{
  "schema_version": "1.1",
  "id": "s4:https://vercel.com/changelog/gemini-3-8-flash-now-available-on-ai-gateway",
  "slug": "gemini-3-8-flash-now-available-on-ai-gateway-0cun77q",
  "url": "https://feed7.dev/p/gemini-3-8-flash-now-available-on-ai-gateway-0cun77q",
  "title": "Gemini 3.8 Flash now available on AI Gateway",
  "why_included": "Gemini 3.8 Flash brings multimodal input, tool calling, web search, and default reasoning to coding agents through Vercel. Its temporary 50% discount runs through December 31.",
  "summary": "Google’s **Gemini 3.8 Flash** is now on Vercel AI Gateway with a **1M-token context window** and a 65,536-token output limit. It accepts text, images, PDFs, and video, and supports tool calls and web search.",
  "practical_implication": "Builders can select google/gemini-3.8-flash in connected coding agents. Thinking is enabled by default, so test its latency, token use, and behavior on representative repository tasks before changing a production routing default.",
  "agent_context": "Google’s **Gemini 3.8 Flash** is now on Vercel AI Gateway with a **1M-token context window** and a 65,536-token output limit. It accepts text, images, PDFs, and video, and supports tool calls and web search.\n\nBuilders can select google/gemini-3.8-flash in connected coding agents. Thinking is enabled by default, so test its latency, token use, and behavior on representative repository tasks before changing a production routing default.\n\nVercel says it improves software engineering, agent work, and multi-step reasoning at the prior Flash model’s speed and cost, but supplies no evaluation results here. The advertised **50% discount ends December 31**.",
  "source": {
    "name": "Vercel",
    "url": "https://vercel.com/changelog/gemini-3-8-flash-now-available-on-ai-gateway",
    "published_at": "2026-09-02T00:00:00.000Z"
  },
  "source_class": "blog_post",
  "content_type": "Engineering Post",
  "layer": "model",
  "domains": [
    "coding"
  ],
  "topics": [
    "coding-agents",
    "reasoning",
    "model-selection"
  ],
  "verification": {
    "status": "official_source",
    "label": "Official Source",
    "method": "source_feed",
    "verified_at": null
  },
  "uncertainty": [
    "Vercel says it improves software engineering, agent work, and multi-step reasoning at the prior Flash model’s speed and cost, but supplies no evaluation results here. The advertised **50% discount ends December 31**."
  ],
  "connected_context": {
    "meaning": "This adds another million-token, multimodal Flash route to the same coding-agent model pool, but does not resolve selection among neighboring models. Its tool calls, web search, and large output allowance broaden the test surface; the absent evaluation data and default thinking make repository-matched measurements of quality, latency, and token use the deciding evidence, with the temporary discount separated from durable routing economics.",
    "corpus_size": 669,
    "generated_at": "2026-09-03T09:59:54.952Z",
    "connections": [
      {
        "title": "Qwen 3.8 Flash now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/qwen-3-8-flash-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/qwen-3-8-flash-now-available-on-ai-gateway-1skoa7y",
        "reason": "Qwen 3.8 Flash has the same stated context and output limits, making it a close matched candidate for testing multimodal coding, tool use, latency, price, and reliability rather than choosing by capacity."
      },
      {
        "title": "GLM 5.3 now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/glm-5-3-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/glm-5-3-now-available-on-ai-gateway-0s7o9zv",
        "reason": "GLM 5.3 reinforces that a 1M-token window and provider claims do not distinguish repository-scale routes without controlled workload evaluation."
      },
      {
        "title": "The Base Model Is Dead — Varun Singh, Arcee AI",
        "source_name": "AI Engineer",
        "source_url": "https://www.youtube.com/watch?v=xbPriQWXtWM",
        "feed7_url": "https://feed7.dev/p/the-base-model-is-dead-varun-singh-arcee-ai-02hts76",
        "reason": "The discussion of capability differences established during model training explains why similar gateway specifications may still yield materially different agent behavior."
      }
    ]
  },
  "lifecycle": "Current",
  "published_at": "2026-09-02T00:00:00.000Z",
  "modified_at": "2026-09-02T00:00:00.000Z",
  "supersedes": [],
  "expires_at": null,
  "formats": {
    "html": "https://feed7.dev/p/gemini-3-8-flash-now-available-on-ai-gateway-0cun77q",
    "json": "https://feed7.dev/p/gemini-3-8-flash-now-available-on-ai-gateway-0cun77q.json",
    "markdown": "https://feed7.dev/p/gemini-3-8-flash-now-available-on-ai-gateway-0cun77q.md"
  }
}