{
  "schema_version": "1.1",
  "id": "s4:https://vercel.com/changelog/deepseek-v4-1-flash-now-available-on-ai-gateway",
  "slug": "deepseek-v4-1-flash-now-available-on-ai-gateway-11mmjw1",
  "url": "https://feed7.dev/p/deepseek-v4-1-flash-now-available-on-ai-gateway-11mmjw1",
  "title": "DeepSeek V4.1 Flash now available on AI Gateway",
  "why_included": "DeepSeek V4.1 Flash brings vision, tool use, reasoning, and prompt caching to Vercel AI Gateway, with direct setup paths for Claude Code, Codex, and Cursor.",
  "summary": "**DeepSeek V4.1 Flash** is available through Vercel AI Gateway with mixed text-and-image input, reasoning, tool use, and prompt caching. It has a **1 million-token context window** and supports outputs up to **384,000 tokens**.",
  "practical_implication": "Builders can select deepseek/deepseek-v4.1-flash after running the latest Vercel CLI setup, then test screenshot reading or chart extraction inside Claude Code, Codex, or Cursor. Gateway controls cover usage, cost, retries, failover, budgets, and routing.",
  "agent_context": "**DeepSeek V4.1 Flash** is available through Vercel AI Gateway with mixed text-and-image input, reasoning, tool use, and prompt caching. It has a **1 million-token context window** and supports outputs up to **384,000 tokens**.\n\nBuilders can select deepseek/deepseek-v4.1-flash after running the latest Vercel CLI setup, then test screenshot reading or chart extraction inside Claude Code, Codex, or Cursor. Gateway controls cover usage, cost, retries, failover, budgets, and routing.\n\nThe material gives architecture and capacity claims but no coding-agent benchmarks, latency measurements, or workload-specific quality results. Large stated limits do not establish reliable performance at those limits.",
  "source": {
    "name": "Vercel",
    "url": "https://vercel.com/changelog/deepseek-v4-1-flash-now-available-on-ai-gateway",
    "published_at": "2026-09-09T00:00:00.000Z"
  },
  "source_class": "blog_post",
  "content_type": "Engineering Post",
  "layer": "infra",
  "domains": [
    "coding",
    "image"
  ],
  "topics": [
    "gateways",
    "model-selection",
    "coding-agents"
  ],
  "verification": {
    "status": "official_source",
    "label": "Official Source",
    "method": "source_feed",
    "verified_at": null
  },
  "uncertainty": [
    "The material gives architecture and capacity claims but no coding-agent benchmarks, latency measurements, or workload-specific quality results. Large stated limits do not establish reliable performance at those limits."
  ],
  "connected_context": {
    "meaning": "This expands the gateway’s multimodal coding-agent pool with unusually large stated context and output limits, prompt caching, and existing routing controls. It strengthens the case for testing consolidated screenshot, chart, and code workflows through one client, but does not distinguish DeepSeek on quality, latency, reliability, or cost; capacity claims remain evaluation inputs, not routing evidence.",
    "corpus_size": 732,
    "generated_at": "2026-09-10T10:07:15.852Z",
    "connections": [
      {
        "title": "Qwen 3.8 Max now available on Vercel AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/qwen-3-8-max-now-available-on-vercel-ai-gateway",
        "feed7_url": "https://feed7.dev/p/qwen-3-8-max-now-available-on-vercel-ai-gateway-1ikih0e",
        "reason": "Both consolidate text, vision, and long-context work behind the same gateway, making matched workload tests—not feature availability—the useful basis for choosing between them."
      },
      {
        "title": "Qwen 3.8 Max 0902 now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/qwen-3-8-max-0902-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/qwen-3-8-max-0902-now-available-on-ai-gateway-0zxju42",
        "reason": "The pinned Qwen snapshot highlights a reproducibility option absent from the supplied DeepSeek description, which matters when evaluating behavior over time."
      },
      {
        "title": "AI Gateway adds unified fast mode support",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/ai-gateway-adds-unified-fast-mode-support",
        "feed7_url": "https://feed7.dev/p/ai-gateway-adds-unified-fast-mode-support-144dq26",
        "reason": "Fast mode is an orthogonal gateway control that could test DeepSeek’s latency behavior, while routing metadata is needed to verify whether accelerated serving was actually used."
      }
    ]
  },
  "lifecycle": "Current",
  "published_at": "2026-09-09T00:00:00.000Z",
  "modified_at": "2026-09-09T00:00:00.000Z",
  "supersedes": [],
  "expires_at": null,
  "formats": {
    "html": "https://feed7.dev/p/deepseek-v4-1-flash-now-available-on-ai-gateway-11mmjw1",
    "json": "https://feed7.dev/p/deepseek-v4-1-flash-now-available-on-ai-gateway-11mmjw1.json",
    "markdown": "https://feed7.dev/p/deepseek-v4-1-flash-now-available-on-ai-gateway-11mmjw1.md"
  }
}