{
  "schema_version": "1.1",
  "id": "s4:https://vercel.com/changelog/claude-opus-5-5-now-available-on-ai-gateway",
  "slug": "claude-opus-5-5-now-available-on-ai-gateway-0eh7dtl",
  "url": "https://feed7.dev/p/claude-opus-5-5-now-available-on-ai-gateway-0eh7dtl",
  "title": "Claude Opus 5.5 now available on AI Gateway",
  "why_included": "Claude Opus 5.5 reaches Vercel AI Gateway with stronger long-run agent positioning, but adaptive thinking and retired forced tool use can break existing requests with HTTP 400s.",
  "summary": "Claude **Opus 5.5** is available on AI Gateway with a **1M-token context window** and up to **128K output tokens**. Anthropic says it matches Fable 5.1 while running about 30% faster and costing about 40% less per task than Opus 5.",
  "practical_implication": "Audit harnesses before switching: fixed or disabled thinking is rejected, and forced tool selection is retired. Steer thinking with effort and prompts; use structured outputs for JSON, and add retry handling when the model skips a requested tool.",
  "agent_context": "Claude **Opus 5.5** is available on AI Gateway with a **1M-token context window** and up to **128K output tokens**. Anthropic says it matches Fable 5.1 while running about 30% faster and costing about 40% less per task than Opus 5.\n\nAudit harnesses before switching: fixed or disabled thinking is rejected, and forced tool selection is retired. Steer thinking with effort and prompts; use structured outputs for JSON, and add retry handling when the model skips a requested tool.\n\nThe speed and cost figures are provider claims, not results from the supplied material. Adaptive thinking and non-forced tools also reduce deterministic control, so production migrations need error-path and tool-selection tests.",
  "source": {
    "name": "Vercel",
    "url": "https://vercel.com/changelog/claude-opus-5-5-now-available-on-ai-gateway",
    "published_at": "2026-09-22T00:00:00.000Z"
  },
  "source_class": "blog_post",
  "content_type": "Engineering Post",
  "layer": "model",
  "domains": [
    "coding"
  ],
  "topics": [
    "coding-agents",
    "model-selection",
    "tool-use"
  ],
  "verification": {
    "status": "official_source",
    "label": "Official Source",
    "method": "source_feed",
    "verified_at": null
  },
  "uncertainty": [
    "The speed and cost figures are provider claims, not results from the supplied material. Adaptive thinking and non-forced tools also reduce deterministic control, so production migrations need error-path and tool-selection tests."
  ],
  "connected_context": {
    "meaning": "This adds a high-capacity long-run agent option but makes migration behavior, not the provider’s speed and cost claims, the immediate engineering concern. Existing harnesses may fail because fixed thinking is rejected and forced tool selection is gone; production adoption therefore depends on request compatibility, structured-output validation, skipped-tool retries, and matched workload results.",
    "corpus_size": 856,
    "generated_at": "2026-09-23T09:04:48.264Z",
    "connections": [
      {
        "title": "Inkling Small from Thinking Machines is now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/inkling-small-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/inkling-small-now-available-on-ai-gateway-1a9781l",
        "reason": "Both expose adjustable reasoning for tool-heavy work, but Opus 5.5 removes fixed thinking and forced tool selection, making control-path compatibility a sharper evaluation criterion than Inkling Small’s efficiency positioning alone."
      },
      {
        "title": "Qwen 3.8 Flash now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/qwen-3-8-flash-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/qwen-3-8-flash-now-available-on-ai-gateway-1skoa7y",
        "reason": "Qwen 3.8 Flash provides another 1M-context agent route, but with a 65K output ceiling versus Opus 5.5’s 128K; capacity still does not resolve comparative tool reliability, latency, or cost."
      },
      {
        "title": "GLM 5.3 Flash now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/glm-5-3-flash-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/glm-5-3-flash-now-available-on-ai-gateway-1u37q78",
        "reason": "GLM 5.3 Flash reinforces that million-token context, function calling, and structured output are available from adjacent routes, so Opus 5.5’s migration risks and workload performance should drive selection rather than feature presence."
      }
    ]
  },
  "lifecycle": "Current",
  "published_at": "2026-09-22T00:00:00.000Z",
  "modified_at": "2026-09-22T00:00:00.000Z",
  "supersedes": [],
  "expires_at": null,
  "formats": {
    "html": "https://feed7.dev/p/claude-opus-5-5-now-available-on-ai-gateway-0eh7dtl",
    "json": "https://feed7.dev/p/claude-opus-5-5-now-available-on-ai-gateway-0eh7dtl.json",
    "markdown": "https://feed7.dev/p/claude-opus-5-5-now-available-on-ai-gateway-0eh7dtl.md"
  }
}