{
  "schema_version": "1.1",
  "id": "s4:https://vercel.com/changelog/ai-gateway-spend-budgets-and-alerts",
  "slug": "ai-gateway-spend-budgets-and-alerts-1969711",
  "url": "https://feed7.dev/p/ai-gateway-spend-budgets-and-alerts-1969711",
  "title": "AI Gateway now supports team and project spend budgets",
  "why_included": "AI Gateway can now enforce spend caps across a team, project, or API key, giving agent workloads layered cost controls instead of relying on per-key limits alone.",
  "summary": "AI Gateway budgets now cover **team, project, and API key** scopes. Every applicable budget must have room; hitting any limit rejects further requests until the cap resets or is raised.",
  "practical_implication": "Set defaults for projects or keys, then add explicit overrides for expensive agents. Use **50%, 75%, and 100% alerts** to catch rising spend before a hard limit interrupts a run.",
  "agent_context": "AI Gateway budgets now cover **team, project, and API key** scopes. Every applicable budget must have room; hitting any limit rejects further requests until the cap resets or is raised.\n\nSet defaults for projects or keys, then add explicit overrides for expensive agents. Use **50%, 75%, and 100% alerts** to catch rising spend before a hard limit interrupts a run.\n\nAlerts are informational and disabled by default. **BYOK spend is excluded by default**, so enable or account for it separately if the budget must reflect total model usage.",
  "source": {
    "name": "Vercel",
    "url": "https://vercel.com/changelog/ai-gateway-spend-budgets-and-alerts",
    "published_at": "2026-07-31T17:00:00.000Z"
  },
  "source_class": "blog_post",
  "content_type": "Engineering Post",
  "layer": "infra",
  "domains": [
    "coding"
  ],
  "topics": [
    "gateways",
    "observability"
  ],
  "verification": {
    "status": "official_source",
    "label": "Official Source",
    "method": "source_feed",
    "verified_at": null
  },
  "uncertainty": [
    "Alerts are informational and disabled by default. **BYOK spend is excluded by default**, so enable or account for it separately if the budget must reflect total model usage."
  ],
  "connected_context": {
    "meaning": "This adds enforceable cost governance above routing: overlapping team, project, and key budgets can stop an agent even when its model route remains healthy. It makes budget hierarchy, alert enablement, and BYOK inclusion part of production reliability planning, since alerts do not prevent overruns and default accounting may omit external-key spend.",
    "corpus_size": 318,
    "generated_at": "2026-08-01T10:07:06.513Z",
    "connections": [
      {
        "title": "AI Gateway logs now have a dedicated page",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/ai-gateway-logs",
        "feed7_url": "https://feed7.dev/p/ai-gateway-logs-1272t5j",
        "reason": "Per-request cost and routing logs provide the evidence needed to investigate budget consumption and decide which project or API-key limits require adjustment."
      },
      {
        "title": "AI Gateway: GPT-5.6 pricing and speed updates",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/ai-gateway-gpt-5-6-pricing-speed-updates",
        "feed7_url": "https://feed7.dev/p/ai-gateway-gpt-5-6-pricing-speed-updates-06cxxee",
        "reason": "Price changes alter how quickly existing workloads consume fixed budgets even without code or model-ID changes, so caps should be reviewed when gateway pricing shifts."
      },
      {
        "title": "AI Gateway adds unified fast mode support",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/ai-gateway-adds-unified-fast-mode-support",
        "feed7_url": "https://feed7.dev/p/ai-gateway-adds-unified-fast-mode-support-144dq26",
        "reason": "Fast tiers usually cost more, making explicit budget overrides and alerts an operational consequence of enabling latency-focused routing for selected agents."
      },
      {
        "title": "WebSocket support for OpenAI Responses API live on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/websocket-support-for-openai-responses-api-live-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/websocket-support-for-openai-responses-api-live-on-ai-gateway-0oj5npd",
        "reason": "Persistent sessions support long tool-heavy runs, while hard budget exhaustion can interrupt those runs; together they require cost headroom appropriate to extended execution."
      }
    ]
  },
  "lifecycle": "Current",
  "published_at": "2026-07-31T17:00:00.000Z",
  "modified_at": "2026-07-31T17:00:00.000Z",
  "supersedes": [],
  "expires_at": null,
  "formats": {
    "html": "https://feed7.dev/p/ai-gateway-spend-budgets-and-alerts-1969711",
    "json": "https://feed7.dev/p/ai-gateway-spend-budgets-and-alerts-1969711.json",
    "markdown": "https://feed7.dev/p/ai-gateway-spend-budgets-and-alerts-1969711.md"
  }
}