{
  "schema_version": "1.1",
  "id": "s4:https://vercel.com/changelog/mimo-v2-6-models-now-available-on-ai-gateway",
  "slug": "mimo-v2-6-models-now-available-on-ai-gateway-13qk4zg",
  "url": "https://feed7.dev/p/mimo-v2-6-models-now-available-on-ai-gateway-13qk4zg",
  "title": "MiMo V2.6 models now available on AI Gateway",
  "why_included": "Vercel AI Gateway adds three MiMo V2.6 variants spanning heavier agent work, efficient multimodal automation, and lower-latency Pro inference.",
  "summary": "AI Gateway now offers MiMo V2.6 Pro, Flash, and Pro UltraSpeed. All have a **1M-token context** and up to **128K output tokens**; Pro uses 1.02T total and 42B active parameters, while Flash uses 309B total and 15B active.",
  "practical_implication": "Choose Pro for complex or long-running software work, Flash for more efficient everyday automation, and **Pro UltraSpeed** when interactive latency matters; Vercel says it serves Pro at **up to 20× output speed**.",
  "agent_context": "AI Gateway now offers MiMo V2.6 Pro, Flash, and Pro UltraSpeed. All have a **1M-token context** and up to **128K output tokens**; Pro uses 1.02T total and 42B active parameters, while Flash uses 309B total and 15B active.\n\nChoose Pro for complex or long-running software work, Flash for more efficient everyday automation, and **Pro UltraSpeed** when interactive latency matters; Vercel says it serves Pro at **up to 20× output speed**.\n\nThe material lists architecture and serving claims but no quality, latency, or cost comparisons, so model selection still needs workload-specific evaluation.",
  "source": {
    "name": "Vercel",
    "url": "https://vercel.com/changelog/mimo-v2-6-models-now-available-on-ai-gateway",
    "published_at": "2026-09-21T00:00:00.000Z"
  },
  "source_class": "blog_post",
  "content_type": "Engineering Post",
  "layer": "model",
  "domains": [
    "coding"
  ],
  "topics": [
    "model-selection",
    "reasoning",
    "tool-use"
  ],
  "verification": {
    "status": "official_source",
    "label": "Official Source",
    "method": "source_feed",
    "verified_at": null
  },
  "uncertainty": [
    "The material lists architecture and serving claims but no quality, latency, or cost comparisons, so model selection still needs workload-specific evaluation."
  ],
  "connected_context": {
    "meaning": "MiMo V2.6 adds three routing points around one large-context family: capability-oriented Pro, efficiency-oriented Flash, and latency-oriented Pro UltraSpeed. This expands rather than resolves model selection; the architecture, context, output, and serving claims define useful trial dimensions, but the prior candidates reinforce that matched tests of quality, tool use, latency, reliability, and total cost remain necessary.",
    "corpus_size": 843,
    "generated_at": "2026-09-22T09:07:46.298Z",
    "connections": [
      {
        "title": "GPT 5.6 Sol, Luna, and Terra now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/gpt-5-6-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/gpt-5-6-now-available-on-ai-gateway-106pgsr",
        "reason": "Both families expose flagship, balanced, and faster or lower-cost routing choices behind the same gateway, reinforcing tiered model selection by workload rather than one universal default."
      },
      {
        "title": "Inkling Small from Thinking Machines is now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/inkling-small-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/inkling-small-now-available-on-ai-gateway-1a9781l",
        "reason": "Inkling Small provides a contrasting lower-compute candidate for coding and tool use, making MiMo Flash’s efficiency positioning testable against another compact route rather than against Pro alone."
      },
      {
        "title": "Qwen 3.8 Flash now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/qwen-3-8-flash-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/qwen-3-8-flash-now-available-on-ai-gateway-1skoa7y",
        "reason": "Qwen 3.8 Flash shares a 1M-token context and coding-agent positioning, showing that capacity alone cannot distinguish it from MiMo Flash; matched quality, output-limit, latency, price, and reliability tests are required."
      },
      {
        "title": "GLM 5.3 now available on AI Gateway",
        "source_name": "Vercel",
        "source_url": "https://vercel.com/changelog/glm-5-3-now-available-on-ai-gateway",
        "feed7_url": "https://feed7.dev/p/glm-5-3-now-available-on-ai-gateway-0s7o9zv",
        "reason": "GLM 5.3’s similar million-token, long-horizon engineering positioning reinforces that repository-scale claims and token capacity are evaluation inputs, not sufficient grounds for changing production defaults."
      }
    ]
  },
  "lifecycle": "Current",
  "published_at": "2026-09-21T00:00:00.000Z",
  "modified_at": "2026-09-21T00:00:00.000Z",
  "supersedes": [],
  "expires_at": null,
  "formats": {
    "html": "https://feed7.dev/p/mimo-v2-6-models-now-available-on-ai-gateway-13qk4zg",
    "json": "https://feed7.dev/p/mimo-v2-6-models-now-available-on-ai-gateway-13qk4zg.json",
    "markdown": "https://feed7.dev/p/mimo-v2-6-models-now-available-on-ai-gateway-13qk4zg.md"
  }
}