{
  "asOf": "2026-09-04",
  "rows": [
    {
      "model": "DeepSeek V4 Flash 0731",
      "provider": "DeepSeek",
      "apiId": "deepseek-v4-flash",
      "releaseDate": "2026-07-31",
      "category": "efficient",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1048576,
      "launchInputPerMtok": 0.14,
      "launchOutputPerMtok": 0.28,
      "aaIntelligence": 50,
      "oneLiner": "A re-post-trained V4 Flash at the same $0.14/$0.28 price, scoring 50 on the Artificial Analysis Intelligence Index, 6 points above the pricier V4 Pro. API public beta; 0731 weights not yet posted.",
      "sourceUrl": "https://api-docs.deepseek.com/updates",
      "sourceType": "provider",
      "verifiedDate": "2026-07-31"
    },
    {
      "model": "Claude Opus 5",
      "provider": "Anthropic",
      "apiId": "claude-opus-5",
      "releaseDate": "2026-07-24",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 5,
      "launchOutputPerMtok": 25,
      "aaIntelligence": 61,
      "oneLiner": "Took the top spot on the Artificial Analysis Intelligence Index at 61, and did it at half the per-token price of Claude Fable 5.",
      "sourceUrl": "https://www.anthropic.com/news/claude-opus-5",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "FLUX 3",
      "provider": "Black Forest Labs",
      "apiId": "",
      "releaseDate": "2026-07-23",
      "category": "image-video",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "oneLiner": "One network spanning image, video, audio and robot action prediction. Video runs in early access; the open-weight Dev build is promised later in 2026.",
      "sourceUrl": "https://www.globenewswire.com/news-release/2026/07/23/3332364/0/en/black-forest-labs-unveils-flux-3-a-new-multimodal-frontier-model-for-visual-intelligence.html",
      "sourceType": "provider-press-release",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Ling-3.0-flash",
      "provider": "Ant Group",
      "apiId": "",
      "releaseDate": "2026-07-23",
      "category": "efficient",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 256000,
      "oneLiner": "A 124B mixture-of-experts firing only 5.1B parameters per token, which Ant says matches its own 1T flagship. Announced as open weight, but shipped API-only with the weights still unpublished.",
      "sourceUrl": "https://www.businesswire.com/news/home/20260726584441/en/Ant-Group-Unveils-Ling-3.0-Flash-Delivering-Top-Tier-Performance-at-a-Fraction-of-the-Parameter-Scale",
      "sourceType": "provider-press-release",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Gemini 3.6 Flash",
      "provider": "Google",
      "apiId": "",
      "releaseDate": "2026-07-21",
      "category": "efficient",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 1.5,
      "launchOutputPerMtok": 7.5,
      "oneLiner": "Google's workhorse tier, cut to $7.50 output from the $9.00 that Gemini 3.5 Flash charged, on a claimed 17% drop in output tokens per task.",
      "sourceUrl": "https://blog.google/innovation-and-ai/models-and-research/gemini-models/gemini-3-6-flash-3-5-flash-lite-3-5-flash-cyber/",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Gemini 3.5 Flash-Lite",
      "provider": "Google",
      "apiId": "",
      "releaseDate": "2026-07-21",
      "category": "efficient",
      "significance": "incremental",
      "openWeight": false,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 0.3,
      "launchOutputPerMtok": 2.5,
      "oneLiner": "The cheapest launch price of any July model at $0.30 in, aimed at classification, extraction and routing rather than reasoning.",
      "sourceUrl": "https://blog.google/innovation-and-ai/models-and-research/gemini-models/gemini-3-6-flash-3-5-flash-lite-3-5-flash-cyber/",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Gemini 3.5 Flash Cyber",
      "provider": "Google",
      "apiId": "",
      "releaseDate": "2026-07-21",
      "category": "specialist",
      "significance": "incremental",
      "openWeight": false,
      "license": "",
      "oneLiner": "Tuned to find and patch software vulnerabilities, and restricted to governments and trusted partners through the CodeMender pilot.",
      "sourceUrl": "https://blog.google/innovation-and-ai/models-and-research/gemini-models/gemini-3-6-flash-3-5-flash-lite-3-5-flash-cyber/",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Laguna S 2.1",
      "provider": "poolside",
      "apiId": "",
      "releaseDate": "2026-07-21",
      "category": "open-weight",
      "significance": "notable",
      "openWeight": true,
      "license": "OpenMDW-1.1",
      "contextWindow": 1000000,
      "launchInputPerMtok": 0.1,
      "launchOutputPerMtok": 0.2,
      "oneLiner": "118B total and 8B active, released under OpenMDW-1.1 at $0.10 in and $0.20 out, the cheapest agentic coding model to ship in July.",
      "sourceUrl": "https://poolside.ai/blog/introducing-laguna-s-2-1",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Qwen-Image-3.0",
      "provider": "Alibaba",
      "apiId": "",
      "releaseDate": "2026-07-21",
      "category": "image-video",
      "significance": "incremental",
      "openWeight": false,
      "license": "",
      "oneLiner": "Closed and invite-only at launch, with no model card, weights or published benchmarks to check the multi-panel layout claims against.",
      "sourceUrl": "https://www.digitalapplied.com/blog/seven-days-seven-releases-july-2026-model-wave",
      "sourceType": "secondary",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Qwen-Audio-3.0-TTS Plus",
      "provider": "Alibaba",
      "apiId": "",
      "releaseDate": "2026-07-20",
      "category": "speech",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "oneLiner": "Took first place on the Artificial Analysis text-to-speech arena. Billed per character, not per token, at roughly a third of what ElevenLabs charges.",
      "sourceUrl": "https://www.marktechpost.com/2026/07/20/alibabas-tongyi-lab-releases-qwen-audio-3-0-tts-a-hosted-text-to-speech-model-in-flash-and-plus-tiers-across-16-languages/",
      "sourceType": "secondary",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Qwen-Audio-3.0-TTS Flash",
      "provider": "Alibaba",
      "apiId": "",
      "releaseDate": "2026-07-20",
      "category": "speech",
      "significance": "incremental",
      "openWeight": false,
      "license": "",
      "oneLiner": "The real-time tier of the same text-to-speech release, covering 16 languages and 20 Chinese dialect regions, hosted only on Alibaba Cloud Model Studio.",
      "sourceUrl": "https://www.marktechpost.com/2026/07/20/alibabas-tongyi-lab-releases-qwen-audio-3-0-tts-a-hosted-text-to-speech-model-in-flash-and-plus-tiers-across-16-languages/",
      "sourceType": "secondary",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Qwen3.8-Max-Preview",
      "provider": "Alibaba",
      "apiId": "",
      "releaseDate": "2026-07-19",
      "category": "frontier",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "oneLiner": "A 2.4T-parameter preview shown at the World AI Conference with no model card, no license and no independent benchmarks published alongside it.",
      "sourceUrl": "https://www.marktechpost.com/2026/07/19/alibaba-previews-qwen3-8-max-a-2-4-trillion-parameter-multimodal-model-days-after-moonshots-kimi-k3-open-weight-launch/",
      "sourceType": "secondary",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Kimi K3",
      "provider": "Moonshot",
      "apiId": "kimi-k3",
      "releaseDate": "2026-07-16",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": true,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 3,
      "launchOutputPerMtok": 15,
      "aaIntelligence": 57,
      "oneLiner": "2.8T parameters and the highest-scoring open-weight model of the month. Weights followed on July 26, a day inside Moonshot's own deadline.",
      "sourceUrl": "https://platform.kimi.ai/docs/pricing/chat-k3",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Inkling",
      "provider": "Thinking Machines",
      "apiId": "thinkingmachines/Inkling",
      "releaseDate": "2026-07-15",
      "category": "open-weight",
      "significance": "notable",
      "openWeight": true,
      "license": "Apache 2.0",
      "contextWindow": 1000000,
      "launchInputPerMtok": 1.87,
      "launchOutputPerMtok": 4.68,
      "aaIntelligence": 41,
      "oneLiner": "975B total and 41B active under Apache 2.0, the most permissive license anyone has attached to a model at that parameter scale.",
      "sourceUrl": "https://tinker-docs.thinkingmachines.ai/tinker/models/",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "GPT-5.6 Sol",
      "provider": "OpenAI",
      "apiId": "",
      "releaseDate": "2026-07-09",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 5,
      "launchOutputPerMtok": 30,
      "aaIntelligence": 59,
      "oneLiner": "The reasoning tier of the GPT-5.6 line. Launched at $5/$30, then cut to $4/$20 on Aug 21, 2026 (promotional through at least Nov 21) after being excluded from the July 30 Terra and Luna cuts.",
      "sourceUrl": "https://developers.openai.com/api/docs/pricing",
      "sourceType": "provider",
      "verifiedDate": "2026-08-25"
    },
    {
      "model": "GPT-5.6 Terra",
      "provider": "OpenAI",
      "apiId": "",
      "releaseDate": "2026-07-09",
      "category": "frontier",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 2.5,
      "launchOutputPerMtok": 15,
      "aaIntelligence": 55,
      "oneLiner": "The middle tier, launched at half of Sol and within four AA index points of it. Cut 20% on July 30, 2026 to $2/$12, half of Sol's promotional $4/$20.",
      "sourceUrl": "https://developers.openai.com/api/docs/pricing",
      "sourceType": "provider",
      "verifiedDate": "2026-07-31"
    },
    {
      "model": "GPT-5.6 Luna",
      "provider": "OpenAI",
      "apiId": "",
      "releaseDate": "2026-07-09",
      "category": "efficient",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 1,
      "launchOutputPerMtok": 6,
      "aaIntelligence": 51,
      "oneLiner": "The high-volume tier, launched at $1 in and $6 out, the best intelligence-per-dollar OpenAI shipped in July. Cut 80% on July 30 to $0.20/$1.20, under the retired GPT-5.4 nano floor.",
      "sourceUrl": "https://developers.openai.com/api/docs/pricing",
      "sourceType": "provider",
      "verifiedDate": "2026-07-31"
    },
    {
      "model": "Muse Spark 1.1",
      "provider": "Meta",
      "apiId": "muse-spark-1.1",
      "releaseDate": "2026-07-09",
      "category": "multimodal",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "oneLiner": "Meta Superintelligence Labs shipped its first paid model, breaking the open-weight-only posture Meta had held since Llama.",
      "sourceUrl": "https://ai.meta.com/blog/introducing-muse-spark-meta-model-api/",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Grok 4.5",
      "provider": "xAI",
      "apiId": "",
      "releaseDate": "2026-07-08",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 500000,
      "launchInputPerMtok": 2,
      "launchOutputPerMtok": 6,
      "aaIntelligence": 54,
      "oneLiner": "Trained on real Cursor session data and priced at $2 in and $6 out, well under half of what Opus 4.8 and GPT-5.5 charged at the time.",
      "sourceUrl": "https://x.ai/news/grok-4-5",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "SWE-1.7",
      "provider": "Cognition",
      "apiId": "",
      "releaseDate": "2026-07-08",
      "category": "specialist",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "oneLiner": "Post-trained on top of Moonshot's already RL-heavy Kimi K2.7 Code, and served inside Devin only. Not sold as a standalone API.",
      "sourceUrl": "https://cognition.com/blog/swe-1-7",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Cohere Transcribe Arabic",
      "provider": "Cohere",
      "apiId": "",
      "releaseDate": "2026-07-07",
      "category": "speech",
      "significance": "incremental",
      "openWeight": true,
      "license": "Apache 2.0",
      "oneLiner": "A 2B open-weight speech recognition model built for Arabic dialect variation and Arabic-English code-switching.",
      "sourceUrl": "https://cohere.com/blog/transcribe-arabic",
      "sourceType": "provider",
      "verifiedDate": "2026-07-29"
    },
    {
      "model": "Gemini 3.7 Flash",
      "provider": "Google",
      "apiId": "gemini-3.7-flash",
      "releaseDate": "2026-08-13",
      "category": "efficient",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "launchInputPerMtok": 0.75,
      "launchOutputPerMtok": 3.75,
      "aaIntelligence": 56,
      "oneLiner": "Launched at an introductory $0.75/$3.75 that Google states doubles on 1 January 2027, with DeepSWE v1.1 up to 65.3% from 3.6 Flash 49.0%.",
      "sourceUrl": "https://blog.google/innovation-and-ai/models-and-research/gemini-models/introducing-gemini-3-7-flash/",
      "sourceType": "provider",
      "verifiedDate": "2026-08-14"
    },
    {
      "model": "Grok 4.6",
      "provider": "xAI",
      "apiId": "grok-4-6",
      "releaseDate": "2026-08-12",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 500000,
      "launchInputPerMtok": 2,
      "launchOutputPerMtok": 6,
      "aaIntelligence": 61,
      "oneLiner": "Arrived 35 days after Grok 4.5 at the same $2/$6, but with cached input raised from $0.30 to $0.50 per Mtok. Ties GPT-5.6 Sol at 61 on the AA index.",
      "sourceUrl": "https://x.ai/news/grok-4-6",
      "sourceType": "provider",
      "verifiedDate": "2026-08-14"
    },
    {
      "model": "Muse Glimmer",
      "provider": "Meta",
      "apiId": "muse-glimmer-30b",
      "releaseDate": "2026-08-10",
      "category": "open-weight",
      "significance": "flagship",
      "openWeight": true,
      "license": "Apache 2.0",
      "contextWindow": 131072,
      "aaIntelligence": 35,
      "oneLiner": "Meta returns to a plain Apache 2.0 licence with a 29.6B dense multimodal model that fits one consumer GPU at 4-bit. Built for local agents, not the leaderboard.",
      "sourceUrl": "https://venturebeat.com/technology/meta-returns-to-open-source-with-muse-glimmer-an-apache-2-0-licensed-30b-parameter-ai-model-optimized-for-agents-available-now",
      "sourceType": "secondary",
      "verifiedDate": "2026-08-14"
    },
    {
      "model": "Muse Spark 1.2",
      "provider": "Meta",
      "apiId": "muse-spark-1.2",
      "releaseDate": "2026-08-05",
      "category": "frontier",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 1.25,
      "launchOutputPerMtok": 4.25,
      "aaIntelligence": 57,
      "oneLiner": "Meta holds the 1.1 rate of $1.25/$4.25 and adds a $0.10/$0.20 contributor tier, plus Muse Code, a terminal coding agent, in beta.",
      "sourceUrl": "https://developer.meta.com/ai/models/muse-spark/",
      "sourceType": "secondary",
      "verifiedDate": "2026-08-14"
    },
    {
      "model": "Qwen3.8-Max",
      "provider": "Alibaba",
      "apiId": "qwen3.8-max",
      "releaseDate": "2026-08-03",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": true,
      "license": "qwen3.8-max (custom, revenue-gated)",
      "contextWindow": 262144,
      "launchInputPerMtok": 2,
      "launchOutputPerMtok": 6,
      "aaIntelligence": 58,
      "oneLiner": "The first Max-class Qwen with published weights (13 August) and, at AA 58, the highest-scoring open-weights model. The licence is custom, not Apache 2.0.",
      "sourceUrl": "https://huggingface.co/Qwen/Qwen3.8-2.4T-A95B",
      "sourceType": "provider",
      "verifiedDate": "2026-08-14"
    },
    {
      "model": "GLM-5.3",
      "provider": "Zhipu",
      "apiId": "glm-5.3",
      "releaseDate": "2026-08-14",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1048576,
      "launchInputPerMtok": 1.4,
      "launchOutputPerMtok": 4.4,
      "oneLiner": "Same base as GLM-5.2, all gains from post-training; Terminal-Bench 3.0 jumps 4.6 to 28.3. Unpriced at launch, then listed at the GLM-5.2 rate. Weights still held back.",
      "sourceUrl": "https://docs.z.ai/guides/llm/glm-5.3",
      "sourceType": "provider",
      "verifiedDate": "2026-08-21"
    },
    {
      "model": "Qwen3.8-27B",
      "provider": "Alibaba",
      "apiId": "qwen3.8-27b",
      "releaseDate": "2026-08-14",
      "category": "open-weight",
      "significance": "flagship",
      "openWeight": true,
      "license": "Apache 2.0",
      "contextWindow": 262144,
      "oneLiner": "The smaller Qwen3.8 sibling, delivered on the promised week: dense 27B native vision-language, Apache 2.0 with no revenue gate, unlike Qwen3.8-Max.",
      "sourceUrl": "https://huggingface.co/Qwen/Qwen3.8-27B",
      "sourceType": "provider",
      "verifiedDate": "2026-08-21"
    },
    {
      "model": "Hy-MT2-30B-A3B",
      "provider": "Tencent",
      "apiId": "hy-mt2-30b-a3b",
      "releaseDate": "2026-08-20",
      "category": "specialist",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 8192,
      "launchInputPerMtok": 0.074,
      "launchOutputPerMtok": 0.295,
      "oneLiner": "Translation flagship, about 3B active params, 33 language pairs plus five Chinese dialect pairs. An 8K context window is the point: a tool, not an agent.",
      "sourceUrl": "https://openrouter.ai/models?order=newest",
      "sourceType": "secondary",
      "verifiedDate": "2026-08-21"
    },
    {
      "model": "Hy-MT2-1.8B",
      "provider": "Tencent",
      "apiId": "hy-mt2-1.8b",
      "releaseDate": "2026-08-20",
      "category": "specialist",
      "significance": "incremental",
      "openWeight": false,
      "license": "",
      "contextWindow": 8192,
      "launchInputPerMtok": 0.044,
      "launchOutputPerMtok": 0.177,
      "oneLiner": "Compact sibling of Hy-MT2-30B-A3B with the same language coverage, at the cheapest published output rate of any model released in August 2026.",
      "sourceUrl": "https://openrouter.ai/models?order=newest",
      "sourceType": "secondary",
      "verifiedDate": "2026-08-21"
    },
    {
      "model": "GLM-5.3-Flash",
      "provider": "Zhipu",
      "apiId": "glm-5.3-flash",
      "releaseDate": "2026-08-26",
      "category": "open-weight",
      "significance": "flagship",
      "openWeight": true,
      "license": "MIT",
      "contextWindow": 1048576,
      "launchInputPerMtok": 0.15,
      "launchOutputPerMtok": 0.5,
      "aaIntelligence": 57,
      "oneLiner": "Z.ai's reveal of the Ox Alpha stealth model: a 320B-total, 18B-active MoE, MIT-licensed with weights on Hugging Face, scoring 57 on the Artificial Analysis Intelligence Index.",
      "sourceUrl": "https://docs.z.ai/guides/llm/glm-5.3-flash",
      "sourceType": "provider",
      "verifiedDate": "2026-08-31"
    },
    {
      "model": "Qwen3.8-Flash-Next",
      "provider": "Alibaba",
      "apiId": "",
      "releaseDate": "2026-08-26",
      "category": "open-weight",
      "significance": "flagship",
      "openWeight": true,
      "license": "Qwen Community 1.0",
      "contextWindow": 262144,
      "launchInputPerMtok": 0.15,
      "launchOutputPerMtok": 0.47,
      "oneLiner": "An open-weight preview of the Qwen4 architecture: 125B total, 6B active, plus a 51B n-gram embedding layer. Multimodal in, text out, 262K native context extensible to 1M.",
      "sourceUrl": "https://huggingface.co/Qwen/Qwen3.8-Flash-Next",
      "sourceType": "provider",
      "verifiedDate": "2026-08-31"
    },
    {
      "model": "DeepSeek V4 Flash Vision Exp",
      "provider": "DeepSeek",
      "apiId": "",
      "releaseDate": "2026-08-21",
      "category": "multimodal",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1048576,
      "launchInputPerMtok": 0.44,
      "launchOutputPerMtok": 1.32,
      "oneLiner": "DeepSeek's first multimodal vision model, experimental, priced exactly as V4 Flash: $0.44/$1.32 per Mtok at peak and half that off-peak. Images bill as input tokens by dimension.",
      "sourceUrl": "https://api-docs.deepseek.com/updates",
      "sourceType": "provider",
      "verifiedDate": "2026-08-31"
    },
    {
      "model": "Claude Fable 5",
      "provider": "Anthropic",
      "apiId": "claude-fable-5",
      "releaseDate": "2026-06-09",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 10,
      "launchOutputPerMtok": 50,
      "oneLiner": "Anthropic's Mythos-class flagship made safe for general use, at $10/$50 per Mtok. Suspended three days later by a US export-control directive, restored after that order lifted on June 30.",
      "sourceUrl": "https://www.anthropic.com/news/claude-fable-5-mythos-5",
      "sourceType": "provider",
      "verifiedDate": "2026-08-31"
    },
    {
      "model": "Claude Mythos 5",
      "provider": "Anthropic",
      "apiId": "",
      "releaseDate": "2026-06-09",
      "category": "frontier",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "launchInputPerMtok": 10,
      "launchOutputPerMtok": 50,
      "oneLiner": "The same underlying model as Fable 5 with safeguards lifted in some areas, at the same $10/$50. Restricted to Project Glasswing partners and select researchers, never generally available.",
      "sourceUrl": "https://www.anthropic.com/news/claude-fable-5-mythos-5",
      "sourceType": "provider",
      "verifiedDate": "2026-08-31"
    },
    {
      "model": "North Mini Code",
      "provider": "Cohere",
      "apiId": "north-mini-code",
      "releaseDate": "2026-06-09",
      "category": "open-weight",
      "significance": "notable",
      "openWeight": true,
      "license": "Apache 2.0",
      "contextWindow": 262144,
      "oneLiner": "Cohere's first developer model: a 30B-total, 3B-active sparse MoE coding agent, Apache 2.0, running on one H100 at FP8. Free on hosted endpoints, so it launched with no rate card.",
      "sourceUrl": "https://cohere.com/blog/north-mini-code",
      "sourceType": "provider",
      "verifiedDate": "2026-06-23"
    },
    {
      "model": "Claude Sonnet 5",
      "provider": "Anthropic",
      "apiId": "claude-sonnet-5",
      "releaseDate": "2026-06-30",
      "category": "efficient",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 2,
      "launchOutputPerMtok": 10,
      "oneLiner": "Anthropic's mid-tier Claude at a $2/$10 introductory rate billed as expiring on August 31. Anthropic cancelled that rise in August and made $2/$10 the standard price.",
      "sourceUrl": "https://platform.claude.com/docs/en/about-claude/pricing",
      "sourceType": "provider",
      "verifiedDate": "2026-08-14"
    },
    {
      "model": "GLM-5.2",
      "provider": "Zhipu",
      "apiId": "",
      "releaseDate": "2026-06-16",
      "category": "open-weight",
      "significance": "flagship",
      "openWeight": true,
      "license": "MIT",
      "launchInputPerMtok": 1.4,
      "launchOutputPerMtok": 4.4,
      "oneLiner": "Z.ai's coding flagship, MIT-licensed with weights on Hugging Face, at $1.40/$4.40 per Mtok. The same rate GLM-5.1 carried and the same rate GLM-5.3 would carry three months later.",
      "sourceUrl": "https://huggingface.co/zai-org/GLM-5.2",
      "sourceType": "provider",
      "verifiedDate": "2026-08-31"
    },
    {
      "model": "Claude Fable 5.1",
      "provider": "Anthropic",
      "apiId": "claude-fable-5-1",
      "releaseDate": "2026-09-01",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 10,
      "launchOutputPerMtok": 50,
      "oneLiner": "Held Fable 5's $10/$50 per Mtok and cut only the cache read, $1.00 to $0.25, a 0.025x multiplier that undercuts Opus 5's cache rate at twice the base input price.",
      "sourceUrl": "https://www.anthropic.com/claude-fable-and-mythos-5-1",
      "sourceType": "provider",
      "verifiedDate": "2026-09-02"
    },
    {
      "model": "Claude Mythos 5.1",
      "provider": "Anthropic",
      "apiId": "",
      "releaseDate": "2026-09-01",
      "category": "frontier",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 10,
      "launchOutputPerMtok": 50,
      "oneLiner": "The same model as Fable 5.1 with safeguards lifted in some areas, at the same $10/$50 and the same $0.25 cache read. Project Glasswing invitation only, never generally available.",
      "sourceUrl": "https://www.anthropic.com/claude-fable-and-mythos-5-1",
      "sourceType": "provider",
      "verifiedDate": "2026-09-02"
    },
    {
      "model": "Gemini 3.8 Flash",
      "provider": "Google",
      "apiId": "gemini-3.8-flash",
      "releaseDate": "2026-09-02",
      "category": "efficient",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1000000,
      "launchInputPerMtok": 0.75,
      "launchOutputPerMtok": 3.75,
      "aaIntelligence": 59,
      "oneLiner": "Third straight Flash generation at $0.75/$3.75 per Mtok, and the third sharing one expiry: the rate doubles on 2027-01-01. AA Intelligence Index 59 at high effort, against 56 for 3.7 Flash.",
      "sourceUrl": "https://blog.google/innovation-and-ai/models-and-research/gemini-models/3-8-flash-and-3-8-flash-cyber/",
      "sourceType": "provider",
      "verifiedDate": "2026-09-03"
    },
    {
      "model": "Gemini 3.8 Flash Cyber",
      "provider": "Google",
      "apiId": "",
      "releaseDate": "2026-09-02",
      "category": "specialist",
      "significance": "notable",
      "openWeight": false,
      "license": "",
      "oneLiner": "Vulnerability-hunting twin of 3.8 Flash, restricted to the Fairwind Program for governments, critical infrastructure operators and software maintainers. No published token rate.",
      "sourceUrl": "https://blog.google/innovation-and-ai/models-and-research/gemini-models/3-8-flash-and-3-8-flash-cyber/",
      "sourceType": "provider",
      "verifiedDate": "2026-09-03"
    },
    {
      "model": "Muse Spark 1.3",
      "provider": "Meta",
      "apiId": "muse-spark-1.3",
      "releaseDate": "2026-09-02",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1048576,
      "launchInputPerMtok": 1.25,
      "launchOutputPerMtok": 4.25,
      "aaIntelligence": 61,
      "oneLiner": "Meta held the 1.2 rate card exactly at $1.25/$4.25 per Mtok and moved AA Intelligence Index 57 to 61 at xhigh. The $0.10/$0.20 contributor tier buys its discount with training rights.",
      "sourceUrl": "https://research.meta.ai/blog/introducing-muse-spark-1-3",
      "sourceType": "provider",
      "verifiedDate": "2026-09-03"
    },
    {
      "model": "GPT-6 Astra",
      "provider": "OpenAI",
      "apiId": "gpt-6-astra",
      "releaseDate": "2026-09-03",
      "category": "frontier",
      "significance": "flagship",
      "openWeight": false,
      "license": "",
      "contextWindow": 1050000,
      "launchInputPerMtok": 10,
      "launchOutputPerMtok": 50,
      "aaIntelligence": 61,
      "oneLiner": "OpenAI's new flagship at $10/$50 per Mtok, 2.5x GPT-5.6 Sol and level with Claude Fable 5.1. Tops Terminal-Bench 4.0 at 58.18% while burning 1.53B tokens to Fable 5.1's 2.75B.",
      "sourceUrl": "https://developers.openai.com/api/docs/pricing",
      "sourceType": "provider",
      "verifiedDate": "2026-09-04"
    }
  ]
}