From 9374a8c4a9e5533f91f401e13b6a8e45719f02ad Mon Sep 17 00:00:00 2001 From: Pan YANG Date: Thu, 24 Sep 2026 21:36:10 +0800 Subject: [PATCH] feat(data): refresh AA index to v4.3.2 --- data/artificial-analysis-index.json | 216 ++++++---- data/data-health.json | 396 +++++++++++------- docs/DATA-HEALTH.md | 222 +++++----- manifests/models/claude-opus-5-5.json | 120 ++++++ manifests/models/deepseek-v4-1-flash.json | 117 ++++++ .../models/deepseek-v4-flash-vision.json | 10 +- manifests/models/deepseek-v4-flash.json | 19 +- manifests/models/deepseek-v4-pro.json | 10 +- manifests/models/gpt-6-luna.json | 124 ++++++ manifests/models/gpt-6-sol.json | 124 ++++++ manifests/models/grok-4-7.json | 123 ++++++ manifests/models/mimo-v2-6-pro.json | 124 ++++++ manifests/models/qwen3-8-max-0902.json | 121 ++++++ manifests/models/qwen3-8-max.json | 6 +- manifests/vendors/alibaba.json | 9 +- manifests/vendors/anthropic.json | 5 +- manifests/vendors/deepseek.json | 11 +- manifests/vendors/openai.json | 11 +- manifests/vendors/xai.json | 8 +- manifests/vendors/xiaomi.json | 9 +- src/lib/generated/metadata.ts | 2 +- src/lib/generated/models.ts | 14 + tests/model-intelligence-index.test.ts | 46 +- tests/model-price-intelligence-index.test.ts | 2 +- 24 files changed, 1446 insertions(+), 403 deletions(-) create mode 100644 manifests/models/claude-opus-5-5.json create mode 100644 manifests/models/deepseek-v4-1-flash.json create mode 100644 manifests/models/gpt-6-luna.json create mode 100644 manifests/models/gpt-6-sol.json create mode 100644 manifests/models/grok-4-7.json create mode 100644 manifests/models/mimo-v2-6-pro.json create mode 100644 manifests/models/qwen3-8-max-0902.json diff --git a/data/artificial-analysis-index.json b/data/artificial-analysis-index.json index 564506b4..398449aa 100644 --- a/data/artificial-analysis-index.json +++ b/data/artificial-analysis-index.json @@ -3,21 +3,21 @@ "source": "Artificial Analysis", "sourceUrl": "https://artificialanalysis.ai/leaderboards/models", "methodologyUrl": "https://artificialanalysis.ai/methodology/intelligence-benchmarking", - "indexVersion": "4.3", - "observedAt": "2026-09-09", + "indexVersion": "4.3.2", + "observedAt": "2026-09-24", "legacyMissingModelIds": ["cursor-composer-2", "cursor-composer-2-5"], "entries": [ { "modelId": "claude-fable-5", "score": 50, "estimated": false, - "configuration": "Claude Fable 5 (Adaptive Reasoning, Max Effort, Opus 4.8 Fallback)" + "configuration": "Claude Fable 5 (with fallback)" }, { "modelId": "claude-fable-5-1", "score": 53, "estimated": false, - "configuration": "Claude Fable 5.1 (Adaptive Reasoning, Max Effort, Default Fallback)" + "configuration": "Claude Fable 5.1 (max with fallback)" }, { "modelId": "claude-haiku-3", @@ -33,9 +33,9 @@ }, { "modelId": "claude-haiku-4-5", - "score": 18, + "score": 17, "estimated": false, - "configuration": "Claude 4.5 Haiku (Reasoning)" + "configuration": "Claude 4.5 Haiku" }, { "modelId": "claude-opus-3", @@ -47,43 +47,49 @@ "modelId": "claude-opus-4", "score": 21, "estimated": true, - "configuration": "Claude 4 Opus (Reasoning)" + "configuration": "Claude 4 Opus" }, { "modelId": "claude-opus-4-1", "score": 23, "estimated": true, - "configuration": "Claude 4.1 Opus (Reasoning)" + "configuration": "Claude 4.1 Opus" }, { "modelId": "claude-opus-4-5", "score": 29, "estimated": true, - "configuration": "Claude Opus 4.5 (Reasoning)" + "configuration": "Claude Opus 4.5" }, { "modelId": "claude-opus-4-6", "score": 32, "estimated": true, - "configuration": "Claude Opus 4.6 (Adaptive Reasoning, Max Effort)" + "configuration": "Claude Opus 4.6 (max)" }, { "modelId": "claude-opus-4-7", "score": 41, "estimated": true, - "configuration": "Claude Opus 4.7 (Adaptive Reasoning, Max Effort)" + "configuration": "Claude Opus 4.7 (max)" }, { "modelId": "claude-opus-4-8", "score": 42, "estimated": false, - "configuration": "Claude Opus 4.8 (Adaptive Reasoning, Max Effort)" + "configuration": "Claude Opus 4.8 (max)" }, { "modelId": "claude-opus-5", "score": 51, "estimated": false, - "configuration": "Claude Opus 5 (Adaptive Reasoning, Max Effort)" + "configuration": "Claude Opus 5 (max)" + }, + { + "modelId": "claude-opus-5-5", + "score": 58, + "estimated": false, + "configuration": "Claude Opus 5.5 (max with fallback)" }, { "modelId": "claude-sonnet-3", @@ -95,67 +101,67 @@ "modelId": "claude-sonnet-3-5-20240620", "score": 7, "estimated": true, - "configuration": "Claude 3.5 Sonnet (June '24)" + "configuration": "Claude 3.5 Sonnet (June)" }, { "modelId": "claude-sonnet-3-5-20241022", "score": 8, "estimated": true, - "configuration": "Claude 3.5 Sonnet (Oct '24)" + "configuration": "Claude 3.5 Sonnet (Oct)" }, { "modelId": "claude-sonnet-3-7", "score": 18, "estimated": true, - "configuration": "Claude 3.7 Sonnet (Reasoning)" + "configuration": "Claude 3.7 Sonnet" }, { "modelId": "claude-sonnet-4", "score": 19, "estimated": true, - "configuration": "Claude 4 Sonnet (Reasoning)" + "configuration": "Claude 4 Sonnet" }, { "modelId": "claude-sonnet-4-5", "score": 21, "estimated": false, - "configuration": "Claude 4.5 Sonnet (Reasoning)" + "configuration": "Claude 4.5 Sonnet" }, { "modelId": "claude-sonnet-4-6", "score": 30, "estimated": false, - "configuration": "Claude Sonnet 4.6 (Adaptive Reasoning, Max Effort)" + "configuration": "Claude Sonnet 4.6 (max)" }, { "modelId": "claude-sonnet-5", "score": 38, "estimated": false, - "configuration": "Claude Sonnet 5 (Adaptive Reasoning, Max Effort)" + "configuration": "Claude Sonnet 5 (max)" }, { "modelId": "deepseek-3-2", "score": 21, "estimated": true, - "configuration": "DeepSeek V3.2 (Reasoning)" + "configuration": "DeepSeek V3.2" }, { "modelId": "deepseek-r1", "score": 11, "estimated": false, - "configuration": "DeepSeek R1 (Jan '25)" + "configuration": "DeepSeek R1 (Jan)" }, { "modelId": "deepseek-r1-0528", "score": 13, "estimated": true, - "configuration": "DeepSeek R1 0528 (May '25)" + "configuration": "DeepSeek R1 0528" }, { "modelId": "deepseek-v3", "score": 8, "estimated": false, - "configuration": "DeepSeek V3 (Dec '24)" + "configuration": "DeepSeek V3 (Dec)" }, { "modelId": "deepseek-v3-1", @@ -167,43 +173,49 @@ "modelId": "deepseek-v3-2-exp", "score": 17, "estimated": true, - "configuration": "DeepSeek V3.2 Exp (Reasoning)" + "configuration": "DeepSeek V3.2 Exp" }, { "modelId": "deepseek-v3-terminus", "score": 15, "estimated": false, - "configuration": "DeepSeek V3.1 Terminus (Reasoning)" + "configuration": "DeepSeek V3.1 Terminus" }, { "modelId": "deepseek-v4-flash", - "score": 35, + "score": 34, "estimated": false, - "configuration": "DeepSeek V4 Flash 0731 (Reasoning, Max Effort)" + "configuration": "DeepSeek V4 Flash 0731 (max)" }, { "modelId": "deepseek-v4-flash-vision", "score": 35, "estimated": false, - "configuration": "DeepSeek V4 Flash Vision (Reasoning, Max Effort)" + "configuration": "DeepSeek V4 Flash Vision (max)" + }, + { + "modelId": "deepseek-v4-1-flash", + "score": 39, + "estimated": false, + "configuration": "DeepSeek V4.1 Flash (max)" }, { "modelId": "deepseek-v4-flash-preview", - "score": 25, + "score": 26, "estimated": false, - "configuration": "DeepSeek V4 Flash (Reasoning, High Effort)" + "configuration": "DeepSeek V4 Flash (high)" }, { "modelId": "deepseek-v4-pro", "score": 36, "estimated": false, - "configuration": "DeepSeek V4 Pro 0813 (Reasoning, Max Effort)" + "configuration": "DeepSeek V4 Pro 0813 (max)" }, { "modelId": "deepseek-v4-pro-preview", - "score": 31, + "score": 30, "estimated": false, - "configuration": "DeepSeek V4 Pro (Reasoning, Max Effort)" + "configuration": "DeepSeek V4 Pro (max)" }, { "modelId": "devstral-2", @@ -221,23 +233,23 @@ "modelId": "gemini-2-0-flash", "score": 9, "estimated": true, - "configuration": "Gemini 2.0 Flash (Feb '25)" + "configuration": "Gemini 2.0 Flash" }, { "modelId": "gemini-2-5-flash", "score": 13, "estimated": true, - "configuration": "Gemini 2.5 Flash (Reasoning)" + "configuration": "Gemini 2.5 Flash" }, { "modelId": "gemini-2-5-flash-lite", "score": 9, "estimated": true, - "configuration": "Gemini 2.5 Flash-Lite (Reasoning)" + "configuration": "Gemini 2.5 Flash-Lite" }, { "modelId": "gemini-2-5-pro", - "score": 17, + "score": 16, "estimated": false, "configuration": "Gemini 2.5 Pro" }, @@ -251,11 +263,11 @@ "modelId": "gemini-3-5-flash", "score": 33, "estimated": false, - "configuration": "Gemini 3.5 Flash (high)" + "configuration": "Gemini 3.5 Flash" }, { "modelId": "gemini-3-5-flash-lite", - "score": 23, + "score": 22, "estimated": false, "configuration": "Gemini 3.5 Flash-Lite" }, @@ -263,7 +275,7 @@ "modelId": "gemini-3-6-flash", "score": 34, "estimated": false, - "configuration": "Gemini 3.6 Flash (high)" + "configuration": "Gemini 3.6 Flash" }, { "modelId": "gemini-3-7-flash", @@ -281,7 +293,7 @@ "modelId": "gemini-3-flash", "score": 26, "estimated": true, - "configuration": "Gemini 3 Flash Preview (Reasoning)" + "configuration": "Gemini 3 Flash" }, { "modelId": "gemini-3-pro", @@ -293,19 +305,19 @@ "modelId": "gemma-4-26b-a4b", "score": 17, "estimated": true, - "configuration": "Gemma 4 26B A4B (Reasoning)" + "configuration": "Gemma 4 26B A4B" }, { "modelId": "gemma-4-31b", - "score": 15, + "score": 19, "estimated": false, - "configuration": "Gemma 4 31B (Reasoning)" + "configuration": "Gemma 4 31B" }, { "modelId": "glm-4-5", "score": 13, "estimated": true, - "configuration": "GLM-4.5 (Reasoning)" + "configuration": "GLM-4.5" }, { "modelId": "glm-4-5-air", @@ -317,37 +329,37 @@ "modelId": "glm-4-5v", "score": 8, "estimated": true, - "configuration": "GLM-4.5V (Reasoning)" + "configuration": "GLM-4.5V" }, { "modelId": "glm-4-6", "score": 19, "estimated": true, - "configuration": "GLM-4.6 (Reasoning)" + "configuration": "GLM-4.6" }, { "modelId": "glm-4-6v", "score": 11, "estimated": true, - "configuration": "GLM-4.6V (Reasoning)" + "configuration": "GLM-4.6V" }, { "modelId": "glm-4-7", "score": 22, "estimated": true, - "configuration": "GLM-4.7 (Reasoning)" + "configuration": "GLM-4.7" }, { "modelId": "glm-4-7-flash", "score": 15, "estimated": true, - "configuration": "GLM-4.7-Flash (Reasoning)" + "configuration": "GLM-4.7-Flash" }, { "modelId": "glm-5", "score": 28, "estimated": true, - "configuration": "GLM-5 (Reasoning)" + "configuration": "GLM-5" }, { "modelId": "glm-5-turbo", @@ -359,17 +371,17 @@ "modelId": "glm-5v-turbo", "score": 23, "estimated": true, - "configuration": "GLM 5V Turbo (Reasoning)" + "configuration": "GLM 5V Turbo" }, { "modelId": "glm-5-1", - "score": 27, + "score": 26, "estimated": true, - "configuration": "GLM-5.1 (Reasoning)" + "configuration": "GLM-5.1" }, { "modelId": "glm-5-2", - "score": 39, + "score": 34, "estimated": true, "configuration": "GLM-5.2 (max)" }, @@ -407,7 +419,7 @@ "modelId": "gpt-4o", "score": 8, "estimated": true, - "configuration": "GPT-4o (Nov '24)" + "configuration": "GPT-4o (Nov)" }, { "modelId": "gpt-4o-mini", @@ -465,7 +477,7 @@ }, { "modelId": "gpt-5-4-mini", - "score": 25, + "score": 24, "estimated": false, "configuration": "GPT-5.4 mini (xhigh)" }, @@ -477,13 +489,13 @@ }, { "modelId": "gpt-5-5", - "score": 39, + "score": 38, "estimated": false, "configuration": "GPT-5.5 (xhigh)" }, { "modelId": "gpt-5-6-luna", - "score": 38, + "score": 37, "estimated": false, "configuration": "GPT-5.6 Luna (max)" }, @@ -499,6 +511,18 @@ "estimated": false, "configuration": "GPT-6 Astra (max)" }, + { + "modelId": "gpt-6-luna", + "score": 37, + "estimated": false, + "configuration": "GPT-6 Luna (max)" + }, + { + "modelId": "gpt-6-sol", + "score": 48, + "estimated": false, + "configuration": "GPT-6 Sol (max)" + }, { "modelId": "gpt-5-6-terra", "score": 42, @@ -533,13 +557,13 @@ "modelId": "grok-4-1-fast", "score": 20, "estimated": true, - "configuration": "Grok 4.1 Fast (Reasoning)" + "configuration": "Grok 4.1 Fast" }, { "modelId": "grok-4-20", "score": 26, "estimated": true, - "configuration": "Grok 4.20 0309 v2 (Reasoning)" + "configuration": "Grok 4.20 0309 v2" }, { "modelId": "grok-4-3", @@ -559,11 +583,17 @@ "estimated": false, "configuration": "Grok 4.6 (high)" }, + { + "modelId": "grok-4-7", + "score": 46, + "estimated": false, + "configuration": "Grok 4.7 (xhigh)" + }, { "modelId": "grok-4-fast", "score": 18, "estimated": true, - "configuration": "Grok 4 Fast (Reasoning)" + "configuration": "Grok 4 Fast" }, { "modelId": "grok-code-fast-1", @@ -573,7 +603,7 @@ }, { "modelId": "hy3", - "score": 26, + "score": 25, "estimated": false, "configuration": "Hy3" }, @@ -593,11 +623,11 @@ "modelId": "kimi-k2-5", "score": 23, "estimated": true, - "configuration": "Kimi K2.5 (Reasoning)" + "configuration": "Kimi K2.5" }, { "modelId": "kimi-k2-6", - "score": 31, + "score": 27, "estimated": true, "configuration": "Kimi K2.6" }, @@ -621,13 +651,13 @@ }, { "modelId": "llama-4-maverick", - "score": 9, + "score": 10, "estimated": false, "configuration": "Llama 4 Maverick" }, { "modelId": "llama-4-scout", - "score": 6, + "score": 8, "estimated": false, "configuration": "Llama 4 Scout" }, @@ -657,13 +687,13 @@ }, { "modelId": "minimax-m3", - "score": 30, + "score": 29, "estimated": false, "configuration": "MiniMax-M3" }, { "modelId": "mimo-v2-5", - "score": 22, + "score": 25, "estimated": false, "configuration": "MiMo-V2.5" }, @@ -673,15 +703,21 @@ "estimated": false, "configuration": "MiMo-V2.5-Pro" }, + { + "modelId": "mimo-v2-6-pro", + "score": 46, + "estimated": false, + "configuration": "MiMo-V2.6-Pro" + }, { "modelId": "mimo-v2-flash", "score": 21, "estimated": true, - "configuration": "MiMo-V2-Flash (Reasoning)" + "configuration": "MiMo-V2-Flash" }, { "modelId": "mistral-medium-3-5", - "score": 15, + "score": 14, "estimated": false, "configuration": "Mistral Medium 3.5" }, @@ -689,7 +725,7 @@ "modelId": "mistral-small-4", "score": 11, "estimated": false, - "configuration": "Mistral Small 4 (Reasoning)" + "configuration": "Mistral Small 4" }, { "modelId": "muse-spark-1-1", @@ -731,31 +767,31 @@ "modelId": "qwen3-5-122b-a10b", "score": 16, "estimated": false, - "configuration": "Qwen3.5 122B A10B (Reasoning)" + "configuration": "Qwen3.5 122B A10B" }, { "modelId": "qwen3-5-35b-a3b", "score": 19, "estimated": true, - "configuration": "Qwen3.5 35B A3B (Reasoning)" + "configuration": "Qwen3.5 35B A3B" }, { "modelId": "qwen3-5-397b-a17b", - "score": 19, + "score": 18, "estimated": false, - "configuration": "Qwen3.5 397B A17B (Reasoning)" + "configuration": "Qwen3.5 397B A17B" }, { "modelId": "qwen3-coder-30b-a3b", "score": 10, "estimated": true, - "configuration": "Qwen3 Coder 30B A3B Instruct" + "configuration": "Qwen3 Coder 30B A3B" }, { "modelId": "qwen3-coder-480b-a35b", "score": 12, "estimated": true, - "configuration": "Qwen3 Coder 480B A35B Instruct" + "configuration": "Qwen3 Coder 480B" }, { "modelId": "qwen3-6-plus", @@ -765,15 +801,15 @@ }, { "modelId": "qwen3-6-27b", - "score": 22, + "score": 21, "estimated": false, - "configuration": "Qwen3.6 27B (Reasoning)" + "configuration": "Qwen3.6 27B" }, { "modelId": "qwen3-6-35b-a3b", - "score": 22, + "score": 18, "estimated": true, - "configuration": "Qwen3.6 35B A3B (Reasoning)" + "configuration": "Qwen3.6 35B A3B" }, { "modelId": "qwen3-6-max-preview", @@ -783,13 +819,13 @@ }, { "modelId": "qwen3-7-max", - "score": 30, + "score": 29, "estimated": false, "configuration": "Qwen3.7 Max" }, { "modelId": "qwen3-7-plus", - "score": 26, + "score": 25, "estimated": false, "configuration": "Qwen3.7 Plus" }, @@ -811,15 +847,21 @@ "estimated": false, "configuration": "Qwen3.8 Max" }, + { + "modelId": "qwen3-8-max-0902", + "score": 45, + "estimated": false, + "configuration": "Qwen3.8 Max (0902)" + }, { "modelId": "qwen3-8-flash-next", - "score": 42, + "score": 40, "estimated": true, "configuration": "Qwen3.8-Flash-Next" }, { "modelId": "qwen3-coder-next", - "score": 10, + "score": 9, "estimated": false, "configuration": "Qwen3 Coder Next" } diff --git a/data/data-health.json b/data/data-health.json index 9b83ba87..3210a85d 100644 --- a/data/data-health.json +++ b/data/data-health.json @@ -1,5 +1,5 @@ { - "asOf": "2026-09-09", + "asOf": "2026-09-24", "thresholds": { "models": 30, "providers": 30, @@ -10,21 +10,21 @@ "vendors": 90 }, "summary": { - "totalRecords": 274, - "recordsWithSources": 274, - "verifiedRecords": 274, - "provenanceComplete": 274, - "staleVerifiedRecords": 134, + "totalRecords": 281, + "recordsWithSources": 281, + "verifiedRecords": 281, + "provenanceComplete": 281, + "staleVerifiedRecords": 148, "translationPlaceholderValues": 300, "danglingRelationships": 0, - "modelBenchmarkCoverage": 8.5, + "modelBenchmarkCoverage": 8.1, "productsWithPricing": 70, "productRecords": 71, "communityUrlsPopulated": 336, "communityUrlsWithProvenance": 336, "duplicatedVendorCommunityUrls": 0, "errors": 0, - "warnings": 134, + "warnings": 148, "info": 0 }, "byCategory": { @@ -50,13 +50,13 @@ "total": 19, "verified": 19, "provenanceComplete": 19, - "stale": 0 + "stale": 2 }, "models": { - "total": 138, - "verified": 138, - "provenanceComplete": 138, - "stale": 117 + "total": 145, + "verified": 145, + "provenanceComplete": 145, + "stale": 129 }, "providers": { "total": 17, @@ -129,943 +129,1041 @@ } }, "issues": [ + { + "severity": "warning", + "code": "stale-verification", + "category": "extensions", + "id": "claude-code", + "message": "Last reviewed 67 days ago; threshold is 60 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "extensions", + "id": "gemini-code-assist", + "message": "Last reviewed 67 days ago; threshold is 60 days." + }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-fable-5", - "message": "Last reviewed 51 days ago; threshold is 30 days." + "message": "Last reviewed 66 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-haiku-3", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-haiku-3-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-haiku-4-5", - "message": "Last reviewed 53 days ago; threshold is 30 days." + "message": "Last reviewed 68 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-3", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-4", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-4-1", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-4-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-4-6", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-4-7", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-4-8", - "message": "Last reviewed 51 days ago; threshold is 30 days." + "message": "Last reviewed 66 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-opus-5", - "message": "Last reviewed 44 days ago; threshold is 30 days." + "message": "Last reviewed 59 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-3", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-3-5-20240620", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-3-5-20241022", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-3-7", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-4", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-4-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-4-6", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "claude-sonnet-5", - "message": "Last reviewed 51 days ago; threshold is 30 days." + "message": "Last reviewed 66 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "cursor-composer-2", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "cursor-composer-2-5", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "deepseek-3-2", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "deepseek-r1", + "message": "Last reviewed 37 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "deepseek-r1-0528", + "message": "Last reviewed 37 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "deepseek-v3", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "deepseek-v3-1", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "deepseek-v3-2-exp", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "deepseek-v3-terminus", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", - "id": "deepseek-v4-flash", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "id": "deepseek-v4-flash-preview", + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", - "id": "deepseek-v4-flash-preview", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "id": "deepseek-v4-pro-preview", + "message": "Last reviewed 42 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "devstral-2", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "devstral-small-2", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-2-0-flash", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-2-5-flash", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-2-5-flash-lite", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-2-5-pro", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-3-1-pro-preview", - "message": "Last reviewed 51 days ago; threshold is 30 days." + "message": "Last reviewed 66 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-3-5-flash", - "message": "Last reviewed 49 days ago; threshold is 30 days." + "message": "Last reviewed 64 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-3-5-flash-lite", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-3-6-flash", - "message": "Last reviewed 49 days ago; threshold is 30 days." + "message": "Last reviewed 64 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "gemini-3-7-flash", + "message": "Last reviewed 41 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-3-flash", - "message": "Last reviewed 49 days ago; threshold is 30 days." + "message": "Last reviewed 64 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemini-3-pro", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemma-4-26b-a4b", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gemma-4-31b", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-5-air", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-5v", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-6", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-6v", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-7", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-4-7-flash", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-5-1", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "glm-5-2", + "message": "Last reviewed 36 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "glm-5-3", + "message": "Last reviewed 36 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-5-turbo", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "glm-5v-turbo", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-4-1", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-4-1-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-4-1-nano", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-4o", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-4o-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-1", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-1-codex", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-1-codex-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-2", - "message": "Last reviewed 53 days ago; threshold is 30 days." + "message": "Last reviewed 68 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-2-codex", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-3-codex", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-4", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-4-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-4-nano", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-6-luna", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-6-sol", - "message": "Last reviewed 51 days ago; threshold is 30 days." + "message": "Last reviewed 66 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-6-terra", - "message": "Last reviewed 51 days ago; threshold is 30 days." + "message": "Last reviewed 66 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-codex", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "gpt-5-nano", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "grok-4", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "grok-4-1-fast", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "grok-4-20", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "grok-4-3", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "grok-4-5", + "message": "Last reviewed 44 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "grok-4-6", + "message": "Last reviewed 42 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "grok-4-fast", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "grok-code-fast-1", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "hy3", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k2-0905", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k2-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k2-6", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k2-7-code", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k2-instruct", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k2-thinking", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "kimi-k3", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "llama-4-maverick", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "llama-4-scout", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "mimo-v2-5", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "mimo-v2-5-pro", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "mimo-v2-flash", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "minimax-m2", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "minimax-m2-1", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "minimax-m2-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "minimax-m2-7", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "minimax-m3", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "mistral-medium-3-5", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "mistral-small-4", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "muse-spark-1-1", + "message": "Last reviewed 44 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "muse-spark-1-2", + "message": "Last reviewed 44 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "o3", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "o3-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "o4-mini", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-5-122b-a10b", - "message": "Last reviewed 44 days ago; threshold is 30 days." + "message": "Last reviewed 59 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-5-35b-a3b", - "message": "Last reviewed 44 days ago; threshold is 30 days." + "message": "Last reviewed 59 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-5-397b-a17b", - "message": "Last reviewed 44 days ago; threshold is 30 days." + "message": "Last reviewed 59 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-6-27b", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-6-35b-a3b", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-6-max-preview", - "message": "Last reviewed 45 days ago; threshold is 30 days." + "message": "Last reviewed 60 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-6-plus", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-7-max", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-7-plus", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "qwen3-8-2-4t-a95b", + "message": "Last reviewed 41 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "qwen3-8-27b", + "message": "Last reviewed 37 days ago; threshold is 30 days." + }, + { + "severity": "warning", + "code": "stale-verification", + "category": "models", + "id": "qwen3-8-max", + "message": "Last reviewed 37 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-coder-30b-a3b", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-coder-480b-a35b", - "message": "Last reviewed 43 days ago; threshold is 30 days." + "message": "Last reviewed 58 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "models", "id": "qwen3-coder-next", - "message": "Last reviewed 50 days ago; threshold is 30 days." + "message": "Last reviewed 65 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "alibaba", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "anthropic", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "baseten", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "deepinfra", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "deepseek", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "fireworks-ai", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "google", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "meta", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "minimax", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "moonshot", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "novita-ai", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "openai", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "openrouter", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "siliconflow", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "together-ai", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "xai", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." }, { "severity": "warning", "code": "stale-verification", "category": "providers", "id": "z-ai", - "message": "Last reviewed 40 days ago; threshold is 30 days." + "message": "Last reviewed 55 days ago; threshold is 30 days." } ] } diff --git a/docs/DATA-HEALTH.md b/docs/DATA-HEALTH.md index a0be4ced..a69e6507 100644 --- a/docs/DATA-HEALTH.md +++ b/docs/DATA-HEALTH.md @@ -1,23 +1,23 @@ # Data Health Report -Snapshot date: 2026-09-09. Regenerate with `pnpm data-health:report`. +Snapshot date: 2026-09-24. Regenerate with `pnpm data-health:report`. ## Scorecard | Metric | Value | | --- | ---: | -| Manifest records | 274 | -| Records with structured sources | 274 | -| Verified records | 274 | -| Verified with complete provenance | 274 | -| Stale verified records | 134 | +| Manifest records | 281 | +| Records with structured sources | 281 | +| Verified records | 281 | +| Verified with complete provenance | 281 | +| Stale verified records | 148 | | Non-English values identical to English | 300 | | Dangling product relationships | 0 | -| Model benchmark coverage | 8.5% | +| Model benchmark coverage | 8.1% | | Products with pricing | 70/71 | | Community URLs with provenance | 336/336 | | Duplicated vendor community URLs | 0 | -| Errors / warnings / info | 0 / 134 / 0 | +| Errors / warnings / info | 0 / 148 / 0 | ## Category Breakdown @@ -26,8 +26,8 @@ Snapshot date: 2026-09-09. Regenerate with `pnpm data-health:report`. | ides | 9 | 9 | 9 | 0 | | clis | 30 | 30 | 30 | 0 | | desktops | 13 | 13 | 13 | 0 | -| extensions | 19 | 19 | 19 | 0 | -| models | 138 | 138 | 138 | 117 | +| extensions | 19 | 19 | 19 | 2 | +| models | 145 | 145 | 145 | 129 | | providers | 17 | 17 | 17 | 17 | | vendors | 48 | 48 | 48 | 0 | @@ -53,7 +53,7 @@ Exact English matches are a triage signal; product names and technical terms can | Issue | Count | | --- | ---: | -| stale-verification | 134 | +| stale-verification | 148 | ## Priority Queue @@ -62,106 +62,106 @@ visible in the scorecards and `data/data-health.json`. | Severity | Issue | Record | Detail | | --- | --- | --- | --- | -| warning | stale-verification | models/claude-fable-5 | Last reviewed 51 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-haiku-3 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-haiku-3-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-haiku-4-5 | Last reviewed 53 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-3 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-4 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-4-1 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-4-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-4-6 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-4-7 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-4-8 | Last reviewed 51 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-opus-5 | Last reviewed 44 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-3 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-3-5-20240620 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-3-5-20241022 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-3-7 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-4 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-4-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-4-6 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/claude-sonnet-5 | Last reviewed 51 days ago; threshold is 30 days. | -| warning | stale-verification | models/cursor-composer-2 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/cursor-composer-2-5 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-3-2 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-v3 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-v3-1 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-v3-2-exp | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-v3-terminus | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-v4-flash | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/deepseek-v4-flash-preview | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/devstral-2 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/devstral-small-2 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-2-0-flash | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-2-5-flash | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-2-5-flash-lite | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-2-5-pro | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-3-1-pro-preview | Last reviewed 51 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-3-5-flash | Last reviewed 49 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-3-5-flash-lite | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-3-6-flash | Last reviewed 49 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-3-flash | Last reviewed 49 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemini-3-pro | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemma-4-26b-a4b | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/gemma-4-31b | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-5-air | Last reviewed 45 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-5v | Last reviewed 45 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-6 | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-6v | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-7 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-4-7-flash | Last reviewed 45 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-5-1 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-5-turbo | Last reviewed 45 days ago; threshold is 30 days. | -| warning | stale-verification | models/glm-5v-turbo | Last reviewed 45 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-4-1 | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-4-1-mini | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-4-1-nano | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-4o | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-4o-mini | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5 | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-1 | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-1-codex | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-1-codex-mini | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-2 | Last reviewed 53 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-2-codex | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-3-codex | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-4 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-4-mini | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-4-nano | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-6-luna | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-6-sol | Last reviewed 51 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-6-terra | Last reviewed 51 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-codex | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-mini | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/gpt-5-nano | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/grok-4 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/grok-4-1-fast | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/grok-4-20 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/grok-4-3 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/grok-4-fast | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/grok-code-fast-1 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/hy3 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k2-0905 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k2-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k2-6 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k2-7-code | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k2-instruct | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k2-thinking | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/kimi-k3 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/llama-4-maverick | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/llama-4-scout | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/mimo-v2-5 | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/mimo-v2-5-pro | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/mimo-v2-flash | Last reviewed 40 days ago; threshold is 30 days. | -| warning | stale-verification | models/minimax-m2 | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/minimax-m2-1 | Last reviewed 43 days ago; threshold is 30 days. | -| warning | stale-verification | models/minimax-m2-5 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/minimax-m2-7 | Last reviewed 50 days ago; threshold is 30 days. | -| warning | stale-verification | models/minimax-m3 | Last reviewed 40 days ago; threshold is 30 days. | +| warning | stale-verification | extensions/claude-code | Last reviewed 67 days ago; threshold is 60 days. | +| warning | stale-verification | extensions/gemini-code-assist | Last reviewed 67 days ago; threshold is 60 days. | +| warning | stale-verification | models/claude-fable-5 | Last reviewed 66 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-haiku-3 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-haiku-3-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-haiku-4-5 | Last reviewed 68 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-3 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-4 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-4-1 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-4-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-4-6 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-4-7 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-4-8 | Last reviewed 66 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-opus-5 | Last reviewed 59 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-3 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-3-5-20240620 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-3-5-20241022 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-3-7 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-4 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-4-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-4-6 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/claude-sonnet-5 | Last reviewed 66 days ago; threshold is 30 days. | +| warning | stale-verification | models/cursor-composer-2 | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/cursor-composer-2-5 | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-3-2 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-r1 | Last reviewed 37 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-r1-0528 | Last reviewed 37 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-v3 | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-v3-1 | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-v3-2-exp | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-v3-terminus | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-v4-flash-preview | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/deepseek-v4-pro-preview | Last reviewed 42 days ago; threshold is 30 days. | +| warning | stale-verification | models/devstral-2 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/devstral-small-2 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-2-0-flash | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-2-5-flash | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-2-5-flash-lite | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-2-5-pro | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-1-pro-preview | Last reviewed 66 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-5-flash | Last reviewed 64 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-5-flash-lite | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-6-flash | Last reviewed 64 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-7-flash | Last reviewed 41 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-flash | Last reviewed 64 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemini-3-pro | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemma-4-26b-a4b | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/gemma-4-31b | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-5-air | Last reviewed 60 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-5v | Last reviewed 60 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-6 | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-6v | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-7 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-4-7-flash | Last reviewed 60 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-5-1 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-5-2 | Last reviewed 36 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-5-3 | Last reviewed 36 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-5-turbo | Last reviewed 60 days ago; threshold is 30 days. | +| warning | stale-verification | models/glm-5v-turbo | Last reviewed 60 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-4-1 | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-4-1-mini | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-4-1-nano | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-4o | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-4o-mini | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5 | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-1 | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-1-codex | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-1-codex-mini | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-2 | Last reviewed 68 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-2-codex | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-3-codex | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-4 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-4-mini | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-4-nano | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-6-luna | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-6-sol | Last reviewed 66 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-6-terra | Last reviewed 66 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-codex | Last reviewed 58 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-mini | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/gpt-5-nano | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4-1-fast | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4-20 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4-3 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4-5 | Last reviewed 44 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4-6 | Last reviewed 42 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-4-fast | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/grok-code-fast-1 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/hy3 | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k2-0905 | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k2-5 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k2-6 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k2-7-code | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k2-instruct | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k2-thinking | Last reviewed 55 days ago; threshold is 30 days. | +| warning | stale-verification | models/kimi-k3 | Last reviewed 65 days ago; threshold is 30 days. | +| warning | stale-verification | models/llama-4-maverick | Last reviewed 55 days ago; threshold is 30 days. | ## Freshness Thresholds diff --git a/manifests/models/claude-opus-5-5.json b/manifests/models/claude-opus-5-5.json new file mode 100644 index 00000000..58c94d9a --- /dev/null +++ b/manifests/models/claude-opus-5-5.json @@ -0,0 +1,120 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "claude-opus-5-5", + "name": "Claude Opus 5.5", + "description": "Anthropic's model for long-running agentic coding and knowledge work, with always-on adaptive thinking, a 1M-token context window, and 128K output.", + "translations": { + "de": { + "description": "Anthropics Modell für lang laufendes agentisches Coding und Wissensarbeit mit stets aktivem adaptivem Denken, 1 Mio. Token Kontext und 128K Ausgabe." + }, + "es": { + "description": "Modelo de Anthropic para programación agéntica prolongada y trabajo del conocimiento, con pensamiento adaptativo siempre activo, contexto de 1 M de tokens y salida de 128K." + }, + "fr": { + "description": "Modèle d’Anthropic pour le codage agentique de longue durée et le travail intellectuel, avec raisonnement adaptatif permanent, contexte de 1 M de jetons et sortie de 128K." + }, + "id": { + "description": "Model Anthropic untuk coding agentik jangka panjang dan pekerjaan pengetahuan, dengan pemikiran adaptif yang selalu aktif, konteks 1 juta token, dan keluaran 128K." + }, + "ja": { + "description": "常時有効の適応型思考、100万トークンのコンテキスト、128K出力を備え、長時間のエージェント型コーディングと知識労働に対応するAnthropicのモデル。" + }, + "ko": { + "description": "항상 활성화되는 적응형 사고, 100만 토큰 컨텍스트와 128K 출력을 갖춘 장기 실행 에이전트 코딩 및 지식 업무용 Anthropic 모델입니다." + }, + "pt": { + "description": "Modelo da Anthropic para programação agêntica de longa duração e trabalho de conhecimento, com pensamento adaptativo sempre ativo, contexto de 1 milhão de tokens e saída de 128K." + }, + "ru": { + "description": "Модель Anthropic для длительного агентного программирования и интеллектуальной работы с постоянно активным адаптивным мышлением, контекстом 1 млн токенов и выводом 128K." + }, + "tr": { + "description": "Her zaman açık uyarlanabilir düşünme, 1 milyon token bağlam ve 128K çıktı ile uzun süreli ajan tabanlı kodlama ve bilgi çalışmaları için Anthropic modeli." + }, + "zh-Hans": { + "description": "Anthropic 面向长时间智能体编码和知识工作的模型,具备始终开启的自适应思考、100 万 token 上下文和 128K 输出。" + }, + "zh-Hant": { + "description": "Anthropic 面向長時間智慧體程式設計和知識工作的模型,具備始終開啟的自適應思考、100 萬 token 上下文和 128K 輸出。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://platform.claude.com/docs/en/models/opus-5-5/overview", + "title": "Claude Opus 5.5 model overview", + "fields": [ + "name", + "description", + "docsUrl", + "contextWindow", + "maxOutput", + "tokenPricing", + "knowledgeCutoff", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://platform.claude.com/docs/en/release-notes/overview", + "title": "Claude Platform release notes", + "fields": ["releaseDate", "lifecycle"] + }, + { + "url": "https://artificialanalysis.ai/models/claude-opus-5-5", + "title": "Artificial Analysis Claude Opus 5.5 model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://www.anthropic.com", + "docsUrl": "https://platform.claude.com/docs/en/models/opus-5-5/overview", + "vendor": "Anthropic", + "size": null, + "activeParameters": null, + "contextWindow": 1000000, + "maxOutput": 128000, + "tokenPricing": { + "status": "available", + "primaryOffer": "global-standard", + "offers": [ + { + "id": "global-standard", + "currency": "USD", + "region": "global", + "serviceTier": "standard", + "effectiveFrom": "2026-09-22", + "effectiveTo": null, + "tiers": [ + { + "condition": null, + "rates": { "input": 4, "output": 20, "cacheRead": 0.2, "cacheWrite": null } + } + ] + } + ] + }, + "releaseDate": "2026-09-22", + "lifecycle": "latest", + "knowledgeCutoff": "2026-06", + "inputModalities": ["text", "image", "pdf"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": null, + "artificialAnalysis": "https://artificialanalysis.ai/models/claude-opus-5-5", + "openrouter": null + } +} diff --git a/manifests/models/deepseek-v4-1-flash.json b/manifests/models/deepseek-v4-1-flash.json new file mode 100644 index 00000000..51bf3039 --- /dev/null +++ b/manifests/models/deepseek-v4-1-flash.json @@ -0,0 +1,117 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "deepseek-v4-1-flash", + "name": "DeepSeek-V4.1-Flash", + "description": "DeepSeek's 552B MoE model with native visual understanding, a 1M-token context window, thinking and non-thinking modes, tool calls, and JSON output.", + "translations": { + "de": { + "description": "DeepSeeks 552B-MoE-Modell mit nativer Bildverarbeitung, 1 Mio. Token Kontext, Denk- und Nicht-Denkmodus, Werkzeugaufrufen und JSON-Ausgabe." + }, + "es": { + "description": "Modelo MoE de 552B de DeepSeek con comprensión visual nativa, contexto de 1 M de tokens, modos con y sin razonamiento, herramientas y salida JSON." + }, + "fr": { + "description": "Modèle MoE 552B de DeepSeek avec compréhension visuelle native, contexte de 1 M de jetons, modes avec et sans raisonnement, appels d’outils et sortie JSON." + }, + "id": { + "description": "Model MoE 552B DeepSeek dengan pemahaman visual native, konteks 1 juta token, mode berpikir dan non-berpikir, pemanggilan alat, serta keluaran JSON." + }, + "ja": { + "description": "ネイティブな画像理解、100万トークンのコンテキスト、思考・非思考モード、ツール呼び出し、JSON出力を備えたDeepSeekの552B MoEモデル。" + }, + "ko": { + "description": "네이티브 시각 이해, 100만 토큰 컨텍스트, 사고 및 비사고 모드, 도구 호출과 JSON 출력을 지원하는 DeepSeek의 552B MoE 모델입니다." + }, + "pt": { + "description": "Modelo MoE de 552B da DeepSeek com compreensão visual nativa, contexto de 1 milhão de tokens, modos com e sem raciocínio, ferramentas e saída JSON." + }, + "ru": { + "description": "MoE-модель DeepSeek на 552 млрд параметров с нативным зрением, контекстом 1 млн токенов, режимами рассуждений, вызовами инструментов и выводом JSON." + }, + "tr": { + "description": "Yerel görsel anlama, 1 milyon token bağlam, düşünme ve düşünmeme modları, araç çağrıları ve JSON çıktısı sunan DeepSeek 552B MoE modeli." + }, + "zh-Hans": { + "description": "DeepSeek 的 552B MoE 模型,具备原生视觉理解、100 万 token 上下文、思考与非思考模式、工具调用和 JSON 输出。" + }, + "zh-Hant": { + "description": "DeepSeek 的 552B MoE 模型,具備原生視覺理解、100 萬 token 上下文、思考與非思考模式、工具呼叫和 JSON 輸出。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://deepseek.com/news/deepseek-v4-1-flash/", + "title": "DeepSeek V4.1 Flash release announcement", + "fields": [ + "name", + "description", + "size", + "releaseDate", + "lifecycle", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://api-docs.deepseek.com/quick_start/pricing/", + "title": "DeepSeek models and pricing", + "fields": [ + "docsUrl", + "contextWindow", + "maxOutput", + "tokenPricing", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://huggingface.co/deepseek-ai/DeepSeek-V4.1-Flash", + "title": "DeepSeek-V4.1-Flash official model card", + "fields": ["size", "platformUrls.huggingface"] + }, + { + "url": "https://artificialanalysis.ai/models/deepseek-v4-1-flash", + "title": "Artificial Analysis DeepSeek V4.1 Flash model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://www.deepseek.com", + "docsUrl": "https://api-docs.deepseek.com/quick_start/pricing/", + "vendor": "DeepSeek", + "size": "552B", + "activeParameters": null, + "contextWindow": 1000000, + "maxOutput": 384000, + "tokenPricing": { + "status": "unavailable", + "reason": "unsupported-pricing-structure", + "primaryOffer": null, + "offers": [] + }, + "releaseDate": "2026-09-10", + "lifecycle": "latest", + "knowledgeCutoff": null, + "inputModalities": ["text", "image"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": "https://huggingface.co/deepseek-ai/DeepSeek-V4.1-Flash", + "artificialAnalysis": "https://artificialanalysis.ai/models/deepseek-v4-1-flash", + "openrouter": null + } +} diff --git a/manifests/models/deepseek-v4-flash-vision.json b/manifests/models/deepseek-v4-flash-vision.json index 35779796..05483a66 100644 --- a/manifests/models/deepseek-v4-flash-vision.json +++ b/manifests/models/deepseek-v4-flash-vision.json @@ -49,12 +49,16 @@ "websiteUrl", "docsUrl", "releaseDate", - "lifecycle", "inputModalities", "outputModalities", "capabilities" ] }, + { + "url": "https://deepseek.com/news/deepseek-v4-1-flash/", + "title": "DeepSeek V4.1 Flash release and V4 retirement announcement", + "fields": ["lifecycle"] + }, { "url": "https://api-docs.deepseek.com/quick_start/pricing", "title": "DeepSeek models and pricing", @@ -71,7 +75,7 @@ "fields": ["platformUrls.artificialAnalysis"] } ], - "lastVerifiedAt": "2026-08-28", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://www.deepseek.com", @@ -88,7 +92,7 @@ "offers": [] }, "releaseDate": "2026-08-21", - "lifecycle": "latest", + "lifecycle": "deprecated", "knowledgeCutoff": null, "inputModalities": ["text", "image"], "outputModalities": ["text"], diff --git a/manifests/models/deepseek-v4-flash.json b/manifests/models/deepseek-v4-flash.json index e51bd198..b2b4a12c 100644 --- a/manifests/models/deepseek-v4-flash.json +++ b/manifests/models/deepseek-v4-flash.json @@ -61,15 +61,12 @@ { "url": "https://api-docs.deepseek.com/updates/", "title": "DeepSeek V4 Flash July 31 API update", - "fields": [ - "name", - "description", - "size", - "activeParameters", - "releaseDate", - "lifecycle", - "capabilities" - ] + "fields": ["name", "description", "size", "activeParameters", "releaseDate", "capabilities"] + }, + { + "url": "https://deepseek.com/news/deepseek-v4-1-flash/", + "title": "DeepSeek V4.1 Flash release and V4 retirement announcement", + "fields": ["lifecycle"] }, { "url": "https://artificialanalysis.ai/models/deepseek-v4-flash/", @@ -77,7 +74,7 @@ "fields": ["platformUrls.artificialAnalysis"] } ], - "lastVerifiedAt": "2026-07-31", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://www.deepseek.com", @@ -113,7 +110,7 @@ ] }, "releaseDate": "2026-07-31", - "lifecycle": "latest", + "lifecycle": "deprecated", "knowledgeCutoff": null, "inputModalities": ["text"], "outputModalities": ["text"], diff --git a/manifests/models/deepseek-v4-pro.json b/manifests/models/deepseek-v4-pro.json index 27ab4951..9f461854 100644 --- a/manifests/models/deepseek-v4-pro.json +++ b/manifests/models/deepseek-v4-pro.json @@ -50,19 +50,23 @@ "contextWindow", "maxOutput", "tokenPricing", - "lifecycle", "inputModalities", "outputModalities", "capabilities" ] }, + { + "url": "https://deepseek.com/news/deepseek-v4-1-flash/", + "title": "DeepSeek V4.1 Flash release and V4 retirement announcement", + "fields": ["lifecycle"] + }, { "url": "https://artificialanalysis.ai/models/deepseek-v4-pro", "title": "Artificial Analysis DeepSeek V4 Pro 0813 model page", "fields": ["releaseDate", "platformUrls.artificialAnalysis"] } ], - "lastVerifiedAt": "2026-08-13", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://www.deepseek.com", @@ -98,7 +102,7 @@ ] }, "releaseDate": "2026-08-13", - "lifecycle": "latest", + "lifecycle": "deprecated", "knowledgeCutoff": null, "inputModalities": ["text"], "outputModalities": ["text"], diff --git a/manifests/models/gpt-6-luna.json b/manifests/models/gpt-6-luna.json new file mode 100644 index 00000000..0b313f47 --- /dev/null +++ b/manifests/models/gpt-6-luna.json @@ -0,0 +1,124 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "gpt-6-luna", + "name": "GPT-6 Luna", + "description": "OpenAI's efficient GPT-6 model for focused, high-volume workloads, with adjustable reasoning, a 1.05M-token context window, and 128K output.", + "translations": { + "de": { + "description": "OpenAIs effizientes GPT-6-Modell für fokussierte Aufgaben mit hohem Volumen, einstellbarem Reasoning, 1,05 Mio. Token Kontext und 128K Ausgabe." + }, + "es": { + "description": "Modelo GPT-6 eficiente de OpenAI para cargas enfocadas y de gran volumen, con razonamiento ajustable, contexto de 1,05 M de tokens y salida de 128K." + }, + "fr": { + "description": "Modèle GPT-6 efficace d’OpenAI pour les charges ciblées à fort volume, avec raisonnement réglable, contexte de 1,05 M de jetons et sortie de 128K." + }, + "id": { + "description": "Model GPT-6 efisien OpenAI untuk beban kerja terfokus bervolume tinggi, dengan penalaran yang dapat diatur, konteks 1,05 juta token, dan keluaran 128K." + }, + "ja": { + "description": "集中的で大量の処理向けに、調整可能な推論、105万トークンのコンテキスト、128K出力を備えたOpenAIの効率的なGPT-6モデル。" + }, + "ko": { + "description": "집중적이고 대규모인 작업을 위해 조절 가능한 추론, 105만 토큰 컨텍스트와 128K 출력을 제공하는 OpenAI의 효율적인 GPT-6 모델입니다." + }, + "pt": { + "description": "Modelo GPT-6 eficiente da OpenAI para cargas focadas e de alto volume, com raciocínio ajustável, contexto de 1,05 milhão de tokens e saída de 128K." + }, + "ru": { + "description": "Эффективная модель GPT-6 от OpenAI для сфокусированных массовых задач с настраиваемыми рассуждениями, контекстом 1,05 млн токенов и выводом 128K." + }, + "tr": { + "description": "Odaklı, yüksek hacimli işler için ayarlanabilir akıl yürütme, 1,05 milyon token bağlam ve 128K çıktı sunan verimli OpenAI GPT-6 modeli." + }, + "zh-Hans": { + "description": "OpenAI 面向聚焦、高吞吐工作负载的高效 GPT-6 模型,支持可调推理、105 万 token 上下文和 128K 输出。" + }, + "zh-Hant": { + "description": "OpenAI 面向聚焦、高吞吐工作負載的高效 GPT-6 模型,支援可調推理、105 萬 token 上下文和 128K 輸出。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://developers.openai.com/api/docs/models/gpt-6-luna", + "title": "GPT-6 Luna model documentation", + "fields": [ + "name", + "description", + "docsUrl", + "contextWindow", + "maxOutput", + "tokenPricing", + "knowledgeCutoff", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://developers.openai.com/api/docs/changelog", + "title": "OpenAI API changelog", + "fields": ["releaseDate", "lifecycle"] + }, + { + "url": "https://artificialanalysis.ai/models/gpt-6-luna", + "title": "Artificial Analysis GPT-6 Luna model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://openai.com", + "docsUrl": "https://developers.openai.com/api/docs/models/gpt-6-luna", + "vendor": "OpenAI", + "size": null, + "activeParameters": null, + "contextWindow": 1050000, + "maxOutput": 128000, + "tokenPricing": { + "status": "available", + "primaryOffer": "global-standard", + "offers": [ + { + "id": "global-standard", + "currency": "USD", + "region": "global", + "serviceTier": "standard", + "effectiveFrom": "2026-09-22", + "effectiveTo": null, + "tiers": [ + { + "condition": { "metric": "inputTokens", "min": 1, "max": 272000 }, + "rates": { "input": 0.1, "output": 0.5, "cacheRead": 0.01, "cacheWrite": 0.125 } + }, + { + "condition": { "metric": "inputTokens", "min": 272001, "max": 922000 }, + "rates": { "input": 0.2, "output": 0.75, "cacheRead": 0.02, "cacheWrite": 0.25 } + } + ] + } + ] + }, + "releaseDate": "2026-09-22", + "lifecycle": "latest", + "knowledgeCutoff": "2026-05-18", + "inputModalities": ["text", "image"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": null, + "artificialAnalysis": "https://artificialanalysis.ai/models/gpt-6-luna", + "openrouter": null + } +} diff --git a/manifests/models/gpt-6-sol.json b/manifests/models/gpt-6-sol.json new file mode 100644 index 00000000..35c40789 --- /dev/null +++ b/manifests/models/gpt-6-sol.json @@ -0,0 +1,124 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "gpt-6-sol", + "name": "GPT-6 Sol", + "description": "OpenAI's GPT-6 model for complex coding and agentic workflows, balancing intelligence and cost with a 1.05M-token context window and 128K output.", + "translations": { + "de": { + "description": "OpenAIs GPT-6-Modell für komplexe Programmier- und Agentenabläufe, das Intelligenz und Kosten mit 1,05 Mio. Token Kontext und 128K Ausgabe ausbalanciert." + }, + "es": { + "description": "Modelo GPT-6 de OpenAI para programación compleja y flujos agénticos, que equilibra inteligencia y coste con contexto de 1,05 M de tokens y salida de 128K." + }, + "fr": { + "description": "Modèle GPT-6 d’OpenAI pour le codage complexe et les flux agentiques, équilibrant intelligence et coût avec 1,05 M de jetons de contexte et 128K en sortie." + }, + "id": { + "description": "Model GPT-6 OpenAI untuk coding kompleks dan alur kerja agentik, menyeimbangkan kecerdasan dan biaya dengan konteks 1,05 juta token serta keluaran 128K." + }, + "ja": { + "description": "複雑なコーディングとエージェント型ワークフロー向けに、知能とコストのバランスを取り、105万トークンのコンテキストと128K出力を備えたOpenAIのGPT-6モデル。" + }, + "ko": { + "description": "복잡한 코딩과 에이전트 워크플로를 위해 지능과 비용의 균형을 맞추고 105만 토큰 컨텍스트와 128K 출력을 제공하는 OpenAI GPT-6 모델입니다." + }, + "pt": { + "description": "Modelo GPT-6 da OpenAI para programação complexa e fluxos agênticos, equilibrando inteligência e custo com contexto de 1,05 milhão de tokens e saída de 128K." + }, + "ru": { + "description": "Модель GPT-6 от OpenAI для сложного программирования и агентных процессов, сочетающая интеллект и стоимость с контекстом 1,05 млн токенов и выводом 128K." + }, + "tr": { + "description": "Karmaşık kodlama ve ajan iş akışları için zeka ile maliyeti dengeleyen, 1,05 milyon token bağlam ve 128K çıktı sunan OpenAI GPT-6 modeli." + }, + "zh-Hans": { + "description": "OpenAI 面向复杂编码和智能体工作流的 GPT-6 模型,以 105 万 token 上下文和 128K 输出平衡能力与成本。" + }, + "zh-Hant": { + "description": "OpenAI 面向複雜程式設計和智慧體工作流程的 GPT-6 模型,以 105 萬 token 上下文和 128K 輸出平衡能力與成本。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://developers.openai.com/api/docs/models/gpt-6-sol", + "title": "GPT-6 Sol model documentation", + "fields": [ + "name", + "description", + "docsUrl", + "contextWindow", + "maxOutput", + "tokenPricing", + "knowledgeCutoff", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://developers.openai.com/api/docs/changelog", + "title": "OpenAI API changelog", + "fields": ["releaseDate", "lifecycle"] + }, + { + "url": "https://artificialanalysis.ai/models/gpt-6-sol", + "title": "Artificial Analysis GPT-6 Sol model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://openai.com", + "docsUrl": "https://developers.openai.com/api/docs/models/gpt-6-sol", + "vendor": "OpenAI", + "size": null, + "activeParameters": null, + "contextWindow": 1050000, + "maxOutput": 128000, + "tokenPricing": { + "status": "available", + "primaryOffer": "global-standard", + "offers": [ + { + "id": "global-standard", + "currency": "USD", + "region": "global", + "serviceTier": "standard", + "effectiveFrom": "2026-09-22", + "effectiveTo": null, + "tiers": [ + { + "condition": { "metric": "inputTokens", "min": 1, "max": 272000 }, + "rates": { "input": 2, "output": 10, "cacheRead": 0.2, "cacheWrite": 2.5 } + }, + { + "condition": { "metric": "inputTokens", "min": 272001, "max": 922000 }, + "rates": { "input": 4, "output": 15, "cacheRead": 0.4, "cacheWrite": 5 } + } + ] + } + ] + }, + "releaseDate": "2026-09-22", + "lifecycle": "latest", + "knowledgeCutoff": "2026-04-20", + "inputModalities": ["text", "image"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": null, + "artificialAnalysis": "https://artificialanalysis.ai/models/gpt-6-sol", + "openrouter": null + } +} diff --git a/manifests/models/grok-4-7.json b/manifests/models/grok-4-7.json new file mode 100644 index 00000000..eca0b0e7 --- /dev/null +++ b/manifests/models/grok-4-7.json @@ -0,0 +1,123 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "grok-4-7", + "name": "Grok 4.7", + "description": "SpaceXAI's frontier model for coding, agentic tasks, and knowledge work, with configurable reasoning and a 500K-token multimodal context window.", + "translations": { + "de": { + "description": "SpaceXAIs Frontier-Modell für Programmierung, agentische Aufgaben und Wissensarbeit mit konfigurierbarem Reasoning und multimodalem 500K-Token-Kontext." + }, + "es": { + "description": "Modelo de frontera de SpaceXAI para programación, tareas agénticas y trabajo del conocimiento, con razonamiento configurable y contexto multimodal de 500K tokens." + }, + "fr": { + "description": "Modèle de pointe de SpaceXAI pour le codage, les tâches agentiques et le travail intellectuel, avec raisonnement configurable et contexte multimodal de 500K jetons." + }, + "id": { + "description": "Model frontier SpaceXAI untuk coding, tugas agentik, dan pekerjaan pengetahuan, dengan penalaran yang dapat dikonfigurasi serta konteks multimodal 500K token." + }, + "ja": { + "description": "設定可能な推論と50万トークンのマルチモーダルコンテキストを備え、コーディング、エージェント型タスク、知識労働に対応するSpaceXAIのフロンティアモデル。" + }, + "ko": { + "description": "구성 가능한 추론과 50만 토큰 멀티모달 컨텍스트를 갖춘 코딩, 에이전트 작업 및 지식 업무용 SpaceXAI 프런티어 모델입니다." + }, + "pt": { + "description": "Modelo de fronteira da SpaceXAI para programação, tarefas agênticas e trabalho de conhecimento, com raciocínio configurável e contexto multimodal de 500K tokens." + }, + "ru": { + "description": "Передовая модель SpaceXAI для программирования, агентных задач и интеллектуальной работы с настраиваемыми рассуждениями и мультимодальным контекстом 500K токенов." + }, + "tr": { + "description": "Yapılandırılabilir akıl yürütme ve 500K token çok modlu bağlam ile kodlama, ajan görevleri ve bilgi çalışmaları için SpaceXAI'ın öncü modeli." + }, + "zh-Hans": { + "description": "SpaceXAI 面向编码、智能体任务和知识工作的前沿模型,支持可配置推理和 50 万 token 多模态上下文。" + }, + "zh-Hant": { + "description": "SpaceXAI 面向程式設計、智慧體任務和知識工作的前沿模型,支援可設定推理和 50 萬 token 多模態上下文。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://docs.x.ai/developers/models/grok-4.7", + "title": "Grok 4.7 model documentation", + "fields": [ + "name", + "description", + "docsUrl", + "contextWindow", + "tokenPricing", + "knowledgeCutoff", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://docs.x.ai/developers/release-notes", + "title": "SpaceXAI API release notes", + "fields": ["releaseDate", "lifecycle"] + }, + { + "url": "https://artificialanalysis.ai/models/grok-4-7", + "title": "Artificial Analysis Grok 4.7 model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://x.ai", + "docsUrl": "https://docs.x.ai/developers/models/grok-4.7", + "vendor": "xAI", + "size": null, + "activeParameters": null, + "contextWindow": 500000, + "maxOutput": null, + "tokenPricing": { + "status": "available", + "primaryOffer": "global-standard", + "offers": [ + { + "id": "global-standard", + "currency": "USD", + "region": "global", + "serviceTier": "standard", + "effectiveFrom": "2026-09-21", + "effectiveTo": null, + "tiers": [ + { + "condition": { "metric": "inputTokens", "min": 1, "max": 200000 }, + "rates": { "input": 2, "output": 6, "cacheRead": 0.5, "cacheWrite": null } + }, + { + "condition": { "metric": "inputTokens", "min": 200001, "max": 500000 }, + "rates": { "input": 4, "output": 12, "cacheRead": 1, "cacheWrite": null } + } + ] + } + ] + }, + "releaseDate": "2026-09-21", + "lifecycle": "latest", + "knowledgeCutoff": "2026-05", + "inputModalities": ["text", "image"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": null, + "artificialAnalysis": "https://artificialanalysis.ai/models/grok-4-7", + "openrouter": null + } +} diff --git a/manifests/models/mimo-v2-6-pro.json b/manifests/models/mimo-v2-6-pro.json new file mode 100644 index 00000000..16c3e627 --- /dev/null +++ b/manifests/models/mimo-v2-6-pro.json @@ -0,0 +1,124 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "mimo-v2-6-pro", + "name": "MiMo-V2.6-Pro", + "description": "Xiaomi's open-weight trillion-parameter flagship for complex projects, long-horizon agents, cybersecurity, and multimodal research with a 1M-token context.", + "translations": { + "de": { + "description": "Das offene Flaggschiffmodell von Xiaomi mit einer Billion Parametern für komplexe Projekte, Langzeitagenten, Cybersicherheit und multimodale Forschung mit 1 Mio. Token Kontext." + }, + "es": { + "description": "Modelo insignia abierto de Xiaomi con un billón de parámetros para proyectos complejos, agentes prolongados, ciberseguridad e investigación multimodal con contexto de 1 M de tokens." + }, + "fr": { + "description": "Modèle phare ouvert de Xiaomi à mille milliards de paramètres pour les projets complexes, agents de longue durée, cybersécurité et recherche multimodale avec 1 M de jetons de contexte." + }, + "id": { + "description": "Model unggulan terbuka Xiaomi berparameter satu triliun untuk proyek kompleks, agen jangka panjang, keamanan siber, dan riset multimodal dengan konteks 1 juta token." + }, + "ja": { + "description": "複雑なプロジェクト、長期稼働エージェント、サイバーセキュリティ、マルチモーダル研究向けに100万トークンのコンテキストを備えたXiaomiの1兆パラメータ級オープンウェイト旗艦モデル。" + }, + "ko": { + "description": "복잡한 프로젝트, 장기 실행 에이전트, 사이버 보안 및 멀티모달 연구를 위해 100만 토큰 컨텍스트를 제공하는 Xiaomi의 1조 파라미터급 오픈 웨이트 플래그십입니다." + }, + "pt": { + "description": "Modelo aberto principal da Xiaomi com 1T parâmetros para projetos complexos, agentes duradouros, cibersegurança e pesquisa multimodal com contexto de 1 milhão de tokens." + }, + "ru": { + "description": "Флагманская модель Xiaomi с открытыми весами и триллионом параметров для сложных проектов, длительных агентов, кибербезопасности и мультимодальных исследований с контекстом 1 млн токенов." + }, + "tr": { + "description": "Karmaşık projeler, uzun süreli ajanlar, siber güvenlik ve çok modlu araştırma için 1 milyon token bağlam sunan Xiaomi'nin trilyon parametreli açık ağırlıklı amiral modeli." + }, + "zh-Hans": { + "description": "Xiaomi 面向复杂项目、长时程智能体、网络安全和多模态研究的万亿参数开放权重旗舰模型,提供 100 万 token 上下文。" + }, + "zh-Hant": { + "description": "Xiaomi 面向複雜專案、長時程智慧體、網路安全和多模態研究的萬億參數開放權重旗艦模型,提供 100 萬 token 上下文。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://mimo.mi.com/models/en-US/mimo-v2.6-pro", + "title": "MiMo-V2.6-Pro model page", + "fields": [ + "name", + "description", + "docsUrl", + "contextWindow", + "maxOutput", + "tokenPricing", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://mimo.mi.com/docs/en-US/news/latest/v2-6", + "title": "MiMo-V2.6 release announcement", + "fields": ["releaseDate", "lifecycle", "size", "activeParameters"] + }, + { + "url": "https://huggingface.co/XiaomiMiMo/MiMo-V2.6-Pro-RL", + "title": "MiMo-V2.6-Pro-RL official model card", + "fields": ["size", "activeParameters", "platformUrls.huggingface"] + }, + { + "url": "https://artificialanalysis.ai/models/mimo-v2-6-pro", + "title": "Artificial Analysis MiMo-V2.6-Pro model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://mimo.xiaomi.com/", + "docsUrl": "https://mimo.mi.com/models/en-US/mimo-v2.6-pro", + "vendor": "Xiaomi", + "size": "1.02T", + "activeParameters": "42B", + "contextWindow": 1000000, + "maxOutput": 128000, + "tokenPricing": { + "status": "available", + "primaryOffer": "global-standard", + "offers": [ + { + "id": "global-standard", + "currency": "USD", + "region": "global", + "serviceTier": "standard", + "effectiveFrom": "2026-09-22", + "effectiveTo": null, + "tiers": [ + { + "condition": null, + "rates": { "input": 0.435, "output": 0.87, "cacheRead": 0.0036, "cacheWrite": null } + } + ] + } + ] + }, + "releaseDate": "2026-09-22", + "lifecycle": "latest", + "knowledgeCutoff": null, + "inputModalities": ["text", "image", "audio", "video"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": "https://huggingface.co/XiaomiMiMo/MiMo-V2.6-Pro-RL", + "artificialAnalysis": "https://artificialanalysis.ai/models/mimo-v2-6-pro", + "openrouter": null + } +} diff --git a/manifests/models/qwen3-8-max-0902.json b/manifests/models/qwen3-8-max-0902.json new file mode 100644 index 00000000..cf8936ac --- /dev/null +++ b/manifests/models/qwen3-8-max-0902.json @@ -0,0 +1,121 @@ +{ + "$schema": "../$schemas/model.schema.json", + "id": "qwen3-8-max-0902", + "name": "Qwen3.8-Max-0902", + "description": "Alibaba's upgraded Qwen3.8-Max snapshot for engineering-scale coding, long-horizon autonomous development, multi-tool agents, and multimodal understanding.", + "translations": { + "de": { + "description": "Alibabas aktualisierter Qwen3.8-Max-Snapshot für Programmierung im Projektmaßstab, langfristige autonome Entwicklung, Multi-Tool-Agenten und multimodales Verständnis." + }, + "es": { + "description": "Versión actualizada de Qwen3.8-Max de Alibaba para programación a escala de ingeniería, desarrollo autónomo prolongado, agentes multiherramienta y comprensión multimodal." + }, + "fr": { + "description": "Version actualisée de Qwen3.8-Max d’Alibaba pour le codage à l’échelle de projets, le développement autonome prolongé, les agents multi-outils et la compréhension multimodale." + }, + "id": { + "description": "Snapshot Qwen3.8-Max Alibaba yang ditingkatkan untuk coding skala rekayasa, pengembangan otonom jangka panjang, agen multi-alat, dan pemahaman multimodal." + }, + "ja": { + "description": "エンジニアリング規模のコーディング、長期自律開発、複数ツールのエージェント、マルチモーダル理解を強化したAlibabaのQwen3.8-Maxスナップショット。" + }, + "ko": { + "description": "엔지니어링 규모 코딩, 장기 자율 개발, 다중 도구 에이전트와 멀티모달 이해를 강화한 Alibaba의 Qwen3.8-Max 스냅샷입니다." + }, + "pt": { + "description": "Versão atualizada do Qwen3.8-Max da Alibaba para programação em escala de engenharia, desenvolvimento autônomo prolongado, agentes com múltiplas ferramentas e compreensão multimodal." + }, + "ru": { + "description": "Обновлённая версия Qwen3.8-Max от Alibaba для программирования инженерного масштаба, длительной автономной разработки, агентов с множеством инструментов и мультимодального понимания." + }, + "tr": { + "description": "Mühendislik ölçeğinde kodlama, uzun süreli otonom geliştirme, çok araçlı ajanlar ve çok modlu anlama için Alibaba'nın yükseltilmiş Qwen3.8-Max sürümü." + }, + "zh-Hans": { + "description": "阿里巴巴升级版 Qwen3.8-Max 快照,面向工程规模编码、长时程自主开发、多工具智能体和多模态理解。" + }, + "zh-Hant": { + "description": "阿里巴巴升級版 Qwen3.8-Max 快照,面向工程規模程式設計、長時程自主開發、多工具智慧體和多模態理解。" + } + }, + "verified": true, + "sources": [ + { + "url": "https://www.qwencloud.com/models/qwen3.8-max-0902", + "title": "Qwen3.8-Max-0902 model page", + "fields": [ + "name", + "description", + "docsUrl", + "contextWindow", + "maxOutput", + "tokenPricing", + "releaseDate", + "lifecycle", + "inputModalities", + "outputModalities", + "capabilities" + ] + }, + { + "url": "https://qwen.ai/blog?id=qwen3.8", + "title": "Qwen3.8-Max: A New Bar for Coding and Cowork", + "fields": ["size", "activeParameters"] + }, + { + "url": "https://artificialanalysis.ai/models/qwen3-8-max", + "title": "Artificial Analysis Qwen3.8 Max 0902 model page", + "fields": ["platformUrls.artificialAnalysis"] + } + ], + "lastVerifiedAt": "2026-09-24", + "verifiedBy": "codex-agent", + "confidence": "high", + "websiteUrl": "https://qwen.ai", + "docsUrl": "https://www.qwencloud.com/models/qwen3.8-max-0902", + "vendor": "Alibaba", + "size": "2.4T", + "activeParameters": "95B", + "contextWindow": 1000000, + "maxOutput": 131072, + "tokenPricing": { + "status": "available", + "primaryOffer": "global-standard", + "offers": [ + { + "id": "global-standard", + "currency": "USD", + "region": "global", + "serviceTier": "standard", + "effectiveFrom": "2026-09-02", + "effectiveTo": null, + "tiers": [ + { + "condition": null, + "rates": { "input": 2, "output": 6, "cacheRead": 0.25, "cacheWrite": 2.5 } + } + ] + } + ] + }, + "releaseDate": "2026-09-02", + "lifecycle": "latest", + "knowledgeCutoff": null, + "inputModalities": ["text", "image", "video"], + "outputModalities": ["text"], + "capabilities": ["function-calling", "tool-choice", "structured-outputs", "reasoning"], + "benchmarks": { + "sweBench": null, + "terminalBench": null, + "mmmu": null, + "mmmuPro": null, + "webDevArena": null, + "sciCode": null, + "liveCodeBench": null + }, + "platformUrls": { + "huggingface": null, + "artificialAnalysis": "https://artificialanalysis.ai/models/qwen3-8-max", + "openrouter": null + } +} diff --git a/manifests/models/qwen3-8-max.json b/manifests/models/qwen3-8-max.json index af13795a..9c458f82 100644 --- a/manifests/models/qwen3-8-max.json +++ b/manifests/models/qwen3-8-max.json @@ -78,8 +78,8 @@ "fields": ["description"] }, { - "url": "https://artificialanalysis.ai/models/qwen3-8-max", - "title": "Artificial Analysis Qwen3.8 Max model page", + "url": "https://artificialanalysis.ai/models/qwen3-8-max-0803", + "title": "Artificial Analysis Qwen3.8 Max 0803 model page", "fields": ["releaseDate", "platformUrls.artificialAnalysis"] } ], @@ -135,7 +135,7 @@ }, "platformUrls": { "huggingface": null, - "artificialAnalysis": "https://artificialanalysis.ai/models/qwen3-8-max", + "artificialAnalysis": "https://artificialanalysis.ai/models/qwen3-8-max-0803", "openrouter": null } } diff --git a/manifests/vendors/alibaba.json b/manifests/vendors/alibaba.json index 5c7640d1..3024a4e8 100644 --- a/manifests/vendors/alibaba.json +++ b/manifests/vendors/alibaba.json @@ -10,7 +10,7 @@ { "id": "qwen-max", "name": "Qwen Max", - "modelIds": ["qwen3-6-max-preview", "qwen3-7-max", "qwen3-8-max"] + "modelIds": ["qwen3-6-max-preview", "qwen3-7-max", "qwen3-8-max", "qwen3-8-max-0902"] }, { "id": "qwen-plus", @@ -100,6 +100,11 @@ "title": "Qwen3.8-Max: A New Bar for Coding and Cowork", "fields": ["modelSeries"] }, + { + "url": "https://www.qwencloud.com/models/qwen3.8-max-0902", + "title": "Qwen3.8-Max-0902 model page", + "fields": ["modelSeries"] + }, { "url": "https://huggingface.co/Qwen/Qwen3.8-2.4T-A95B", "title": "Qwen3.8-2.4T-A95B official model card", @@ -116,7 +121,7 @@ "fields": ["modelSeries"] } ], - "lastVerifiedAt": "2026-08-28", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://www.alibabagroup.com", diff --git a/manifests/vendors/anthropic.json b/manifests/vendors/anthropic.json index 60819ef8..882ff6c0 100644 --- a/manifests/vendors/anthropic.json +++ b/manifests/vendors/anthropic.json @@ -18,7 +18,8 @@ "claude-opus-4-6", "claude-opus-4-7", "claude-opus-4-8", - "claude-opus-5" + "claude-opus-5", + "claude-opus-5-5" ] }, { @@ -115,7 +116,7 @@ "fields": ["modelSeries"] } ], - "lastVerifiedAt": "2026-09-04", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://www.anthropic.com", diff --git a/manifests/vendors/deepseek.json b/manifests/vendors/deepseek.json index 57dd7162..9caebda4 100644 --- a/manifests/vendors/deepseek.json +++ b/manifests/vendors/deepseek.json @@ -25,7 +25,12 @@ { "id": "deepseek-flash", "name": "DeepSeek Flash", - "modelIds": ["deepseek-v4-flash-preview", "deepseek-v4-flash", "deepseek-v4-flash-vision"] + "modelIds": [ + "deepseek-v4-flash-preview", + "deepseek-v4-flash", + "deepseek-v4-flash-vision", + "deepseek-v4-1-flash" + ] } ], "description": "A leading AI research company focused on developing advanced language models and AI technology for coding and general intelligence.", @@ -98,11 +103,11 @@ }, { "url": "https://api-docs.deepseek.com/updates/", - "title": "DeepSeek V4 Flash Vision release", + "title": "DeepSeek model release changelog", "fields": ["modelSeries"] } ], - "lastVerifiedAt": "2026-08-28", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://www.deepseek.com", diff --git a/manifests/vendors/openai.json b/manifests/vendors/openai.json index a988325b..e99cfbb9 100644 --- a/manifests/vendors/openai.json +++ b/manifests/vendors/openai.json @@ -18,7 +18,8 @@ "gpt-5-2", "gpt-5-4", "gpt-5-5", - "gpt-5-6-sol" + "gpt-5-6-sol", + "gpt-6-sol" ] }, { @@ -36,7 +37,7 @@ { "id": "gpt-nano-luna", "name": "GPT nano(Luna)", - "modelIds": ["gpt-4-1-nano", "gpt-5-nano", "gpt-5-4-nano", "gpt-5-6-luna"] + "modelIds": ["gpt-4-1-nano", "gpt-5-nano", "gpt-5-4-nano", "gpt-5-6-luna", "gpt-6-luna"] }, { "id": "gpt-codex", @@ -133,12 +134,12 @@ "fields": ["communityUrls.reddit"] }, { - "url": "https://developers.openai.com/api/docs/models/gpt-6-astra", - "title": "GPT-6 Astra model documentation", + "url": "https://developers.openai.com/api/docs/models", + "title": "OpenAI model catalog", "fields": ["modelSeries"] } ], - "lastVerifiedAt": "2026-09-04", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://openai.com", diff --git a/manifests/vendors/xai.json b/manifests/vendors/xai.json index 5685bff0..bf23497c 100644 --- a/manifests/vendors/xai.json +++ b/manifests/vendors/xai.json @@ -10,7 +10,7 @@ { "id": "grok", "name": "Grok", - "modelIds": ["grok-4", "grok-4-20", "grok-4-3", "grok-4-5", "grok-4-6"] + "modelIds": ["grok-4", "grok-4-20", "grok-4-3", "grok-4-5", "grok-4-6", "grok-4-7"] }, { "id": "grok-code", @@ -92,12 +92,12 @@ "fields": ["communityUrls.blog"] }, { - "url": "https://x.ai/news/grok-4-6", - "title": "Introducing Grok 4.6", + "url": "https://docs.x.ai/developers/models/grok-4.7", + "title": "Grok 4.7 model documentation", "fields": ["modelSeries"] } ], - "lastVerifiedAt": "2026-08-13", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://x.ai", diff --git a/manifests/vendors/xiaomi.json b/manifests/vendors/xiaomi.json index 0595cb4f..695a3ccd 100644 --- a/manifests/vendors/xiaomi.json +++ b/manifests/vendors/xiaomi.json @@ -10,7 +10,7 @@ { "id": "mimo-pro", "name": "MiMo Pro", - "modelIds": ["mimo-v2-5-pro"] + "modelIds": ["mimo-v2-5-pro", "mimo-v2-6-pro"] }, { "id": "mimo", @@ -70,9 +70,14 @@ "url": "https://github.com/XiaomiMiMo", "title": "Xiaomi MiMo official GitHub organization", "fields": ["communityUrls.github"] + }, + { + "url": "https://mimo.mi.com/docs/en-US/news/latest/v2-6", + "title": "MiMo-V2.6 release announcement", + "fields": ["modelSeries"] } ], - "lastVerifiedAt": "2026-07-27", + "lastVerifiedAt": "2026-09-24", "verifiedBy": "codex-agent", "confidence": "high", "websiteUrl": "https://mimo.xiaomi.com/", diff --git a/src/lib/generated/metadata.ts b/src/lib/generated/metadata.ts index 408ef1e8..a3243514 100644 --- a/src/lib/generated/metadata.ts +++ b/src/lib/generated/metadata.ts @@ -3639,7 +3639,7 @@ export const stackCounts: Record = { clis: 30, desktops: 13, extensions: 19, - models: 138, + models: 145, 'model-providers': 17, vendors: 48, } diff --git a/src/lib/generated/models.ts b/src/lib/generated/models.ts index 0efcb2e3..4e959ed2 100644 --- a/src/lib/generated/models.ts +++ b/src/lib/generated/models.ts @@ -17,6 +17,7 @@ import ClaudeOpus46 from '../../../manifests/models/claude-opus-4-6.json' import ClaudeOpus47 from '../../../manifests/models/claude-opus-4-7.json' import ClaudeOpus48 from '../../../manifests/models/claude-opus-4-8.json' import ClaudeOpus5 from '../../../manifests/models/claude-opus-5.json' +import ClaudeOpus55 from '../../../manifests/models/claude-opus-5-5.json' import ClaudeSonnet3 from '../../../manifests/models/claude-sonnet-3.json' import ClaudeSonnet3520240620 from '../../../manifests/models/claude-sonnet-3-5-20240620.json' import ClaudeSonnet3520241022 from '../../../manifests/models/claude-sonnet-3-5-20241022.json' @@ -34,6 +35,7 @@ import DeepseekV3 from '../../../manifests/models/deepseek-v3.json' import DeepseekV31 from '../../../manifests/models/deepseek-v3-1.json' import DeepseekV32Exp from '../../../manifests/models/deepseek-v3-2-exp.json' import DeepseekV3Terminus from '../../../manifests/models/deepseek-v3-terminus.json' +import DeepseekV41Flash from '../../../manifests/models/deepseek-v4-1-flash.json' import DeepseekV4Flash from '../../../manifests/models/deepseek-v4-flash.json' import DeepseekV4FlashPreview from '../../../manifests/models/deepseek-v4-flash-preview.json' import DeepseekV4FlashVision from '../../../manifests/models/deepseek-v4-flash-vision.json' @@ -92,11 +94,14 @@ import Gpt5Codex from '../../../manifests/models/gpt-5-codex.json' import Gpt5Mini from '../../../manifests/models/gpt-5-mini.json' import Gpt5Nano from '../../../manifests/models/gpt-5-nano.json' import Gpt6Astra from '../../../manifests/models/gpt-6-astra.json' +import Gpt6Luna from '../../../manifests/models/gpt-6-luna.json' +import Gpt6Sol from '../../../manifests/models/gpt-6-sol.json' import Grok4 from '../../../manifests/models/grok-4.json' import Grok41Fast from '../../../manifests/models/grok-4-1-fast.json' import Grok43 from '../../../manifests/models/grok-4-3.json' import Grok45 from '../../../manifests/models/grok-4-5.json' import Grok46 from '../../../manifests/models/grok-4-6.json' +import Grok47 from '../../../manifests/models/grok-4-7.json' import Grok420 from '../../../manifests/models/grok-4-20.json' import Grok4Fast from '../../../manifests/models/grok-4-fast.json' import GrokCodeFast1 from '../../../manifests/models/grok-code-fast-1.json' @@ -112,6 +117,7 @@ import Llama4Maverick from '../../../manifests/models/llama-4-maverick.json' import Llama4Scout from '../../../manifests/models/llama-4-scout.json' import MimoV25 from '../../../manifests/models/mimo-v2-5.json' import MimoV25Pro from '../../../manifests/models/mimo-v2-5-pro.json' +import MimoV26Pro from '../../../manifests/models/mimo-v2-6-pro.json' import MimoV2Flash from '../../../manifests/models/mimo-v2-flash.json' import MinimaxM2 from '../../../manifests/models/minimax-m2.json' import MinimaxM21 from '../../../manifests/models/minimax-m2-1.json' @@ -139,6 +145,7 @@ import Qwen3824tA95b from '../../../manifests/models/qwen3-8-2-4t-a95b.json' import Qwen3827b from '../../../manifests/models/qwen3-8-27b.json' import Qwen38FlashNext from '../../../manifests/models/qwen3-8-flash-next.json' import Qwen38Max from '../../../manifests/models/qwen3-8-max.json' +import Qwen38Max0902 from '../../../manifests/models/qwen3-8-max-0902.json' import Qwen3Coder30bA3b from '../../../manifests/models/qwen3-coder-30b-a3b.json' import Qwen3Coder480bA35b from '../../../manifests/models/qwen3-coder-480b-a35b.json' import Qwen3CoderNext from '../../../manifests/models/qwen3-coder-next.json' @@ -157,6 +164,7 @@ export const modelsData = [ ClaudeOpus47, ClaudeOpus48, ClaudeOpus4, + ClaudeOpus55, ClaudeOpus5, ClaudeSonnet3520240620, ClaudeSonnet3520241022, @@ -175,6 +183,7 @@ export const modelsData = [ DeepseekV32Exp, DeepseekV3Terminus, DeepseekV3, + DeepseekV41Flash, DeepseekV4FlashPreview, DeepseekV4FlashVision, DeepseekV4Flash, @@ -233,10 +242,13 @@ export const modelsData = [ Gpt5Nano, Gpt5, Gpt6Astra, + Gpt6Luna, + Gpt6Sol, Grok41Fast, Grok43, Grok45, Grok46, + Grok47, Grok420, Grok4Fast, Grok4, @@ -253,6 +265,7 @@ export const modelsData = [ Llama4Scout, MimoV25Pro, MimoV25, + MimoV26Pro, MimoV2Flash, MinimaxM21, MinimaxM25, @@ -279,6 +292,7 @@ export const modelsData = [ Qwen3824tA95b, Qwen3827b, Qwen38FlashNext, + Qwen38Max0902, Qwen38Max, Qwen3Coder30bA3b, Qwen3Coder480bA35b, diff --git a/tests/model-intelligence-index.test.ts b/tests/model-intelligence-index.test.ts index 49b47f37..d012bfc2 100644 --- a/tests/model-intelligence-index.test.ts +++ b/tests/model-intelligence-index.test.ts @@ -198,8 +198,9 @@ describe('model intelligence index', () => { ]) ).toEqual([ ['qwen3-6-max-preview', 28, true, 'Qwen3.6 Max Preview'], - ['qwen3-7-max', 30, false, 'Qwen3.7 Max'], + ['qwen3-7-max', 29, false, 'Qwen3.7 Max'], ['qwen3-8-max', 40, false, 'Qwen3.8 Max'], + ['qwen3-8-max-0902', 45, false, 'Qwen3.8 Max (0902)'], ]) expect(qwenSeries[3]?.points.find(point => point.modelId === 'qwen3-8-27b')).toMatchObject({ modelId: 'qwen3-8-27b', @@ -209,22 +210,22 @@ describe('model intelligence index', () => { }) expect(qwenSeries[3]?.points.at(-1)).toMatchObject({ modelId: 'qwen3-8-flash-next', - score: 42, + score: 40, estimated: true, configuration: 'Qwen3.8-Flash-Next', }) }) - it('connects Grok 4.6 to the flagship Grok line', () => { + it('connects Grok 4.7 to the flagship Grok line', () => { const grokSeries = modelIntelligenceSeries.find( series => series.vendor === 'xAI' && series.name === 'Grok' ) expect(grokSeries?.points.at(-1)).toMatchObject({ - modelId: 'grok-4-6', - score: 44, + modelId: 'grok-4-7', + score: 46, estimated: false, - configuration: 'Grok 4.6 (high)', + configuration: 'Grok 4.7 (xhigh)', }) }) @@ -247,7 +248,7 @@ describe('model intelligence index', () => { modelId: 'claude-fable-5-1', score: 53, estimated: false, - configuration: 'Claude Fable 5.1 (Adaptive Reasoning, Max Effort, Default Fallback)', + configuration: 'Claude Fable 5.1 (max with fallback)', }) }) @@ -272,6 +273,18 @@ describe('model intelligence index', () => { configuration: 'GPT-6 Astra (max)', }), ]) + expect(openAISeries[0]?.points.at(-1)).toMatchObject({ + modelId: 'gpt-6-sol', + score: 48, + estimated: false, + configuration: 'GPT-6 Sol (max)', + }) + expect(openAISeries[2]?.points.at(-1)).toMatchObject({ + modelId: 'gpt-6-luna', + score: 37, + estimated: false, + configuration: 'GPT-6 Luna (max)', + }) }) it('connects Muse Spark 1.3 to the existing Muse series', () => { @@ -307,7 +320,7 @@ describe('model intelligence index', () => { 'deepseek-v4-pro', ]) expect(deepSeekSeries[0]?.points.slice(-2).map(point => [point.modelId, point.score])).toEqual([ - ['deepseek-v4-pro-preview', 31], + ['deepseek-v4-pro-preview', 30], ['deepseek-v4-pro', 36], ]) expect( @@ -318,9 +331,10 @@ describe('model intelligence index', () => { point.configuration, ]) ).toEqual([ - ['deepseek-v4-flash-preview', 25, false, 'DeepSeek V4 Flash (Reasoning, High Effort)'], - ['deepseek-v4-flash', 35, false, 'DeepSeek V4 Flash 0731 (Reasoning, Max Effort)'], - ['deepseek-v4-flash-vision', 35, false, 'DeepSeek V4 Flash Vision (Reasoning, Max Effort)'], + ['deepseek-v4-flash-preview', 26, false, 'DeepSeek V4 Flash (high)'], + ['deepseek-v4-flash', 34, false, 'DeepSeek V4 Flash 0731 (max)'], + ['deepseek-v4-flash-vision', 35, false, 'DeepSeek V4 Flash Vision (max)'], + ['deepseek-v4-1-flash', 39, false, 'DeepSeek V4.1 Flash (max)'], ]) }) @@ -352,7 +366,7 @@ describe('model intelligence index', () => { expect(mimoSeries.map(series => series.name)).toEqual(['MiMo Pro', 'MiMo', 'MiMo Flash']) expect(mimoSeries.map(series => series.points.map(point => point.modelId))).toEqual([ - ['mimo-v2-5-pro'], + ['mimo-v2-5-pro', 'mimo-v2-6-pro'], ['mimo-v2-5'], ['mimo-v2-flash'], ]) @@ -505,7 +519,7 @@ describe('model intelligence index', () => { expect(haikuPoints).toEqual([ ['claude-haiku-3', 6, true], ['claude-haiku-3-5', 9, true], - ['claude-haiku-4-5', 18, false], + ['claude-haiku-4-5', 17, false], ]) expect( @@ -532,8 +546,8 @@ describe('model intelligence index', () => { ['glm-5', 'GLM', 28], ['glm-5-turbo', 'GLM Turbo', 27], ['glm-5v-turbo', 'GLM Vision', 23], - ['glm-5-1', 'GLM', 27], - ['glm-5-2', 'GLM', 39], + ['glm-5-1', 'GLM', 26], + ['glm-5-2', 'GLM', 34], ['glm-5-3', 'GLM', 45], ['glm-5-3-flash', 'GLM Air / Flash', 42], ]) @@ -546,7 +560,7 @@ describe('model intelligence index', () => { expect(modelIntelligenceMeta.methodologyUrl).toBe( 'https://artificialanalysis.ai/methodology/intelligence-benchmarking' ) - expect(modelIntelligenceMeta.indexVersion).toBe('4.3') + expect(modelIntelligenceMeta.indexVersion).toBe('4.3.2') expect(modelIntelligenceMeta.observedAt).toMatch(/^\d{4}-\d{2}-\d{2}$/) }) }) diff --git a/tests/model-price-intelligence-index.test.ts b/tests/model-price-intelligence-index.test.ts index d1e2df08..38c2b0ea 100644 --- a/tests/model-price-intelligence-index.test.ts +++ b/tests/model-price-intelligence-index.test.ts @@ -91,7 +91,7 @@ describe('model price-intelligence index', () => { ) const hy3 = modelPriceIntelligencePoints.find(point => point.modelId === 'hy3') - expect(hy3).toMatchObject({ inputPrice: 0.14, outputPrice: 0.58, score: 26 }) + expect(hy3).toMatchObject({ inputPrice: 0.14, outputPrice: 0.58, score: 25 }) expect(hy3?.blendedPrice).toBeCloseTo(0.184) })