@yanlinglabs/winter-provider-catalog 0.0.23 → 0.0.27

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -7046,6 +7046,7 @@
7046
7046
  "id": "xai",
7047
7047
  "displayName": "xAI (Grok)",
7048
7048
  "protocols": [
7049
+ "openai-responses",
7049
7050
  "openai-chat-completions"
7050
7051
  ],
7051
7052
  "authKinds": [
@@ -7056,7 +7057,7 @@
7056
7057
  },
7057
7058
  "modelDiscovery": "openai-models",
7058
7059
  "liveCatalogAuthority": "unknown",
7059
- "adapterId": "winter.openai-chat-completions",
7060
+ "adapterId": "winter.openai-responses",
7060
7061
  "family": "openai",
7061
7062
  "upstream": {
7062
7063
  "project": "winter",
@@ -10001,13 +10002,13 @@
10001
10002
  },
10002
10003
  "pricing": {
10003
10004
  "value": {
10004
- "inputPerMTokUsd": 0.144,
10005
- "outputPerMTokUsd": 0.574
10005
+ "inputPerMTokUsd": 0.149,
10006
+ "outputPerMTokUsd": 0.596
10006
10007
  },
10007
10008
  "source": "official-doc",
10008
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-coder-next — China (Beijing) qwen3-coder-next: input $0.144, output $0.574 per million tokens; ≤32K input; >32K–128K $0.216/$0.861; >128K–256K $0.359/$1.434",
10009
- "confidence": "declared",
10010
- "observedAt": "2026-09-25T09:40:00Z"
10009
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3-coder-next, ≤32K; 32K–128K: ¥1.5/¥6; 128K–256K: ¥2.5/¥10 input tier: ¥1 input/¥4 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10010
+ "confidence": "inferred",
10011
+ "observedAt": "2026-09-25T12:31:56.381Z"
10011
10012
  },
10012
10013
  "unsupportedParameters": [],
10013
10014
  "status": "candidate",
@@ -10109,13 +10110,13 @@
10109
10110
  },
10110
10111
  "pricing": {
10111
10112
  "value": {
10112
- "inputPerMTokUsd": 0.574,
10113
- "outputPerMTokUsd": 2.294
10113
+ "inputPerMTokUsd": 0.596,
10114
+ "outputPerMTokUsd": 2.384
10114
10115
  },
10115
10116
  "source": "official-doc",
10116
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-coder-plus — China (Beijing) qwen3-coder-plus: input $0.574, output $2.294 per million tokens; ≤32K input; >32K–128K $0.861/$3.441; >128K–256K $1.434/$5.735; >256K–1M $2.868/$28.671",
10117
- "confidence": "declared",
10118
- "observedAt": "2026-09-25T09:40:00Z"
10117
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3-coder-plus, ≤32K; 32K–128K: ¥6/¥24; 128K–256K: ¥10/¥40; 256K–1M: ¥20/¥200 input tier: ¥4 input/¥16 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10118
+ "confidence": "inferred",
10119
+ "observedAt": "2026-09-25T12:31:56.381Z"
10119
10120
  },
10120
10121
  "unsupportedParameters": [],
10121
10122
  "status": "candidate",
@@ -10219,13 +10220,13 @@
10219
10220
  },
10220
10221
  "pricing": {
10221
10222
  "value": {
10222
- "inputPerMTokUsd": 0.115,
10223
- "outputPerMTokUsd": 0.917
10223
+ "inputPerMTokUsd": 0.1192,
10224
+ "outputPerMTokUsd": 0.9536
10224
10225
  },
10225
10226
  "source": "official-doc",
10226
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-5-122b-a10b — China (Beijing) qwen3.5-122b-a10b: input $0.115, output $0.917 per million tokens; ≤128K input; >128K–256K $0.287/$2.294",
10227
- "confidence": "declared",
10228
- "observedAt": "2026-09-25T09:40:00Z"
10227
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.5-122b-a10b, ≤128K; >128K–256K: ¥2/¥16 input tier: ¥0.8 input/¥6.4 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10228
+ "confidence": "inferred",
10229
+ "observedAt": "2026-09-25T12:31:56.381Z"
10229
10230
  },
10230
10231
  "unsupportedParameters": [],
10231
10232
  "status": "candidate",
@@ -10329,13 +10330,13 @@
10329
10330
  },
10330
10331
  "pricing": {
10331
10332
  "value": {
10332
- "inputPerMTokUsd": 0.172,
10333
- "outputPerMTokUsd": 1.032
10333
+ "inputPerMTokUsd": 0.1788,
10334
+ "outputPerMTokUsd": 1.0728
10334
10335
  },
10335
10336
  "source": "official-doc",
10336
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-5-397b-a17b — China (Beijing) qwen3.5-397b-a17b: input $0.172, output $1.032 per million tokens; ≤128K input; >128K–256K $0.43/$2.58",
10337
- "confidence": "declared",
10338
- "observedAt": "2026-09-25T09:40:00Z"
10337
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.5-397b-a17b, ≤128K; >128K–256K: ¥3/¥18 input tier: ¥1.2 input/¥7.2 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10338
+ "confidence": "inferred",
10339
+ "observedAt": "2026-09-25T12:31:56.381Z"
10339
10340
  },
10340
10341
  "unsupportedParameters": [],
10341
10342
  "status": "candidate",
@@ -10425,13 +10426,13 @@
10425
10426
  },
10426
10427
  "pricing": {
10427
10428
  "value": {
10428
- "inputPerMTokUsd": 0.115,
10429
- "outputPerMTokUsd": 0.688
10429
+ "inputPerMTokUsd": 0.1192,
10430
+ "outputPerMTokUsd": 0.7152
10430
10431
  },
10431
10432
  "source": "official-doc",
10432
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.5-plus: input $0.115, output $0.688 per million tokens; ≤128K input; >128K–256K $0.287/$1.72; >256K–1M $0.573/$3.44",
10433
- "confidence": "declared",
10434
- "observedAt": "2026-09-25T09:40:00Z"
10433
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.5-plus, ≤128K; 128K–256K: ¥2/¥12; 256K–1M: ¥4/¥24 input tier: ¥0.8 input/¥4.8 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10434
+ "confidence": "inferred",
10435
+ "observedAt": "2026-09-25T12:31:56.381Z"
10435
10436
  },
10436
10437
  "unsupportedParameters": [],
10437
10438
  "status": "candidate",
@@ -10535,13 +10536,13 @@
10535
10536
  },
10536
10537
  "pricing": {
10537
10538
  "value": {
10538
- "inputPerMTokUsd": 0.412564,
10539
- "outputPerMTokUsd": 2.475384
10539
+ "inputPerMTokUsd": 0.447,
10540
+ "outputPerMTokUsd": 2.682
10540
10541
  },
10541
10542
  "source": "official-doc",
10542
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-6-27b — China (Beijing) qwen3.6-27b: input $0.412564, output $2.475384 per million tokens; all input",
10543
- "confidence": "declared",
10544
- "observedAt": "2026-09-25T09:40:00Z"
10543
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.6-27b, ≤256K input tier: ¥3 input/¥18 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10544
+ "confidence": "inferred",
10545
+ "observedAt": "2026-09-25T12:31:56.381Z"
10545
10546
  },
10546
10547
  "unsupportedParameters": [],
10547
10548
  "status": "candidate",
@@ -10638,13 +10639,13 @@
10638
10639
  },
10639
10640
  "pricing": {
10640
10641
  "value": {
10641
- "inputPerMTokUsd": 0.165,
10642
- "outputPerMTokUsd": 0.99
10642
+ "inputPerMTokUsd": 0.1788,
10643
+ "outputPerMTokUsd": 1.0728
10643
10644
  },
10644
10645
  "source": "official-doc",
10645
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.6-flash: input $0.165, output $0.99 per million tokens; ≤256K input; >256K–1M $0.66/$3.961",
10646
- "confidence": "declared",
10647
- "observedAt": "2026-09-25T09:40:00Z"
10646
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.6-flash, ≤256K; 256K–1M: ¥4.8/¥28.8 input tier: ¥1.2 input/¥7.2 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10647
+ "confidence": "inferred",
10648
+ "observedAt": "2026-09-25T12:31:56.381Z"
10648
10649
  },
10649
10650
  "unsupportedParameters": [],
10650
10651
  "status": "candidate",
@@ -10734,13 +10735,13 @@
10734
10735
  },
10735
10736
  "pricing": {
10736
10737
  "value": {
10737
- "inputPerMTokUsd": 0.276,
10738
- "outputPerMTokUsd": 1.651
10738
+ "inputPerMTokUsd": 0.298,
10739
+ "outputPerMTokUsd": 1.788
10739
10740
  },
10740
10741
  "source": "official-doc",
10741
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.6-plus: input $0.276, output $1.651 per million tokens; ≤256K input; >256K–1M $1.101/$6.602",
10742
- "confidence": "declared",
10743
- "observedAt": "2026-09-25T09:40:00Z"
10742
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.6-plus, ≤256K; >256K–1M: ¥8/¥48 input tier: ¥2 input/¥12 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10743
+ "confidence": "inferred",
10744
+ "observedAt": "2026-09-25T12:31:56.381Z"
10744
10745
  },
10745
10746
  "unsupportedParameters": [],
10746
10747
  "status": "candidate",
@@ -10828,13 +10829,13 @@
10828
10829
  },
10829
10830
  "pricing": {
10830
10831
  "value": {
10831
- "inputPerMTokUsd": 1.65,
10832
- "outputPerMTokUsd": 4.951
10832
+ "inputPerMTokUsd": 1.788,
10833
+ "outputPerMTokUsd": 5.364
10833
10834
  },
10834
10835
  "source": "official-doc",
10835
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.7-max: input $1.65, output $4.951 per million tokens; 0–1M input tokens",
10836
- "confidence": "declared",
10837
- "observedAt": "2026-09-25T09:40:00Z"
10836
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.7-max, 0–1M input tier: ¥12 input/¥36 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10837
+ "confidence": "inferred",
10838
+ "observedAt": "2026-09-25T12:31:56.381Z"
10838
10839
  },
10839
10840
  "unsupportedParameters": [],
10840
10841
  "status": "candidate",
@@ -10931,13 +10932,13 @@
10931
10932
  },
10932
10933
  "pricing": {
10933
10934
  "value": {
10934
- "inputPerMTokUsd": 0.276,
10935
- "outputPerMTokUsd": 1.101
10935
+ "inputPerMTokUsd": 0.298,
10936
+ "outputPerMTokUsd": 1.192
10936
10937
  },
10937
10938
  "source": "official-doc",
10938
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.7-plus: input $0.276, output $1.101 per million tokens; ≤256K input; >256K–1M $0.826/$3.301; list price before limited-time 20% discount",
10939
- "confidence": "declared",
10940
- "observedAt": "2026-09-25T09:40:00Z"
10939
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.7-plus, ≤256K; >256K–1M: ¥6/¥24; base list price before temporary 20% promotion input tier: ¥2 input/¥8 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
10940
+ "confidence": "inferred",
10941
+ "observedAt": "2026-09-25T12:31:56.381Z"
10941
10942
  },
10942
10943
  "unsupportedParameters": [],
10943
10944
  "status": "candidate",
@@ -11027,13 +11028,13 @@
11027
11028
  },
11028
11029
  "pricing": {
11029
11030
  "value": {
11030
- "inputPerMTokUsd": 0.113,
11031
- "outputPerMTokUsd": 0.382
11031
+ "inputPerMTokUsd": 0.1192,
11032
+ "outputPerMTokUsd": 0.4023
11032
11033
  },
11033
11034
  "source": "official-doc",
11034
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.8-flash: input $0.113, output $0.382 per million tokens; 0–1M input tokens",
11035
- "confidence": "declared",
11036
- "observedAt": "2026-09-25T09:40:00Z"
11035
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.8-flash, 0–1M input tier: ¥0.8 input/¥2.7 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
11036
+ "confidence": "inferred",
11037
+ "observedAt": "2026-09-25T12:31:56.381Z"
11037
11038
  },
11038
11039
  "unsupportedParameters": [],
11039
11040
  "status": "candidate",
@@ -11123,13 +11124,13 @@
11123
11124
  },
11124
11125
  "pricing": {
11125
11126
  "value": {
11126
- "inputPerMTokUsd": 1.65,
11127
- "outputPerMTokUsd": 4.951
11127
+ "inputPerMTokUsd": 1.788,
11128
+ "outputPerMTokUsd": 5.364
11128
11129
  },
11129
11130
  "source": "official-doc",
11130
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.8-max: input $1.65, output $4.951 per million tokens; 0–1M input tokens",
11131
- "confidence": "declared",
11132
- "observedAt": "2026-09-25T09:40:00Z"
11131
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.8-max, 0–1M input tier: ¥12 input/¥36 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
11132
+ "confidence": "inferred",
11133
+ "observedAt": "2026-09-25T12:31:56.381Z"
11133
11134
  },
11134
11135
  "unsupportedParameters": [],
11135
11136
  "status": "candidate",
@@ -11455,13 +11456,13 @@
11455
11456
  },
11456
11457
  "pricing": {
11457
11458
  "value": {
11458
- "inputPerMTokUsd": 0.144,
11459
- "outputPerMTokUsd": 0.574
11459
+ "inputPerMTokUsd": 0.149,
11460
+ "outputPerMTokUsd": 0.596
11460
11461
  },
11461
11462
  "source": "official-doc",
11462
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-coder-next — China (Beijing) qwen3-coder-next: input $0.144, output $0.574 per million tokens; ≤32K input; >32K–128K $0.216/$0.861; >128K–256K $0.359/$1.434",
11463
- "confidence": "declared",
11464
- "observedAt": "2026-09-25T09:40:00Z"
11463
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3-coder-next, ≤32K; 32K–128K: ¥1.5/¥6; 128K–256K: ¥2.5/¥10 input tier: ¥1 input/¥4 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
11464
+ "confidence": "inferred",
11465
+ "observedAt": "2026-09-25T12:31:56.381Z"
11465
11466
  },
11466
11467
  "unsupportedParameters": [],
11467
11468
  "status": "candidate",
@@ -11557,13 +11558,13 @@
11557
11558
  },
11558
11559
  "pricing": {
11559
11560
  "value": {
11560
- "inputPerMTokUsd": 0.574,
11561
- "outputPerMTokUsd": 2.294
11561
+ "inputPerMTokUsd": 0.596,
11562
+ "outputPerMTokUsd": 2.384
11562
11563
  },
11563
11564
  "source": "official-doc",
11564
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/qwen3-coder-plus — China (Beijing) qwen3-coder-plus: input $0.574, output $2.294 per million tokens; ≤32K input; >32K–128K $0.861/$3.441; >128K–256K $1.434/$5.735; >256K–1M $2.868/$28.671",
11565
- "confidence": "declared",
11566
- "observedAt": "2026-09-25T09:40:00Z"
11565
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3-coder-plus, ≤32K; 32K–128K: ¥6/¥24; 128K–256K: ¥10/¥40; 256K–1M: ¥20/¥200 input tier: ¥4 input/¥16 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
11566
+ "confidence": "inferred",
11567
+ "observedAt": "2026-09-25T12:31:56.381Z"
11567
11568
  },
11568
11569
  "unsupportedParameters": [],
11569
11570
  "status": "candidate",
@@ -12160,13 +12161,13 @@
12160
12161
  },
12161
12162
  "pricing": {
12162
12163
  "value": {
12163
- "inputPerMTokUsd": 0.165,
12164
- "outputPerMTokUsd": 0.99
12164
+ "inputPerMTokUsd": 0.1788,
12165
+ "outputPerMTokUsd": 1.0728
12165
12166
  },
12166
12167
  "source": "official-doc",
12167
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.6-flash: input $0.165, output $0.99 per million tokens; ≤256K input; >256K–1M $0.66/$3.961",
12168
- "confidence": "declared",
12169
- "observedAt": "2026-09-25T09:40:00Z"
12168
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.6-flash, ≤256K; 256K–1M: ¥4.8/¥28.8 input tier: ¥1.2 input/¥7.2 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
12169
+ "confidence": "inferred",
12170
+ "observedAt": "2026-09-25T12:31:56.381Z"
12170
12171
  },
12171
12172
  "unsupportedParameters": [],
12172
12173
  "status": "candidate",
@@ -12531,13 +12532,13 @@
12531
12532
  },
12532
12533
  "pricing": {
12533
12534
  "value": {
12534
- "inputPerMTokUsd": 0.113,
12535
- "outputPerMTokUsd": 0.382
12535
+ "inputPerMTokUsd": 0.1192,
12536
+ "outputPerMTokUsd": 0.4023
12536
12537
  },
12537
12538
  "source": "official-doc",
12538
- "sourceRef": "https://www.alibabacloud.com/help/en/model-studio/model-pricing — China (Beijing) qwen3.8-flash: input $0.113, output $0.382 per million tokens; 0–1M input tokens",
12539
- "confidence": "declared",
12540
- "observedAt": "2026-09-25T09:40:00Z"
12539
+ "sourceRef": "https://help.aliyun.com/zh/model-studio/model-pricing — Beijing CNY list, qwen3.8-flash, 0–1M input tier: ¥0.8 input/¥2.7 output per million; USD = CNY × 0.1490 (CNY→USD reference rate for 2026-09-25, https://www.investing.com/currencies/cny-usd-historical-data) — the one rate every CNY-derived price in this catalog uses; not an Alibaba FX quote.",
12540
+ "confidence": "inferred",
12541
+ "observedAt": "2026-09-25T12:31:56.381Z"
12541
12542
  },
12542
12543
  "unsupportedParameters": [],
12543
12544
  "status": "candidate",
@@ -14375,7 +14376,16 @@
14375
14376
  "max"
14376
14377
  ],
14377
14378
  "continuation": "opaque-provider-state",
14378
- "defaultEffort": "high"
14379
+ "defaultEffort": "high",
14380
+ "effortRequest": {
14381
+ "value": {
14382
+ "field": "output_config.effort"
14383
+ },
14384
+ "source": "official-doc",
14385
+ "confidence": "declared",
14386
+ "observedAt": "2026-09-25T12:30:00Z",
14387
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
14388
+ }
14379
14389
  },
14380
14390
  "pricing": {
14381
14391
  "value": {
@@ -14392,9 +14402,50 @@
14392
14402
  "unsupportedParameters": [
14393
14403
  "temperature",
14394
14404
  "top_p",
14395
- "top_k"
14405
+ "top_k",
14406
+ "thinking.type.enabled",
14407
+ "thinking.type.disabled"
14396
14408
  ],
14397
14409
  "status": "candidate",
14410
+ "deferredToolLoading": {
14411
+ "value": true,
14412
+ "source": "official-doc",
14413
+ "confidence": "declared",
14414
+ "observedAt": "2026-09-25T18:30:00Z",
14415
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
14416
+ },
14417
+ "midConversationSystem": {
14418
+ "value": true,
14419
+ "source": "official-doc",
14420
+ "confidence": "declared",
14421
+ "observedAt": "2026-09-25T19:00:00Z",
14422
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
14423
+ },
14424
+ "midConversationToolChanges": {
14425
+ "value": {
14426
+ "beta": "mid-conversation-tool-changes-2026-07-01"
14427
+ },
14428
+ "source": "official-doc",
14429
+ "confidence": "declared",
14430
+ "observedAt": "2026-09-26T00:00:00Z",
14431
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
14432
+ },
14433
+ "inlineToolDefinitions": {
14434
+ "value": {
14435
+ "beta": "inline-tools-2026-09-15"
14436
+ },
14437
+ "source": "official-doc",
14438
+ "confidence": "declared",
14439
+ "observedAt": "2026-09-26T00:00:00Z",
14440
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
14441
+ },
14442
+ "assistantPrefill": {
14443
+ "value": false,
14444
+ "source": "official-doc",
14445
+ "confidence": "declared",
14446
+ "observedAt": "2026-09-26T00:00:00Z",
14447
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
14448
+ },
14398
14449
  "canonicalModelId": "claude-fable-5",
14399
14450
  "modelFamily": "claude"
14400
14451
  },
@@ -14527,6 +14578,33 @@
14527
14578
  "sourceRef": "continuity report §4.4, BOTH halves, applied per anthropic/claude-opus-5's own evidence. Own-state acceptance: a Claude model's own thinking blocks are replayed to it unchanged, in order, with signatures intact. Why the domain is NARROW: prior thinking/redacted_thinking blocks are tied to the model that produced them, so \"same provider\" is not automatically \"same continuation domain\" — the domain is this model alone.",
14528
14579
  "confidence": "declared",
14529
14580
  "observedAt": "2026-09-07T00:00:00Z"
14581
+ },
14582
+ "effortRequest": {
14583
+ "value": {
14584
+ "field": "output_config.effort"
14585
+ },
14586
+ "source": "official-doc",
14587
+ "confidence": "declared",
14588
+ "observedAt": "2026-09-25T12:30:00Z",
14589
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
14590
+ },
14591
+ "blockBinding": {
14592
+ "value": {
14593
+ "beta": "thinking-binding-controls-2026-08-01"
14594
+ },
14595
+ "source": "official-doc",
14596
+ "confidence": "declared",
14597
+ "observedAt": "2026-09-25T13:00:00Z",
14598
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting — \"A 400 error says a thinking block signature is invalid\": on Claude Fable 5.1 and Opus 5.5 a replayed thinking block is bound to the conversation (system, tools and earlier messages) and is rejected once that prefix changes (enforced for accounts created on or after 2026-08-31); the documented escape is the `thinking-binding-controls-2026-08-01` beta with `thinking.block_binding.prefix_mismatch_behavior: \"drop_block\"`. See also https://platform.claude.com/docs/en/models/opus-5-5/whats-new-opus-5-5 (\"Thinking blocks are tied to the model and the conversation\")."
14599
+ },
14600
+ "perMessageEffort": {
14601
+ "value": {
14602
+ "beta": "mid-conversation-output-config-2026-07-01"
14603
+ },
14604
+ "source": "official-doc",
14605
+ "confidence": "declared",
14606
+ "observedAt": "2026-09-25T18:00:00Z",
14607
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort#change-effort-mid-conversation-beta — \"On Claude Fable 5.1, Claude Mythos 5.1, Claude Opus 5.5, and Claude Opus 5, use a per-message effort change, which keeps the prompt cache\": a role:\"system\" message with empty content and output_config.effort, beta header mid-conversation-output-config-2026-07-01; \"Models without per-message effort, including Claude Fable 5, return a 400 error\". claude 2.1.282 sends the older alias per-turn-control-2026-07-01, which the page does not document."
14530
14608
  }
14531
14609
  },
14532
14610
  "pricing": {
@@ -14541,8 +14619,52 @@
14541
14619
  "confidence": "declared",
14542
14620
  "observedAt": "2026-09-08T00:00:00Z"
14543
14621
  },
14544
- "unsupportedParameters": [],
14622
+ "unsupportedParameters": [
14623
+ "thinking.type.enabled",
14624
+ "thinking.type.disabled",
14625
+ "tool_choice.any",
14626
+ "tool_choice.tool"
14627
+ ],
14545
14628
  "status": "candidate",
14629
+ "deferredToolLoading": {
14630
+ "value": true,
14631
+ "source": "official-doc",
14632
+ "confidence": "declared",
14633
+ "observedAt": "2026-09-25T18:30:00Z",
14634
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
14635
+ },
14636
+ "midConversationSystem": {
14637
+ "value": true,
14638
+ "source": "official-doc",
14639
+ "confidence": "declared",
14640
+ "observedAt": "2026-09-25T19:00:00Z",
14641
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
14642
+ },
14643
+ "midConversationToolChanges": {
14644
+ "value": {
14645
+ "beta": "mid-conversation-tool-changes-2026-07-01"
14646
+ },
14647
+ "source": "official-doc",
14648
+ "confidence": "declared",
14649
+ "observedAt": "2026-09-26T00:00:00Z",
14650
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
14651
+ },
14652
+ "inlineToolDefinitions": {
14653
+ "value": {
14654
+ "beta": "inline-tools-2026-09-15"
14655
+ },
14656
+ "source": "official-doc",
14657
+ "confidence": "declared",
14658
+ "observedAt": "2026-09-26T00:00:00Z",
14659
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
14660
+ },
14661
+ "assistantPrefill": {
14662
+ "value": false,
14663
+ "source": "official-doc",
14664
+ "confidence": "declared",
14665
+ "observedAt": "2026-09-26T00:00:00Z",
14666
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
14667
+ },
14546
14668
  "canonicalModelId": "claude-fable-5.1",
14547
14669
  "modelFamily": "claude"
14548
14670
  },
@@ -14637,8 +14759,17 @@
14637
14759
  "confidence": "declared",
14638
14760
  "observedAt": "2026-09-05T00:00:00Z"
14639
14761
  },
14640
- "unsupportedParameters": [],
14762
+ "unsupportedParameters": [
14763
+ "thinking.type.adaptive"
14764
+ ],
14641
14765
  "status": "candidate",
14766
+ "deferredToolLoading": {
14767
+ "value": true,
14768
+ "source": "official-doc",
14769
+ "confidence": "declared",
14770
+ "observedAt": "2026-09-25T18:30:00Z",
14771
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
14772
+ },
14642
14773
  "canonicalModelId": "claude-haiku-4.5-20251001",
14643
14774
  "modelFamily": "claude"
14644
14775
  },
@@ -14733,8 +14864,17 @@
14733
14864
  "confidence": "declared",
14734
14865
  "observedAt": "2026-09-19T00:00:00Z"
14735
14866
  },
14736
- "unsupportedParameters": [],
14867
+ "unsupportedParameters": [
14868
+ "thinking.type.adaptive"
14869
+ ],
14737
14870
  "status": "candidate",
14871
+ "deferredToolLoading": {
14872
+ "value": true,
14873
+ "source": "official-doc",
14874
+ "confidence": "declared",
14875
+ "observedAt": "2026-09-25T18:30:00Z",
14876
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
14877
+ },
14738
14878
  "canonicalModelId": "claude-haiku-4.5",
14739
14879
  "modelFamily": "claude"
14740
14880
  },
@@ -14832,7 +14972,16 @@
14832
14972
  "high"
14833
14973
  ],
14834
14974
  "continuation": "opaque-provider-state",
14835
- "defaultEffort": "high"
14975
+ "defaultEffort": "high",
14976
+ "effortRequest": {
14977
+ "value": {
14978
+ "field": "output_config.effort"
14979
+ },
14980
+ "source": "official-doc",
14981
+ "confidence": "declared",
14982
+ "observedAt": "2026-09-25T12:30:00Z",
14983
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
14984
+ }
14836
14985
  },
14837
14986
  "pricing": {
14838
14987
  "value": {
@@ -14846,8 +14995,17 @@
14846
14995
  "confidence": "declared",
14847
14996
  "observedAt": "2026-09-19T00:00:00Z"
14848
14997
  },
14849
- "unsupportedParameters": [],
14998
+ "unsupportedParameters": [
14999
+ "thinking.type.adaptive"
15000
+ ],
14850
15001
  "status": "candidate",
15002
+ "deferredToolLoading": {
15003
+ "value": true,
15004
+ "source": "official-doc",
15005
+ "confidence": "declared",
15006
+ "observedAt": "2026-09-25T18:30:00Z",
15007
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
15008
+ },
14851
15009
  "canonicalModelId": "claude-opus-4.5",
14852
15010
  "modelFamily": "claude"
14853
15011
  },
@@ -14856,7 +15014,9 @@
14856
15014
  "providerId": "anthropic",
14857
15015
  "upstreamId": "claude-opus-4.6",
14858
15016
  "displayName": "Claude Opus 4.6",
14859
- "aliases": [],
15017
+ "aliases": [
15018
+ "claude-opus-4-6"
15019
+ ],
14860
15020
  "endpoints": [
14861
15021
  "chat"
14862
15022
  ],
@@ -14943,7 +15103,16 @@
14943
15103
  "max"
14944
15104
  ],
14945
15105
  "continuation": "opaque-provider-state",
14946
- "defaultEffort": "high"
15106
+ "defaultEffort": "high",
15107
+ "effortRequest": {
15108
+ "value": {
15109
+ "field": "output_config.effort"
15110
+ },
15111
+ "source": "official-doc",
15112
+ "confidence": "declared",
15113
+ "observedAt": "2026-09-25T12:30:00Z",
15114
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
15115
+ }
14947
15116
  },
14948
15117
  "pricing": {
14949
15118
  "value": {
@@ -14959,6 +15128,20 @@
14959
15128
  },
14960
15129
  "unsupportedParameters": [],
14961
15130
  "status": "candidate",
15131
+ "deferredToolLoading": {
15132
+ "value": true,
15133
+ "source": "official-doc",
15134
+ "confidence": "declared",
15135
+ "observedAt": "2026-09-25T18:30:00Z",
15136
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
15137
+ },
15138
+ "assistantPrefill": {
15139
+ "value": false,
15140
+ "source": "official-doc",
15141
+ "confidence": "declared",
15142
+ "observedAt": "2026-09-26T00:00:00Z",
15143
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
15144
+ },
14962
15145
  "canonicalModelId": "claude-opus-4.6",
14963
15146
  "modelFamily": "claude"
14964
15147
  },
@@ -14967,7 +15150,9 @@
14967
15150
  "providerId": "anthropic",
14968
15151
  "upstreamId": "claude-opus-4.7",
14969
15152
  "displayName": "Claude Opus 4.7",
14970
- "aliases": [],
15153
+ "aliases": [
15154
+ "claude-opus-4-7"
15155
+ ],
14971
15156
  "endpoints": [
14972
15157
  "chat"
14973
15158
  ],
@@ -15055,7 +15240,16 @@
15055
15240
  "max"
15056
15241
  ],
15057
15242
  "continuation": "opaque-provider-state",
15058
- "defaultEffort": "high"
15243
+ "defaultEffort": "high",
15244
+ "effortRequest": {
15245
+ "value": {
15246
+ "field": "output_config.effort"
15247
+ },
15248
+ "source": "official-doc",
15249
+ "confidence": "declared",
15250
+ "observedAt": "2026-09-25T12:30:00Z",
15251
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
15252
+ }
15059
15253
  },
15060
15254
  "pricing": {
15061
15255
  "value": {
@@ -15072,9 +15266,24 @@
15072
15266
  "unsupportedParameters": [
15073
15267
  "temperature",
15074
15268
  "top_p",
15075
- "top_k"
15269
+ "top_k",
15270
+ "thinking.type.enabled"
15076
15271
  ],
15077
15272
  "status": "candidate",
15273
+ "deferredToolLoading": {
15274
+ "value": true,
15275
+ "source": "official-doc",
15276
+ "confidence": "declared",
15277
+ "observedAt": "2026-09-25T18:30:00Z",
15278
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
15279
+ },
15280
+ "assistantPrefill": {
15281
+ "value": false,
15282
+ "source": "official-doc",
15283
+ "confidence": "declared",
15284
+ "observedAt": "2026-09-26T00:00:00Z",
15285
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
15286
+ },
15078
15287
  "canonicalModelId": "claude-opus-4.7",
15079
15288
  "modelFamily": "claude"
15080
15289
  },
@@ -15083,7 +15292,9 @@
15083
15292
  "providerId": "anthropic",
15084
15293
  "upstreamId": "claude-opus-4.8",
15085
15294
  "displayName": "Claude Opus 4.8",
15086
- "aliases": [],
15295
+ "aliases": [
15296
+ "claude-opus-4-8"
15297
+ ],
15087
15298
  "endpoints": [
15088
15299
  "chat"
15089
15300
  ],
@@ -15171,7 +15382,16 @@
15171
15382
  "max"
15172
15383
  ],
15173
15384
  "continuation": "opaque-provider-state",
15174
- "defaultEffort": "high"
15385
+ "defaultEffort": "high",
15386
+ "effortRequest": {
15387
+ "value": {
15388
+ "field": "output_config.effort"
15389
+ },
15390
+ "source": "official-doc",
15391
+ "confidence": "declared",
15392
+ "observedAt": "2026-09-25T12:30:00Z",
15393
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
15394
+ }
15175
15395
  },
15176
15396
  "pricing": {
15177
15397
  "value": {
@@ -15188,9 +15408,49 @@
15188
15408
  "unsupportedParameters": [
15189
15409
  "temperature",
15190
15410
  "top_p",
15191
- "top_k"
15411
+ "top_k",
15412
+ "thinking.type.enabled"
15192
15413
  ],
15193
15414
  "status": "candidate",
15415
+ "deferredToolLoading": {
15416
+ "value": true,
15417
+ "source": "official-doc",
15418
+ "confidence": "declared",
15419
+ "observedAt": "2026-09-25T18:30:00Z",
15420
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
15421
+ },
15422
+ "midConversationSystem": {
15423
+ "value": true,
15424
+ "source": "official-doc",
15425
+ "confidence": "declared",
15426
+ "observedAt": "2026-09-25T19:00:00Z",
15427
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
15428
+ },
15429
+ "midConversationToolChanges": {
15430
+ "value": {
15431
+ "beta": "mid-conversation-tool-changes-2026-07-01"
15432
+ },
15433
+ "source": "official-doc",
15434
+ "confidence": "declared",
15435
+ "observedAt": "2026-09-26T00:00:00Z",
15436
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
15437
+ },
15438
+ "inlineToolDefinitions": {
15439
+ "value": {
15440
+ "beta": "inline-tools-2026-09-15"
15441
+ },
15442
+ "source": "official-doc",
15443
+ "confidence": "declared",
15444
+ "observedAt": "2026-09-26T00:00:00Z",
15445
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
15446
+ },
15447
+ "assistantPrefill": {
15448
+ "value": false,
15449
+ "source": "official-doc",
15450
+ "confidence": "declared",
15451
+ "observedAt": "2026-09-26T00:00:00Z",
15452
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
15453
+ },
15194
15454
  "canonicalModelId": "claude-opus-4.8",
15195
15455
  "modelFamily": "claude"
15196
15456
  },
@@ -15330,6 +15590,24 @@
15330
15590
  "sourceRef": "continuity report §4.4, BOTH halves. Own-state acceptance: the model's own `thinking` blocks are replayed to it unchanged, in order, with their signatures intact (§4.4's replay rules and its worked request). Why the domain is NARROW: \"when changing Claude models, prior `thinking` and `redacted_thinking` blocks should be stripped because they are tied to the model that produced them. Therefore 'same provider' is not automatically 'same continuation domain.'\" The domain is this model alone, never the Anthropic provider.",
15331
15591
  "confidence": "declared",
15332
15592
  "observedAt": "2026-09-05T00:00:00Z"
15593
+ },
15594
+ "effortRequest": {
15595
+ "value": {
15596
+ "field": "output_config.effort"
15597
+ },
15598
+ "source": "official-doc",
15599
+ "confidence": "declared",
15600
+ "observedAt": "2026-09-25T12:30:00Z",
15601
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
15602
+ },
15603
+ "perMessageEffort": {
15604
+ "value": {
15605
+ "beta": "mid-conversation-output-config-2026-07-01"
15606
+ },
15607
+ "source": "official-doc",
15608
+ "confidence": "declared",
15609
+ "observedAt": "2026-09-25T18:00:00Z",
15610
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort#change-effort-mid-conversation-beta — \"On Claude Fable 5.1, Claude Mythos 5.1, Claude Opus 5.5, and Claude Opus 5, use a per-message effort change, which keeps the prompt cache\": a role:\"system\" message with empty content and output_config.effort, beta header mid-conversation-output-config-2026-07-01; \"Models without per-message effort, including Claude Fable 5, return a 400 error\". claude 2.1.282 sends the older alias per-turn-control-2026-07-01, which the page does not document."
15333
15611
  }
15334
15612
  },
15335
15613
  "pricing": {
@@ -15347,9 +15625,51 @@
15347
15625
  "unsupportedParameters": [
15348
15626
  "temperature",
15349
15627
  "top_p",
15350
- "top_k"
15628
+ "top_k",
15629
+ "thinking.type.enabled",
15630
+ "thinking.type.disabled+output_config.effort.xhigh",
15631
+ "thinking.type.disabled+output_config.effort.max"
15351
15632
  ],
15352
15633
  "status": "candidate",
15634
+ "deferredToolLoading": {
15635
+ "value": true,
15636
+ "source": "official-doc",
15637
+ "confidence": "declared",
15638
+ "observedAt": "2026-09-25T18:30:00Z",
15639
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
15640
+ },
15641
+ "midConversationSystem": {
15642
+ "value": true,
15643
+ "source": "official-doc",
15644
+ "confidence": "declared",
15645
+ "observedAt": "2026-09-25T19:00:00Z",
15646
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
15647
+ },
15648
+ "midConversationToolChanges": {
15649
+ "value": {
15650
+ "beta": "mid-conversation-tool-changes-2026-07-01"
15651
+ },
15652
+ "source": "official-doc",
15653
+ "confidence": "declared",
15654
+ "observedAt": "2026-09-26T00:00:00Z",
15655
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
15656
+ },
15657
+ "inlineToolDefinitions": {
15658
+ "value": {
15659
+ "beta": "inline-tools-2026-09-15"
15660
+ },
15661
+ "source": "official-doc",
15662
+ "confidence": "declared",
15663
+ "observedAt": "2026-09-26T00:00:00Z",
15664
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
15665
+ },
15666
+ "assistantPrefill": {
15667
+ "value": false,
15668
+ "source": "official-doc",
15669
+ "confidence": "declared",
15670
+ "observedAt": "2026-09-26T00:00:00Z",
15671
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
15672
+ },
15353
15673
  "canonicalModelId": "claude-opus-5",
15354
15674
  "modelFamily": "claude"
15355
15675
  },
@@ -15491,6 +15811,33 @@
15491
15811
  "sourceRef": "continuity report §3 (wire probe of the pinned Claude Agent SDK, 0.3.250→0.3.258) — every uncompacted historical thinking/redacted block is replayed on every later request",
15492
15812
  "confidence": "declared",
15493
15813
  "observedAt": "2026-09-05T00:00:00Z"
15814
+ },
15815
+ "effortRequest": {
15816
+ "value": {
15817
+ "field": "output_config.effort"
15818
+ },
15819
+ "source": "official-doc",
15820
+ "confidence": "declared",
15821
+ "observedAt": "2026-09-25T12:30:00Z",
15822
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
15823
+ },
15824
+ "blockBinding": {
15825
+ "value": {
15826
+ "beta": "thinking-binding-controls-2026-08-01"
15827
+ },
15828
+ "source": "official-doc",
15829
+ "confidence": "declared",
15830
+ "observedAt": "2026-09-25T13:00:00Z",
15831
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting — \"A 400 error says a thinking block signature is invalid\": on Claude Fable 5.1 and Opus 5.5 a replayed thinking block is bound to the conversation (system, tools and earlier messages) and is rejected once that prefix changes (enforced for accounts created on or after 2026-08-31); the documented escape is the `thinking-binding-controls-2026-08-01` beta with `thinking.block_binding.prefix_mismatch_behavior: \"drop_block\"`. See also https://platform.claude.com/docs/en/models/opus-5-5/whats-new-opus-5-5 (\"Thinking blocks are tied to the model and the conversation\")."
15832
+ },
15833
+ "perMessageEffort": {
15834
+ "value": {
15835
+ "beta": "mid-conversation-output-config-2026-07-01"
15836
+ },
15837
+ "source": "official-doc",
15838
+ "confidence": "declared",
15839
+ "observedAt": "2026-09-25T18:00:00Z",
15840
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort#change-effort-mid-conversation-beta — \"On Claude Fable 5.1, Claude Mythos 5.1, Claude Opus 5.5, and Claude Opus 5, use a per-message effort change, which keeps the prompt cache\": a role:\"system\" message with empty content and output_config.effort, beta header mid-conversation-output-config-2026-07-01; \"Models without per-message effort, including Claude Fable 5, return a 400 error\". claude 2.1.282 sends the older alias per-turn-control-2026-07-01, which the page does not document."
15494
15841
  }
15495
15842
  },
15496
15843
  "pricing": {
@@ -15505,8 +15852,52 @@
15505
15852
  "confidence": "declared",
15506
15853
  "observedAt": "2026-09-25T10:25:20Z"
15507
15854
  },
15508
- "unsupportedParameters": [],
15855
+ "unsupportedParameters": [
15856
+ "thinking.type.enabled",
15857
+ "thinking.type.disabled",
15858
+ "tool_choice.any",
15859
+ "tool_choice.tool"
15860
+ ],
15509
15861
  "status": "candidate",
15862
+ "deferredToolLoading": {
15863
+ "value": true,
15864
+ "source": "official-doc",
15865
+ "confidence": "declared",
15866
+ "observedAt": "2026-09-25T18:30:00Z",
15867
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
15868
+ },
15869
+ "midConversationSystem": {
15870
+ "value": true,
15871
+ "source": "official-doc",
15872
+ "confidence": "declared",
15873
+ "observedAt": "2026-09-25T19:00:00Z",
15874
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
15875
+ },
15876
+ "midConversationToolChanges": {
15877
+ "value": {
15878
+ "beta": "mid-conversation-tool-changes-2026-07-01"
15879
+ },
15880
+ "source": "official-doc",
15881
+ "confidence": "declared",
15882
+ "observedAt": "2026-09-26T00:00:00Z",
15883
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
15884
+ },
15885
+ "inlineToolDefinitions": {
15886
+ "value": {
15887
+ "beta": "inline-tools-2026-09-15"
15888
+ },
15889
+ "source": "official-doc",
15890
+ "confidence": "declared",
15891
+ "observedAt": "2026-09-26T00:00:00Z",
15892
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
15893
+ },
15894
+ "assistantPrefill": {
15895
+ "value": false,
15896
+ "source": "official-doc",
15897
+ "confidence": "declared",
15898
+ "observedAt": "2026-09-26T00:00:00Z",
15899
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
15900
+ },
15510
15901
  "canonicalModelId": "claude-opus-5.5",
15511
15902
  "modelFamily": "claude"
15512
15903
  },
@@ -15618,8 +16009,17 @@
15618
16009
  "confidence": "declared",
15619
16010
  "observedAt": "2026-09-19T00:00:00Z"
15620
16011
  },
15621
- "unsupportedParameters": [],
16012
+ "unsupportedParameters": [
16013
+ "thinking.type.adaptive"
16014
+ ],
15622
16015
  "status": "candidate",
16016
+ "deferredToolLoading": {
16017
+ "value": true,
16018
+ "source": "official-doc",
16019
+ "confidence": "declared",
16020
+ "observedAt": "2026-09-25T18:30:00Z",
16021
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
16022
+ },
15623
16023
  "canonicalModelId": "claude-sonnet-4.5",
15624
16024
  "modelFamily": "claude"
15625
16025
  },
@@ -15628,7 +16028,9 @@
15628
16028
  "providerId": "anthropic",
15629
16029
  "upstreamId": "claude-sonnet-4.6",
15630
16030
  "displayName": "Claude Sonnet 4.6",
15631
- "aliases": [],
16031
+ "aliases": [
16032
+ "claude-sonnet-4-6"
16033
+ ],
15632
16034
  "endpoints": [
15633
16035
  "chat"
15634
16036
  ],
@@ -15715,7 +16117,16 @@
15715
16117
  "max"
15716
16118
  ],
15717
16119
  "continuation": "opaque-provider-state",
15718
- "defaultEffort": "high"
16120
+ "defaultEffort": "high",
16121
+ "effortRequest": {
16122
+ "value": {
16123
+ "field": "output_config.effort"
16124
+ },
16125
+ "source": "official-doc",
16126
+ "confidence": "declared",
16127
+ "observedAt": "2026-09-25T12:30:00Z",
16128
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
16129
+ }
15719
16130
  },
15720
16131
  "pricing": {
15721
16132
  "value": {
@@ -15731,6 +16142,13 @@
15731
16142
  },
15732
16143
  "unsupportedParameters": [],
15733
16144
  "status": "candidate",
16145
+ "deferredToolLoading": {
16146
+ "value": true,
16147
+ "source": "official-doc",
16148
+ "confidence": "declared",
16149
+ "observedAt": "2026-09-25T18:30:00Z",
16150
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
16151
+ },
15734
16152
  "canonicalModelId": "claude-sonnet-4.6",
15735
16153
  "modelFamily": "claude"
15736
16154
  },
@@ -15872,6 +16290,15 @@
15872
16290
  "sourceRef": "continuity report §4.4, BOTH halves. Own-state acceptance: §4.4 requires the model's own assistant blocks — `thinking` with its signature, `redacted_thinking` — replayed unchanged and in order across a tool loop. Why the domain is NARROW: Anthropic documents model switching as a boundary at which those blocks are STRIPPED, because they are tied to the producing model; \"same provider\" is therefore not \"same continuation domain\", and this domain is the model alone.",
15873
16291
  "confidence": "declared",
15874
16292
  "observedAt": "2026-09-05T00:00:00Z"
16293
+ },
16294
+ "effortRequest": {
16295
+ "value": {
16296
+ "field": "output_config.effort"
16297
+ },
16298
+ "source": "official-doc",
16299
+ "confidence": "declared",
16300
+ "observedAt": "2026-09-25T12:30:00Z",
16301
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
15875
16302
  }
15876
16303
  },
15877
16304
  "pricing": {
@@ -15889,9 +16316,17 @@
15889
16316
  "unsupportedParameters": [
15890
16317
  "temperature",
15891
16318
  "top_p",
15892
- "top_k"
16319
+ "top_k",
16320
+ "thinking.type.enabled"
15893
16321
  ],
15894
16322
  "status": "candidate",
16323
+ "assistantPrefill": {
16324
+ "value": false,
16325
+ "source": "official-doc",
16326
+ "confidence": "declared",
16327
+ "observedAt": "2026-09-26T00:00:00Z",
16328
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
16329
+ },
15895
16330
  "canonicalModelId": "claude-sonnet-5",
15896
16331
  "modelFamily": "claude"
15897
16332
  },
@@ -21736,13 +22171,12 @@
21736
22171
  "value": [
21737
22172
  "text",
21738
22173
  "image",
21739
- "video",
21740
- "pdf"
22174
+ "video"
21741
22175
  ],
21742
22176
  "source": "official-doc",
21743
- "sourceRef": "https://docs.aws.amazon.com/nova/latest/userguide/what-is-nova.html — Amazon Nova Lite supports text, image, video, pdf inputs; PDF is mapped from document support",
21744
- "confidence": "inferred",
21745
- "observedAt": "2026-09-25T11:00:00Z"
22177
+ "sourceRef": "https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-amazon-nova-lite.html — per-model Input Modalities table marks Text, Image, Video supported; PDF is not listed as a separate modality",
22178
+ "confidence": "declared",
22179
+ "observedAt": "2026-09-25T12:31:56.381Z"
21746
22180
  },
21747
22181
  "outputModalities": {
21748
22182
  "value": [
@@ -21887,13 +22321,12 @@
21887
22321
  "value": [
21888
22322
  "text",
21889
22323
  "image",
21890
- "video",
21891
- "pdf"
22324
+ "video"
21892
22325
  ],
21893
22326
  "source": "official-doc",
21894
- "sourceRef": "https://docs.aws.amazon.com/nova/latest/userguide/what-is-nova.html — Amazon Nova Premier supports text, image, video, pdf inputs; PDF is mapped from document support",
21895
- "confidence": "inferred",
21896
- "observedAt": "2026-09-25T11:00:00Z"
22327
+ "sourceRef": "https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-amazon-nova-premier.html — per-model Input Modalities table marks Text, Image, Video supported; PDF is not listed as a separate modality",
22328
+ "confidence": "declared",
22329
+ "observedAt": "2026-09-25T12:31:56.381Z"
21897
22330
  },
21898
22331
  "outputModalities": {
21899
22332
  "value": [
@@ -21975,13 +22408,12 @@
21975
22408
  "value": [
21976
22409
  "text",
21977
22410
  "image",
21978
- "video",
21979
- "pdf"
22411
+ "video"
21980
22412
  ],
21981
22413
  "source": "official-doc",
21982
- "sourceRef": "https://docs.aws.amazon.com/nova/latest/userguide/what-is-nova.html — Amazon Nova Pro supports text, image, video, pdf inputs; PDF is mapped from document support",
21983
- "confidence": "inferred",
21984
- "observedAt": "2026-09-25T11:00:00Z"
22414
+ "sourceRef": "https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-amazon-nova-pro.html — per-model Input Modalities table marks Text, Image, Video supported; PDF is not listed as a separate modality",
22415
+ "confidence": "declared",
22416
+ "observedAt": "2026-09-25T12:31:56.381Z"
21985
22417
  },
21986
22418
  "outputModalities": {
21987
22419
  "value": [
@@ -22419,13 +22851,12 @@
22419
22851
  "value": [
22420
22852
  "text",
22421
22853
  "image",
22422
- "video",
22423
- "pdf"
22854
+ "video"
22424
22855
  ],
22425
22856
  "source": "official-doc",
22426
- "sourceRef": "https://docs.aws.amazon.com/nova/latest/nova2-userguide/what-is-nova-2.html — Amazon Nova 2 Lite (US Geo) supports text, image, video, pdf inputs; PDF is mapped from document support",
22427
- "confidence": "inferred",
22428
- "observedAt": "2026-09-25T11:00:00Z"
22857
+ "sourceRef": "https://docs.aws.amazon.com/bedrock/latest/userguide/model-card-amazon-nova-2-lite.html — per-model Input Modalities table marks Text, Image, Video supported; PDF is not listed as a separate modality",
22858
+ "confidence": "declared",
22859
+ "observedAt": "2026-09-25T12:31:56.381Z"
22429
22860
  },
22430
22861
  "outputModalities": {
22431
22862
  "value": [
@@ -25099,6 +25530,20 @@
25099
25530
  },
25100
25531
  "unsupportedParameters": [],
25101
25532
  "status": "candidate",
25533
+ "promptCacheKey": {
25534
+ "value": true,
25535
+ "source": "official-doc",
25536
+ "confidence": "declared",
25537
+ "observedAt": "2026-09-25T19:30:00Z",
25538
+ "sourceRef": "https://github.com/openai/codex/blob/main/codex-rs/core/src/client.rs — Codex (the vendor's own client for this backend) sets prompt_cache_key on every Responses request, to the session id (or source:parent_thread_id for internal sessions); https://developers.openai.com/api/docs/guides/prompt-caching documents the field."
25539
+ },
25540
+ "clientToolSearch": {
25541
+ "value": true,
25542
+ "source": "upstream-static",
25543
+ "confidence": "declared",
25544
+ "observedAt": "2026-09-26T00:00:00Z",
25545
+ "sourceRef": "codex-rs models-manager/models.json (cloned 25270df, 2026-09-26): supports_search_tool is true on every bundled model; core/src/tools/spec_plan.rs:653-666 turns on the client tool_search (tools/src/tool_search.rs:84-110, core/src/tools/handlers/tool_search_spec.rs:98) when the provider's namespace capability is on (default true) -- the vendor's own client sends it to chatgpt.com/backend-api/codex."
25546
+ },
25102
25547
  "canonicalModelId": "gpt-5.6-luna",
25103
25548
  "modelFamily": "gpt"
25104
25549
  },
@@ -25206,6 +25651,20 @@
25206
25651
  },
25207
25652
  "unsupportedParameters": [],
25208
25653
  "status": "candidate",
25654
+ "promptCacheKey": {
25655
+ "value": true,
25656
+ "source": "official-doc",
25657
+ "confidence": "declared",
25658
+ "observedAt": "2026-09-25T19:30:00Z",
25659
+ "sourceRef": "https://github.com/openai/codex/blob/main/codex-rs/core/src/client.rs — Codex (the vendor's own client for this backend) sets prompt_cache_key on every Responses request, to the session id (or source:parent_thread_id for internal sessions); https://developers.openai.com/api/docs/guides/prompt-caching documents the field."
25660
+ },
25661
+ "clientToolSearch": {
25662
+ "value": true,
25663
+ "source": "upstream-static",
25664
+ "confidence": "declared",
25665
+ "observedAt": "2026-09-26T00:00:00Z",
25666
+ "sourceRef": "codex-rs models-manager/models.json (cloned 25270df, 2026-09-26): supports_search_tool is true on every bundled model; core/src/tools/spec_plan.rs:653-666 turns on the client tool_search (tools/src/tool_search.rs:84-110, core/src/tools/handlers/tool_search_spec.rs:98) when the provider's namespace capability is on (default true) -- the vendor's own client sends it to chatgpt.com/backend-api/codex."
25667
+ },
25209
25668
  "canonicalModelId": "gpt-5.6-sol",
25210
25669
  "modelFamily": "gpt"
25211
25670
  },
@@ -25313,6 +25772,20 @@
25313
25772
  },
25314
25773
  "unsupportedParameters": [],
25315
25774
  "status": "candidate",
25775
+ "promptCacheKey": {
25776
+ "value": true,
25777
+ "source": "official-doc",
25778
+ "confidence": "declared",
25779
+ "observedAt": "2026-09-25T19:30:00Z",
25780
+ "sourceRef": "https://github.com/openai/codex/blob/main/codex-rs/core/src/client.rs — Codex (the vendor's own client for this backend) sets prompt_cache_key on every Responses request, to the session id (or source:parent_thread_id for internal sessions); https://developers.openai.com/api/docs/guides/prompt-caching documents the field."
25781
+ },
25782
+ "clientToolSearch": {
25783
+ "value": true,
25784
+ "source": "upstream-static",
25785
+ "confidence": "declared",
25786
+ "observedAt": "2026-09-26T00:00:00Z",
25787
+ "sourceRef": "codex-rs models-manager/models.json (cloned 25270df, 2026-09-26): supports_search_tool is true on every bundled model; core/src/tools/spec_plan.rs:653-666 turns on the client tool_search (tools/src/tool_search.rs:84-110, core/src/tools/handlers/tool_search_spec.rs:98) when the provider's namespace capability is on (default true) -- the vendor's own client sends it to chatgpt.com/backend-api/codex."
25788
+ },
25316
25789
  "canonicalModelId": "gpt-5.6-terra",
25317
25790
  "modelFamily": "gpt"
25318
25791
  },
@@ -25422,6 +25895,20 @@
25422
25895
  },
25423
25896
  "unsupportedParameters": [],
25424
25897
  "status": "candidate",
25898
+ "promptCacheKey": {
25899
+ "value": true,
25900
+ "source": "official-doc",
25901
+ "confidence": "declared",
25902
+ "observedAt": "2026-09-25T19:30:00Z",
25903
+ "sourceRef": "https://github.com/openai/codex/blob/main/codex-rs/core/src/client.rs — Codex (the vendor's own client for this backend) sets prompt_cache_key on every Responses request, to the session id (or source:parent_thread_id for internal sessions); https://developers.openai.com/api/docs/guides/prompt-caching documents the field."
25904
+ },
25905
+ "clientToolSearch": {
25906
+ "value": true,
25907
+ "source": "upstream-static",
25908
+ "confidence": "declared",
25909
+ "observedAt": "2026-09-26T00:00:00Z",
25910
+ "sourceRef": "codex-rs models-manager/models.json (cloned 25270df, 2026-09-26): supports_search_tool is true on every bundled model; core/src/tools/spec_plan.rs:653-666 turns on the client tool_search (tools/src/tool_search.rs:84-110, core/src/tools/handlers/tool_search_spec.rs:98) when the provider's namespace capability is on (default true) -- the vendor's own client sends it to chatgpt.com/backend-api/codex."
25911
+ },
25425
25912
  "canonicalModelId": "gpt-6-astra",
25426
25913
  "modelFamily": "gpt"
25427
25914
  },
@@ -25531,6 +26018,20 @@
25531
26018
  },
25532
26019
  "unsupportedParameters": [],
25533
26020
  "status": "candidate",
26021
+ "promptCacheKey": {
26022
+ "value": true,
26023
+ "source": "official-doc",
26024
+ "confidence": "declared",
26025
+ "observedAt": "2026-09-25T19:30:00Z",
26026
+ "sourceRef": "https://github.com/openai/codex/blob/main/codex-rs/core/src/client.rs — Codex (the vendor's own client for this backend) sets prompt_cache_key on every Responses request, to the session id (or source:parent_thread_id for internal sessions); https://developers.openai.com/api/docs/guides/prompt-caching documents the field."
26027
+ },
26028
+ "clientToolSearch": {
26029
+ "value": true,
26030
+ "source": "upstream-static",
26031
+ "confidence": "declared",
26032
+ "observedAt": "2026-09-26T00:00:00Z",
26033
+ "sourceRef": "codex-rs models-manager/models.json (cloned 25270df, 2026-09-26): supports_search_tool is true on every bundled model; core/src/tools/spec_plan.rs:653-666 turns on the client tool_search (tools/src/tool_search.rs:84-110, core/src/tools/handlers/tool_search_spec.rs:98) when the provider's namespace capability is on (default true) -- the vendor's own client sends it to chatgpt.com/backend-api/codex."
26034
+ },
25534
26035
  "canonicalModelId": "gpt-6-luna",
25535
26036
  "modelFamily": "gpt"
25536
26037
  },
@@ -25640,6 +26141,20 @@
25640
26141
  },
25641
26142
  "unsupportedParameters": [],
25642
26143
  "status": "candidate",
26144
+ "promptCacheKey": {
26145
+ "value": true,
26146
+ "source": "official-doc",
26147
+ "confidence": "declared",
26148
+ "observedAt": "2026-09-25T19:30:00Z",
26149
+ "sourceRef": "https://github.com/openai/codex/blob/main/codex-rs/core/src/client.rs — Codex (the vendor's own client for this backend) sets prompt_cache_key on every Responses request, to the session id (or source:parent_thread_id for internal sessions); https://developers.openai.com/api/docs/guides/prompt-caching documents the field."
26150
+ },
26151
+ "clientToolSearch": {
26152
+ "value": true,
26153
+ "source": "upstream-static",
26154
+ "confidence": "declared",
26155
+ "observedAt": "2026-09-26T00:00:00Z",
26156
+ "sourceRef": "codex-rs models-manager/models.json (cloned 25270df, 2026-09-26): supports_search_tool is true on every bundled model; core/src/tools/spec_plan.rs:653-666 turns on the client tool_search (tools/src/tool_search.rs:84-110, core/src/tools/handlers/tool_search_spec.rs:98) when the provider's namespace capability is on (default true) -- the vendor's own client sends it to chatgpt.com/backend-api/codex."
26157
+ },
25643
26158
  "canonicalModelId": "gpt-6-sol",
25644
26159
  "modelFamily": "gpt"
25645
26160
  },
@@ -26352,7 +26867,16 @@
26352
26867
  "max"
26353
26868
  ],
26354
26869
  "continuation": "opaque-provider-state",
26355
- "defaultEffort": "high"
26870
+ "defaultEffort": "high",
26871
+ "effortRequest": {
26872
+ "value": {
26873
+ "field": "output_config.effort"
26874
+ },
26875
+ "source": "official-doc",
26876
+ "confidence": "declared",
26877
+ "observedAt": "2026-09-25T12:30:00Z",
26878
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
26879
+ }
26356
26880
  },
26357
26881
  "pricing": {
26358
26882
  "value": {
@@ -26369,9 +26893,50 @@
26369
26893
  "unsupportedParameters": [
26370
26894
  "temperature",
26371
26895
  "top_p",
26372
- "top_k"
26896
+ "top_k",
26897
+ "thinking.type.enabled",
26898
+ "thinking.type.disabled"
26373
26899
  ],
26374
26900
  "status": "candidate",
26901
+ "deferredToolLoading": {
26902
+ "value": true,
26903
+ "source": "official-doc",
26904
+ "confidence": "declared",
26905
+ "observedAt": "2026-09-25T18:30:00Z",
26906
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
26907
+ },
26908
+ "midConversationSystem": {
26909
+ "value": true,
26910
+ "source": "official-doc",
26911
+ "confidence": "declared",
26912
+ "observedAt": "2026-09-25T19:00:00Z",
26913
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
26914
+ },
26915
+ "midConversationToolChanges": {
26916
+ "value": {
26917
+ "beta": "mid-conversation-tool-changes-2026-07-01"
26918
+ },
26919
+ "source": "official-doc",
26920
+ "confidence": "declared",
26921
+ "observedAt": "2026-09-26T00:00:00Z",
26922
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
26923
+ },
26924
+ "inlineToolDefinitions": {
26925
+ "value": {
26926
+ "beta": "inline-tools-2026-09-15"
26927
+ },
26928
+ "source": "official-doc",
26929
+ "confidence": "declared",
26930
+ "observedAt": "2026-09-26T00:00:00Z",
26931
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
26932
+ },
26933
+ "assistantPrefill": {
26934
+ "value": false,
26935
+ "source": "official-doc",
26936
+ "confidence": "declared",
26937
+ "observedAt": "2026-09-26T00:00:00Z",
26938
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
26939
+ },
26375
26940
  "canonicalModelId": "claude-fable-5",
26376
26941
  "modelFamily": "claude"
26377
26942
  },
@@ -26504,6 +27069,33 @@
26504
27069
  "sourceRef": "continuity report §4.4, BOTH halves, applied per anthropic/claude-opus-5's own evidence. Own-state acceptance: a Claude model's own thinking blocks are replayed to it unchanged, in order, with signatures intact. Why the domain is NARROW: prior thinking/redacted_thinking blocks are tied to the model that produced them, so \"same provider\" is not automatically \"same continuation domain\" — the domain is this model alone.",
26505
27070
  "confidence": "declared",
26506
27071
  "observedAt": "2026-09-16T00:00:00Z"
27072
+ },
27073
+ "effortRequest": {
27074
+ "value": {
27075
+ "field": "output_config.effort"
27076
+ },
27077
+ "source": "official-doc",
27078
+ "confidence": "declared",
27079
+ "observedAt": "2026-09-25T12:30:00Z",
27080
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
27081
+ },
27082
+ "blockBinding": {
27083
+ "value": {
27084
+ "beta": "thinking-binding-controls-2026-08-01"
27085
+ },
27086
+ "source": "official-doc",
27087
+ "confidence": "declared",
27088
+ "observedAt": "2026-09-25T13:00:00Z",
27089
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting — \"A 400 error says a thinking block signature is invalid\": on Claude Fable 5.1 and Opus 5.5 a replayed thinking block is bound to the conversation (system, tools and earlier messages) and is rejected once that prefix changes (enforced for accounts created on or after 2026-08-31); the documented escape is the `thinking-binding-controls-2026-08-01` beta with `thinking.block_binding.prefix_mismatch_behavior: \"drop_block\"`. See also https://platform.claude.com/docs/en/models/opus-5-5/whats-new-opus-5-5 (\"Thinking blocks are tied to the model and the conversation\")."
27090
+ },
27091
+ "perMessageEffort": {
27092
+ "value": {
27093
+ "beta": "mid-conversation-output-config-2026-07-01"
27094
+ },
27095
+ "source": "official-doc",
27096
+ "confidence": "declared",
27097
+ "observedAt": "2026-09-25T18:00:00Z",
27098
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort#change-effort-mid-conversation-beta — \"On Claude Fable 5.1, Claude Mythos 5.1, Claude Opus 5.5, and Claude Opus 5, use a per-message effort change, which keeps the prompt cache\": a role:\"system\" message with empty content and output_config.effort, beta header mid-conversation-output-config-2026-07-01; \"Models without per-message effort, including Claude Fable 5, return a 400 error\". claude 2.1.282 sends the older alias per-turn-control-2026-07-01, which the page does not document."
26507
27099
  }
26508
27100
  },
26509
27101
  "pricing": {
@@ -26518,8 +27110,52 @@
26518
27110
  "confidence": "declared",
26519
27111
  "observedAt": "2026-09-16T00:00:00Z"
26520
27112
  },
26521
- "unsupportedParameters": [],
27113
+ "unsupportedParameters": [
27114
+ "thinking.type.enabled",
27115
+ "thinking.type.disabled",
27116
+ "tool_choice.any",
27117
+ "tool_choice.tool"
27118
+ ],
26522
27119
  "status": "candidate",
27120
+ "deferredToolLoading": {
27121
+ "value": true,
27122
+ "source": "official-doc",
27123
+ "confidence": "declared",
27124
+ "observedAt": "2026-09-25T18:30:00Z",
27125
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27126
+ },
27127
+ "midConversationSystem": {
27128
+ "value": true,
27129
+ "source": "official-doc",
27130
+ "confidence": "declared",
27131
+ "observedAt": "2026-09-25T19:00:00Z",
27132
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
27133
+ },
27134
+ "midConversationToolChanges": {
27135
+ "value": {
27136
+ "beta": "mid-conversation-tool-changes-2026-07-01"
27137
+ },
27138
+ "source": "official-doc",
27139
+ "confidence": "declared",
27140
+ "observedAt": "2026-09-26T00:00:00Z",
27141
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
27142
+ },
27143
+ "inlineToolDefinitions": {
27144
+ "value": {
27145
+ "beta": "inline-tools-2026-09-15"
27146
+ },
27147
+ "source": "official-doc",
27148
+ "confidence": "declared",
27149
+ "observedAt": "2026-09-26T00:00:00Z",
27150
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
27151
+ },
27152
+ "assistantPrefill": {
27153
+ "value": false,
27154
+ "source": "official-doc",
27155
+ "confidence": "declared",
27156
+ "observedAt": "2026-09-26T00:00:00Z",
27157
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
27158
+ },
26523
27159
  "canonicalModelId": "claude-fable-5.1",
26524
27160
  "modelFamily": "claude"
26525
27161
  },
@@ -26614,8 +27250,17 @@
26614
27250
  "confidence": "declared",
26615
27251
  "observedAt": "2026-09-16T00:00:00Z"
26616
27252
  },
26617
- "unsupportedParameters": [],
27253
+ "unsupportedParameters": [
27254
+ "thinking.type.adaptive"
27255
+ ],
26618
27256
  "status": "candidate",
27257
+ "deferredToolLoading": {
27258
+ "value": true,
27259
+ "source": "official-doc",
27260
+ "confidence": "declared",
27261
+ "observedAt": "2026-09-25T18:30:00Z",
27262
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27263
+ },
26619
27264
  "canonicalModelId": "claude-haiku-4.5-20251001",
26620
27265
  "modelFamily": "claude"
26621
27266
  },
@@ -26710,8 +27355,17 @@
26710
27355
  "confidence": "declared",
26711
27356
  "observedAt": "2026-09-19T00:00:00Z"
26712
27357
  },
26713
- "unsupportedParameters": [],
27358
+ "unsupportedParameters": [
27359
+ "thinking.type.adaptive"
27360
+ ],
26714
27361
  "status": "candidate",
27362
+ "deferredToolLoading": {
27363
+ "value": true,
27364
+ "source": "official-doc",
27365
+ "confidence": "declared",
27366
+ "observedAt": "2026-09-25T18:30:00Z",
27367
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27368
+ },
26715
27369
  "canonicalModelId": "claude-haiku-4.5",
26716
27370
  "modelFamily": "claude"
26717
27371
  },
@@ -26809,7 +27463,16 @@
26809
27463
  "high"
26810
27464
  ],
26811
27465
  "continuation": "opaque-provider-state",
26812
- "defaultEffort": "high"
27466
+ "defaultEffort": "high",
27467
+ "effortRequest": {
27468
+ "value": {
27469
+ "field": "output_config.effort"
27470
+ },
27471
+ "source": "official-doc",
27472
+ "confidence": "declared",
27473
+ "observedAt": "2026-09-25T12:30:00Z",
27474
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
27475
+ }
26813
27476
  },
26814
27477
  "pricing": {
26815
27478
  "value": {
@@ -26823,8 +27486,17 @@
26823
27486
  "confidence": "declared",
26824
27487
  "observedAt": "2026-09-19T00:00:00Z"
26825
27488
  },
26826
- "unsupportedParameters": [],
27489
+ "unsupportedParameters": [
27490
+ "thinking.type.adaptive"
27491
+ ],
26827
27492
  "status": "candidate",
27493
+ "deferredToolLoading": {
27494
+ "value": true,
27495
+ "source": "official-doc",
27496
+ "confidence": "declared",
27497
+ "observedAt": "2026-09-25T18:30:00Z",
27498
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27499
+ },
26828
27500
  "canonicalModelId": "claude-opus-4.5",
26829
27501
  "modelFamily": "claude"
26830
27502
  },
@@ -26833,7 +27505,9 @@
26833
27505
  "providerId": "console",
26834
27506
  "upstreamId": "claude-opus-4.6",
26835
27507
  "displayName": "Claude Opus 4.6",
26836
- "aliases": [],
27508
+ "aliases": [
27509
+ "claude-opus-4-6"
27510
+ ],
26837
27511
  "endpoints": [
26838
27512
  "chat"
26839
27513
  ],
@@ -26920,7 +27594,16 @@
26920
27594
  "max"
26921
27595
  ],
26922
27596
  "continuation": "opaque-provider-state",
26923
- "defaultEffort": "high"
27597
+ "defaultEffort": "high",
27598
+ "effortRequest": {
27599
+ "value": {
27600
+ "field": "output_config.effort"
27601
+ },
27602
+ "source": "official-doc",
27603
+ "confidence": "declared",
27604
+ "observedAt": "2026-09-25T12:30:00Z",
27605
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
27606
+ }
26924
27607
  },
26925
27608
  "pricing": {
26926
27609
  "value": {
@@ -26936,6 +27619,20 @@
26936
27619
  },
26937
27620
  "unsupportedParameters": [],
26938
27621
  "status": "candidate",
27622
+ "deferredToolLoading": {
27623
+ "value": true,
27624
+ "source": "official-doc",
27625
+ "confidence": "declared",
27626
+ "observedAt": "2026-09-25T18:30:00Z",
27627
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27628
+ },
27629
+ "assistantPrefill": {
27630
+ "value": false,
27631
+ "source": "official-doc",
27632
+ "confidence": "declared",
27633
+ "observedAt": "2026-09-26T00:00:00Z",
27634
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
27635
+ },
26939
27636
  "canonicalModelId": "claude-opus-4.6",
26940
27637
  "modelFamily": "claude"
26941
27638
  },
@@ -26944,7 +27641,9 @@
26944
27641
  "providerId": "console",
26945
27642
  "upstreamId": "claude-opus-4.7",
26946
27643
  "displayName": "Claude Opus 4.7",
26947
- "aliases": [],
27644
+ "aliases": [
27645
+ "claude-opus-4-7"
27646
+ ],
26948
27647
  "endpoints": [
26949
27648
  "chat"
26950
27649
  ],
@@ -27032,7 +27731,16 @@
27032
27731
  "max"
27033
27732
  ],
27034
27733
  "continuation": "opaque-provider-state",
27035
- "defaultEffort": "high"
27734
+ "defaultEffort": "high",
27735
+ "effortRequest": {
27736
+ "value": {
27737
+ "field": "output_config.effort"
27738
+ },
27739
+ "source": "official-doc",
27740
+ "confidence": "declared",
27741
+ "observedAt": "2026-09-25T12:30:00Z",
27742
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
27743
+ }
27036
27744
  },
27037
27745
  "pricing": {
27038
27746
  "value": {
@@ -27049,9 +27757,24 @@
27049
27757
  "unsupportedParameters": [
27050
27758
  "temperature",
27051
27759
  "top_p",
27052
- "top_k"
27760
+ "top_k",
27761
+ "thinking.type.enabled"
27053
27762
  ],
27054
27763
  "status": "candidate",
27764
+ "deferredToolLoading": {
27765
+ "value": true,
27766
+ "source": "official-doc",
27767
+ "confidence": "declared",
27768
+ "observedAt": "2026-09-25T18:30:00Z",
27769
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27770
+ },
27771
+ "assistantPrefill": {
27772
+ "value": false,
27773
+ "source": "official-doc",
27774
+ "confidence": "declared",
27775
+ "observedAt": "2026-09-26T00:00:00Z",
27776
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
27777
+ },
27055
27778
  "canonicalModelId": "claude-opus-4.7",
27056
27779
  "modelFamily": "claude"
27057
27780
  },
@@ -27060,7 +27783,9 @@
27060
27783
  "providerId": "console",
27061
27784
  "upstreamId": "claude-opus-4.8",
27062
27785
  "displayName": "Claude Opus 4.8",
27063
- "aliases": [],
27786
+ "aliases": [
27787
+ "claude-opus-4-8"
27788
+ ],
27064
27789
  "endpoints": [
27065
27790
  "chat"
27066
27791
  ],
@@ -27148,7 +27873,16 @@
27148
27873
  "max"
27149
27874
  ],
27150
27875
  "continuation": "opaque-provider-state",
27151
- "defaultEffort": "high"
27876
+ "defaultEffort": "high",
27877
+ "effortRequest": {
27878
+ "value": {
27879
+ "field": "output_config.effort"
27880
+ },
27881
+ "source": "official-doc",
27882
+ "confidence": "declared",
27883
+ "observedAt": "2026-09-25T12:30:00Z",
27884
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
27885
+ }
27152
27886
  },
27153
27887
  "pricing": {
27154
27888
  "value": {
@@ -27165,9 +27899,49 @@
27165
27899
  "unsupportedParameters": [
27166
27900
  "temperature",
27167
27901
  "top_p",
27168
- "top_k"
27902
+ "top_k",
27903
+ "thinking.type.enabled"
27169
27904
  ],
27170
27905
  "status": "candidate",
27906
+ "deferredToolLoading": {
27907
+ "value": true,
27908
+ "source": "official-doc",
27909
+ "confidence": "declared",
27910
+ "observedAt": "2026-09-25T18:30:00Z",
27911
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
27912
+ },
27913
+ "midConversationSystem": {
27914
+ "value": true,
27915
+ "source": "official-doc",
27916
+ "confidence": "declared",
27917
+ "observedAt": "2026-09-25T19:00:00Z",
27918
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
27919
+ },
27920
+ "midConversationToolChanges": {
27921
+ "value": {
27922
+ "beta": "mid-conversation-tool-changes-2026-07-01"
27923
+ },
27924
+ "source": "official-doc",
27925
+ "confidence": "declared",
27926
+ "observedAt": "2026-09-26T00:00:00Z",
27927
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
27928
+ },
27929
+ "inlineToolDefinitions": {
27930
+ "value": {
27931
+ "beta": "inline-tools-2026-09-15"
27932
+ },
27933
+ "source": "official-doc",
27934
+ "confidence": "declared",
27935
+ "observedAt": "2026-09-26T00:00:00Z",
27936
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
27937
+ },
27938
+ "assistantPrefill": {
27939
+ "value": false,
27940
+ "source": "official-doc",
27941
+ "confidence": "declared",
27942
+ "observedAt": "2026-09-26T00:00:00Z",
27943
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
27944
+ },
27171
27945
  "canonicalModelId": "claude-opus-4.8",
27172
27946
  "modelFamily": "claude"
27173
27947
  },
@@ -27307,6 +28081,24 @@
27307
28081
  "sourceRef": "continuity report §4.4, BOTH halves. Own-state acceptance: the model's own `thinking` blocks are replayed to it unchanged, in order, with their signatures intact (§4.4's replay rules and its worked request). Why the domain is NARROW: \"when changing Claude models, prior `thinking` and `redacted_thinking` blocks should be stripped because they are tied to the model that produced them. Therefore 'same provider' is not automatically 'same continuation domain.'\" The domain is this model alone, never the Anthropic provider.",
27308
28082
  "confidence": "declared",
27309
28083
  "observedAt": "2026-09-16T00:00:00Z"
28084
+ },
28085
+ "effortRequest": {
28086
+ "value": {
28087
+ "field": "output_config.effort"
28088
+ },
28089
+ "source": "official-doc",
28090
+ "confidence": "declared",
28091
+ "observedAt": "2026-09-25T12:30:00Z",
28092
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
28093
+ },
28094
+ "perMessageEffort": {
28095
+ "value": {
28096
+ "beta": "mid-conversation-output-config-2026-07-01"
28097
+ },
28098
+ "source": "official-doc",
28099
+ "confidence": "declared",
28100
+ "observedAt": "2026-09-25T18:00:00Z",
28101
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort#change-effort-mid-conversation-beta — \"On Claude Fable 5.1, Claude Mythos 5.1, Claude Opus 5.5, and Claude Opus 5, use a per-message effort change, which keeps the prompt cache\": a role:\"system\" message with empty content and output_config.effort, beta header mid-conversation-output-config-2026-07-01; \"Models without per-message effort, including Claude Fable 5, return a 400 error\". claude 2.1.282 sends the older alias per-turn-control-2026-07-01, which the page does not document."
27310
28102
  }
27311
28103
  },
27312
28104
  "pricing": {
@@ -27324,9 +28116,51 @@
27324
28116
  "unsupportedParameters": [
27325
28117
  "temperature",
27326
28118
  "top_p",
27327
- "top_k"
28119
+ "top_k",
28120
+ "thinking.type.enabled",
28121
+ "thinking.type.disabled+output_config.effort.xhigh",
28122
+ "thinking.type.disabled+output_config.effort.max"
27328
28123
  ],
27329
28124
  "status": "candidate",
28125
+ "deferredToolLoading": {
28126
+ "value": true,
28127
+ "source": "official-doc",
28128
+ "confidence": "declared",
28129
+ "observedAt": "2026-09-25T18:30:00Z",
28130
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
28131
+ },
28132
+ "midConversationSystem": {
28133
+ "value": true,
28134
+ "source": "official-doc",
28135
+ "confidence": "declared",
28136
+ "observedAt": "2026-09-25T19:00:00Z",
28137
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
28138
+ },
28139
+ "midConversationToolChanges": {
28140
+ "value": {
28141
+ "beta": "mid-conversation-tool-changes-2026-07-01"
28142
+ },
28143
+ "source": "official-doc",
28144
+ "confidence": "declared",
28145
+ "observedAt": "2026-09-26T00:00:00Z",
28146
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
28147
+ },
28148
+ "inlineToolDefinitions": {
28149
+ "value": {
28150
+ "beta": "inline-tools-2026-09-15"
28151
+ },
28152
+ "source": "official-doc",
28153
+ "confidence": "declared",
28154
+ "observedAt": "2026-09-26T00:00:00Z",
28155
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
28156
+ },
28157
+ "assistantPrefill": {
28158
+ "value": false,
28159
+ "source": "official-doc",
28160
+ "confidence": "declared",
28161
+ "observedAt": "2026-09-26T00:00:00Z",
28162
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
28163
+ },
27330
28164
  "canonicalModelId": "claude-opus-5",
27331
28165
  "modelFamily": "claude"
27332
28166
  },
@@ -27468,6 +28302,33 @@
27468
28302
  "sourceRef": "continuity report §3 (wire probe of the pinned Claude Agent SDK, 0.3.250→0.3.258) — every uncompacted historical thinking/redacted block is replayed on every later request",
27469
28303
  "confidence": "declared",
27470
28304
  "observedAt": "2026-09-05T00:00:00Z"
28305
+ },
28306
+ "effortRequest": {
28307
+ "value": {
28308
+ "field": "output_config.effort"
28309
+ },
28310
+ "source": "official-doc",
28311
+ "confidence": "declared",
28312
+ "observedAt": "2026-09-25T12:30:00Z",
28313
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
28314
+ },
28315
+ "blockBinding": {
28316
+ "value": {
28317
+ "beta": "thinking-binding-controls-2026-08-01"
28318
+ },
28319
+ "source": "official-doc",
28320
+ "confidence": "declared",
28321
+ "observedAt": "2026-09-25T13:00:00Z",
28322
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting — \"A 400 error says a thinking block signature is invalid\": on Claude Fable 5.1 and Opus 5.5 a replayed thinking block is bound to the conversation (system, tools and earlier messages) and is rejected once that prefix changes (enforced for accounts created on or after 2026-08-31); the documented escape is the `thinking-binding-controls-2026-08-01` beta with `thinking.block_binding.prefix_mismatch_behavior: \"drop_block\"`. See also https://platform.claude.com/docs/en/models/opus-5-5/whats-new-opus-5-5 (\"Thinking blocks are tied to the model and the conversation\")."
28323
+ },
28324
+ "perMessageEffort": {
28325
+ "value": {
28326
+ "beta": "mid-conversation-output-config-2026-07-01"
28327
+ },
28328
+ "source": "official-doc",
28329
+ "confidence": "declared",
28330
+ "observedAt": "2026-09-25T18:00:00Z",
28331
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort#change-effort-mid-conversation-beta — \"On Claude Fable 5.1, Claude Mythos 5.1, Claude Opus 5.5, and Claude Opus 5, use a per-message effort change, which keeps the prompt cache\": a role:\"system\" message with empty content and output_config.effort, beta header mid-conversation-output-config-2026-07-01; \"Models without per-message effort, including Claude Fable 5, return a 400 error\". claude 2.1.282 sends the older alias per-turn-control-2026-07-01, which the page does not document."
27471
28332
  }
27472
28333
  },
27473
28334
  "pricing": {
@@ -27482,8 +28343,52 @@
27482
28343
  "confidence": "declared",
27483
28344
  "observedAt": "2026-09-25T10:25:20Z"
27484
28345
  },
27485
- "unsupportedParameters": [],
28346
+ "unsupportedParameters": [
28347
+ "thinking.type.enabled",
28348
+ "thinking.type.disabled",
28349
+ "tool_choice.any",
28350
+ "tool_choice.tool"
28351
+ ],
27486
28352
  "status": "candidate",
28353
+ "deferredToolLoading": {
28354
+ "value": true,
28355
+ "source": "official-doc",
28356
+ "confidence": "declared",
28357
+ "observedAt": "2026-09-25T18:30:00Z",
28358
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
28359
+ },
28360
+ "midConversationSystem": {
28361
+ "value": true,
28362
+ "source": "official-doc",
28363
+ "confidence": "declared",
28364
+ "observedAt": "2026-09-25T19:00:00Z",
28365
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — \"This feature is available on Claude Fable 5.1, Claude Mythos 5.1, Claude Fable 5, Claude Mythos 5, Claude Opus 5.5, Claude Opus 4.8, and Claude Opus 5. No beta header is required for mid-conversation system messages. This feature is not available on Claude Sonnet 5.\""
28366
+ },
28367
+ "midConversationToolChanges": {
28368
+ "value": {
28369
+ "beta": "mid-conversation-tool-changes-2026-07-01"
28370
+ },
28371
+ "source": "official-doc",
28372
+ "confidence": "declared",
28373
+ "observedAt": "2026-09-26T00:00:00Z",
28374
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition / tool_removal blocks in a role:\"system\" message, by reference (\"Add/remove by reference: mid-conversation-tool-changes-2026-07-01 — Claude API, Amazon Bedrock, Google Cloud\"); available on Claude Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 4.8 and Opus 5, \"Not available on Claude Sonnet 5\". Retrieved 2026-09-26; release note 2026-07-24 (https://platform.claude.com/docs/en/release-notes/overview)."
28375
+ },
28376
+ "inlineToolDefinitions": {
28377
+ "value": {
28378
+ "beta": "inline-tools-2026-09-15"
28379
+ },
28380
+ "source": "official-doc",
28381
+ "confidence": "declared",
28382
+ "observedAt": "2026-09-26T00:00:00Z",
28383
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/mid-conversation-system-messages — tool_addition with a tool_definition (\"Add/remove by reference + define by value: inline-tools-2026-09-15 — Claude API\"); \"The inline-tools-2026-09-15 header covers all reference-based changes\"; redefinition: \"send a different definition under the same name ... The new definition replaces the earlier one from that position onward\" (no removal first, a tool declared in tools included); a definition reusing the name of a different type of tool is a 400 tool_name_conflict. Same model list as midConversationToolChanges. Claude API only, so set on the anthropic and console providers (both api.anthropic.com), never on a cloud reseller. Retrieved 2026-09-26; release note 2026-09-22."
28384
+ },
28385
+ "assistantPrefill": {
28386
+ "value": false,
28387
+ "source": "official-doc",
28388
+ "confidence": "declared",
28389
+ "observedAt": "2026-09-26T00:00:00Z",
28390
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
28391
+ },
27487
28392
  "canonicalModelId": "claude-opus-5.5",
27488
28393
  "modelFamily": "claude"
27489
28394
  },
@@ -27595,8 +28500,17 @@
27595
28500
  "confidence": "declared",
27596
28501
  "observedAt": "2026-09-19T00:00:00Z"
27597
28502
  },
27598
- "unsupportedParameters": [],
28503
+ "unsupportedParameters": [
28504
+ "thinking.type.adaptive"
28505
+ ],
27599
28506
  "status": "candidate",
28507
+ "deferredToolLoading": {
28508
+ "value": true,
28509
+ "source": "official-doc",
28510
+ "confidence": "declared",
28511
+ "observedAt": "2026-09-25T18:30:00Z",
28512
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
28513
+ },
27600
28514
  "canonicalModelId": "claude-sonnet-4.5",
27601
28515
  "modelFamily": "claude"
27602
28516
  },
@@ -27605,7 +28519,9 @@
27605
28519
  "providerId": "console",
27606
28520
  "upstreamId": "claude-sonnet-4.6",
27607
28521
  "displayName": "Claude Sonnet 4.6",
27608
- "aliases": [],
28522
+ "aliases": [
28523
+ "claude-sonnet-4-6"
28524
+ ],
27609
28525
  "endpoints": [
27610
28526
  "chat"
27611
28527
  ],
@@ -27692,7 +28608,16 @@
27692
28608
  "max"
27693
28609
  ],
27694
28610
  "continuation": "opaque-provider-state",
27695
- "defaultEffort": "high"
28611
+ "defaultEffort": "high",
28612
+ "effortRequest": {
28613
+ "value": {
28614
+ "field": "output_config.effort"
28615
+ },
28616
+ "source": "official-doc",
28617
+ "confidence": "declared",
28618
+ "observedAt": "2026-09-25T12:30:00Z",
28619
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
28620
+ }
27696
28621
  },
27697
28622
  "pricing": {
27698
28623
  "value": {
@@ -27708,6 +28633,13 @@
27708
28633
  },
27709
28634
  "unsupportedParameters": [],
27710
28635
  "status": "candidate",
28636
+ "deferredToolLoading": {
28637
+ "value": true,
28638
+ "source": "official-doc",
28639
+ "confidence": "declared",
28640
+ "observedAt": "2026-09-25T18:30:00Z",
28641
+ "sourceRef": "https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-search-tool — the Model compatibility table lists Fable 5.1, Mythos 5.1, Fable 5, Mythos 5, Opus 5.5, Opus 5, Opus 4.8, 4.7, 4.6, Sonnet 4.6, Opus 4.5, Sonnet 4.5 and Haiku 4.5 (not Sonnet 5); \"Custom tool search implementation\": a client tool returns a tool_result whose content carries {type:\"tool_reference\", tool_name} blocks, every referenced tool declared in tools with defer_loading: true, and the API expands them; \"The prefix is untouched, so prompt caching is preserved\" (https://platform.claude.com/docs/en/agents-and-tools/tool-use/tool-use-with-prompt-caching)."
28642
+ },
27711
28643
  "canonicalModelId": "claude-sonnet-4.6",
27712
28644
  "modelFamily": "claude"
27713
28645
  },
@@ -27849,6 +28781,15 @@
27849
28781
  "sourceRef": "continuity report §4.4, BOTH halves. Own-state acceptance: §4.4 requires the model's own assistant blocks — `thinking` with its signature, `redacted_thinking` — replayed unchanged and in order across a tool loop. Why the domain is NARROW: Anthropic documents model switching as a boundary at which those blocks are STRIPPED, because they are tied to the producing model; \"same provider\" is therefore not \"same continuation domain\", and this domain is the model alone.",
27850
28782
  "confidence": "declared",
27851
28783
  "observedAt": "2026-09-16T00:00:00Z"
28784
+ },
28785
+ "effortRequest": {
28786
+ "value": {
28787
+ "field": "output_config.effort"
28788
+ },
28789
+ "source": "official-doc",
28790
+ "confidence": "declared",
28791
+ "observedAt": "2026-09-25T12:30:00Z",
28792
+ "sourceRef": "https://platform.claude.com/docs/en/build-with-claude/effort — the effort page's supportedModels list names claude-fable-5-1, claude-fable-5, claude-opus-5-5, claude-opus-5, claude-opus-4-8, claude-opus-4-7, claude-opus-4-6, claude-opus-4-5-20251101, claude-sonnet-5 and claude-sonnet-4-6: effort is requested through `output_config.effort` (Claude 4.7 and later reject a manual `thinking.budget_tokens`, https://platform.claude.com/docs/en/build-with-claude/thinking-troubleshooting#rejected-configurations)."
27852
28793
  }
27853
28794
  },
27854
28795
  "pricing": {
@@ -27866,9 +28807,17 @@
27866
28807
  "unsupportedParameters": [
27867
28808
  "temperature",
27868
28809
  "top_p",
27869
- "top_k"
28810
+ "top_k",
28811
+ "thinking.type.enabled"
27870
28812
  ],
27871
28813
  "status": "candidate",
28814
+ "assistantPrefill": {
28815
+ "value": false,
28816
+ "source": "official-doc",
28817
+ "confidence": "declared",
28818
+ "observedAt": "2026-09-26T00:00:00Z",
28819
+ "sourceRef": "https://platform.claude.com/docs/en/models/opus-5-5/migration-guide — \"Don't end messages with a prefilled assistant turn: it is rejected\"; \"Prefilling assistant messages returns a 400 error on Claude Opus 4.6 and later Opus models, including Claude Opus 5.5\"; Sonnet 5 and Opus 5: \"Manual extended thinking, non-default sampling parameters, and assistant prefill return a 400 error on both models\"; the Fable 5.1 guide: prefill \"returns a 400 error, unchanged from Claude Fable 5\". Live gate 2026-09-26 on claude-opus-5-5: \"This model does not support assistant message prefill. The conversation must end with a user message.\""
28820
+ },
27872
28821
  "canonicalModelId": "claude-sonnet-5",
27873
28822
  "modelFamily": "claude"
27874
28823
  },
@@ -29142,12 +30091,11 @@
29142
30091
  "supported": {
29143
30092
  "value": true,
29144
30093
  "source": "official-doc",
29145
- "sourceRef": "https://api-docs.deepseek.com/guides/thinking_mode/ — Chat reasoning_effort and Responses reasoning.effort allow none/low/high/max; default high; none disables",
30094
+ "sourceRef": "https://api-docs.deepseek.com/guides/thinking_mode/ — Chat Completions thinking toggle uses thinking.type enabled/disabled; reasoning_effort low/high/max; default high. Responses reasoning.effort additionally accepts none",
29146
30095
  "confidence": "declared",
29147
- "observedAt": "2026-09-25T08:45:47Z"
30096
+ "observedAt": "2026-09-25T12:31:56.381Z"
29148
30097
  },
29149
30098
  "efforts": [
29150
- "none",
29151
30099
  "low",
29152
30100
  "high",
29153
30101
  "max"
@@ -29272,12 +30220,11 @@
29272
30220
  "supported": {
29273
30221
  "value": true,
29274
30222
  "source": "official-doc",
29275
- "sourceRef": "https://api-docs.deepseek.com/guides/thinking_mode/ — Chat reasoning_effort and Responses reasoning.effort allow none/low/high/max; default high; none disables",
30223
+ "sourceRef": "https://api-docs.deepseek.com/guides/thinking_mode/ — Chat Completions thinking toggle uses thinking.type enabled/disabled; reasoning_effort low/high/max; default high. Responses reasoning.effort additionally accepts none",
29276
30224
  "confidence": "declared",
29277
- "observedAt": "2026-09-25T08:45:47Z"
30225
+ "observedAt": "2026-09-25T12:31:56.381Z"
29278
30226
  },
29279
30227
  "efforts": [
29280
- "none",
29281
30228
  "low",
29282
30229
  "high",
29283
30230
  "max"
@@ -42469,9 +43416,9 @@
42469
43416
  "cacheWritePerMTokUsd": 0.375
42470
43417
  },
42471
43418
  "source": "official-doc",
42472
- "sourceRef": "https://platform.minimax.io/subscribe/token-plan?tab=api-enterprise — current MiniMax API pay-as-you-go Token Plan price table for this M2.7 model, per million tokens (retrieved 2026-09-19).",
43419
+ "sourceRef": "https://platform.minimax.io/docs/guides/pricing-paygo — MiniMax-M2.7 Pay as You Go Standard table: $0.3 input/$1.2 output/$0.06 prompt-cache read/$0.375 prompt-cache write per million tokens",
42473
43420
  "confidence": "declared",
42474
- "observedAt": "2026-09-19T00:00:00Z"
43421
+ "observedAt": "2026-09-25T12:31:56.381Z"
42475
43422
  },
42476
43423
  "unsupportedParameters": [
42477
43424
  "function_call"
@@ -42561,9 +43508,9 @@
42561
43508
  "cacheWritePerMTokUsd": 0.375
42562
43509
  },
42563
43510
  "source": "official-doc",
42564
- "sourceRef": "https://platform.minimax.io/subscribe/token-plan?tab=api-enterprise — current MiniMax API pay-as-you-go Token Plan price table for this M2.7 model, per million tokens (retrieved 2026-09-19).",
43511
+ "sourceRef": "https://platform.minimax.io/docs/guides/pricing-paygo — MiniMax-M2.7-highspeed Pay as You Go Standard table: $0.6 input/$2.4 output/$0.06 prompt-cache read/$0.375 prompt-cache write per million tokens",
42565
43512
  "confidence": "declared",
42566
- "observedAt": "2026-09-19T00:00:00Z"
43513
+ "observedAt": "2026-09-25T12:31:56.381Z"
42567
43514
  },
42568
43515
  "unsupportedParameters": [
42569
43516
  "function_call"
@@ -50171,6 +51118,13 @@
50171
51118
  },
50172
51119
  "unsupportedParameters": [],
50173
51120
  "status": "candidate",
51121
+ "promptCacheKey": {
51122
+ "value": true,
51123
+ "source": "official-doc",
51124
+ "confidence": "declared",
51125
+ "observedAt": "2026-09-25T19:30:00Z",
51126
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51127
+ },
50174
51128
  "canonicalModelId": "gpt-4.1",
50175
51129
  "modelFamily": "gpt"
50176
51130
  },
@@ -50251,6 +51205,13 @@
50251
51205
  },
50252
51206
  "unsupportedParameters": [],
50253
51207
  "status": "candidate",
51208
+ "promptCacheKey": {
51209
+ "value": true,
51210
+ "source": "official-doc",
51211
+ "confidence": "declared",
51212
+ "observedAt": "2026-09-25T19:30:00Z",
51213
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51214
+ },
50254
51215
  "canonicalModelId": "gpt-4.1-mini",
50255
51216
  "modelFamily": "gpt"
50256
51217
  },
@@ -50338,6 +51299,13 @@
50338
51299
  },
50339
51300
  "unsupportedParameters": [],
50340
51301
  "status": "candidate",
51302
+ "promptCacheKey": {
51303
+ "value": true,
51304
+ "source": "official-doc",
51305
+ "confidence": "declared",
51306
+ "observedAt": "2026-09-25T19:30:00Z",
51307
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51308
+ },
50341
51309
  "canonicalModelId": "gpt-4.1-nano",
50342
51310
  "modelFamily": "gpt"
50343
51311
  },
@@ -50418,6 +51386,13 @@
50418
51386
  },
50419
51387
  "unsupportedParameters": [],
50420
51388
  "status": "candidate",
51389
+ "promptCacheKey": {
51390
+ "value": true,
51391
+ "source": "official-doc",
51392
+ "confidence": "declared",
51393
+ "observedAt": "2026-09-25T19:30:00Z",
51394
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51395
+ },
50421
51396
  "canonicalModelId": "gpt-4o",
50422
51397
  "modelFamily": "gpt"
50423
51398
  },
@@ -50498,6 +51473,13 @@
50498
51473
  },
50499
51474
  "unsupportedParameters": [],
50500
51475
  "status": "candidate",
51476
+ "promptCacheKey": {
51477
+ "value": true,
51478
+ "source": "official-doc",
51479
+ "confidence": "declared",
51480
+ "observedAt": "2026-09-25T19:30:00Z",
51481
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51482
+ },
50501
51483
  "canonicalModelId": "gpt-4o-2024-11-20",
50502
51484
  "modelFamily": "gpt"
50503
51485
  },
@@ -50578,6 +51560,13 @@
50578
51560
  },
50579
51561
  "unsupportedParameters": [],
50580
51562
  "status": "candidate",
51563
+ "promptCacheKey": {
51564
+ "value": true,
51565
+ "source": "official-doc",
51566
+ "confidence": "declared",
51567
+ "observedAt": "2026-09-25T19:30:00Z",
51568
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51569
+ },
50581
51570
  "canonicalModelId": "gpt-4o-mini",
50582
51571
  "modelFamily": "gpt"
50583
51572
  },
@@ -50683,6 +51672,34 @@
50683
51672
  },
50684
51673
  "unsupportedParameters": [],
50685
51674
  "status": "candidate",
51675
+ "promptCacheKey": {
51676
+ "value": true,
51677
+ "source": "official-doc",
51678
+ "confidence": "declared",
51679
+ "observedAt": "2026-09-25T19:30:00Z",
51680
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51681
+ },
51682
+ "clientToolSearch": {
51683
+ "value": true,
51684
+ "source": "official-doc",
51685
+ "confidence": "declared",
51686
+ "observedAt": "2026-09-26T00:00:00Z",
51687
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
51688
+ },
51689
+ "additionalToolsItem": {
51690
+ "value": true,
51691
+ "source": "official-doc",
51692
+ "confidence": "declared",
51693
+ "observedAt": "2026-09-26T00:00:00Z",
51694
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
51695
+ },
51696
+ "allowedToolsChoice": {
51697
+ "value": true,
51698
+ "source": "official-doc",
51699
+ "confidence": "declared",
51700
+ "observedAt": "2026-09-26T00:00:00Z",
51701
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
51702
+ },
50686
51703
  "canonicalModelId": "gpt-5.4",
50687
51704
  "modelFamily": "gpt"
50688
51705
  },
@@ -50795,6 +51812,34 @@
50795
51812
  },
50796
51813
  "unsupportedParameters": [],
50797
51814
  "status": "candidate",
51815
+ "promptCacheKey": {
51816
+ "value": true,
51817
+ "source": "official-doc",
51818
+ "confidence": "declared",
51819
+ "observedAt": "2026-09-25T19:30:00Z",
51820
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51821
+ },
51822
+ "clientToolSearch": {
51823
+ "value": true,
51824
+ "source": "official-doc",
51825
+ "confidence": "declared",
51826
+ "observedAt": "2026-09-26T00:00:00Z",
51827
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
51828
+ },
51829
+ "additionalToolsItem": {
51830
+ "value": true,
51831
+ "source": "official-doc",
51832
+ "confidence": "declared",
51833
+ "observedAt": "2026-09-26T00:00:00Z",
51834
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
51835
+ },
51836
+ "allowedToolsChoice": {
51837
+ "value": true,
51838
+ "source": "official-doc",
51839
+ "confidence": "declared",
51840
+ "observedAt": "2026-09-26T00:00:00Z",
51841
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
51842
+ },
50798
51843
  "canonicalModelId": "gpt-5.4-mini",
50799
51844
  "modelFamily": "gpt"
50800
51845
  },
@@ -50907,6 +51952,34 @@
50907
51952
  },
50908
51953
  "unsupportedParameters": [],
50909
51954
  "status": "candidate",
51955
+ "promptCacheKey": {
51956
+ "value": true,
51957
+ "source": "official-doc",
51958
+ "confidence": "declared",
51959
+ "observedAt": "2026-09-25T19:30:00Z",
51960
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
51961
+ },
51962
+ "clientToolSearch": {
51963
+ "value": true,
51964
+ "source": "official-doc",
51965
+ "confidence": "declared",
51966
+ "observedAt": "2026-09-26T00:00:00Z",
51967
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
51968
+ },
51969
+ "additionalToolsItem": {
51970
+ "value": true,
51971
+ "source": "official-doc",
51972
+ "confidence": "declared",
51973
+ "observedAt": "2026-09-26T00:00:00Z",
51974
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
51975
+ },
51976
+ "allowedToolsChoice": {
51977
+ "value": true,
51978
+ "source": "official-doc",
51979
+ "confidence": "declared",
51980
+ "observedAt": "2026-09-26T00:00:00Z",
51981
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
51982
+ },
50910
51983
  "canonicalModelId": "gpt-5.4-nano",
50911
51984
  "modelFamily": "gpt"
50912
51985
  },
@@ -50994,6 +52067,34 @@
50994
52067
  },
50995
52068
  "unsupportedParameters": [],
50996
52069
  "status": "candidate",
52070
+ "promptCacheKey": {
52071
+ "value": true,
52072
+ "source": "official-doc",
52073
+ "confidence": "declared",
52074
+ "observedAt": "2026-09-25T19:30:00Z",
52075
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52076
+ },
52077
+ "clientToolSearch": {
52078
+ "value": true,
52079
+ "source": "official-doc",
52080
+ "confidence": "declared",
52081
+ "observedAt": "2026-09-26T00:00:00Z",
52082
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52083
+ },
52084
+ "additionalToolsItem": {
52085
+ "value": true,
52086
+ "source": "official-doc",
52087
+ "confidence": "declared",
52088
+ "observedAt": "2026-09-26T00:00:00Z",
52089
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52090
+ },
52091
+ "allowedToolsChoice": {
52092
+ "value": true,
52093
+ "source": "official-doc",
52094
+ "confidence": "declared",
52095
+ "observedAt": "2026-09-26T00:00:00Z",
52096
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52097
+ },
50997
52098
  "canonicalModelId": "gpt-5.4-pro",
50998
52099
  "modelFamily": "gpt"
50999
52100
  },
@@ -51099,6 +52200,34 @@
51099
52200
  },
51100
52201
  "unsupportedParameters": [],
51101
52202
  "status": "candidate",
52203
+ "promptCacheKey": {
52204
+ "value": true,
52205
+ "source": "official-doc",
52206
+ "confidence": "declared",
52207
+ "observedAt": "2026-09-25T19:30:00Z",
52208
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52209
+ },
52210
+ "clientToolSearch": {
52211
+ "value": true,
52212
+ "source": "official-doc",
52213
+ "confidence": "declared",
52214
+ "observedAt": "2026-09-26T00:00:00Z",
52215
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52216
+ },
52217
+ "additionalToolsItem": {
52218
+ "value": true,
52219
+ "source": "official-doc",
52220
+ "confidence": "declared",
52221
+ "observedAt": "2026-09-26T00:00:00Z",
52222
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52223
+ },
52224
+ "allowedToolsChoice": {
52225
+ "value": true,
52226
+ "source": "official-doc",
52227
+ "confidence": "declared",
52228
+ "observedAt": "2026-09-26T00:00:00Z",
52229
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52230
+ },
51102
52231
  "canonicalModelId": "gpt-5.5",
51103
52232
  "modelFamily": "gpt"
51104
52233
  },
@@ -51193,6 +52322,34 @@
51193
52322
  },
51194
52323
  "unsupportedParameters": [],
51195
52324
  "status": "candidate",
52325
+ "promptCacheKey": {
52326
+ "value": true,
52327
+ "source": "official-doc",
52328
+ "confidence": "declared",
52329
+ "observedAt": "2026-09-25T19:30:00Z",
52330
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52331
+ },
52332
+ "clientToolSearch": {
52333
+ "value": true,
52334
+ "source": "official-doc",
52335
+ "confidence": "declared",
52336
+ "observedAt": "2026-09-26T00:00:00Z",
52337
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52338
+ },
52339
+ "additionalToolsItem": {
52340
+ "value": true,
52341
+ "source": "official-doc",
52342
+ "confidence": "declared",
52343
+ "observedAt": "2026-09-26T00:00:00Z",
52344
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52345
+ },
52346
+ "allowedToolsChoice": {
52347
+ "value": true,
52348
+ "source": "official-doc",
52349
+ "confidence": "declared",
52350
+ "observedAt": "2026-09-26T00:00:00Z",
52351
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52352
+ },
51196
52353
  "canonicalModelId": "gpt-5.5-pro",
51197
52354
  "modelFamily": "gpt"
51198
52355
  },
@@ -51307,6 +52464,34 @@
51307
52464
  },
51308
52465
  "unsupportedParameters": [],
51309
52466
  "status": "candidate",
52467
+ "promptCacheKey": {
52468
+ "value": true,
52469
+ "source": "official-doc",
52470
+ "confidence": "declared",
52471
+ "observedAt": "2026-09-25T19:30:00Z",
52472
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52473
+ },
52474
+ "clientToolSearch": {
52475
+ "value": true,
52476
+ "source": "official-doc",
52477
+ "confidence": "declared",
52478
+ "observedAt": "2026-09-26T00:00:00Z",
52479
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52480
+ },
52481
+ "additionalToolsItem": {
52482
+ "value": true,
52483
+ "source": "official-doc",
52484
+ "confidence": "declared",
52485
+ "observedAt": "2026-09-26T00:00:00Z",
52486
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52487
+ },
52488
+ "allowedToolsChoice": {
52489
+ "value": true,
52490
+ "source": "official-doc",
52491
+ "confidence": "declared",
52492
+ "observedAt": "2026-09-26T00:00:00Z",
52493
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52494
+ },
51310
52495
  "canonicalModelId": "gpt-5.6",
51311
52496
  "modelFamily": "gpt"
51312
52497
  },
@@ -51421,6 +52606,34 @@
51421
52606
  },
51422
52607
  "unsupportedParameters": [],
51423
52608
  "status": "candidate",
52609
+ "promptCacheKey": {
52610
+ "value": true,
52611
+ "source": "official-doc",
52612
+ "confidence": "declared",
52613
+ "observedAt": "2026-09-25T19:30:00Z",
52614
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52615
+ },
52616
+ "clientToolSearch": {
52617
+ "value": true,
52618
+ "source": "official-doc",
52619
+ "confidence": "declared",
52620
+ "observedAt": "2026-09-26T00:00:00Z",
52621
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52622
+ },
52623
+ "additionalToolsItem": {
52624
+ "value": true,
52625
+ "source": "official-doc",
52626
+ "confidence": "declared",
52627
+ "observedAt": "2026-09-26T00:00:00Z",
52628
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52629
+ },
52630
+ "allowedToolsChoice": {
52631
+ "value": true,
52632
+ "source": "official-doc",
52633
+ "confidence": "declared",
52634
+ "observedAt": "2026-09-26T00:00:00Z",
52635
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52636
+ },
51424
52637
  "canonicalModelId": "gpt-5.6-luna",
51425
52638
  "modelFamily": "gpt"
51426
52639
  },
@@ -51535,6 +52748,34 @@
51535
52748
  },
51536
52749
  "unsupportedParameters": [],
51537
52750
  "status": "candidate",
52751
+ "promptCacheKey": {
52752
+ "value": true,
52753
+ "source": "official-doc",
52754
+ "confidence": "declared",
52755
+ "observedAt": "2026-09-25T19:30:00Z",
52756
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52757
+ },
52758
+ "clientToolSearch": {
52759
+ "value": true,
52760
+ "source": "official-doc",
52761
+ "confidence": "declared",
52762
+ "observedAt": "2026-09-26T00:00:00Z",
52763
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52764
+ },
52765
+ "additionalToolsItem": {
52766
+ "value": true,
52767
+ "source": "official-doc",
52768
+ "confidence": "declared",
52769
+ "observedAt": "2026-09-26T00:00:00Z",
52770
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52771
+ },
52772
+ "allowedToolsChoice": {
52773
+ "value": true,
52774
+ "source": "official-doc",
52775
+ "confidence": "declared",
52776
+ "observedAt": "2026-09-26T00:00:00Z",
52777
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52778
+ },
51538
52779
  "canonicalModelId": "gpt-5.6-sol",
51539
52780
  "modelFamily": "gpt"
51540
52781
  },
@@ -51649,6 +52890,34 @@
51649
52890
  },
51650
52891
  "unsupportedParameters": [],
51651
52892
  "status": "candidate",
52893
+ "promptCacheKey": {
52894
+ "value": true,
52895
+ "source": "official-doc",
52896
+ "confidence": "declared",
52897
+ "observedAt": "2026-09-25T19:30:00Z",
52898
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
52899
+ },
52900
+ "clientToolSearch": {
52901
+ "value": true,
52902
+ "source": "official-doc",
52903
+ "confidence": "declared",
52904
+ "observedAt": "2026-09-26T00:00:00Z",
52905
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
52906
+ },
52907
+ "additionalToolsItem": {
52908
+ "value": true,
52909
+ "source": "official-doc",
52910
+ "confidence": "declared",
52911
+ "observedAt": "2026-09-26T00:00:00Z",
52912
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
52913
+ },
52914
+ "allowedToolsChoice": {
52915
+ "value": true,
52916
+ "source": "official-doc",
52917
+ "confidence": "declared",
52918
+ "observedAt": "2026-09-26T00:00:00Z",
52919
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
52920
+ },
51652
52921
  "canonicalModelId": "gpt-5.6-terra",
51653
52922
  "modelFamily": "gpt"
51654
52923
  },
@@ -51746,7 +53015,7 @@
51746
53015
  "max"
51747
53016
  ],
51748
53017
  "continuation": "opaque-provider-state",
51749
- "defaultEffort": "low",
53018
+ "defaultEffort": "medium",
51750
53019
  "readableState": {
51751
53020
  "value": "summary",
51752
53021
  "source": "official-doc",
@@ -51783,6 +53052,15 @@
51783
53052
  "sourceRef": "continuity report §5.1/§5.3 for own-state acceptance — the Responses surface replays its own completed `reasoning.encrypted_content` items to the same model. WS-13 §8.2: \"OpenAI-compatible\" is never the capability test, so this key shares NO domain with any other row.",
51784
53053
  "confidence": "declared",
51785
53054
  "observedAt": "2026-09-06T00:00:00Z"
53055
+ },
53056
+ "perMessageEffort": {
53057
+ "value": {
53058
+ "item": "configuration_update"
53059
+ },
53060
+ "source": "official-doc",
53061
+ "confidence": "declared",
53062
+ "observedAt": "2026-09-26T00:00:00Z",
53063
+ "sourceRef": "https://developers.openai.com/api/docs/guides/reasoning — \"Configuration updates are supported by the GPT-6 model family in standard, single-agent mode. They change only reasoning effort.\" Shape {\"type\":\"configuration_update\",\"reasoning\":{\"effort\":…}}, placed \"before the next user message in the input array\"; \"the API rejects adjacent updates\"; \"The response's reasoning.effort continues to report the request-level setting\"; with store:false, replay updates \"in their original positions\". Retrieved 2026-09-26. NOT set on codex-oauth/gpt-6-* (codex-rs enables it only behind an off-by-default feature, backend acceptance unprobed — scripts/probe-openai-midconv.ts decides) nor on azure-openai (undocumented there)."
51786
53064
  }
51787
53065
  },
51788
53066
  "pricing": {
@@ -51799,6 +53077,34 @@
51799
53077
  },
51800
53078
  "unsupportedParameters": [],
51801
53079
  "status": "candidate",
53080
+ "promptCacheKey": {
53081
+ "value": true,
53082
+ "source": "official-doc",
53083
+ "confidence": "declared",
53084
+ "observedAt": "2026-09-25T19:30:00Z",
53085
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
53086
+ },
53087
+ "clientToolSearch": {
53088
+ "value": true,
53089
+ "source": "official-doc",
53090
+ "confidence": "declared",
53091
+ "observedAt": "2026-09-26T00:00:00Z",
53092
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
53093
+ },
53094
+ "additionalToolsItem": {
53095
+ "value": true,
53096
+ "source": "official-doc",
53097
+ "confidence": "declared",
53098
+ "observedAt": "2026-09-26T00:00:00Z",
53099
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
53100
+ },
53101
+ "allowedToolsChoice": {
53102
+ "value": true,
53103
+ "source": "official-doc",
53104
+ "confidence": "declared",
53105
+ "observedAt": "2026-09-26T00:00:00Z",
53106
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
53107
+ },
51802
53108
  "canonicalModelId": "gpt-6-astra",
51803
53109
  "modelFamily": "gpt"
51804
53110
  },
@@ -51934,6 +53240,15 @@
51934
53240
  "sourceRef": "continuity report §5.1/§5.3 for own-state acceptance — the Responses surface replays its own completed `reasoning.encrypted_content` items to the same model. WS-13 §8.2: \"OpenAI-compatible\" is never the capability test, so this key shares NO domain with any other row.",
51935
53241
  "confidence": "declared",
51936
53242
  "observedAt": "2026-09-06T00:00:00Z"
53243
+ },
53244
+ "perMessageEffort": {
53245
+ "value": {
53246
+ "item": "configuration_update"
53247
+ },
53248
+ "source": "official-doc",
53249
+ "confidence": "declared",
53250
+ "observedAt": "2026-09-26T00:00:00Z",
53251
+ "sourceRef": "https://developers.openai.com/api/docs/guides/reasoning — \"Configuration updates are supported by the GPT-6 model family in standard, single-agent mode. They change only reasoning effort.\" Shape {\"type\":\"configuration_update\",\"reasoning\":{\"effort\":…}}, placed \"before the next user message in the input array\"; \"the API rejects adjacent updates\"; \"The response's reasoning.effort continues to report the request-level setting\"; with store:false, replay updates \"in their original positions\". Retrieved 2026-09-26. NOT set on codex-oauth/gpt-6-* (codex-rs enables it only behind an off-by-default feature, backend acceptance unprobed — scripts/probe-openai-midconv.ts decides) nor on azure-openai (undocumented there)."
51937
53252
  }
51938
53253
  },
51939
53254
  "pricing": {
@@ -51950,6 +53265,34 @@
51950
53265
  },
51951
53266
  "unsupportedParameters": [],
51952
53267
  "status": "candidate",
53268
+ "promptCacheKey": {
53269
+ "value": true,
53270
+ "source": "official-doc",
53271
+ "confidence": "declared",
53272
+ "observedAt": "2026-09-25T19:30:00Z",
53273
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
53274
+ },
53275
+ "clientToolSearch": {
53276
+ "value": true,
53277
+ "source": "official-doc",
53278
+ "confidence": "declared",
53279
+ "observedAt": "2026-09-26T00:00:00Z",
53280
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
53281
+ },
53282
+ "additionalToolsItem": {
53283
+ "value": true,
53284
+ "source": "official-doc",
53285
+ "confidence": "declared",
53286
+ "observedAt": "2026-09-26T00:00:00Z",
53287
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
53288
+ },
53289
+ "allowedToolsChoice": {
53290
+ "value": true,
53291
+ "source": "official-doc",
53292
+ "confidence": "declared",
53293
+ "observedAt": "2026-09-26T00:00:00Z",
53294
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
53295
+ },
51953
53296
  "canonicalModelId": "gpt-6-luna",
51954
53297
  "modelFamily": "gpt"
51955
53298
  },
@@ -52085,6 +53428,15 @@
52085
53428
  "sourceRef": "continuity report §5.1/§5.3 for own-state acceptance — the Responses surface replays its own completed `reasoning.encrypted_content` items to the same model. WS-13 §8.2: \"OpenAI-compatible\" is never the capability test, so this key shares NO domain with any other row.",
52086
53429
  "confidence": "declared",
52087
53430
  "observedAt": "2026-09-06T00:00:00Z"
53431
+ },
53432
+ "perMessageEffort": {
53433
+ "value": {
53434
+ "item": "configuration_update"
53435
+ },
53436
+ "source": "official-doc",
53437
+ "confidence": "declared",
53438
+ "observedAt": "2026-09-26T00:00:00Z",
53439
+ "sourceRef": "https://developers.openai.com/api/docs/guides/reasoning — \"Configuration updates are supported by the GPT-6 model family in standard, single-agent mode. They change only reasoning effort.\" Shape {\"type\":\"configuration_update\",\"reasoning\":{\"effort\":…}}, placed \"before the next user message in the input array\"; \"the API rejects adjacent updates\"; \"The response's reasoning.effort continues to report the request-level setting\"; with store:false, replay updates \"in their original positions\". Retrieved 2026-09-26. NOT set on codex-oauth/gpt-6-* (codex-rs enables it only behind an off-by-default feature, backend acceptance unprobed — scripts/probe-openai-midconv.ts decides) nor on azure-openai (undocumented there)."
52088
53440
  }
52089
53441
  },
52090
53442
  "pricing": {
@@ -52101,6 +53453,34 @@
52101
53453
  },
52102
53454
  "unsupportedParameters": [],
52103
53455
  "status": "candidate",
53456
+ "promptCacheKey": {
53457
+ "value": true,
53458
+ "source": "official-doc",
53459
+ "confidence": "declared",
53460
+ "observedAt": "2026-09-25T19:30:00Z",
53461
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
53462
+ },
53463
+ "clientToolSearch": {
53464
+ "value": true,
53465
+ "source": "official-doc",
53466
+ "confidence": "declared",
53467
+ "observedAt": "2026-09-26T00:00:00Z",
53468
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — \"Only gpt-5.4 and later models support tool_search\" in the Responses API; client execution: {\"type\":\"tool_search\",\"execution\":\"client\",description,parameters}, tool_search_call -> tool_search_output {call_id, execution:\"client\", status, tools}, deferred functions keep defer_loading:true, optionally inside {\"type\":\"namespace\",name,description,tools}. Retrieved 2026-09-26."
53469
+ },
53470
+ "additionalToolsItem": {
53471
+ "value": true,
53472
+ "source": "official-doc",
53473
+ "confidence": "declared",
53474
+ "observedAt": "2026-09-26T00:00:00Z",
53475
+ "sourceRef": "https://developers.openai.com/api/docs/guides/tools-tool-search — {\"type\":\"additional_tools\",\"role\":\"developer\",\"tools\":[...]}: \"Tools in an additional_tools item become available only after that item appears in the input\"; its position is kept on replay. The page names no model list, so it is set only on the rows the same page documents tool_search for (GPT-5.4 and later). NOT on codex-oauth: codex-rs sends one only at input[0] (core/src/client.rs:910-941), never mid-conversation -- the live probe decides. Retrieved 2026-09-26."
53476
+ },
53477
+ "allowedToolsChoice": {
53478
+ "value": true,
53479
+ "source": "official-doc",
53480
+ "confidence": "declared",
53481
+ "observedAt": "2026-09-26T00:00:00Z",
53482
+ "sourceRef": "https://developers.openai.com/api/docs/guides/function-calling — tool_choice {\"type\":\"allowed_tools\",\"mode\":\"auto\"|\"required\",\"tools\":[{\"type\":\"function\",\"name\":…}]}: \"make only a subset of tools available across model requests, but not modify the list of tools you pass in, so you can maximize savings from prompt caching\"; \"When you use tool search, tool_choice still applies to the tools that are currently callable in the turn.\" No model list on the page; set on the same GPT-5.4+ rows. NOT on codex-oauth: codex-rs always sends tool_choice \"auto\" (core/src/client.rs:1004) and its allowed_tools is a client-side filter (ext/extension-api/src/tool_policy.rs:14-36) -- the live probe decides. Retrieved 2026-09-26."
53483
+ },
52104
53484
  "canonicalModelId": "gpt-6-sol",
52105
53485
  "modelFamily": "gpt"
52106
53486
  },
@@ -52203,6 +53583,13 @@
52203
53583
  },
52204
53584
  "unsupportedParameters": [],
52205
53585
  "status": "candidate",
53586
+ "promptCacheKey": {
53587
+ "value": true,
53588
+ "source": "official-doc",
53589
+ "confidence": "declared",
53590
+ "observedAt": "2026-09-25T19:30:00Z",
53591
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
53592
+ },
52206
53593
  "canonicalModelId": "o3",
52207
53594
  "modelFamily": "o-series"
52208
53595
  },
@@ -52297,6 +53684,13 @@
52297
53684
  },
52298
53685
  "unsupportedParameters": [],
52299
53686
  "status": "candidate",
53687
+ "promptCacheKey": {
53688
+ "value": true,
53689
+ "source": "official-doc",
53690
+ "confidence": "declared",
53691
+ "observedAt": "2026-09-25T19:30:00Z",
53692
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
53693
+ },
52300
53694
  "canonicalModelId": "o3-mini",
52301
53695
  "modelFamily": "o-series"
52302
53696
  },
@@ -52385,6 +53779,7 @@
52385
53779
  "high"
52386
53780
  ],
52387
53781
  "continuation": "opaque-provider-state",
53782
+ "defaultEffort": "medium",
52388
53783
  "readableState": {
52389
53784
  "value": "summary",
52390
53785
  "source": "official-doc",
@@ -52443,6 +53838,13 @@
52443
53838
  },
52444
53839
  "unsupportedParameters": [],
52445
53840
  "status": "candidate",
53841
+ "promptCacheKey": {
53842
+ "value": true,
53843
+ "source": "official-doc",
53844
+ "confidence": "declared",
53845
+ "observedAt": "2026-09-25T19:30:00Z",
53846
+ "sourceRef": "https://developers.openai.com/api/docs/guides/prompt-caching — the Responses API takes prompt_cache_key: \"Use a stable prompt_cache_key to optimize cache routing for requests that share a reusable prefix\" (pre-GPT-5.6); on GPT-5.6+ routing is automatic and the key is \"optional for separate cache accounting\"."
53847
+ },
52446
53848
  "canonicalModelId": "o4-mini",
52447
53849
  "modelFamily": "o-series"
52448
53850
  },
@@ -65599,6 +67001,27 @@
65599
67001
  "endpoints": [
65600
67002
  "chat"
65601
67003
  ],
67004
+ "contextWindow": {
67005
+ "value": 1000000,
67006
+ "source": "official-doc",
67007
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-flash-202605 row lists context window 1M tokens; k/M expanded as decimal",
67008
+ "confidence": "inferred",
67009
+ "observedAt": "2026-09-25T12:31:56.381Z"
67010
+ },
67011
+ "maxInputTokens": {
67012
+ "value": 1000000,
67013
+ "source": "official-doc",
67014
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-flash-202605 row lists maximum input 1M tokens; k/M expanded as decimal",
67015
+ "confidence": "inferred",
67016
+ "observedAt": "2026-09-25T12:31:56.381Z"
67017
+ },
67018
+ "maxOutputTokens": {
67019
+ "value": 384000,
67020
+ "source": "official-doc",
67021
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-flash-202605 row lists maximum output 384k tokens; k/M expanded as decimal",
67022
+ "confidence": "inferred",
67023
+ "observedAt": "2026-09-25T12:31:56.381Z"
67024
+ },
65602
67025
  "inputModalities": {
65603
67026
  "value": [
65604
67027
  "text"
@@ -65659,6 +67082,27 @@
65659
67082
  "endpoints": [
65660
67083
  "chat"
65661
67084
  ],
67085
+ "contextWindow": {
67086
+ "value": 1000000,
67087
+ "source": "official-doc",
67088
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-pro-202606 row lists context window 1M tokens; k/M expanded as decimal",
67089
+ "confidence": "inferred",
67090
+ "observedAt": "2026-09-25T12:31:56.381Z"
67091
+ },
67092
+ "maxInputTokens": {
67093
+ "value": 1000000,
67094
+ "source": "official-doc",
67095
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-pro-202606 row lists maximum input 1M tokens; k/M expanded as decimal",
67096
+ "confidence": "inferred",
67097
+ "observedAt": "2026-09-25T12:31:56.381Z"
67098
+ },
67099
+ "maxOutputTokens": {
67100
+ "value": 384000,
67101
+ "source": "official-doc",
67102
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-pro-202606 row lists maximum output 384k tokens; k/M expanded as decimal",
67103
+ "confidence": "inferred",
67104
+ "observedAt": "2026-09-25T12:31:56.381Z"
67105
+ },
65662
67106
  "inputModalities": {
65663
67107
  "value": [
65664
67108
  "text"
@@ -65718,6 +67162,27 @@
65718
67162
  "endpoints": [
65719
67163
  "chat"
65720
67164
  ],
67165
+ "contextWindow": {
67166
+ "value": 200000,
67167
+ "source": "official-doc",
67168
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5 row lists context window 200k tokens; k/M expanded as decimal",
67169
+ "confidence": "inferred",
67170
+ "observedAt": "2026-09-25T12:31:56.381Z"
67171
+ },
67172
+ "maxInputTokens": {
67173
+ "value": 200000,
67174
+ "source": "official-doc",
67175
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5 row lists maximum input 200k tokens; k/M expanded as decimal",
67176
+ "confidence": "inferred",
67177
+ "observedAt": "2026-09-25T12:31:56.381Z"
67178
+ },
67179
+ "maxOutputTokens": {
67180
+ "value": 128000,
67181
+ "source": "official-doc",
67182
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5 row lists maximum output 128k tokens; k/M expanded as decimal",
67183
+ "confidence": "inferred",
67184
+ "observedAt": "2026-09-25T12:31:56.381Z"
67185
+ },
65721
67186
  "inputModalities": {
65722
67187
  "value": [
65723
67188
  "text"
@@ -65777,6 +67242,27 @@
65777
67242
  "endpoints": [
65778
67243
  "chat"
65779
67244
  ],
67245
+ "contextWindow": {
67246
+ "value": 200000,
67247
+ "source": "official-doc",
67248
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.1 row lists context window 200k tokens; k/M expanded as decimal",
67249
+ "confidence": "inferred",
67250
+ "observedAt": "2026-09-25T12:31:56.381Z"
67251
+ },
67252
+ "maxInputTokens": {
67253
+ "value": 200000,
67254
+ "source": "official-doc",
67255
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.1 row lists maximum input 200k tokens; k/M expanded as decimal",
67256
+ "confidence": "inferred",
67257
+ "observedAt": "2026-09-25T12:31:56.381Z"
67258
+ },
67259
+ "maxOutputTokens": {
67260
+ "value": 128000,
67261
+ "source": "official-doc",
67262
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.1 row lists maximum output 128k tokens; k/M expanded as decimal",
67263
+ "confidence": "inferred",
67264
+ "observedAt": "2026-09-25T12:31:56.381Z"
67265
+ },
65780
67266
  "inputModalities": {
65781
67267
  "value": [
65782
67268
  "text"
@@ -65836,6 +67322,27 @@
65836
67322
  "endpoints": [
65837
67323
  "chat"
65838
67324
  ],
67325
+ "contextWindow": {
67326
+ "value": 1000000,
67327
+ "source": "official-doc",
67328
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.2 row lists context window 1M tokens; k/M expanded as decimal",
67329
+ "confidence": "inferred",
67330
+ "observedAt": "2026-09-25T12:31:56.381Z"
67331
+ },
67332
+ "maxInputTokens": {
67333
+ "value": 1000000,
67334
+ "source": "official-doc",
67335
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.2 row lists maximum input 1M tokens; k/M expanded as decimal",
67336
+ "confidence": "inferred",
67337
+ "observedAt": "2026-09-25T12:31:56.381Z"
67338
+ },
67339
+ "maxOutputTokens": {
67340
+ "value": 128000,
67341
+ "source": "official-doc",
67342
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.2 row lists maximum output 128k tokens; k/M expanded as decimal",
67343
+ "confidence": "inferred",
67344
+ "observedAt": "2026-09-25T12:31:56.381Z"
67345
+ },
65839
67346
  "inputModalities": {
65840
67347
  "value": [
65841
67348
  "text"
@@ -65895,6 +67402,27 @@
65895
67402
  "endpoints": [
65896
67403
  "chat"
65897
67404
  ],
67405
+ "contextWindow": {
67406
+ "value": 1000000,
67407
+ "source": "official-doc",
67408
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3 row lists context window 1M tokens; k/M expanded as decimal",
67409
+ "confidence": "inferred",
67410
+ "observedAt": "2026-09-25T12:31:56.381Z"
67411
+ },
67412
+ "maxInputTokens": {
67413
+ "value": 1000000,
67414
+ "source": "official-doc",
67415
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3 row lists maximum input 1M tokens; k/M expanded as decimal",
67416
+ "confidence": "inferred",
67417
+ "observedAt": "2026-09-25T12:31:56.381Z"
67418
+ },
67419
+ "maxOutputTokens": {
67420
+ "value": 128000,
67421
+ "source": "official-doc",
67422
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3 row lists maximum output 128k tokens; k/M expanded as decimal",
67423
+ "confidence": "inferred",
67424
+ "observedAt": "2026-09-25T12:31:56.381Z"
67425
+ },
65898
67426
  "inputModalities": {
65899
67427
  "value": [
65900
67428
  "text"
@@ -65952,6 +67480,27 @@
65952
67480
  "endpoints": [
65953
67481
  "chat"
65954
67482
  ],
67483
+ "contextWindow": {
67484
+ "value": 1000000,
67485
+ "source": "official-doc",
67486
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3-flash row lists context window 1M tokens; k/M expanded as decimal",
67487
+ "confidence": "inferred",
67488
+ "observedAt": "2026-09-25T12:31:56.381Z"
67489
+ },
67490
+ "maxInputTokens": {
67491
+ "value": 1000000,
67492
+ "source": "official-doc",
67493
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3-flash row lists maximum input 1M tokens; k/M expanded as decimal",
67494
+ "confidence": "inferred",
67495
+ "observedAt": "2026-09-25T12:31:56.381Z"
67496
+ },
67497
+ "maxOutputTokens": {
67498
+ "value": 128000,
67499
+ "source": "official-doc",
67500
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3-flash row lists maximum output 128k tokens; k/M expanded as decimal",
67501
+ "confidence": "inferred",
67502
+ "observedAt": "2026-09-25T12:31:56.381Z"
67503
+ },
65955
67504
  "inputModalities": {
65956
67505
  "value": [
65957
67506
  "text"
@@ -66196,6 +67745,27 @@
66196
67745
  "endpoints": [
66197
67746
  "chat"
66198
67747
  ],
67748
+ "contextWindow": {
67749
+ "value": 256000,
67750
+ "source": "official-doc",
67751
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k2.7-code row lists context window 256k tokens; k/M expanded as decimal",
67752
+ "confidence": "inferred",
67753
+ "observedAt": "2026-09-25T12:31:56.381Z"
67754
+ },
67755
+ "maxInputTokens": {
67756
+ "value": 256000,
67757
+ "source": "official-doc",
67758
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k2.7-code row lists maximum input 256k tokens; k/M expanded as decimal",
67759
+ "confidence": "inferred",
67760
+ "observedAt": "2026-09-25T12:31:56.381Z"
67761
+ },
67762
+ "maxOutputTokens": {
67763
+ "value": 256000,
67764
+ "source": "official-doc",
67765
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k2.7-code row lists maximum output 256k tokens; k/M expanded as decimal",
67766
+ "confidence": "inferred",
67767
+ "observedAt": "2026-09-25T12:31:56.381Z"
67768
+ },
66199
67769
  "inputModalities": {
66200
67770
  "value": [
66201
67771
  "text"
@@ -66253,6 +67823,27 @@
66253
67823
  "endpoints": [
66254
67824
  "chat"
66255
67825
  ],
67826
+ "contextWindow": {
67827
+ "value": 1000000,
67828
+ "source": "official-doc",
67829
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k3 row lists context window 1M tokens; k/M expanded as decimal",
67830
+ "confidence": "inferred",
67831
+ "observedAt": "2026-09-25T12:31:56.381Z"
67832
+ },
67833
+ "maxInputTokens": {
67834
+ "value": 1000000,
67835
+ "source": "official-doc",
67836
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k3 row lists maximum input 1M tokens; k/M expanded as decimal",
67837
+ "confidence": "inferred",
67838
+ "observedAt": "2026-09-25T12:31:56.381Z"
67839
+ },
67840
+ "maxOutputTokens": {
67841
+ "value": 1000000,
67842
+ "source": "official-doc",
67843
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k3 row lists maximum output 1M tokens; k/M expanded as decimal",
67844
+ "confidence": "inferred",
67845
+ "observedAt": "2026-09-25T12:31:56.381Z"
67846
+ },
66256
67847
  "inputModalities": {
66257
67848
  "value": [
66258
67849
  "text"
@@ -66312,6 +67903,27 @@
66312
67903
  "endpoints": [
66313
67904
  "chat"
66314
67905
  ],
67906
+ "contextWindow": {
67907
+ "value": 200000,
67908
+ "source": "official-doc",
67909
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m2.7 row lists context window 200k tokens; k/M expanded as decimal",
67910
+ "confidence": "inferred",
67911
+ "observedAt": "2026-09-25T12:31:56.381Z"
67912
+ },
67913
+ "maxInputTokens": {
67914
+ "value": 200000,
67915
+ "source": "official-doc",
67916
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m2.7 row lists maximum input 200k tokens; k/M expanded as decimal",
67917
+ "confidence": "inferred",
67918
+ "observedAt": "2026-09-25T12:31:56.381Z"
67919
+ },
67920
+ "maxOutputTokens": {
67921
+ "value": 128000,
67922
+ "source": "official-doc",
67923
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m2.7 row lists maximum output 128k tokens; k/M expanded as decimal",
67924
+ "confidence": "inferred",
67925
+ "observedAt": "2026-09-25T12:31:56.381Z"
67926
+ },
66315
67927
  "inputModalities": {
66316
67928
  "value": [
66317
67929
  "text"
@@ -66371,6 +67983,20 @@
66371
67983
  "endpoints": [
66372
67984
  "chat"
66373
67985
  ],
67986
+ "contextWindow": {
67987
+ "value": 1000000,
67988
+ "source": "official-doc",
67989
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m3 row lists context window 1M tokens; k/M expanded as decimal",
67990
+ "confidence": "inferred",
67991
+ "observedAt": "2026-09-25T12:31:56.381Z"
67992
+ },
67993
+ "maxInputTokens": {
67994
+ "value": 1000000,
67995
+ "source": "official-doc",
67996
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m3 row lists maximum input 1M tokens; k/M expanded as decimal",
67997
+ "confidence": "inferred",
67998
+ "observedAt": "2026-09-25T12:31:56.381Z"
67999
+ },
66374
68000
  "inputModalities": {
66375
68001
  "value": [
66376
68002
  "text"
@@ -66477,6 +68103,27 @@
66477
68103
  "endpoints": [
66478
68104
  "chat"
66479
68105
  ],
68106
+ "contextWindow": {
68107
+ "value": 1000000,
68108
+ "source": "official-doc",
68109
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-flash-202605 row lists context window 1M tokens; k/M expanded as decimal",
68110
+ "confidence": "inferred",
68111
+ "observedAt": "2026-09-25T12:31:56.381Z"
68112
+ },
68113
+ "maxInputTokens": {
68114
+ "value": 1000000,
68115
+ "source": "official-doc",
68116
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-flash-202605 row lists maximum input 1M tokens; k/M expanded as decimal",
68117
+ "confidence": "inferred",
68118
+ "observedAt": "2026-09-25T12:31:56.381Z"
68119
+ },
68120
+ "maxOutputTokens": {
68121
+ "value": 384000,
68122
+ "source": "official-doc",
68123
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-flash-202605 row lists maximum output 384k tokens; k/M expanded as decimal",
68124
+ "confidence": "inferred",
68125
+ "observedAt": "2026-09-25T12:31:56.381Z"
68126
+ },
66480
68127
  "inputModalities": {
66481
68128
  "value": [
66482
68129
  "text"
@@ -66541,6 +68188,27 @@
66541
68188
  "endpoints": [
66542
68189
  "chat"
66543
68190
  ],
68191
+ "contextWindow": {
68192
+ "value": 1000000,
68193
+ "source": "official-doc",
68194
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-pro-202606 row lists context window 1M tokens; k/M expanded as decimal",
68195
+ "confidence": "inferred",
68196
+ "observedAt": "2026-09-25T12:31:56.381Z"
68197
+ },
68198
+ "maxInputTokens": {
68199
+ "value": 1000000,
68200
+ "source": "official-doc",
68201
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-pro-202606 row lists maximum input 1M tokens; k/M expanded as decimal",
68202
+ "confidence": "inferred",
68203
+ "observedAt": "2026-09-25T12:31:56.381Z"
68204
+ },
68205
+ "maxOutputTokens": {
68206
+ "value": 384000,
68207
+ "source": "official-doc",
68208
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — deepseek-v4-pro-202606 row lists maximum output 384k tokens; k/M expanded as decimal",
68209
+ "confidence": "inferred",
68210
+ "observedAt": "2026-09-25T12:31:56.381Z"
68211
+ },
66544
68212
  "inputModalities": {
66545
68213
  "value": [
66546
68214
  "text"
@@ -66604,6 +68272,27 @@
66604
68272
  "endpoints": [
66605
68273
  "chat"
66606
68274
  ],
68275
+ "contextWindow": {
68276
+ "value": 200000,
68277
+ "source": "official-doc",
68278
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5 row lists context window 200k tokens; k/M expanded as decimal",
68279
+ "confidence": "inferred",
68280
+ "observedAt": "2026-09-25T12:31:56.381Z"
68281
+ },
68282
+ "maxInputTokens": {
68283
+ "value": 200000,
68284
+ "source": "official-doc",
68285
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5 row lists maximum input 200k tokens; k/M expanded as decimal",
68286
+ "confidence": "inferred",
68287
+ "observedAt": "2026-09-25T12:31:56.381Z"
68288
+ },
68289
+ "maxOutputTokens": {
68290
+ "value": 128000,
68291
+ "source": "official-doc",
68292
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5 row lists maximum output 128k tokens; k/M expanded as decimal",
68293
+ "confidence": "inferred",
68294
+ "observedAt": "2026-09-25T12:31:56.381Z"
68295
+ },
66607
68296
  "inputModalities": {
66608
68297
  "value": [
66609
68298
  "text"
@@ -66667,6 +68356,27 @@
66667
68356
  "endpoints": [
66668
68357
  "chat"
66669
68358
  ],
68359
+ "contextWindow": {
68360
+ "value": 200000,
68361
+ "source": "official-doc",
68362
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.1 row lists context window 200k tokens; k/M expanded as decimal",
68363
+ "confidence": "inferred",
68364
+ "observedAt": "2026-09-25T12:31:56.381Z"
68365
+ },
68366
+ "maxInputTokens": {
68367
+ "value": 200000,
68368
+ "source": "official-doc",
68369
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.1 row lists maximum input 200k tokens; k/M expanded as decimal",
68370
+ "confidence": "inferred",
68371
+ "observedAt": "2026-09-25T12:31:56.381Z"
68372
+ },
68373
+ "maxOutputTokens": {
68374
+ "value": 128000,
68375
+ "source": "official-doc",
68376
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.1 row lists maximum output 128k tokens; k/M expanded as decimal",
68377
+ "confidence": "inferred",
68378
+ "observedAt": "2026-09-25T12:31:56.381Z"
68379
+ },
66670
68380
  "inputModalities": {
66671
68381
  "value": [
66672
68382
  "text"
@@ -66730,6 +68440,27 @@
66730
68440
  "endpoints": [
66731
68441
  "chat"
66732
68442
  ],
68443
+ "contextWindow": {
68444
+ "value": 1000000,
68445
+ "source": "official-doc",
68446
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.2 row lists context window 1M tokens; k/M expanded as decimal",
68447
+ "confidence": "inferred",
68448
+ "observedAt": "2026-09-25T12:31:56.381Z"
68449
+ },
68450
+ "maxInputTokens": {
68451
+ "value": 1000000,
68452
+ "source": "official-doc",
68453
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.2 row lists maximum input 1M tokens; k/M expanded as decimal",
68454
+ "confidence": "inferred",
68455
+ "observedAt": "2026-09-25T12:31:56.381Z"
68456
+ },
68457
+ "maxOutputTokens": {
68458
+ "value": 128000,
68459
+ "source": "official-doc",
68460
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.2 row lists maximum output 128k tokens; k/M expanded as decimal",
68461
+ "confidence": "inferred",
68462
+ "observedAt": "2026-09-25T12:31:56.381Z"
68463
+ },
66733
68464
  "inputModalities": {
66734
68465
  "value": [
66735
68466
  "text"
@@ -66793,6 +68524,27 @@
66793
68524
  "endpoints": [
66794
68525
  "chat"
66795
68526
  ],
68527
+ "contextWindow": {
68528
+ "value": 1000000,
68529
+ "source": "official-doc",
68530
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3 row lists context window 1M tokens; k/M expanded as decimal",
68531
+ "confidence": "inferred",
68532
+ "observedAt": "2026-09-25T12:31:56.381Z"
68533
+ },
68534
+ "maxInputTokens": {
68535
+ "value": 1000000,
68536
+ "source": "official-doc",
68537
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3 row lists maximum input 1M tokens; k/M expanded as decimal",
68538
+ "confidence": "inferred",
68539
+ "observedAt": "2026-09-25T12:31:56.381Z"
68540
+ },
68541
+ "maxOutputTokens": {
68542
+ "value": 128000,
68543
+ "source": "official-doc",
68544
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3 row lists maximum output 128k tokens; k/M expanded as decimal",
68545
+ "confidence": "inferred",
68546
+ "observedAt": "2026-09-25T12:31:56.381Z"
68547
+ },
66796
68548
  "inputModalities": {
66797
68549
  "value": [
66798
68550
  "text"
@@ -66854,6 +68606,27 @@
66854
68606
  "endpoints": [
66855
68607
  "chat"
66856
68608
  ],
68609
+ "contextWindow": {
68610
+ "value": 1000000,
68611
+ "source": "official-doc",
68612
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3-flash row lists context window 1M tokens; k/M expanded as decimal",
68613
+ "confidence": "inferred",
68614
+ "observedAt": "2026-09-25T12:31:56.381Z"
68615
+ },
68616
+ "maxInputTokens": {
68617
+ "value": 1000000,
68618
+ "source": "official-doc",
68619
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3-flash row lists maximum input 1M tokens; k/M expanded as decimal",
68620
+ "confidence": "inferred",
68621
+ "observedAt": "2026-09-25T12:31:56.381Z"
68622
+ },
68623
+ "maxOutputTokens": {
68624
+ "value": 128000,
68625
+ "source": "official-doc",
68626
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — glm-5.3-flash row lists maximum output 128k tokens; k/M expanded as decimal",
68627
+ "confidence": "inferred",
68628
+ "observedAt": "2026-09-25T12:31:56.381Z"
68629
+ },
66857
68630
  "inputModalities": {
66858
68631
  "value": [
66859
68632
  "text"
@@ -67110,6 +68883,27 @@
67110
68883
  "endpoints": [
67111
68884
  "chat"
67112
68885
  ],
68886
+ "contextWindow": {
68887
+ "value": 256000,
68888
+ "source": "official-doc",
68889
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k2.7-code row lists context window 256k tokens; k/M expanded as decimal",
68890
+ "confidence": "inferred",
68891
+ "observedAt": "2026-09-25T12:31:56.381Z"
68892
+ },
68893
+ "maxInputTokens": {
68894
+ "value": 256000,
68895
+ "source": "official-doc",
68896
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k2.7-code row lists maximum input 256k tokens; k/M expanded as decimal",
68897
+ "confidence": "inferred",
68898
+ "observedAt": "2026-09-25T12:31:56.381Z"
68899
+ },
68900
+ "maxOutputTokens": {
68901
+ "value": 256000,
68902
+ "source": "official-doc",
68903
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k2.7-code row lists maximum output 256k tokens; k/M expanded as decimal",
68904
+ "confidence": "inferred",
68905
+ "observedAt": "2026-09-25T12:31:56.381Z"
68906
+ },
67113
68907
  "inputModalities": {
67114
68908
  "value": [
67115
68909
  "text"
@@ -67171,6 +68965,27 @@
67171
68965
  "endpoints": [
67172
68966
  "chat"
67173
68967
  ],
68968
+ "contextWindow": {
68969
+ "value": 1000000,
68970
+ "source": "official-doc",
68971
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k3 row lists context window 1M tokens; k/M expanded as decimal",
68972
+ "confidence": "inferred",
68973
+ "observedAt": "2026-09-25T12:31:56.381Z"
68974
+ },
68975
+ "maxInputTokens": {
68976
+ "value": 1000000,
68977
+ "source": "official-doc",
68978
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k3 row lists maximum input 1M tokens; k/M expanded as decimal",
68979
+ "confidence": "inferred",
68980
+ "observedAt": "2026-09-25T12:31:56.381Z"
68981
+ },
68982
+ "maxOutputTokens": {
68983
+ "value": 1000000,
68984
+ "source": "official-doc",
68985
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — kimi-k3 row lists maximum output 1M tokens; k/M expanded as decimal",
68986
+ "confidence": "inferred",
68987
+ "observedAt": "2026-09-25T12:31:56.381Z"
68988
+ },
67174
68989
  "inputModalities": {
67175
68990
  "value": [
67176
68991
  "text"
@@ -67234,6 +69049,27 @@
67234
69049
  "endpoints": [
67235
69050
  "chat"
67236
69051
  ],
69052
+ "contextWindow": {
69053
+ "value": 200000,
69054
+ "source": "official-doc",
69055
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m2.7 row lists context window 200k tokens; k/M expanded as decimal",
69056
+ "confidence": "inferred",
69057
+ "observedAt": "2026-09-25T12:31:56.381Z"
69058
+ },
69059
+ "maxInputTokens": {
69060
+ "value": 200000,
69061
+ "source": "official-doc",
69062
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m2.7 row lists maximum input 200k tokens; k/M expanded as decimal",
69063
+ "confidence": "inferred",
69064
+ "observedAt": "2026-09-25T12:31:56.381Z"
69065
+ },
69066
+ "maxOutputTokens": {
69067
+ "value": 128000,
69068
+ "source": "official-doc",
69069
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m2.7 row lists maximum output 128k tokens; k/M expanded as decimal",
69070
+ "confidence": "inferred",
69071
+ "observedAt": "2026-09-25T12:31:56.381Z"
69072
+ },
67237
69073
  "inputModalities": {
67238
69074
  "value": [
67239
69075
  "text"
@@ -67297,6 +69133,20 @@
67297
69133
  "endpoints": [
67298
69134
  "chat"
67299
69135
  ],
69136
+ "contextWindow": {
69137
+ "value": 1000000,
69138
+ "source": "official-doc",
69139
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m3 row lists context window 1M tokens; k/M expanded as decimal",
69140
+ "confidence": "inferred",
69141
+ "observedAt": "2026-09-25T12:31:56.381Z"
69142
+ },
69143
+ "maxInputTokens": {
69144
+ "value": 1000000,
69145
+ "source": "official-doc",
69146
+ "sourceRef": "https://cloud.tencent.com/document/product/1823/130051 — minimax-m3 row lists maximum input 1M tokens; k/M expanded as decimal",
69147
+ "confidence": "inferred",
69148
+ "observedAt": "2026-09-25T12:31:56.381Z"
69149
+ },
67300
69150
  "inputModalities": {
67301
69151
  "value": [
67302
69152
  "text"
@@ -67844,14 +69694,14 @@
67844
69694
  },
67845
69695
  "pricing": {
67846
69696
  "value": {
67847
- "inputPerMTokUsd": 0.149,
67848
- "outputPerMTokUsd": 0.596,
67849
- "cacheReadPerMTokUsd": 0.03725
69697
+ "inputPerMTokUsd": 0.132,
69698
+ "outputPerMTokUsd": 0.528,
69699
+ "cacheReadPerMTokUsd": 0.033
67850
69700
  },
67851
69701
  "source": "official-doc",
67852
- "sourceRef": "https://cloud.tencent.com/document/product/1823/130055 — hy3: CNY 1 input, 4 output, 0.25 cache hit per million tokens. Converted at 2026-09-25 1 CNY = $0.1490 (https://www.investing.com/currencies/cny-usd-historical-data).",
67853
- "confidence": "inferred",
67854
- "observedAt": "2026-09-25T10:06:30Z"
69702
+ "sourceRef": "https://proxy-hk.tencentcloud.com/pt/document/product/1300/78937 — Singapore international USD list, hy3: $0.132 input/$0.528 output/$0.033 cache-hit per million tokens",
69703
+ "confidence": "declared",
69704
+ "observedAt": "2026-09-25T12:31:56.381Z"
67855
69705
  },
67856
69706
  "unsupportedParameters": [],
67857
69707
  "status": "candidate",
@@ -67947,14 +69797,14 @@
67947
69797
  },
67948
69798
  "pricing": {
67949
69799
  "value": {
67950
- "inputPerMTokUsd": 0.894,
67951
- "outputPerMTokUsd": 2.682,
67952
- "cacheReadPerMTokUsd": 0.0447
69800
+ "inputPerMTokUsd": 0.834,
69801
+ "outputPerMTokUsd": 2.501,
69802
+ "cacheReadPerMTokUsd": 0.042
67953
69803
  },
67954
69804
  "source": "official-doc",
67955
- "sourceRef": "https://cloud.tencent.com/document/product/1823/130055 — hy4-preview: CNY 6 input, 18 output, 0.3 cache hit per million tokens. Converted at 2026-09-25 1 CNY = $0.1490 (https://www.investing.com/currencies/cny-usd-historical-data).",
67956
- "confidence": "inferred",
67957
- "observedAt": "2026-09-25T10:06:30Z"
69805
+ "sourceRef": "https://proxy-hk.tencentcloud.com/pt/document/product/1300/78937 — Singapore international USD list, hy4-preview: $0.834 input/$2.501 output/$0.042 cache-hit per million tokens",
69806
+ "confidence": "declared",
69807
+ "observedAt": "2026-09-25T12:31:56.381Z"
67958
69808
  },
67959
69809
  "unsupportedParameters": [],
67960
69810
  "status": "candidate",
@@ -68054,14 +69904,14 @@
68054
69904
  },
68055
69905
  "pricing": {
68056
69906
  "value": {
68057
- "inputPerMTokUsd": 0.149,
68058
- "outputPerMTokUsd": 0.596,
68059
- "cacheReadPerMTokUsd": 0.03725
69907
+ "inputPerMTokUsd": 0.132,
69908
+ "outputPerMTokUsd": 0.528,
69909
+ "cacheReadPerMTokUsd": 0.033
68060
69910
  },
68061
69911
  "source": "official-doc",
68062
- "sourceRef": "https://cloud.tencent.com/document/product/1823/130055 — hy3: CNY 1 input, 4 output, 0.25 cache hit per million tokens. Converted at 2026-09-25 1 CNY = $0.1490 (https://www.investing.com/currencies/cny-usd-historical-data).",
68063
- "confidence": "inferred",
68064
- "observedAt": "2026-09-25T10:06:30Z"
69912
+ "sourceRef": "https://proxy-hk.tencentcloud.com/pt/document/product/1300/78937 — Singapore international USD list, hy3: $0.132 input/$0.528 output/$0.033 cache-hit per million tokens",
69913
+ "confidence": "declared",
69914
+ "observedAt": "2026-09-25T12:31:56.381Z"
68065
69915
  },
68066
69916
  "unsupportedParameters": [],
68067
69917
  "status": "candidate",
@@ -68161,14 +70011,14 @@
68161
70011
  },
68162
70012
  "pricing": {
68163
70013
  "value": {
68164
- "inputPerMTokUsd": 0.894,
68165
- "outputPerMTokUsd": 2.682,
68166
- "cacheReadPerMTokUsd": 0.0447
70014
+ "inputPerMTokUsd": 0.834,
70015
+ "outputPerMTokUsd": 2.501,
70016
+ "cacheReadPerMTokUsd": 0.042
68167
70017
  },
68168
70018
  "source": "official-doc",
68169
- "sourceRef": "https://cloud.tencent.com/document/product/1823/130055 — hy4-preview: CNY 6 input, 18 output, 0.3 cache hit per million tokens. Converted at 2026-09-25 1 CNY = $0.1490 (https://www.investing.com/currencies/cny-usd-historical-data).",
68170
- "confidence": "inferred",
68171
- "observedAt": "2026-09-25T10:06:30Z"
70019
+ "sourceRef": "https://proxy-hk.tencentcloud.com/pt/document/product/1300/78937 — Singapore international USD list, hy4-preview: $0.834 input/$2.501 output/$0.042 cache-hit per million tokens",
70020
+ "confidence": "declared",
70021
+ "observedAt": "2026-09-25T12:31:56.381Z"
68172
70022
  },
68173
70023
  "unsupportedParameters": [],
68174
70024
  "status": "candidate",
@@ -72265,7 +74115,8 @@
72265
74115
  "grok-4.20-non-reasoning-gv2"
72266
74116
  ],
72267
74117
  "endpoints": [
72268
- "chat"
74118
+ "chat",
74119
+ "responses"
72269
74120
  ],
72270
74121
  "contextWindow": {
72271
74122
  "value": 1000000,
@@ -72360,7 +74211,8 @@
72360
74211
  "grok-4.20-reasoning-gv2"
72361
74212
  ],
72362
74213
  "endpoints": [
72363
- "chat"
74214
+ "chat",
74215
+ "responses"
72364
74216
  ],
72365
74217
  "contextWindow": {
72366
74218
  "value": 1000000,
@@ -72425,7 +74277,7 @@
72425
74277
  "observedAt": "2026-09-25T08:40:00Z"
72426
74278
  },
72427
74279
  "efforts": [],
72428
- "continuation": "none"
74280
+ "continuation": "opaque-provider-state"
72429
74281
  },
72430
74282
  "pricing": {
72431
74283
  "value": {
@@ -72443,6 +74295,111 @@
72443
74295
  "canonicalModelId": "grok-4.20-0309-reasoning",
72444
74296
  "modelFamily": "grok"
72445
74297
  },
74298
+ {
74299
+ "key": "xai/grok-4.20-multi-agent-0309",
74300
+ "providerId": "xai",
74301
+ "upstreamId": "grok-4.20-multi-agent-0309",
74302
+ "displayName": "Grok 4.20 Multi-Agent Beta",
74303
+ "aliases": [
74304
+ "grok-4.20-multi-agent",
74305
+ "grok-4.20-multi-agent-latest",
74306
+ "grok-4.20-multi-agent-beta-latest",
74307
+ "grok-4.20-multi-agent-experimental-beta-0304",
74308
+ "grok-4.20-multi-agent-experimental-beta-latest",
74309
+ "grok-4.20-multi-agent-beta-0309"
74310
+ ],
74311
+ "endpoints": [
74312
+ "responses"
74313
+ ],
74314
+ "contextWindow": {
74315
+ "value": 1000000,
74316
+ "source": "official-doc",
74317
+ "sourceRef": "https://docs.x.ai/developers/models/grok-4.20-multi-agent-0309 — model page context window 1,000,000 tokens",
74318
+ "confidence": "declared",
74319
+ "observedAt": "2026-09-25T08:40:00Z"
74320
+ },
74321
+ "inputModalities": {
74322
+ "value": [
74323
+ "text",
74324
+ "image"
74325
+ ],
74326
+ "source": "official-doc",
74327
+ "sourceRef": "https://docs.x.ai/developers/models/grok-4.20-multi-agent-0309 — model page Modalities: Text, Image input",
74328
+ "confidence": "declared",
74329
+ "observedAt": "2026-09-25T08:40:00Z"
74330
+ },
74331
+ "outputModalities": {
74332
+ "value": [
74333
+ "text"
74334
+ ],
74335
+ "source": "official-doc",
74336
+ "sourceRef": "https://docs.x.ai/developers/models/grok-4.20-multi-agent-0309 — model page Modalities: Text output",
74337
+ "confidence": "declared",
74338
+ "observedAt": "2026-09-25T08:40:00Z"
74339
+ },
74340
+ "toolCalling": {
74341
+ "value": "none",
74342
+ "source": "official-doc",
74343
+ "sourceRef": "https://docs.x.ai/developers/model-capabilities/text/multi-agent — multi-agent limitations: client-side/custom function calling unsupported; built-in server tools only",
74344
+ "confidence": "declared",
74345
+ "observedAt": "2026-09-25T08:40:00Z"
74346
+ },
74347
+ "nativeTools": {
74348
+ "value": false,
74349
+ "source": "official-doc",
74350
+ "sourceRef": "https://docs.x.ai/developers/model-capabilities/text/multi-agent — client-side/custom function tools unsupported",
74351
+ "confidence": "declared",
74352
+ "observedAt": "2026-09-25T08:40:00Z"
74353
+ },
74354
+ "structuredOutput": {
74355
+ "value": true,
74356
+ "source": "official-doc",
74357
+ "sourceRef": "https://docs.x.ai/developers/models/grok-4.20-multi-agent-0309 — model page Structured outputs capability",
74358
+ "confidence": "declared",
74359
+ "observedAt": "2026-09-25T08:40:00Z"
74360
+ },
74361
+ "promptCaching": {
74362
+ "value": true,
74363
+ "source": "official-doc",
74364
+ "sourceRef": "https://docs.x.ai/developers/models/grok-4.20-multi-agent-0309 — model page lists Cached tokens input rate; prompt caching available",
74365
+ "confidence": "declared",
74366
+ "observedAt": "2026-09-25T08:40:00Z"
74367
+ },
74368
+ "reasoning": {
74369
+ "supported": {
74370
+ "value": true,
74371
+ "source": "official-doc",
74372
+ "sourceRef": "https://docs.x.ai/developers/model-capabilities/text/multi-agent — reasoning.effort low/medium selects 4 agents, high/xhigh selects 16; previous_response_id supports multi-turn",
74373
+ "confidence": "declared",
74374
+ "observedAt": "2026-09-25T08:40:00Z"
74375
+ },
74376
+ "efforts": [
74377
+ "low",
74378
+ "medium",
74379
+ "high",
74380
+ "xhigh"
74381
+ ],
74382
+ "continuation": "opaque-provider-state"
74383
+ },
74384
+ "pricing": {
74385
+ "value": {
74386
+ "inputPerMTokUsd": 1.25,
74387
+ "outputPerMTokUsd": 2.5,
74388
+ "cacheReadPerMTokUsd": 0.2
74389
+ },
74390
+ "source": "official-doc",
74391
+ "sourceRef": "https://docs.x.ai/developers/pricing — grok-4.20-multi-agent-0309 Standard global short-context rate (<200k prompt tokens): $1.25 input, $0.20 cached input, $2.50 output per 1M; >=200k the entire request is charged $2.50/$0.40/$5.00. US-regional 1.1x applies only to models currently on that endpoint, listed as grok-4.7 and grok-4.6; not this model (re-read 2026-09-25, WS-23).",
74392
+ "confidence": "declared",
74393
+ "observedAt": "2026-09-25T17:09:48Z"
74394
+ },
74395
+ "unsupportedParameters": [
74396
+ "max_output_tokens",
74397
+ "tools"
74398
+ ],
74399
+ "status": "candidate",
74400
+ "canonicalModelId": "grok-4.20-multi-agent-0309",
74401
+ "modelFamily": "grok"
74402
+ },
72446
74403
  {
72447
74404
  "key": "xai/grok-4.3",
72448
74405
  "providerId": "xai",
@@ -72452,7 +74409,8 @@
72452
74409
  "grok-4.3-latest"
72453
74410
  ],
72454
74411
  "endpoints": [
72455
- "chat"
74412
+ "chat",
74413
+ "responses"
72456
74414
  ],
72457
74415
  "contextWindow": {
72458
74416
  "value": 1000000,
@@ -72523,7 +74481,7 @@
72523
74481
  "high",
72524
74482
  "xhigh"
72525
74483
  ],
72526
- "continuation": "plaintext",
74484
+ "continuation": "opaque-provider-state",
72527
74485
  "defaultEffort": "low"
72528
74486
  },
72529
74487
  "pricing": {
@@ -72533,9 +74491,9 @@
72533
74491
  "cacheReadPerMTokUsd": 0.2
72534
74492
  },
72535
74493
  "source": "official-doc",
72536
- "sourceRef": "https://docs.x.ai/developers/pricing — current global Standard Text API price table for grok-4.3; it charges the listed long-context rate for every token after the prompt reaches 200K and applies a 10% US-regional premium. This catalog records the standard short-context rate (retrieved 2026-09-19).",
74494
+ "sourceRef": "https://docs.x.ai/developers/pricing — grok-4.3 Standard global short-context rate: $1.25 input/$2.5 output/$0.2 cached per million; long-context threshold 200K. US-regional 1.1x applies only to models currently on that endpoint, listed as grok-4.7 and grok-4.6; not grok-4.3.",
72537
74495
  "confidence": "declared",
72538
- "observedAt": "2026-09-19T00:00:00Z"
74496
+ "observedAt": "2026-09-25T12:31:56.381Z"
72539
74497
  },
72540
74498
  "unsupportedParameters": [],
72541
74499
  "status": "candidate",
@@ -72552,7 +74510,8 @@
72552
74510
  "grok-build-latest"
72553
74511
  ],
72554
74512
  "endpoints": [
72555
- "chat"
74513
+ "chat",
74514
+ "responses"
72556
74515
  ],
72557
74516
  "contextWindow": {
72558
74517
  "value": 500000,
@@ -72621,7 +74580,7 @@
72621
74580
  "medium",
72622
74581
  "high"
72623
74582
  ],
72624
- "continuation": "plaintext",
74583
+ "continuation": "opaque-provider-state",
72625
74584
  "defaultEffort": "high"
72626
74585
  },
72627
74586
  "pricing": {
@@ -72651,7 +74610,8 @@
72651
74610
  "displayName": "Grok 4.6",
72652
74611
  "aliases": [],
72653
74612
  "endpoints": [
72654
- "chat"
74613
+ "chat",
74614
+ "responses"
72655
74615
  ],
72656
74616
  "contextWindow": {
72657
74617
  "value": 500000,
@@ -72721,7 +74681,7 @@
72721
74681
  "high",
72722
74682
  "xhigh"
72723
74683
  ],
72724
- "continuation": "plaintext",
74684
+ "continuation": "opaque-provider-state",
72725
74685
  "defaultEffort": "high"
72726
74686
  },
72727
74687
  "pricing": {
@@ -72751,7 +74711,8 @@
72751
74711
  "displayName": "Grok 4.7",
72752
74712
  "aliases": [],
72753
74713
  "endpoints": [
72754
- "chat"
74714
+ "chat",
74715
+ "responses"
72755
74716
  ],
72756
74717
  "contextWindow": {
72757
74718
  "value": 500000,
@@ -72821,8 +74782,29 @@
72821
74782
  "high",
72822
74783
  "xhigh"
72823
74784
  ],
72824
- "continuation": "plaintext",
72825
- "defaultEffort": "high"
74785
+ "continuation": "opaque-provider-state",
74786
+ "defaultEffort": "high",
74787
+ "readableState": {
74788
+ "value": "summary",
74789
+ "source": "official-doc",
74790
+ "sourceRef": "https://docs.x.ai/developers/model-capabilities/text/reasoning — \"For `grok-4.7`, we expose summarizations of the model's internal reasoning\" (Summarized Reasoning Content); https://docs.x.ai/developers/rest-api-reference/inference/responses — the example reasoning output item carries `summary: [{ type: \"summary_text\", … }]`",
74791
+ "confidence": "declared",
74792
+ "observedAt": "2026-09-25T17:09:48Z"
74793
+ },
74794
+ "summaryRequest": {
74795
+ "value": {
74796
+ "field": "reasoning.summary",
74797
+ "values": [
74798
+ "detailed",
74799
+ "auto",
74800
+ "concise"
74801
+ ]
74802
+ },
74803
+ "source": "official-doc",
74804
+ "sourceRef": "https://docs.x.ai/developers/rest-api-reference/inference/responses — `reasoning.summary`: \"Possible values are `auto`, `concise` and `detailed`. Only included for compatibility. The model shall always return `detailed`.\" `detailed` is listed FIRST because the adapter sends the first value, and it is the one the model returns regardless",
74805
+ "confidence": "declared",
74806
+ "observedAt": "2026-09-25T17:09:48Z"
74807
+ }
72826
74808
  },
72827
74809
  "pricing": {
72828
74810
  "value": {
@@ -72855,7 +74837,8 @@
72855
74837
  "grok-code-fast-1-0825"
72856
74838
  ],
72857
74839
  "endpoints": [
72858
- "chat"
74840
+ "chat",
74841
+ "responses"
72859
74842
  ],
72860
74843
  "contextWindow": {
72861
74844
  "value": 256000,
@@ -72920,7 +74903,7 @@
72920
74903
  "observedAt": "2026-09-25T08:40:00Z"
72921
74904
  },
72922
74905
  "efforts": [],
72923
- "continuation": "none"
74906
+ "continuation": "opaque-provider-state"
72924
74907
  },
72925
74908
  "pricing": {
72926
74909
  "value": {
@@ -72929,9 +74912,9 @@
72929
74912
  "cacheReadPerMTokUsd": 0.2
72930
74913
  },
72931
74914
  "source": "official-doc",
72932
- "sourceRef": "https://docs.x.ai/developers/pricing — current global Standard Text API price table for grok-build-0.1; it charges the listed long-context rate for every token after the prompt reaches 200K and applies a 10% US-regional premium. This catalog records the standard short-context rate (retrieved 2026-09-19).",
74915
+ "sourceRef": "https://docs.x.ai/developers/pricing — grok-build-0.1 Standard global short-context rate: $1 input/$2 output/$0.2 cached per million; long-context threshold 200K. US-regional 1.1x applies only to models currently on that endpoint, listed as grok-4.7 and grok-4.6; not grok-build-0.1.",
72933
74916
  "confidence": "declared",
72934
- "observedAt": "2026-09-19T00:00:00Z"
74917
+ "observedAt": "2026-09-25T12:31:56.381Z"
72935
74918
  },
72936
74919
  "unsupportedParameters": [],
72937
74920
  "status": "candidate",
@@ -73033,7 +75016,7 @@
73033
75016
  "observedAt": "2026-09-25T09:52:37Z"
73034
75017
  },
73035
75018
  "unsupportedParameters": [],
73036
- "status": "deprecated",
75019
+ "status": "candidate",
73037
75020
  "canonicalModelId": "mimo-v2.5",
73038
75021
  "modelFamily": "mimo"
73039
75022
  },
@@ -73129,7 +75112,7 @@
73129
75112
  "observedAt": "2026-09-25T09:52:37Z"
73130
75113
  },
73131
75114
  "unsupportedParameters": [],
73132
- "status": "deprecated",
75115
+ "status": "candidate",
73133
75116
  "canonicalModelId": "mimo-v2.5-pro",
73134
75117
  "modelFamily": "mimo"
73135
75118
  },
@@ -73514,7 +75497,7 @@
73514
75497
  "continuation": "none"
73515
75498
  },
73516
75499
  "unsupportedParameters": [],
73517
- "status": "deprecated",
75500
+ "status": "candidate",
73518
75501
  "canonicalModelId": "mimo-v2.5",
73519
75502
  "modelFamily": "mimo"
73520
75503
  },
@@ -73599,7 +75582,7 @@
73599
75582
  "continuation": "none"
73600
75583
  },
73601
75584
  "unsupportedParameters": [],
73602
- "status": "deprecated",
75585
+ "status": "candidate",
73603
75586
  "canonicalModelId": "mimo-v2.5-pro",
73604
75587
  "modelFamily": "mimo"
73605
75588
  },
@@ -73863,7 +75846,7 @@
73863
75846
  "continuation": "none"
73864
75847
  },
73865
75848
  "unsupportedParameters": [],
73866
- "status": "deprecated",
75849
+ "status": "candidate",
73867
75850
  "canonicalModelId": "mimo-v2.5",
73868
75851
  "modelFamily": "mimo"
73869
75852
  },
@@ -73948,7 +75931,7 @@
73948
75931
  "continuation": "none"
73949
75932
  },
73950
75933
  "unsupportedParameters": [],
73951
- "status": "deprecated",
75934
+ "status": "candidate",
73952
75935
  "canonicalModelId": "mimo-v2.5-pro",
73953
75936
  "modelFamily": "mimo"
73954
75937
  },
@@ -74212,7 +76195,7 @@
74212
76195
  "continuation": "plaintext"
74213
76196
  },
74214
76197
  "unsupportedParameters": [],
74215
- "status": "deprecated",
76198
+ "status": "candidate",
74216
76199
  "canonicalModelId": "mimo-v2.5",
74217
76200
  "modelFamily": "mimo"
74218
76201
  },
@@ -74297,7 +76280,7 @@
74297
76280
  "continuation": "plaintext"
74298
76281
  },
74299
76282
  "unsupportedParameters": [],
74300
- "status": "deprecated",
76283
+ "status": "candidate",
74301
76284
  "canonicalModelId": "mimo-v2.5-pro",
74302
76285
  "modelFamily": "mimo"
74303
76286
  },
@@ -74561,7 +76544,7 @@
74561
76544
  "continuation": "none"
74562
76545
  },
74563
76546
  "unsupportedParameters": [],
74564
- "status": "deprecated",
76547
+ "status": "candidate",
74565
76548
  "canonicalModelId": "mimo-v2.5",
74566
76549
  "modelFamily": "mimo"
74567
76550
  },
@@ -74646,7 +76629,7 @@
74646
76629
  "continuation": "none"
74647
76630
  },
74648
76631
  "unsupportedParameters": [],
74649
- "status": "deprecated",
76632
+ "status": "candidate",
74650
76633
  "canonicalModelId": "mimo-v2.5-pro",
74651
76634
  "modelFamily": "mimo"
74652
76635
  },
@@ -74910,7 +76893,7 @@
74910
76893
  "continuation": "plaintext"
74911
76894
  },
74912
76895
  "unsupportedParameters": [],
74913
- "status": "deprecated",
76896
+ "status": "candidate",
74914
76897
  "canonicalModelId": "mimo-v2.5",
74915
76898
  "modelFamily": "mimo"
74916
76899
  },
@@ -74995,7 +76978,7 @@
74995
76978
  "continuation": "plaintext"
74996
76979
  },
74997
76980
  "unsupportedParameters": [],
74998
- "status": "deprecated",
76981
+ "status": "candidate",
74999
76982
  "canonicalModelId": "mimo-v2.5-pro",
75000
76983
  "modelFamily": "mimo"
75001
76984
  },
@@ -75259,7 +77242,7 @@
75259
77242
  "continuation": "plaintext"
75260
77243
  },
75261
77244
  "unsupportedParameters": [],
75262
- "status": "deprecated",
77245
+ "status": "candidate",
75263
77246
  "canonicalModelId": "mimo-v2.5",
75264
77247
  "modelFamily": "mimo"
75265
77248
  },
@@ -75344,7 +77327,7 @@
75344
77327
  "continuation": "plaintext"
75345
77328
  },
75346
77329
  "unsupportedParameters": [],
75347
- "status": "deprecated",
77330
+ "status": "candidate",
75348
77331
  "canonicalModelId": "mimo-v2.5-pro",
75349
77332
  "modelFamily": "mimo"
75350
77333
  },
@@ -75619,7 +77602,7 @@
75619
77602
  "observedAt": "2026-09-25T09:52:37Z"
75620
77603
  },
75621
77604
  "unsupportedParameters": [],
75622
- "status": "deprecated",
77605
+ "status": "candidate",
75623
77606
  "canonicalModelId": "mimo-v2.5",
75624
77607
  "modelFamily": "mimo"
75625
77608
  },
@@ -75715,7 +77698,7 @@
75715
77698
  "observedAt": "2026-09-25T09:52:37Z"
75716
77699
  },
75717
77700
  "unsupportedParameters": [],
75718
- "status": "deprecated",
77701
+ "status": "candidate",
75719
77702
  "canonicalModelId": "mimo-v2.5-pro",
75720
77703
  "modelFamily": "mimo"
75721
77704
  },
@@ -76983,6 +78966,13 @@
76983
78966
  "endpoints": [
76984
78967
  "chat"
76985
78968
  ],
78969
+ "contextWindow": {
78970
+ "value": 128000,
78971
+ "source": "official-doc",
78972
+ "sourceRef": "https://docs.z.ai/guides/llm/glm-4.5 — glm-4.5 model guide/table gives 128K context length; K expanded as 1,000 tokens",
78973
+ "confidence": "inferred",
78974
+ "observedAt": "2026-09-25T12:31:56.381Z"
78975
+ },
76986
78976
  "maxOutputTokens": {
76987
78977
  "value": 98304,
76988
78978
  "source": "official-doc",
@@ -77065,6 +79055,13 @@
77065
79055
  "endpoints": [
77066
79056
  "chat"
77067
79057
  ],
79058
+ "contextWindow": {
79059
+ "value": 128000,
79060
+ "source": "official-doc",
79061
+ "sourceRef": "https://docs.z.ai/llms-full.txt — glm-4.5-air model guide/table gives 128K context length; K expanded as 1,000 tokens",
79062
+ "confidence": "inferred",
79063
+ "observedAt": "2026-09-25T12:31:56.381Z"
79064
+ },
77068
79065
  "maxOutputTokens": {
77069
79066
  "value": 98304,
77070
79067
  "source": "official-doc",
@@ -77147,6 +79144,13 @@
77147
79144
  "endpoints": [
77148
79145
  "chat"
77149
79146
  ],
79147
+ "contextWindow": {
79148
+ "value": 200000,
79149
+ "source": "official-doc",
79150
+ "sourceRef": "https://docs.z.ai/guides/llm/glm-4.6 — glm-4.6 model guide/table gives 200K context length; K expanded as 1,000 tokens",
79151
+ "confidence": "inferred",
79152
+ "observedAt": "2026-09-25T12:31:56.381Z"
79153
+ },
77150
79154
  "maxOutputTokens": {
77151
79155
  "value": 131072,
77152
79156
  "source": "official-doc",
@@ -77229,6 +79233,13 @@
77229
79233
  "endpoints": [
77230
79234
  "chat"
77231
79235
  ],
79236
+ "contextWindow": {
79237
+ "value": 200000,
79238
+ "source": "official-doc",
79239
+ "sourceRef": "https://docs.z.ai/guides/llm/glm-4.7 — glm-4.7 model guide/table gives 200K context length; K expanded as 1,000 tokens",
79240
+ "confidence": "inferred",
79241
+ "observedAt": "2026-09-25T12:31:56.381Z"
79242
+ },
77232
79243
  "maxOutputTokens": {
77233
79244
  "value": 131072,
77234
79245
  "source": "official-doc",
@@ -77332,6 +79343,13 @@
77332
79343
  "endpoints": [
77333
79344
  "chat"
77334
79345
  ],
79346
+ "contextWindow": {
79347
+ "value": 200000,
79348
+ "source": "official-doc",
79349
+ "sourceRef": "https://docs.z.ai/llms-full.txt — glm-4.7-flash model guide/table gives 200K context length; K expanded as 1,000 tokens",
79350
+ "confidence": "inferred",
79351
+ "observedAt": "2026-09-25T12:31:56.381Z"
79352
+ },
77335
79353
  "inputModalities": {
77336
79354
  "value": [
77337
79355
  "text"
@@ -77421,6 +79439,13 @@
77421
79439
  "endpoints": [
77422
79440
  "chat"
77423
79441
  ],
79442
+ "contextWindow": {
79443
+ "value": 200000,
79444
+ "source": "official-doc",
79445
+ "sourceRef": "https://docs.z.ai/llms-full.txt — glm-4.7-flashx model guide/table gives 200K context length; K expanded as 1,000 tokens",
79446
+ "confidence": "inferred",
79447
+ "observedAt": "2026-09-25T12:31:56.381Z"
79448
+ },
77424
79449
  "inputModalities": {
77425
79450
  "value": [
77426
79451
  "text"
@@ -77496,6 +79521,13 @@
77496
79521
  "endpoints": [
77497
79522
  "chat"
77498
79523
  ],
79524
+ "contextWindow": {
79525
+ "value": 200000,
79526
+ "source": "official-doc",
79527
+ "sourceRef": "https://docs.z.ai/guides/llm/glm-5 — glm-5 model guide/table gives 200K context length; K expanded as 1,000 tokens",
79528
+ "confidence": "inferred",
79529
+ "observedAt": "2026-09-25T12:31:56.381Z"
79530
+ },
77499
79531
  "maxOutputTokens": {
77500
79532
  "value": 131072,
77501
79533
  "source": "official-doc",
@@ -77599,6 +79631,13 @@
77599
79631
  "endpoints": [
77600
79632
  "chat"
77601
79633
  ],
79634
+ "contextWindow": {
79635
+ "value": 200000,
79636
+ "source": "official-doc",
79637
+ "sourceRef": "https://docs.z.ai/guides/llm/glm-5-turbo — glm-5-turbo model guide/table gives 200K context length; K expanded as 1,000 tokens",
79638
+ "confidence": "inferred",
79639
+ "observedAt": "2026-09-25T12:31:56.381Z"
79640
+ },
77602
79641
  "inputModalities": {
77603
79642
  "value": [
77604
79643
  "text"
@@ -77677,6 +79716,13 @@
77677
79716
  "endpoints": [
77678
79717
  "chat"
77679
79718
  ],
79719
+ "contextWindow": {
79720
+ "value": 200000,
79721
+ "source": "official-doc",
79722
+ "sourceRef": "https://docs.z.ai/guides/llm/glm-5.1 — glm-5.1 model guide/table gives 200K context length; K expanded as 1,000 tokens",
79723
+ "confidence": "inferred",
79724
+ "observedAt": "2026-09-25T12:31:56.381Z"
79725
+ },
77680
79726
  "maxOutputTokens": {
77681
79727
  "value": 131072,
77682
79728
  "source": "official-doc",