@molecule/api-resource-ai-models 1.10.0 → 1.12.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +23 -6
- package/dist/models.d.ts +22 -5
- package/dist/models.d.ts.map +1 -1
- package/dist/models.js +157 -25
- package/package.json +1 -1
package/README.md
CHANGED
|
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
|
|
|
3
3
|
Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
|
|
4
4
|
Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
|
|
5
5
|
To change this document, edit the module-level JSDoc in src/index.ts.
|
|
6
|
-
Generated: 2026-09-
|
|
6
|
+
Generated: 2026-09-30T04:56:41.977Z
|
|
7
7
|
-->
|
|
8
8
|
|
|
9
9
|
# @molecule/api-resource-ai-models
|
|
@@ -931,6 +931,14 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
931
931
|
> /v1/chat/completions with reasoning_effort 'none'. The 5.6 family is now
|
|
932
932
|
> superseded; every 5.6 price on the page is unchanged and the -sol promo
|
|
933
933
|
> footnote still reads "at least through November 21, 2026".)
|
|
934
|
+
> (re-verified 2026-09-30 on the pricing page + the gpt-6.1-sol model page:
|
|
935
|
+
> ADDED gpt-6.1-sol ($2/$10, cached $0.10 — 5% of input, not 10% — cache
|
|
936
|
+
writes $2.50; >272K band $4/$15, cached $0.20, writes $5 not modeled;
|
|
937
|
+
> 1.05M ctx / 128K out; image input; knowledge cutoff Apr 30 2026; effort
|
|
938
|
+
> low|medium|high|xhigh|max, no none/minimal; tools on /v1/responses only).
|
|
939
|
+
> It supersedes gpt-6-sol. Every other GPT-6 and 5.6 price on the page is
|
|
940
|
+
> unchanged and the 5.6-sol promo footnote still reads "at least through
|
|
941
|
+
> November 21, 2026".)
|
|
934
942
|
- Google: https://ai.google.dev/gemini-api/docs/pricing (gemini-3.6-flash GA
|
|
935
943
|
2026-07-21 $1.50/$7.50 supersedes 3.5-flash as the agentic flagship;
|
|
936
944
|
gemini-3.1-pro-preview still the pro tier — "3.5 Pro" has NOT shipped as
|
|
@@ -958,11 +966,20 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
958
966
|
low|medium|high, vision, tools, caching, search grounding, code execution,
|
|
959
967
|
url context — identical surface to 3.7-flash. Knowledge cutoff is published
|
|
960
968
|
by neither Google nor models.dev for this id)
|
|
961
|
-
- xAI: https://docs.x.ai/developers/models + /developers/grok-4-
|
|
962
|
-
(
|
|
963
|
-
|
|
964
|
-
|
|
965
|
-
|
|
969
|
+
- xAI: https://docs.x.ai/developers/models + /developers/grok-4-7 +
|
|
970
|
+
/developers/model-capabilities/text/reasoning (re-read 2026-09-30:
|
|
971
|
+
grok-4.7 flagship: $2/$6, cached $0.50, 500K ctx, ≥200K prompts bill 2× —
|
|
972
|
+
not modeled; reasoning_effort low|medium|high|xhigh default high, image
|
|
973
|
+
input, knowledge cutoff May 2026, "No text output limit"; live-probed on
|
|
974
|
+
/v1/chat/completions with tools, forced + required tool_choice,
|
|
975
|
+
temperature 0, xhigh and an image — all 200, effort 'none' 400s.
|
|
976
|
+
grok-4.6 (re-read 2026-09-30 from the models page's model list — its own
|
|
977
|
+
page now redirects to grok-4-7; specs from that page as archived
|
|
978
|
+
2026-09-10): $2/$6, cached $0.50, 500K ctx, image input, reasoning_effort
|
|
979
|
+
low|medium|high|xhigh default high, knowledge cutoff February 1, 2026 —
|
|
980
|
+
cataloged superseded by grok-4.7.
|
|
981
|
+
grok-4.5 still served at $2/$6, cached $0.30; grok-4.3 still served at
|
|
982
|
+
$1.25/$2.50 with the bigger 1M window; grok-code-fast-1 no longer listed)
|
|
966
983
|
- Chinese public holidays (DeepSeek's peak excludes them):
|
|
967
984
|
https://www.12371.gov.cn/web/article/web/content_1451614684968525824.html
|
|
968
985
|
(the State Council's 2026 notice; see CHINA_PUBLIC_HOLIDAYS)
|
package/dist/models.d.ts
CHANGED
|
@@ -147,6 +147,14 @@ import type { ModelDefinition } from './types.js';
|
|
|
147
147
|
* /v1/chat/completions with reasoning_effort 'none'. The 5.6 family is now
|
|
148
148
|
* superseded; every 5.6 price on the page is unchanged and the -sol promo
|
|
149
149
|
* footnote still reads "at least through November 21, 2026".)
|
|
150
|
+
* (re-verified 2026-09-30 on the pricing page + the gpt-6.1-sol model page:
|
|
151
|
+
* ADDED gpt-6.1-sol ($2/$10, cached $0.10 — 5% of input, not 10% — cache
|
|
152
|
+
* writes $2.50; >272K band $4/$15, cached $0.20, writes $5 not modeled;
|
|
153
|
+
* 1.05M ctx / 128K out; image input; knowledge cutoff Apr 30 2026; effort
|
|
154
|
+
* low|medium|high|xhigh|max, no none/minimal; tools on /v1/responses only).
|
|
155
|
+
* It supersedes gpt-6-sol. Every other GPT-6 and 5.6 price on the page is
|
|
156
|
+
* unchanged and the 5.6-sol promo footnote still reads "at least through
|
|
157
|
+
* November 21, 2026".)
|
|
150
158
|
* - Google: https://ai.google.dev/gemini-api/docs/pricing (gemini-3.6-flash GA
|
|
151
159
|
* 2026-07-21 $1.50/$7.50 supersedes 3.5-flash as the agentic flagship;
|
|
152
160
|
* gemini-3.1-pro-preview still the pro tier — "3.5 Pro" has NOT shipped as
|
|
@@ -174,11 +182,20 @@ import type { ModelDefinition } from './types.js';
|
|
|
174
182
|
* low|medium|high, vision, tools, caching, search grounding, code execution,
|
|
175
183
|
* url context — identical surface to 3.7-flash. Knowledge cutoff is published
|
|
176
184
|
* by neither Google nor models.dev for this id)
|
|
177
|
-
* - xAI: https://docs.x.ai/developers/models + /developers/grok-4-
|
|
178
|
-
* (
|
|
179
|
-
*
|
|
180
|
-
*
|
|
181
|
-
*
|
|
185
|
+
* - xAI: https://docs.x.ai/developers/models + /developers/grok-4-7 +
|
|
186
|
+
* /developers/model-capabilities/text/reasoning (re-read 2026-09-30:
|
|
187
|
+
* grok-4.7 flagship: $2/$6, cached $0.50, 500K ctx, ≥200K prompts bill 2× —
|
|
188
|
+
* not modeled; reasoning_effort low|medium|high|xhigh default high, image
|
|
189
|
+
* input, knowledge cutoff May 2026, "No text output limit"; live-probed on
|
|
190
|
+
* /v1/chat/completions with tools, forced + required tool_choice,
|
|
191
|
+
* temperature 0, xhigh and an image — all 200, effort 'none' 400s.
|
|
192
|
+
* grok-4.6 (re-read 2026-09-30 from the models page's model list — its own
|
|
193
|
+
* page now redirects to grok-4-7; specs from that page as archived
|
|
194
|
+
* 2026-09-10): $2/$6, cached $0.50, 500K ctx, image input, reasoning_effort
|
|
195
|
+
* low|medium|high|xhigh default high, knowledge cutoff February 1, 2026 —
|
|
196
|
+
* cataloged superseded by grok-4.7.
|
|
197
|
+
* grok-4.5 still served at $2/$6, cached $0.30; grok-4.3 still served at
|
|
198
|
+
* $1.25/$2.50 with the bigger 1M window; grok-code-fast-1 no longer listed)
|
|
182
199
|
* - Chinese public holidays (DeepSeek's peak excludes them):
|
|
183
200
|
* https://www.12371.gov.cn/web/article/web/content_1451614684968525824.html
|
|
184
201
|
* (the State Council's 2026 notice; see CHINA_PUBLIC_HOLIDAYS)
|
package/dist/models.d.ts.map
CHANGED
|
@@ -1 +1 @@
|
|
|
1
|
-
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAwDjD
|
|
1
|
+
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAwDjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAsSG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAgxEnC,CAAA"}
|
package/dist/models.js
CHANGED
|
@@ -196,6 +196,14 @@ const CHINA_PUBLIC_HOLIDAYS = [
|
|
|
196
196
|
* /v1/chat/completions with reasoning_effort 'none'. The 5.6 family is now
|
|
197
197
|
* superseded; every 5.6 price on the page is unchanged and the -sol promo
|
|
198
198
|
* footnote still reads "at least through November 21, 2026".)
|
|
199
|
+
* (re-verified 2026-09-30 on the pricing page + the gpt-6.1-sol model page:
|
|
200
|
+
* ADDED gpt-6.1-sol ($2/$10, cached $0.10 — 5% of input, not 10% — cache
|
|
201
|
+
* writes $2.50; >272K band $4/$15, cached $0.20, writes $5 not modeled;
|
|
202
|
+
* 1.05M ctx / 128K out; image input; knowledge cutoff Apr 30 2026; effort
|
|
203
|
+
* low|medium|high|xhigh|max, no none/minimal; tools on /v1/responses only).
|
|
204
|
+
* It supersedes gpt-6-sol. Every other GPT-6 and 5.6 price on the page is
|
|
205
|
+
* unchanged and the 5.6-sol promo footnote still reads "at least through
|
|
206
|
+
* November 21, 2026".)
|
|
199
207
|
* - Google: https://ai.google.dev/gemini-api/docs/pricing (gemini-3.6-flash GA
|
|
200
208
|
* 2026-07-21 $1.50/$7.50 supersedes 3.5-flash as the agentic flagship;
|
|
201
209
|
* gemini-3.1-pro-preview still the pro tier — "3.5 Pro" has NOT shipped as
|
|
@@ -223,11 +231,20 @@ const CHINA_PUBLIC_HOLIDAYS = [
|
|
|
223
231
|
* low|medium|high, vision, tools, caching, search grounding, code execution,
|
|
224
232
|
* url context — identical surface to 3.7-flash. Knowledge cutoff is published
|
|
225
233
|
* by neither Google nor models.dev for this id)
|
|
226
|
-
* - xAI: https://docs.x.ai/developers/models + /developers/grok-4-
|
|
227
|
-
* (
|
|
228
|
-
*
|
|
229
|
-
*
|
|
230
|
-
*
|
|
234
|
+
* - xAI: https://docs.x.ai/developers/models + /developers/grok-4-7 +
|
|
235
|
+
* /developers/model-capabilities/text/reasoning (re-read 2026-09-30:
|
|
236
|
+
* grok-4.7 flagship: $2/$6, cached $0.50, 500K ctx, ≥200K prompts bill 2× —
|
|
237
|
+
* not modeled; reasoning_effort low|medium|high|xhigh default high, image
|
|
238
|
+
* input, knowledge cutoff May 2026, "No text output limit"; live-probed on
|
|
239
|
+
* /v1/chat/completions with tools, forced + required tool_choice,
|
|
240
|
+
* temperature 0, xhigh and an image — all 200, effort 'none' 400s.
|
|
241
|
+
* grok-4.6 (re-read 2026-09-30 from the models page's model list — its own
|
|
242
|
+
* page now redirects to grok-4-7; specs from that page as archived
|
|
243
|
+
* 2026-09-10): $2/$6, cached $0.50, 500K ctx, image input, reasoning_effort
|
|
244
|
+
* low|medium|high|xhigh default high, knowledge cutoff February 1, 2026 —
|
|
245
|
+
* cataloged superseded by grok-4.7.
|
|
246
|
+
* grok-4.5 still served at $2/$6, cached $0.30; grok-4.3 still served at
|
|
247
|
+
* $1.25/$2.50 with the bigger 1M window; grok-code-fast-1 no longer listed)
|
|
231
248
|
* - Chinese public holidays (DeepSeek's peak excludes them):
|
|
232
249
|
* https://www.12371.gov.cn/web/article/web/content_1451614684968525824.html
|
|
233
250
|
* (the State Council's 2026 notice; see CHINA_PUBLIC_HOLIDAYS)
|
|
@@ -842,6 +859,13 @@ export const MODELS = [
|
|
|
842
859
|
// again despite the docs pages listing it. GPT-6 has no Terra tier: sol at
|
|
843
860
|
// $2/$10 is priced at 5.6-terra's tier and undercuts 5.6-sol, so it supersedes
|
|
844
861
|
// both; luna supersedes 5.6-luna at half the price.
|
|
862
|
+
//
|
|
863
|
+
// gpt-6.1-sol (released 2026-09-29) IS astra's case: its model page lists no
|
|
864
|
+
// 'none'/'minimal' effort and says "Use the Responses API for tool calling.
|
|
865
|
+
// Chat Completions is supported without tool calling." The bond calls
|
|
866
|
+
// /v1/responses on OpenAI's own endpoint, so, like astra, it needs no
|
|
867
|
+
// `toolsRequireReasoningOff` pin. Same list price as gpt-6-sol with half the
|
|
868
|
+
// cached-input rate, so it supersedes gpt-6-sol.
|
|
845
869
|
// ---------------------------------------------------------------------------
|
|
846
870
|
{
|
|
847
871
|
id: 'gpt-6-astra',
|
|
@@ -876,6 +900,37 @@ export const MODELS = [
|
|
|
876
900
|
cacheWritePricePerMTok: 12.5,
|
|
877
901
|
knowledgeCutoff: '2026-04-30',
|
|
878
902
|
},
|
|
903
|
+
{
|
|
904
|
+
id: 'gpt-6.1-sol',
|
|
905
|
+
provider: 'openai',
|
|
906
|
+
label: 'GPT-6.1 Sol',
|
|
907
|
+
description: 'OpenAI for complex coding & agentic work',
|
|
908
|
+
// Documented as 1.05M; floored to 1M like the other OpenAI entries.
|
|
909
|
+
contextWindow: 1_000_000,
|
|
910
|
+
maxOutputTokens: 128_000,
|
|
911
|
+
supportsThinking: true,
|
|
912
|
+
thinkingBudgetTokens: 16_000,
|
|
913
|
+
thinkingConfigurable: true,
|
|
914
|
+
// No 'none' or 'minimal' on this model (model page, 2026-09-30).
|
|
915
|
+
supportedEffortLevels: ['low', 'medium', 'high', 'xhigh', 'max'],
|
|
916
|
+
defaultEffortLevel: 'medium',
|
|
917
|
+
// Omitted like the rest of the reasoning family; the pre-commit dispatch
|
|
918
|
+
// probe sends the temperature-0 shape.
|
|
919
|
+
rejectsTemperature: true,
|
|
920
|
+
supportsVision: true,
|
|
921
|
+
supportsPromptCaching: true,
|
|
922
|
+
supportsTools: true,
|
|
923
|
+
webSearchToolType: 'web_search',
|
|
924
|
+
webSearchPricePer1k: 10,
|
|
925
|
+
codeExecutionToolType: 'code_interpreter',
|
|
926
|
+
// Standard tier. Cached input is 5% of input (not the 10% of GPT-6);
|
|
927
|
+
// >272K band ($4/$15, cached $0.20, writes $5) not modeled.
|
|
928
|
+
inputPricePerMTok: 2,
|
|
929
|
+
outputPricePerMTok: 10,
|
|
930
|
+
cacheReadPricePerMTok: 0.1,
|
|
931
|
+
cacheWritePricePerMTok: 2.5,
|
|
932
|
+
knowledgeCutoff: '2026-04-30',
|
|
933
|
+
},
|
|
879
934
|
{
|
|
880
935
|
id: 'gpt-6-sol',
|
|
881
936
|
provider: 'openai',
|
|
@@ -909,6 +964,10 @@ export const MODELS = [
|
|
|
909
964
|
cacheReadPricePerMTok: 0.2,
|
|
910
965
|
cacheWritePricePerMTok: 2.5,
|
|
911
966
|
knowledgeCutoff: '2026-04-20',
|
|
967
|
+
// Superseded by gpt-6.1-sol (same $2/$10, cached input $0.10 vs $0.20).
|
|
968
|
+
// Still served and priceable — just not offered in the picker.
|
|
969
|
+
deprecatedAt: '2026-09-30',
|
|
970
|
+
supersededBy: 'gpt-6.1-sol',
|
|
912
971
|
},
|
|
913
972
|
{
|
|
914
973
|
id: 'gpt-6-luna',
|
|
@@ -981,10 +1040,11 @@ export const MODELS = [
|
|
|
981
1040
|
cacheWritePricePerMTok: 6.25,
|
|
982
1041
|
// Not published — best-effort estimate.
|
|
983
1042
|
knowledgeCutoff: '2026-03-01',
|
|
984
|
-
// Superseded by gpt-6-sol ($2/$10 vs $5/$30 list)
|
|
1043
|
+
// Superseded by gpt-6-sol ($2/$10 vs $5/$30 list); points one hop to
|
|
1044
|
+
// gpt-6.1-sol since 2026-09-30, when gpt-6-sol was superseded. Still served and
|
|
985
1045
|
// priceable — just not offered in the picker.
|
|
986
1046
|
deprecatedAt: '2026-09-23',
|
|
987
|
-
supersededBy: 'gpt-6-sol',
|
|
1047
|
+
supersededBy: 'gpt-6.1-sol',
|
|
988
1048
|
},
|
|
989
1049
|
{
|
|
990
1050
|
id: 'gpt-5.6-terra',
|
|
@@ -1019,9 +1079,10 @@ export const MODELS = [
|
|
|
1019
1079
|
// Not published — best-effort estimate.
|
|
1020
1080
|
knowledgeCutoff: '2026-03-01',
|
|
1021
1081
|
// Superseded by gpt-6-sol: GPT-6 has no Terra tier, and sol sits at this
|
|
1022
|
-
// tier's price ($2/$10 vs $2/$12).
|
|
1082
|
+
// tier's price ($2/$10 vs $2/$12). Points one hop to gpt-6.1-sol since
|
|
1083
|
+
// 2026-09-30.
|
|
1023
1084
|
deprecatedAt: '2026-09-23',
|
|
1024
|
-
supersededBy: 'gpt-6-sol',
|
|
1085
|
+
supersededBy: 'gpt-6.1-sol',
|
|
1025
1086
|
},
|
|
1026
1087
|
{
|
|
1027
1088
|
id: 'gpt-5.6-luna',
|
|
@@ -1097,10 +1158,11 @@ export const MODELS = [
|
|
|
1097
1158
|
cacheWritePricePerMTok: 5,
|
|
1098
1159
|
knowledgeCutoff: '2025-12-01',
|
|
1099
1160
|
// Superseded by gpt-5.6-sol (same frontier tier, same $5/$30), and since
|
|
1100
|
-
// 2026-09-23 points one hop to gpt-6-sol, 5.6-sol's own successor.
|
|
1161
|
+
// 2026-09-23 points one hop to gpt-6-sol, 5.6-sol's own successor (gpt-6.1-sol
|
|
1162
|
+
// since 2026-09-30). Still
|
|
1101
1163
|
// listed as current by OpenAI, so it stays priceable — just not offered.
|
|
1102
1164
|
deprecatedAt: '2026-07-09',
|
|
1103
|
-
supersededBy: 'gpt-6-sol',
|
|
1165
|
+
supersededBy: 'gpt-6.1-sol',
|
|
1104
1166
|
},
|
|
1105
1167
|
{
|
|
1106
1168
|
id: 'gpt-5.4',
|
|
@@ -1133,9 +1195,10 @@ export const MODELS = [
|
|
|
1133
1195
|
// balanced tier for LESS ($2/$12 vs $2.50/$15) — superseded, so the picker
|
|
1134
1196
|
// offers only the 5.6 generation (this is OUR taxonomy, not OpenAI's
|
|
1135
1197
|
// deprecations page; the model stays priceable). Points one hop to
|
|
1136
|
-
// gpt-6-sol since 2026-09-23, when 5.6-terra was itself superseded
|
|
1198
|
+
// gpt-6-sol since 2026-09-23, when 5.6-terra was itself superseded, and to
|
|
1199
|
+
// gpt-6.1-sol since 2026-09-30.
|
|
1137
1200
|
deprecatedAt: '2026-07-28',
|
|
1138
|
-
supersededBy: 'gpt-6-sol',
|
|
1201
|
+
supersededBy: 'gpt-6.1-sol',
|
|
1139
1202
|
},
|
|
1140
1203
|
{
|
|
1141
1204
|
id: 'gpt-5.4-mini',
|
|
@@ -1381,13 +1444,78 @@ export const MODELS = [
|
|
|
1381
1444
|
},
|
|
1382
1445
|
// ---------------------------------------------------------------------------
|
|
1383
1446
|
// xAI (Grok)
|
|
1384
|
-
// Verified: https://docs.x.ai/developers/models + /developers/grok-4-
|
|
1385
|
-
// (2026-
|
|
1386
|
-
// grok-4.
|
|
1387
|
-
//
|
|
1388
|
-
//
|
|
1389
|
-
//
|
|
1447
|
+
// Verified: https://docs.x.ai/developers/models + /developers/grok-4-7 +
|
|
1448
|
+
// /developers/model-capabilities/text/reasoning (2026-09-30)
|
|
1449
|
+
// grok-4.7 is the flagship: 500K ctx, reasoning_effort low|medium|high|xhigh
|
|
1450
|
+
// default high ("xhigh is available on grok-4.6 and later"), image input,
|
|
1451
|
+
// knowledge cutoff May 2026. grok-4.5 is still served but superseded.
|
|
1452
|
+
// grok-4.3 stays served with the BIGGER 1M window (reasoning_effort
|
|
1453
|
+
// none|low|medium|high, default low). Max output tokens still not documented
|
|
1454
|
+
// for any model (grok-4.7's page says "No text output limit").
|
|
1390
1455
|
// ---------------------------------------------------------------------------
|
|
1456
|
+
{
|
|
1457
|
+
id: 'grok-4.7',
|
|
1458
|
+
provider: 'xai',
|
|
1459
|
+
label: 'Grok 4.7',
|
|
1460
|
+
description: 'xAI frontier — coding, agentic tasks & knowledge work',
|
|
1461
|
+
contextWindow: 500_000,
|
|
1462
|
+
// "No text output limit" per docs.x.ai — conservative cap.
|
|
1463
|
+
maxOutputTokens: 128_000,
|
|
1464
|
+
supportsThinking: true,
|
|
1465
|
+
thinkingBudgetTokens: 16_000,
|
|
1466
|
+
thinkingConfigurable: true,
|
|
1467
|
+
// reasoning_effort low|medium|high|xhigh, default high (docs.x.ai/developers/
|
|
1468
|
+
// grok-4-7). No 'none' tier — the live endpoint 400s on it (2026-09-30).
|
|
1469
|
+
supportedEffortLevels: ['low', 'medium', 'high', 'xhigh'],
|
|
1470
|
+
defaultEffortLevel: 'high',
|
|
1471
|
+
// Multimodal input (text + images), text out.
|
|
1472
|
+
supportsVision: true,
|
|
1473
|
+
supportsPromptCaching: true,
|
|
1474
|
+
supportsTools: true,
|
|
1475
|
+
// <200K-token prompts; ≥200K bills 2× ($4/$1.00/$12 — tiering not
|
|
1476
|
+
// modeled, same as grok-4.5 and the Gemini 3.1 Pro >200K tier).
|
|
1477
|
+
inputPricePerMTok: 2,
|
|
1478
|
+
outputPricePerMTok: 6,
|
|
1479
|
+
// xAI cached input billed at a flat $0.50/M (up from 4.5's $0.30), no
|
|
1480
|
+
// write premium.
|
|
1481
|
+
cacheReadPricePerMTok: 0.5,
|
|
1482
|
+
cacheWritePricePerMTok: 2,
|
|
1483
|
+
// Official (docs.x.ai): May 2026.
|
|
1484
|
+
knowledgeCutoff: '2026-05-01',
|
|
1485
|
+
},
|
|
1486
|
+
{
|
|
1487
|
+
id: 'grok-4.6',
|
|
1488
|
+
provider: 'xai',
|
|
1489
|
+
label: 'Grok 4.6',
|
|
1490
|
+
description: 'xAI frontier — coding, agentic tasks & knowledge work',
|
|
1491
|
+
contextWindow: 500_000,
|
|
1492
|
+
// "No text output limit" per docs.x.ai — conservative cap.
|
|
1493
|
+
maxOutputTokens: 128_000,
|
|
1494
|
+
supportsThinking: true,
|
|
1495
|
+
thinkingBudgetTokens: 16_000,
|
|
1496
|
+
thinkingConfigurable: true,
|
|
1497
|
+
// reasoning_effort low|medium|high|xhigh, default high (xAI's model list
|
|
1498
|
+
// `reasoningEffortOptions` for grok-4.6). Same surface as grok-4.7.
|
|
1499
|
+
supportedEffortLevels: ['low', 'medium', 'high', 'xhigh'],
|
|
1500
|
+
defaultEffortLevel: 'high',
|
|
1501
|
+
// Multimodal input (text + images), text out.
|
|
1502
|
+
supportsVision: true,
|
|
1503
|
+
supportsPromptCaching: true,
|
|
1504
|
+
supportsTools: true,
|
|
1505
|
+
// <200K-token prompts; ≥200K bills 2× ($4/$1.00/$12 — tiering not modeled).
|
|
1506
|
+
inputPricePerMTok: 2,
|
|
1507
|
+
outputPricePerMTok: 6,
|
|
1508
|
+
// xAI cached input billed at a flat $0.50/M, no write premium.
|
|
1509
|
+
cacheReadPricePerMTok: 0.5,
|
|
1510
|
+
cacheWritePricePerMTok: 2,
|
|
1511
|
+
// Official (docs.x.ai/developers/grok-4-6, as of 2026-09-10): February 1, 2026.
|
|
1512
|
+
knowledgeCutoff: '2026-02-01',
|
|
1513
|
+
// Superseded by grok-4.7 on arrival: same general-purpose Grok line, same
|
|
1514
|
+
// list price, window and effort levels, newer weights. Cataloged so a
|
|
1515
|
+
// request naming grok-4.6 is priced; never offered in the picker.
|
|
1516
|
+
deprecatedAt: '2026-09-30',
|
|
1517
|
+
supersededBy: 'grok-4.7',
|
|
1518
|
+
},
|
|
1391
1519
|
{
|
|
1392
1520
|
id: 'grok-4.5',
|
|
1393
1521
|
provider: 'xai',
|
|
@@ -1416,6 +1544,10 @@ export const MODELS = [
|
|
|
1416
1544
|
cacheWritePricePerMTok: 2,
|
|
1417
1545
|
// Official (docs.x.ai): 2026-02-01.
|
|
1418
1546
|
knowledgeCutoff: '2026-02-01',
|
|
1547
|
+
// Superseded by grok-4.7: same general-purpose Grok line, same list price
|
|
1548
|
+
// and window, newer weights. Stays priceable for saved selections.
|
|
1549
|
+
deprecatedAt: '2026-09-30',
|
|
1550
|
+
supersededBy: 'grok-4.7',
|
|
1419
1551
|
},
|
|
1420
1552
|
{
|
|
1421
1553
|
id: 'grok-4.3',
|
|
@@ -1441,13 +1573,13 @@ export const MODELS = [
|
|
|
1441
1573
|
cacheReadPricePerMTok: 0.2,
|
|
1442
1574
|
cacheWritePricePerMTok: 1.25,
|
|
1443
1575
|
knowledgeCutoff: '2025-12-01',
|
|
1444
|
-
// Superseded by grok-4.
|
|
1445
|
-
// Grok line, not a separately-named
|
|
1446
|
-
// 500K) and a lower price, which is
|
|
1447
|
-
//
|
|
1448
|
-
//
|
|
1576
|
+
// Superseded by grok-4.7 (was grok-4.5 until 2026-09-30): the previous
|
|
1577
|
+
// version of the same general-purpose Grok line, not a separately-named
|
|
1578
|
+
// tier. It keeps a bigger window (1M vs 500K) and a lower price, which is
|
|
1579
|
+
// why it was previously left selectable — but offering two generations of
|
|
1580
|
+
// one family is exactly what the picker no longer does. Stays priceable.
|
|
1449
1581
|
deprecatedAt: '2026-07-28',
|
|
1450
|
-
supersededBy: 'grok-4.
|
|
1582
|
+
supersededBy: 'grok-4.7',
|
|
1451
1583
|
},
|
|
1452
1584
|
{
|
|
1453
1585
|
id: 'grok-build-0.1',
|
package/package.json
CHANGED