@caupulican/pi-ai 0.90.6 → 0.90.11
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +1 -1
- package/dist/image-models.generated.d.ts +30 -0
- package/dist/image-models.generated.d.ts.map +1 -1
- package/dist/image-models.generated.js +30 -0
- package/dist/image-models.generated.js.map +1 -1
- package/dist/models.generated.d.ts +262 -162
- package/dist/models.generated.d.ts.map +1 -1
- package/dist/models.generated.js +276 -213
- package/dist/models.generated.js.map +1 -1
- package/dist/providers/openai-responses.js +6 -0
- package/dist/providers/openai-responses.js.map +1 -1
- package/dist/utils/oauth/anthropic.d.ts.map +1 -1
- package/dist/utils/oauth/anthropic.js +1 -0
- package/dist/utils/oauth/anthropic.js.map +1 -1
- package/dist/utils/oauth/device-code.d.ts +3 -0
- package/dist/utils/oauth/device-code.d.ts.map +1 -1
- package/dist/utils/oauth/device-code.js +15 -2
- package/dist/utils/oauth/device-code.js.map +1 -1
- package/dist/utils/oauth/github-copilot.d.ts.map +1 -1
- package/dist/utils/oauth/github-copilot.js +1 -0
- package/dist/utils/oauth/github-copilot.js.map +1 -1
- package/dist/utils/oauth/kimi-coding.d.ts.map +1 -1
- package/dist/utils/oauth/kimi-coding.js +2 -0
- package/dist/utils/oauth/kimi-coding.js.map +1 -1
- package/dist/utils/oauth/openai-codex.d.ts.map +1 -1
- package/dist/utils/oauth/openai-codex.js +1 -0
- package/dist/utils/oauth/openai-codex.js.map +1 -1
- package/dist/utils/oauth/types.d.ts +4 -0
- package/dist/utils/oauth/types.d.ts.map +1 -1
- package/dist/utils/oauth/types.js.map +1 -1
- package/dist/utils/oauth/xai.d.ts.map +1 -1
- package/dist/utils/oauth/xai.js +15 -3
- package/dist/utils/oauth/xai.js.map +1 -1
- package/dist/utils/tool-repair/registry.d.ts +1 -1
- package/dist/utils/tool-repair/registry.d.ts.map +1 -1
- package/dist/utils/tool-repair/registry.js +14 -1
- package/dist/utils/tool-repair/registry.js.map +1 -1
- package/dist/utils/tool-repair/repairer.d.ts.map +1 -1
- package/dist/utils/tool-repair/repairer.js +62 -9
- package/dist/utils/tool-repair/repairer.js.map +1 -1
- package/package.json +1 -1
package/dist/models.generated.js
CHANGED
|
@@ -4148,6 +4148,24 @@ export const MODELS = {
|
|
|
4148
4148
|
contextWindow: 131072,
|
|
4149
4149
|
maxTokens: 32768,
|
|
4150
4150
|
},
|
|
4151
|
+
"accounts/fireworks/models/inkling": {
|
|
4152
|
+
id: "accounts/fireworks/models/inkling",
|
|
4153
|
+
name: "Inkling",
|
|
4154
|
+
api: "anthropic-messages",
|
|
4155
|
+
provider: "fireworks",
|
|
4156
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4157
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4158
|
+
reasoning: true,
|
|
4159
|
+
input: ["text", "image"],
|
|
4160
|
+
cost: {
|
|
4161
|
+
input: 1,
|
|
4162
|
+
output: 4.05,
|
|
4163
|
+
cacheRead: 0.17,
|
|
4164
|
+
cacheWrite: 0,
|
|
4165
|
+
},
|
|
4166
|
+
contextWindow: 1048576,
|
|
4167
|
+
maxTokens: 1048576,
|
|
4168
|
+
},
|
|
4151
4169
|
"accounts/fireworks/models/kimi-k2p6": {
|
|
4152
4170
|
id: "accounts/fireworks/models/kimi-k2p6",
|
|
4153
4171
|
name: "Kimi K2.6",
|
|
@@ -4238,6 +4256,60 @@ export const MODELS = {
|
|
|
4238
4256
|
contextWindow: 512000,
|
|
4239
4257
|
maxTokens: 512000,
|
|
4240
4258
|
},
|
|
4259
|
+
"accounts/fireworks/models/muse-glimmer-30b": {
|
|
4260
|
+
id: "accounts/fireworks/models/muse-glimmer-30b",
|
|
4261
|
+
name: "Muse Glimmer 30B",
|
|
4262
|
+
api: "anthropic-messages",
|
|
4263
|
+
provider: "fireworks",
|
|
4264
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4265
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4266
|
+
reasoning: true,
|
|
4267
|
+
input: ["text", "image"],
|
|
4268
|
+
cost: {
|
|
4269
|
+
input: 0.35,
|
|
4270
|
+
output: 1.5,
|
|
4271
|
+
cacheRead: 0.04,
|
|
4272
|
+
cacheWrite: 0,
|
|
4273
|
+
},
|
|
4274
|
+
contextWindow: 131072,
|
|
4275
|
+
maxTokens: 131072,
|
|
4276
|
+
},
|
|
4277
|
+
"accounts/fireworks/models/nemotron-3-ultra-nvfp4": {
|
|
4278
|
+
id: "accounts/fireworks/models/nemotron-3-ultra-nvfp4",
|
|
4279
|
+
name: "Nemotron 3 Ultra 550B A55B",
|
|
4280
|
+
api: "anthropic-messages",
|
|
4281
|
+
provider: "fireworks",
|
|
4282
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4283
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4284
|
+
reasoning: true,
|
|
4285
|
+
input: ["text"],
|
|
4286
|
+
cost: {
|
|
4287
|
+
input: 0.6,
|
|
4288
|
+
output: 2.4,
|
|
4289
|
+
cacheRead: 0.119,
|
|
4290
|
+
cacheWrite: 0,
|
|
4291
|
+
},
|
|
4292
|
+
contextWindow: 262144,
|
|
4293
|
+
maxTokens: 128000,
|
|
4294
|
+
},
|
|
4295
|
+
"accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b": {
|
|
4296
|
+
id: "accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b",
|
|
4297
|
+
name: "Nemotron 3.5 Lightning 30B A3B",
|
|
4298
|
+
api: "anthropic-messages",
|
|
4299
|
+
provider: "fireworks",
|
|
4300
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4301
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4302
|
+
reasoning: true,
|
|
4303
|
+
input: ["text"],
|
|
4304
|
+
cost: {
|
|
4305
|
+
input: 0.05,
|
|
4306
|
+
output: 0.2,
|
|
4307
|
+
cacheRead: 0.01,
|
|
4308
|
+
cacheWrite: 0,
|
|
4309
|
+
},
|
|
4310
|
+
contextWindow: 262144,
|
|
4311
|
+
maxTokens: 262144,
|
|
4312
|
+
},
|
|
4241
4313
|
"accounts/fireworks/models/qwen3p7-plus": {
|
|
4242
4314
|
id: "accounts/fireworks/models/qwen3p7-plus",
|
|
4243
4315
|
name: "Qwen 3.7 Plus",
|
|
@@ -4256,6 +4328,24 @@ export const MODELS = {
|
|
|
4256
4328
|
contextWindow: 262144,
|
|
4257
4329
|
maxTokens: 65536,
|
|
4258
4330
|
},
|
|
4331
|
+
"accounts/fireworks/models/qwen3p8-max": {
|
|
4332
|
+
id: "accounts/fireworks/models/qwen3p8-max",
|
|
4333
|
+
name: "Qwen3.8 Max",
|
|
4334
|
+
api: "anthropic-messages",
|
|
4335
|
+
provider: "fireworks",
|
|
4336
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4337
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4338
|
+
reasoning: true,
|
|
4339
|
+
input: ["text"],
|
|
4340
|
+
cost: {
|
|
4341
|
+
input: 2,
|
|
4342
|
+
output: 6,
|
|
4343
|
+
cacheRead: 0.25,
|
|
4344
|
+
cacheWrite: 0,
|
|
4345
|
+
},
|
|
4346
|
+
contextWindow: 262144,
|
|
4347
|
+
maxTokens: 131072,
|
|
4348
|
+
},
|
|
4259
4349
|
"accounts/fireworks/routers/glm-5p2-fast": {
|
|
4260
4350
|
id: "accounts/fireworks/routers/glm-5p2-fast",
|
|
4261
4351
|
name: "GLM 5.2 Fast",
|
|
@@ -5284,6 +5374,24 @@ export const MODELS = {
|
|
|
5284
5374
|
contextWindow: 1048576,
|
|
5285
5375
|
maxTokens: 65536,
|
|
5286
5376
|
},
|
|
5377
|
+
"gemini-3.7-flash": {
|
|
5378
|
+
id: "gemini-3.7-flash",
|
|
5379
|
+
name: "Gemini 3.7 Flash",
|
|
5380
|
+
api: "google-generative-ai",
|
|
5381
|
+
provider: "google",
|
|
5382
|
+
baseUrl: "https://generativelanguage.googleapis.com/v1beta",
|
|
5383
|
+
reasoning: true,
|
|
5384
|
+
thinkingLevelMap: { "off": null },
|
|
5385
|
+
input: ["text", "image"],
|
|
5386
|
+
cost: {
|
|
5387
|
+
input: 0.75,
|
|
5388
|
+
output: 3.75,
|
|
5389
|
+
cacheRead: 0.075,
|
|
5390
|
+
cacheWrite: 0,
|
|
5391
|
+
},
|
|
5392
|
+
contextWindow: 1048576,
|
|
5393
|
+
maxTokens: 65536,
|
|
5394
|
+
},
|
|
5287
5395
|
"gemini-flash-latest": {
|
|
5288
5396
|
id: "gemini-flash-latest",
|
|
5289
5397
|
name: "Gemini Flash Latest",
|
|
@@ -6191,6 +6299,24 @@ export const MODELS = {
|
|
|
6191
6299
|
contextWindow: 64000,
|
|
6192
6300
|
maxTokens: 8192,
|
|
6193
6301
|
},
|
|
6302
|
+
"deepseek-ai/DeepSeek-V3-0324": {
|
|
6303
|
+
id: "deepseek-ai/DeepSeek-V3-0324",
|
|
6304
|
+
name: "DeepSeek V3 0324",
|
|
6305
|
+
api: "openai-completions",
|
|
6306
|
+
provider: "huggingface",
|
|
6307
|
+
baseUrl: "https://router.huggingface.co/v1",
|
|
6308
|
+
compat: { "supportsDeveloperRole": false },
|
|
6309
|
+
reasoning: false,
|
|
6310
|
+
input: ["text"],
|
|
6311
|
+
cost: {
|
|
6312
|
+
input: 0.27,
|
|
6313
|
+
output: 1.12,
|
|
6314
|
+
cacheRead: 0,
|
|
6315
|
+
cacheWrite: 0,
|
|
6316
|
+
},
|
|
6317
|
+
contextWindow: 163840,
|
|
6318
|
+
maxTokens: 163840,
|
|
6319
|
+
},
|
|
6194
6320
|
"deepseek-ai/DeepSeek-V3.1": {
|
|
6195
6321
|
id: "deepseek-ai/DeepSeek-V3.1",
|
|
6196
6322
|
name: "DeepSeek-V3.1",
|
|
@@ -9077,6 +9203,24 @@ export const MODELS = {
|
|
|
9077
9203
|
contextWindow: 1048576,
|
|
9078
9204
|
maxTokens: 65536,
|
|
9079
9205
|
},
|
|
9206
|
+
"gemini-3.7-flash": {
|
|
9207
|
+
id: "gemini-3.7-flash",
|
|
9208
|
+
name: "Gemini 3.7 Flash",
|
|
9209
|
+
api: "google-generative-ai",
|
|
9210
|
+
provider: "opencode",
|
|
9211
|
+
baseUrl: "https://opencode.ai/zen/v1",
|
|
9212
|
+
reasoning: true,
|
|
9213
|
+
thinkingLevelMap: { "off": null },
|
|
9214
|
+
input: ["text", "image"],
|
|
9215
|
+
cost: {
|
|
9216
|
+
input: 1.5,
|
|
9217
|
+
output: 7.5,
|
|
9218
|
+
cacheRead: 0.15,
|
|
9219
|
+
cacheWrite: 0,
|
|
9220
|
+
},
|
|
9221
|
+
contextWindow: 1048576,
|
|
9222
|
+
maxTokens: 65536,
|
|
9223
|
+
},
|
|
9080
9224
|
"glm-5": {
|
|
9081
9225
|
id: "glm-5",
|
|
9082
9226
|
name: "GLM-5",
|
|
@@ -9647,23 +9791,6 @@ export const MODELS = {
|
|
|
9647
9791
|
contextWindow: 256000,
|
|
9648
9792
|
maxTokens: 32000,
|
|
9649
9793
|
},
|
|
9650
|
-
"ling-3.0-tiny-free": {
|
|
9651
|
-
id: "ling-3.0-tiny-free",
|
|
9652
|
-
name: "Ling-3.0-tiny Free",
|
|
9653
|
-
api: "openai-completions",
|
|
9654
|
-
provider: "opencode",
|
|
9655
|
-
baseUrl: "https://opencode.ai/zen/v1",
|
|
9656
|
-
reasoning: true,
|
|
9657
|
-
input: ["text"],
|
|
9658
|
-
cost: {
|
|
9659
|
-
input: 0,
|
|
9660
|
-
output: 0,
|
|
9661
|
-
cacheRead: 0,
|
|
9662
|
-
cacheWrite: 0,
|
|
9663
|
-
},
|
|
9664
|
-
contextWindow: 262144,
|
|
9665
|
-
maxTokens: 32768,
|
|
9666
|
-
},
|
|
9667
9794
|
"mimo-v2.5-free": {
|
|
9668
9795
|
id: "mimo-v2.5-free",
|
|
9669
9796
|
name: "MiMo V2.5 Free",
|
|
@@ -11061,7 +11188,7 @@ export const MODELS = {
|
|
|
11061
11188
|
cacheRead: 0,
|
|
11062
11189
|
cacheWrite: 0,
|
|
11063
11190
|
},
|
|
11064
|
-
contextWindow:
|
|
11191
|
+
contextWindow: 64000,
|
|
11065
11192
|
maxTokens: 16000,
|
|
11066
11193
|
},
|
|
11067
11194
|
"deepseek/deepseek-r1-0528": {
|
|
@@ -11162,13 +11289,13 @@ export const MODELS = {
|
|
|
11162
11289
|
thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
|
|
11163
11290
|
input: ["text"],
|
|
11164
11291
|
cost: {
|
|
11165
|
-
input: 0.
|
|
11166
|
-
output: 0.
|
|
11167
|
-
cacheRead: 0.
|
|
11292
|
+
input: 0.14,
|
|
11293
|
+
output: 0.28,
|
|
11294
|
+
cacheRead: 0.028,
|
|
11168
11295
|
cacheWrite: 0,
|
|
11169
11296
|
},
|
|
11170
11297
|
contextWindow: 1048576,
|
|
11171
|
-
maxTokens:
|
|
11298
|
+
maxTokens: 393216,
|
|
11172
11299
|
},
|
|
11173
11300
|
"deepseek/deepseek-v4-pro": {
|
|
11174
11301
|
id: "deepseek/deepseek-v4-pro",
|
|
@@ -11574,10 +11701,10 @@ export const MODELS = {
|
|
|
11574
11701
|
reasoning: true,
|
|
11575
11702
|
input: ["text", "image"],
|
|
11576
11703
|
cost: {
|
|
11577
|
-
input:
|
|
11578
|
-
output:
|
|
11579
|
-
cacheRead: 0.
|
|
11580
|
-
cacheWrite: 0.
|
|
11704
|
+
input: 0.75,
|
|
11705
|
+
output: 3.75,
|
|
11706
|
+
cacheRead: 0.075,
|
|
11707
|
+
cacheWrite: 0.0416666666666667,
|
|
11581
11708
|
},
|
|
11582
11709
|
contextWindow: 1048576,
|
|
11583
11710
|
maxTokens: 65536,
|
|
@@ -11591,10 +11718,44 @@ export const MODELS = {
|
|
|
11591
11718
|
reasoning: true,
|
|
11592
11719
|
input: ["text", "image"],
|
|
11593
11720
|
cost: {
|
|
11594
|
-
input: 0.
|
|
11595
|
-
output:
|
|
11596
|
-
cacheRead: 0.
|
|
11597
|
-
cacheWrite: 0.
|
|
11721
|
+
input: 0.375,
|
|
11722
|
+
output: 1.875,
|
|
11723
|
+
cacheRead: 0.0375,
|
|
11724
|
+
cacheWrite: 0.0416666666666667,
|
|
11725
|
+
},
|
|
11726
|
+
contextWindow: 1048576,
|
|
11727
|
+
maxTokens: 65536,
|
|
11728
|
+
},
|
|
11729
|
+
"google/gemini-3.7-flash": {
|
|
11730
|
+
id: "google/gemini-3.7-flash",
|
|
11731
|
+
name: "Google: Gemini 3.7 Flash",
|
|
11732
|
+
api: "openai-completions",
|
|
11733
|
+
provider: "openrouter",
|
|
11734
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
11735
|
+
reasoning: true,
|
|
11736
|
+
input: ["text", "image"],
|
|
11737
|
+
cost: {
|
|
11738
|
+
input: 0.375,
|
|
11739
|
+
output: 1.875,
|
|
11740
|
+
cacheRead: 0.0375,
|
|
11741
|
+
cacheWrite: 0.0208333333333333,
|
|
11742
|
+
},
|
|
11743
|
+
contextWindow: 1048576,
|
|
11744
|
+
maxTokens: 65536,
|
|
11745
|
+
},
|
|
11746
|
+
"google/gemini-3.7-flash:batch": {
|
|
11747
|
+
id: "google/gemini-3.7-flash:batch",
|
|
11748
|
+
name: "Google: Gemini 3.7 Flash (batch)",
|
|
11749
|
+
api: "openai-completions",
|
|
11750
|
+
provider: "openrouter",
|
|
11751
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
11752
|
+
reasoning: true,
|
|
11753
|
+
input: ["text", "image"],
|
|
11754
|
+
cost: {
|
|
11755
|
+
input: 0.1875,
|
|
11756
|
+
output: 0.9375,
|
|
11757
|
+
cacheRead: 0.01875,
|
|
11758
|
+
cacheWrite: 0.0208333333333333,
|
|
11598
11759
|
},
|
|
11599
11760
|
contextWindow: 1048576,
|
|
11600
11761
|
maxTokens: 65536,
|
|
@@ -11787,23 +11948,6 @@ export const MODELS = {
|
|
|
11787
11948
|
contextWindow: 262144,
|
|
11788
11949
|
maxTokens: 32768,
|
|
11789
11950
|
},
|
|
11790
|
-
"inclusionai/ling-3.0-tiny:free": {
|
|
11791
|
-
id: "inclusionai/ling-3.0-tiny:free",
|
|
11792
|
-
name: "inclusionAI: Ling 3.0 Tiny (free)",
|
|
11793
|
-
api: "openai-completions",
|
|
11794
|
-
provider: "openrouter",
|
|
11795
|
-
baseUrl: "https://openrouter.ai/api/v1",
|
|
11796
|
-
reasoning: true,
|
|
11797
|
-
input: ["text"],
|
|
11798
|
-
cost: {
|
|
11799
|
-
input: 0,
|
|
11800
|
-
output: 0,
|
|
11801
|
-
cacheRead: 0,
|
|
11802
|
-
cacheWrite: 0,
|
|
11803
|
-
},
|
|
11804
|
-
contextWindow: 262144,
|
|
11805
|
-
maxTokens: 32768,
|
|
11806
|
-
},
|
|
11807
11951
|
"inclusionai/ring-2.6-1t": {
|
|
11808
11952
|
id: "inclusionai/ring-2.6-1t",
|
|
11809
11953
|
name: "inclusionAI: Ring-2.6-1T",
|
|
@@ -11887,7 +12031,7 @@ export const MODELS = {
|
|
|
11887
12031
|
cacheWrite: 0,
|
|
11888
12032
|
},
|
|
11889
12033
|
contextWindow: 128000,
|
|
11890
|
-
maxTokens:
|
|
12034
|
+
maxTokens: 8192,
|
|
11891
12035
|
},
|
|
11892
12036
|
"meituan/longcat-2.0": {
|
|
11893
12037
|
id: "meituan/longcat-2.0",
|
|
@@ -12511,9 +12655,9 @@ export const MODELS = {
|
|
|
12511
12655
|
reasoning: true,
|
|
12512
12656
|
input: ["text", "image"],
|
|
12513
12657
|
cost: {
|
|
12514
|
-
input: 0.
|
|
12515
|
-
output:
|
|
12516
|
-
cacheRead: 0.
|
|
12658
|
+
input: 0.95,
|
|
12659
|
+
output: 4,
|
|
12660
|
+
cacheRead: 0.16,
|
|
12517
12661
|
cacheWrite: 0,
|
|
12518
12662
|
},
|
|
12519
12663
|
contextWindow: 262144,
|
|
@@ -12740,6 +12884,23 @@ export const MODELS = {
|
|
|
12740
12884
|
contextWindow: 1000000,
|
|
12741
12885
|
maxTokens: 65536,
|
|
12742
12886
|
},
|
|
12887
|
+
"nvidia/nemotron-3.5-lightning": {
|
|
12888
|
+
id: "nvidia/nemotron-3.5-lightning",
|
|
12889
|
+
name: "NVIDIA: Nemotron 3.5 Lightning",
|
|
12890
|
+
api: "openai-completions",
|
|
12891
|
+
provider: "openrouter",
|
|
12892
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
12893
|
+
reasoning: true,
|
|
12894
|
+
input: ["text"],
|
|
12895
|
+
cost: {
|
|
12896
|
+
input: 0.09999999999999999,
|
|
12897
|
+
output: 0.25,
|
|
12898
|
+
cacheRead: 0.049999999999999996,
|
|
12899
|
+
cacheWrite: 0,
|
|
12900
|
+
},
|
|
12901
|
+
contextWindow: 1000000,
|
|
12902
|
+
maxTokens: 262144,
|
|
12903
|
+
},
|
|
12743
12904
|
"nvidia/nemotron-3.5-lightning:free": {
|
|
12744
12905
|
id: "nvidia/nemotron-3.5-lightning:free",
|
|
12745
12906
|
name: "NVIDIA: Nemotron 3.5 Lightning (free)",
|
|
@@ -14820,13 +14981,13 @@ export const MODELS = {
|
|
|
14820
14981
|
reasoning: false,
|
|
14821
14982
|
input: ["text"],
|
|
14822
14983
|
cost: {
|
|
14823
|
-
input: 0.
|
|
14984
|
+
input: 0.09999999999999999,
|
|
14824
14985
|
output: 1.1,
|
|
14825
|
-
cacheRead: 0,
|
|
14986
|
+
cacheRead: 0.07,
|
|
14826
14987
|
cacheWrite: 0,
|
|
14827
14988
|
},
|
|
14828
14989
|
contextWindow: 262144,
|
|
14829
|
-
maxTokens:
|
|
14990
|
+
maxTokens: 262144,
|
|
14830
14991
|
},
|
|
14831
14992
|
"qwen/qwen3-next-80b-a3b-thinking": {
|
|
14832
14993
|
id: "qwen/qwen3-next-80b-a3b-thinking",
|
|
@@ -14843,7 +15004,7 @@ export const MODELS = {
|
|
|
14843
15004
|
cacheWrite: 0,
|
|
14844
15005
|
},
|
|
14845
15006
|
contextWindow: 262144,
|
|
14846
|
-
maxTokens:
|
|
15007
|
+
maxTokens: 32768,
|
|
14847
15008
|
},
|
|
14848
15009
|
"qwen/qwen3-vl-235b-a22b-instruct": {
|
|
14849
15010
|
id: "qwen/qwen3-vl-235b-a22b-instruct",
|
|
@@ -14888,13 +15049,13 @@ export const MODELS = {
|
|
|
14888
15049
|
reasoning: false,
|
|
14889
15050
|
input: ["text", "image"],
|
|
14890
15051
|
cost: {
|
|
14891
|
-
input: 0.
|
|
14892
|
-
output: 0.
|
|
15052
|
+
input: 0.13,
|
|
15053
|
+
output: 0.52,
|
|
14893
15054
|
cacheRead: 0,
|
|
14894
15055
|
cacheWrite: 0,
|
|
14895
15056
|
},
|
|
14896
15057
|
contextWindow: 262144,
|
|
14897
|
-
maxTokens:
|
|
15058
|
+
maxTokens: 32768,
|
|
14898
15059
|
},
|
|
14899
15060
|
"qwen/qwen3-vl-30b-a3b-thinking": {
|
|
14900
15061
|
id: "qwen/qwen3-vl-30b-a3b-thinking",
|
|
@@ -15007,9 +15168,9 @@ export const MODELS = {
|
|
|
15007
15168
|
reasoning: true,
|
|
15008
15169
|
input: ["text", "image"],
|
|
15009
15170
|
cost: {
|
|
15010
|
-
input: 0.
|
|
15011
|
-
output: 1,
|
|
15012
|
-
cacheRead: 0,
|
|
15171
|
+
input: 0.25,
|
|
15172
|
+
output: 1.25,
|
|
15173
|
+
cacheRead: 0.25,
|
|
15013
15174
|
cacheWrite: 0,
|
|
15014
15175
|
},
|
|
15015
15176
|
contextWindow: 262144,
|
|
@@ -15247,11 +15408,11 @@ export const MODELS = {
|
|
|
15247
15408
|
cost: {
|
|
15248
15409
|
input: 2,
|
|
15249
15410
|
output: 6,
|
|
15250
|
-
cacheRead: 0.
|
|
15411
|
+
cacheRead: 0.25,
|
|
15251
15412
|
cacheWrite: 0,
|
|
15252
15413
|
},
|
|
15253
|
-
contextWindow:
|
|
15254
|
-
maxTokens:
|
|
15414
|
+
contextWindow: 1010000,
|
|
15415
|
+
maxTokens: 262144,
|
|
15255
15416
|
},
|
|
15256
15417
|
"qwen/qwen3.8-max": {
|
|
15257
15418
|
id: "qwen/qwen3.8-max",
|
|
@@ -15704,9 +15865,9 @@ export const MODELS = {
|
|
|
15704
15865
|
reasoning: true,
|
|
15705
15866
|
input: ["text"],
|
|
15706
15867
|
cost: {
|
|
15707
|
-
input: 0.
|
|
15708
|
-
output: 2
|
|
15709
|
-
cacheRead: 0.
|
|
15868
|
+
input: 0.5,
|
|
15869
|
+
output: 2,
|
|
15870
|
+
cacheRead: 0.09999999999999999,
|
|
15710
15871
|
cacheWrite: 0,
|
|
15711
15872
|
},
|
|
15712
15873
|
contextWindow: 204800,
|
|
@@ -15823,13 +15984,13 @@ export const MODELS = {
|
|
|
15823
15984
|
reasoning: true,
|
|
15824
15985
|
input: ["text"],
|
|
15825
15986
|
cost: {
|
|
15826
|
-
input: 0.
|
|
15827
|
-
output: 1.
|
|
15828
|
-
cacheRead: 0.
|
|
15987
|
+
input: 0.63,
|
|
15988
|
+
output: 1.9800000000000002,
|
|
15989
|
+
cacheRead: 0.0945,
|
|
15829
15990
|
cacheWrite: 0,
|
|
15830
15991
|
},
|
|
15831
15992
|
contextWindow: 1048576,
|
|
15832
|
-
maxTokens:
|
|
15993
|
+
maxTokens: 4096,
|
|
15833
15994
|
},
|
|
15834
15995
|
"z-ai/glm-5.2:batch": {
|
|
15835
15996
|
id: "z-ai/glm-5.2:batch",
|
|
@@ -15965,10 +16126,10 @@ export const MODELS = {
|
|
|
15965
16126
|
reasoning: true,
|
|
15966
16127
|
input: ["text", "image"],
|
|
15967
16128
|
cost: {
|
|
15968
|
-
input:
|
|
15969
|
-
output:
|
|
15970
|
-
cacheRead: 0.
|
|
15971
|
-
cacheWrite: 0.
|
|
16129
|
+
input: 0.375,
|
|
16130
|
+
output: 1.875,
|
|
16131
|
+
cacheRead: 0.0375,
|
|
16132
|
+
cacheWrite: 0.0208333333333333,
|
|
15972
16133
|
},
|
|
15973
16134
|
contextWindow: 1048576,
|
|
15974
16135
|
maxTokens: 65536,
|
|
@@ -17448,6 +17609,23 @@ export const MODELS = {
|
|
|
17448
17609
|
contextWindow: 1000000,
|
|
17449
17610
|
maxTokens: 64000,
|
|
17450
17611
|
},
|
|
17612
|
+
"alibaba/qwen3.8-2.4t-a95b": {
|
|
17613
|
+
id: "alibaba/qwen3.8-2.4t-a95b",
|
|
17614
|
+
name: "Qwen3.8 2.4T A95B",
|
|
17615
|
+
api: "anthropic-messages",
|
|
17616
|
+
provider: "vercel-ai-gateway",
|
|
17617
|
+
baseUrl: "https://ai-gateway.vercel.sh",
|
|
17618
|
+
reasoning: true,
|
|
17619
|
+
input: ["text"],
|
|
17620
|
+
cost: {
|
|
17621
|
+
input: 2,
|
|
17622
|
+
output: 6,
|
|
17623
|
+
cacheRead: 0.25,
|
|
17624
|
+
cacheWrite: 0,
|
|
17625
|
+
},
|
|
17626
|
+
contextWindow: 262144,
|
|
17627
|
+
maxTokens: 131072,
|
|
17628
|
+
},
|
|
17451
17629
|
"alibaba/qwen3.8-max": {
|
|
17452
17630
|
id: "alibaba/qwen3.8-max",
|
|
17453
17631
|
name: "Qwen 3.8 Max",
|
|
@@ -18214,6 +18392,23 @@ export const MODELS = {
|
|
|
18214
18392
|
contextWindow: 1000000,
|
|
18215
18393
|
maxTokens: 64000,
|
|
18216
18394
|
},
|
|
18395
|
+
"google/gemini-3.7-flash": {
|
|
18396
|
+
id: "google/gemini-3.7-flash",
|
|
18397
|
+
name: "Gemini 3.7 Flash",
|
|
18398
|
+
api: "anthropic-messages",
|
|
18399
|
+
provider: "vercel-ai-gateway",
|
|
18400
|
+
baseUrl: "https://ai-gateway.vercel.sh",
|
|
18401
|
+
reasoning: true,
|
|
18402
|
+
input: ["text", "image"],
|
|
18403
|
+
cost: {
|
|
18404
|
+
input: 0.75,
|
|
18405
|
+
output: 3.75,
|
|
18406
|
+
cacheRead: 0.075,
|
|
18407
|
+
cacheWrite: 0,
|
|
18408
|
+
},
|
|
18409
|
+
contextWindow: 1000000,
|
|
18410
|
+
maxTokens: 65536,
|
|
18411
|
+
},
|
|
18217
18412
|
"google/gemma-4-26b-a4b-it": {
|
|
18218
18413
|
id: "google/gemma-4-26b-a4b-it",
|
|
18219
18414
|
name: "Google Gemma 4 26B A4B",
|
|
@@ -18245,7 +18440,7 @@ export const MODELS = {
|
|
|
18245
18440
|
cacheRead: 0,
|
|
18246
18441
|
cacheWrite: 0,
|
|
18247
18442
|
},
|
|
18248
|
-
contextWindow:
|
|
18443
|
+
contextWindow: 262144,
|
|
18249
18444
|
maxTokens: 131072,
|
|
18250
18445
|
},
|
|
18251
18446
|
"inception/mercury-2": {
|
|
@@ -18299,23 +18494,6 @@ export const MODELS = {
|
|
|
18299
18494
|
contextWindow: 256000,
|
|
18300
18495
|
maxTokens: 32000,
|
|
18301
18496
|
},
|
|
18302
|
-
"inclusionai/ling-3.0-tiny-free": {
|
|
18303
|
-
id: "inclusionai/ling-3.0-tiny-free",
|
|
18304
|
-
name: "Ling 3.0 Tiny (Free)",
|
|
18305
|
-
api: "anthropic-messages",
|
|
18306
|
-
provider: "vercel-ai-gateway",
|
|
18307
|
-
baseUrl: "https://ai-gateway.vercel.sh",
|
|
18308
|
-
reasoning: true,
|
|
18309
|
-
input: ["text"],
|
|
18310
|
-
cost: {
|
|
18311
|
-
input: 0,
|
|
18312
|
-
output: 0,
|
|
18313
|
-
cacheRead: 0,
|
|
18314
|
-
cacheWrite: 0,
|
|
18315
|
-
},
|
|
18316
|
-
contextWindow: 256000,
|
|
18317
|
-
maxTokens: 32000,
|
|
18318
|
-
},
|
|
18319
18497
|
"interfaze/interfaze-beta": {
|
|
18320
18498
|
id: "interfaze/interfaze-beta",
|
|
18321
18499
|
name: "Interfaze Beta",
|
|
@@ -20456,91 +20634,6 @@ export const MODELS = {
|
|
|
20456
20634
|
},
|
|
20457
20635
|
},
|
|
20458
20636
|
"xai": {
|
|
20459
|
-
"grok-3": {
|
|
20460
|
-
id: "grok-3",
|
|
20461
|
-
name: "Grok 3",
|
|
20462
|
-
api: "openai-completions",
|
|
20463
|
-
provider: "xai",
|
|
20464
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20465
|
-
reasoning: false,
|
|
20466
|
-
input: ["text"],
|
|
20467
|
-
cost: {
|
|
20468
|
-
input: 3,
|
|
20469
|
-
output: 15,
|
|
20470
|
-
cacheRead: 0.75,
|
|
20471
|
-
cacheWrite: 0,
|
|
20472
|
-
},
|
|
20473
|
-
contextWindow: 131072,
|
|
20474
|
-
maxTokens: 8192,
|
|
20475
|
-
},
|
|
20476
|
-
"grok-3-fast": {
|
|
20477
|
-
id: "grok-3-fast",
|
|
20478
|
-
name: "Grok 3 Fast",
|
|
20479
|
-
api: "openai-completions",
|
|
20480
|
-
provider: "xai",
|
|
20481
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20482
|
-
reasoning: false,
|
|
20483
|
-
input: ["text"],
|
|
20484
|
-
cost: {
|
|
20485
|
-
input: 5,
|
|
20486
|
-
output: 25,
|
|
20487
|
-
cacheRead: 1.25,
|
|
20488
|
-
cacheWrite: 0,
|
|
20489
|
-
},
|
|
20490
|
-
contextWindow: 131072,
|
|
20491
|
-
maxTokens: 8192,
|
|
20492
|
-
},
|
|
20493
|
-
"grok-4.20-0309-non-reasoning": {
|
|
20494
|
-
id: "grok-4.20-0309-non-reasoning",
|
|
20495
|
-
name: "Grok 4.20 (Non-Reasoning)",
|
|
20496
|
-
api: "openai-completions",
|
|
20497
|
-
provider: "xai",
|
|
20498
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20499
|
-
reasoning: false,
|
|
20500
|
-
input: ["text", "image"],
|
|
20501
|
-
cost: {
|
|
20502
|
-
input: 1.25,
|
|
20503
|
-
output: 2.5,
|
|
20504
|
-
cacheRead: 0.2,
|
|
20505
|
-
cacheWrite: 0,
|
|
20506
|
-
},
|
|
20507
|
-
contextWindow: 1000000,
|
|
20508
|
-
maxTokens: 30000,
|
|
20509
|
-
},
|
|
20510
|
-
"grok-4.20-0309-reasoning": {
|
|
20511
|
-
id: "grok-4.20-0309-reasoning",
|
|
20512
|
-
name: "Grok 4.20 (Reasoning)",
|
|
20513
|
-
api: "openai-completions",
|
|
20514
|
-
provider: "xai",
|
|
20515
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20516
|
-
reasoning: true,
|
|
20517
|
-
input: ["text", "image"],
|
|
20518
|
-
cost: {
|
|
20519
|
-
input: 1.25,
|
|
20520
|
-
output: 2.5,
|
|
20521
|
-
cacheRead: 0.2,
|
|
20522
|
-
cacheWrite: 0,
|
|
20523
|
-
},
|
|
20524
|
-
contextWindow: 1000000,
|
|
20525
|
-
maxTokens: 30000,
|
|
20526
|
-
},
|
|
20527
|
-
"grok-4.3": {
|
|
20528
|
-
id: "grok-4.3",
|
|
20529
|
-
name: "Grok 4.3",
|
|
20530
|
-
api: "openai-completions",
|
|
20531
|
-
provider: "xai",
|
|
20532
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20533
|
-
reasoning: true,
|
|
20534
|
-
input: ["text", "image"],
|
|
20535
|
-
cost: {
|
|
20536
|
-
input: 1.25,
|
|
20537
|
-
output: 2.5,
|
|
20538
|
-
cacheRead: 0.2,
|
|
20539
|
-
cacheWrite: 0,
|
|
20540
|
-
},
|
|
20541
|
-
contextWindow: 1000000,
|
|
20542
|
-
maxTokens: 30000,
|
|
20543
|
-
},
|
|
20544
20637
|
"grok-4.5": {
|
|
20545
20638
|
id: "grok-4.5",
|
|
20546
20639
|
name: "Grok 4.5",
|
|
@@ -20549,6 +20642,7 @@ export const MODELS = {
|
|
|
20549
20642
|
baseUrl: "https://api.x.ai/v1",
|
|
20550
20643
|
compat: { "supportsLongCacheRetention": false },
|
|
20551
20644
|
reasoning: true,
|
|
20645
|
+
defaultThinkingLevel: "high",
|
|
20552
20646
|
thinkingLevelMap: { "off": null, "minimal": null },
|
|
20553
20647
|
input: ["text", "image"],
|
|
20554
20648
|
cost: {
|
|
@@ -20563,10 +20657,13 @@ export const MODELS = {
|
|
|
20563
20657
|
"grok-4.6": {
|
|
20564
20658
|
id: "grok-4.6",
|
|
20565
20659
|
name: "Grok 4.6",
|
|
20566
|
-
api: "openai-
|
|
20660
|
+
api: "openai-responses",
|
|
20567
20661
|
provider: "xai",
|
|
20568
20662
|
baseUrl: "https://api.x.ai/v1",
|
|
20663
|
+
compat: { "supportsLongCacheRetention": false },
|
|
20569
20664
|
reasoning: true,
|
|
20665
|
+
defaultThinkingLevel: "high",
|
|
20666
|
+
thinkingLevelMap: { "off": null, "minimal": null, "xhigh": "xhigh" },
|
|
20570
20667
|
input: ["text", "image"],
|
|
20571
20668
|
cost: {
|
|
20572
20669
|
input: 2,
|
|
@@ -20577,40 +20674,6 @@ export const MODELS = {
|
|
|
20577
20674
|
contextWindow: 500000,
|
|
20578
20675
|
maxTokens: 500000,
|
|
20579
20676
|
},
|
|
20580
|
-
"grok-build-0.1": {
|
|
20581
|
-
id: "grok-build-0.1",
|
|
20582
|
-
name: "Grok Build 0.1",
|
|
20583
|
-
api: "openai-completions",
|
|
20584
|
-
provider: "xai",
|
|
20585
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20586
|
-
reasoning: true,
|
|
20587
|
-
input: ["text", "image"],
|
|
20588
|
-
cost: {
|
|
20589
|
-
input: 1,
|
|
20590
|
-
output: 2,
|
|
20591
|
-
cacheRead: 0.2,
|
|
20592
|
-
cacheWrite: 0,
|
|
20593
|
-
},
|
|
20594
|
-
contextWindow: 256000,
|
|
20595
|
-
maxTokens: 256000,
|
|
20596
|
-
},
|
|
20597
|
-
"grok-code-fast-1": {
|
|
20598
|
-
id: "grok-code-fast-1",
|
|
20599
|
-
name: "Grok Code Fast 1",
|
|
20600
|
-
api: "openai-completions",
|
|
20601
|
-
provider: "xai",
|
|
20602
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20603
|
-
reasoning: false,
|
|
20604
|
-
input: ["text"],
|
|
20605
|
-
cost: {
|
|
20606
|
-
input: 0.2,
|
|
20607
|
-
output: 1.5,
|
|
20608
|
-
cacheRead: 0.02,
|
|
20609
|
-
cacheWrite: 0,
|
|
20610
|
-
},
|
|
20611
|
-
contextWindow: 32768,
|
|
20612
|
-
maxTokens: 8192,
|
|
20613
|
-
},
|
|
20614
20677
|
},
|
|
20615
20678
|
"xiaomi": {
|
|
20616
20679
|
"mimo-v2-flash": {
|