@caupulican/pi-ai 0.90.7 → 0.90.11
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +1 -1
- package/dist/image-models.generated.d.ts +15 -0
- package/dist/image-models.generated.d.ts.map +1 -1
- package/dist/image-models.generated.js +15 -0
- package/dist/image-models.generated.js.map +1 -1
- package/dist/models.generated.d.ts +245 -162
- package/dist/models.generated.d.ts.map +1 -1
- package/dist/models.generated.js +254 -208
- package/dist/models.generated.js.map +1 -1
- package/dist/providers/openai-responses.js +6 -0
- package/dist/providers/openai-responses.js.map +1 -1
- package/dist/utils/oauth/anthropic.d.ts.map +1 -1
- package/dist/utils/oauth/anthropic.js +1 -0
- package/dist/utils/oauth/anthropic.js.map +1 -1
- package/dist/utils/oauth/device-code.d.ts +3 -0
- package/dist/utils/oauth/device-code.d.ts.map +1 -1
- package/dist/utils/oauth/device-code.js +15 -2
- package/dist/utils/oauth/device-code.js.map +1 -1
- package/dist/utils/oauth/github-copilot.d.ts.map +1 -1
- package/dist/utils/oauth/github-copilot.js +1 -0
- package/dist/utils/oauth/github-copilot.js.map +1 -1
- package/dist/utils/oauth/kimi-coding.d.ts.map +1 -1
- package/dist/utils/oauth/kimi-coding.js +2 -0
- package/dist/utils/oauth/kimi-coding.js.map +1 -1
- package/dist/utils/oauth/openai-codex.d.ts.map +1 -1
- package/dist/utils/oauth/openai-codex.js +1 -0
- package/dist/utils/oauth/openai-codex.js.map +1 -1
- package/dist/utils/oauth/types.d.ts +4 -0
- package/dist/utils/oauth/types.d.ts.map +1 -1
- package/dist/utils/oauth/types.js.map +1 -1
- package/dist/utils/oauth/xai.d.ts.map +1 -1
- package/dist/utils/oauth/xai.js +15 -3
- package/dist/utils/oauth/xai.js.map +1 -1
- package/dist/utils/tool-repair/registry.d.ts +1 -1
- package/dist/utils/tool-repair/registry.d.ts.map +1 -1
- package/dist/utils/tool-repair/registry.js +14 -1
- package/dist/utils/tool-repair/registry.js.map +1 -1
- package/dist/utils/tool-repair/repairer.d.ts.map +1 -1
- package/dist/utils/tool-repair/repairer.js +62 -9
- package/dist/utils/tool-repair/repairer.js.map +1 -1
- package/package.json +1 -1
package/dist/models.generated.js
CHANGED
|
@@ -4148,6 +4148,24 @@ export const MODELS = {
|
|
|
4148
4148
|
contextWindow: 131072,
|
|
4149
4149
|
maxTokens: 32768,
|
|
4150
4150
|
},
|
|
4151
|
+
"accounts/fireworks/models/inkling": {
|
|
4152
|
+
id: "accounts/fireworks/models/inkling",
|
|
4153
|
+
name: "Inkling",
|
|
4154
|
+
api: "anthropic-messages",
|
|
4155
|
+
provider: "fireworks",
|
|
4156
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4157
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4158
|
+
reasoning: true,
|
|
4159
|
+
input: ["text", "image"],
|
|
4160
|
+
cost: {
|
|
4161
|
+
input: 1,
|
|
4162
|
+
output: 4.05,
|
|
4163
|
+
cacheRead: 0.17,
|
|
4164
|
+
cacheWrite: 0,
|
|
4165
|
+
},
|
|
4166
|
+
contextWindow: 1048576,
|
|
4167
|
+
maxTokens: 1048576,
|
|
4168
|
+
},
|
|
4151
4169
|
"accounts/fireworks/models/kimi-k2p6": {
|
|
4152
4170
|
id: "accounts/fireworks/models/kimi-k2p6",
|
|
4153
4171
|
name: "Kimi K2.6",
|
|
@@ -4238,6 +4256,60 @@ export const MODELS = {
|
|
|
4238
4256
|
contextWindow: 512000,
|
|
4239
4257
|
maxTokens: 512000,
|
|
4240
4258
|
},
|
|
4259
|
+
"accounts/fireworks/models/muse-glimmer-30b": {
|
|
4260
|
+
id: "accounts/fireworks/models/muse-glimmer-30b",
|
|
4261
|
+
name: "Muse Glimmer 30B",
|
|
4262
|
+
api: "anthropic-messages",
|
|
4263
|
+
provider: "fireworks",
|
|
4264
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4265
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4266
|
+
reasoning: true,
|
|
4267
|
+
input: ["text", "image"],
|
|
4268
|
+
cost: {
|
|
4269
|
+
input: 0.35,
|
|
4270
|
+
output: 1.5,
|
|
4271
|
+
cacheRead: 0.04,
|
|
4272
|
+
cacheWrite: 0,
|
|
4273
|
+
},
|
|
4274
|
+
contextWindow: 131072,
|
|
4275
|
+
maxTokens: 131072,
|
|
4276
|
+
},
|
|
4277
|
+
"accounts/fireworks/models/nemotron-3-ultra-nvfp4": {
|
|
4278
|
+
id: "accounts/fireworks/models/nemotron-3-ultra-nvfp4",
|
|
4279
|
+
name: "Nemotron 3 Ultra 550B A55B",
|
|
4280
|
+
api: "anthropic-messages",
|
|
4281
|
+
provider: "fireworks",
|
|
4282
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4283
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4284
|
+
reasoning: true,
|
|
4285
|
+
input: ["text"],
|
|
4286
|
+
cost: {
|
|
4287
|
+
input: 0.6,
|
|
4288
|
+
output: 2.4,
|
|
4289
|
+
cacheRead: 0.119,
|
|
4290
|
+
cacheWrite: 0,
|
|
4291
|
+
},
|
|
4292
|
+
contextWindow: 262144,
|
|
4293
|
+
maxTokens: 128000,
|
|
4294
|
+
},
|
|
4295
|
+
"accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b": {
|
|
4296
|
+
id: "accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b",
|
|
4297
|
+
name: "Nemotron 3.5 Lightning 30B A3B",
|
|
4298
|
+
api: "anthropic-messages",
|
|
4299
|
+
provider: "fireworks",
|
|
4300
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4301
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4302
|
+
reasoning: true,
|
|
4303
|
+
input: ["text"],
|
|
4304
|
+
cost: {
|
|
4305
|
+
input: 0.05,
|
|
4306
|
+
output: 0.2,
|
|
4307
|
+
cacheRead: 0.01,
|
|
4308
|
+
cacheWrite: 0,
|
|
4309
|
+
},
|
|
4310
|
+
contextWindow: 262144,
|
|
4311
|
+
maxTokens: 262144,
|
|
4312
|
+
},
|
|
4241
4313
|
"accounts/fireworks/models/qwen3p7-plus": {
|
|
4242
4314
|
id: "accounts/fireworks/models/qwen3p7-plus",
|
|
4243
4315
|
name: "Qwen 3.7 Plus",
|
|
@@ -4256,6 +4328,24 @@ export const MODELS = {
|
|
|
4256
4328
|
contextWindow: 262144,
|
|
4257
4329
|
maxTokens: 65536,
|
|
4258
4330
|
},
|
|
4331
|
+
"accounts/fireworks/models/qwen3p8-max": {
|
|
4332
|
+
id: "accounts/fireworks/models/qwen3p8-max",
|
|
4333
|
+
name: "Qwen3.8 Max",
|
|
4334
|
+
api: "anthropic-messages",
|
|
4335
|
+
provider: "fireworks",
|
|
4336
|
+
baseUrl: "https://api.fireworks.ai/inference",
|
|
4337
|
+
compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
|
|
4338
|
+
reasoning: true,
|
|
4339
|
+
input: ["text"],
|
|
4340
|
+
cost: {
|
|
4341
|
+
input: 2,
|
|
4342
|
+
output: 6,
|
|
4343
|
+
cacheRead: 0.25,
|
|
4344
|
+
cacheWrite: 0,
|
|
4345
|
+
},
|
|
4346
|
+
contextWindow: 262144,
|
|
4347
|
+
maxTokens: 131072,
|
|
4348
|
+
},
|
|
4259
4349
|
"accounts/fireworks/routers/glm-5p2-fast": {
|
|
4260
4350
|
id: "accounts/fireworks/routers/glm-5p2-fast",
|
|
4261
4351
|
name: "GLM 5.2 Fast",
|
|
@@ -5284,6 +5374,24 @@ export const MODELS = {
|
|
|
5284
5374
|
contextWindow: 1048576,
|
|
5285
5375
|
maxTokens: 65536,
|
|
5286
5376
|
},
|
|
5377
|
+
"gemini-3.7-flash": {
|
|
5378
|
+
id: "gemini-3.7-flash",
|
|
5379
|
+
name: "Gemini 3.7 Flash",
|
|
5380
|
+
api: "google-generative-ai",
|
|
5381
|
+
provider: "google",
|
|
5382
|
+
baseUrl: "https://generativelanguage.googleapis.com/v1beta",
|
|
5383
|
+
reasoning: true,
|
|
5384
|
+
thinkingLevelMap: { "off": null },
|
|
5385
|
+
input: ["text", "image"],
|
|
5386
|
+
cost: {
|
|
5387
|
+
input: 0.75,
|
|
5388
|
+
output: 3.75,
|
|
5389
|
+
cacheRead: 0.075,
|
|
5390
|
+
cacheWrite: 0,
|
|
5391
|
+
},
|
|
5392
|
+
contextWindow: 1048576,
|
|
5393
|
+
maxTokens: 65536,
|
|
5394
|
+
},
|
|
5287
5395
|
"gemini-flash-latest": {
|
|
5288
5396
|
id: "gemini-flash-latest",
|
|
5289
5397
|
name: "Gemini Flash Latest",
|
|
@@ -6191,6 +6299,24 @@ export const MODELS = {
|
|
|
6191
6299
|
contextWindow: 64000,
|
|
6192
6300
|
maxTokens: 8192,
|
|
6193
6301
|
},
|
|
6302
|
+
"deepseek-ai/DeepSeek-V3-0324": {
|
|
6303
|
+
id: "deepseek-ai/DeepSeek-V3-0324",
|
|
6304
|
+
name: "DeepSeek V3 0324",
|
|
6305
|
+
api: "openai-completions",
|
|
6306
|
+
provider: "huggingface",
|
|
6307
|
+
baseUrl: "https://router.huggingface.co/v1",
|
|
6308
|
+
compat: { "supportsDeveloperRole": false },
|
|
6309
|
+
reasoning: false,
|
|
6310
|
+
input: ["text"],
|
|
6311
|
+
cost: {
|
|
6312
|
+
input: 0.27,
|
|
6313
|
+
output: 1.12,
|
|
6314
|
+
cacheRead: 0,
|
|
6315
|
+
cacheWrite: 0,
|
|
6316
|
+
},
|
|
6317
|
+
contextWindow: 163840,
|
|
6318
|
+
maxTokens: 163840,
|
|
6319
|
+
},
|
|
6194
6320
|
"deepseek-ai/DeepSeek-V3.1": {
|
|
6195
6321
|
id: "deepseek-ai/DeepSeek-V3.1",
|
|
6196
6322
|
name: "DeepSeek-V3.1",
|
|
@@ -9077,6 +9203,24 @@ export const MODELS = {
|
|
|
9077
9203
|
contextWindow: 1048576,
|
|
9078
9204
|
maxTokens: 65536,
|
|
9079
9205
|
},
|
|
9206
|
+
"gemini-3.7-flash": {
|
|
9207
|
+
id: "gemini-3.7-flash",
|
|
9208
|
+
name: "Gemini 3.7 Flash",
|
|
9209
|
+
api: "google-generative-ai",
|
|
9210
|
+
provider: "opencode",
|
|
9211
|
+
baseUrl: "https://opencode.ai/zen/v1",
|
|
9212
|
+
reasoning: true,
|
|
9213
|
+
thinkingLevelMap: { "off": null },
|
|
9214
|
+
input: ["text", "image"],
|
|
9215
|
+
cost: {
|
|
9216
|
+
input: 1.5,
|
|
9217
|
+
output: 7.5,
|
|
9218
|
+
cacheRead: 0.15,
|
|
9219
|
+
cacheWrite: 0,
|
|
9220
|
+
},
|
|
9221
|
+
contextWindow: 1048576,
|
|
9222
|
+
maxTokens: 65536,
|
|
9223
|
+
},
|
|
9080
9224
|
"glm-5": {
|
|
9081
9225
|
id: "glm-5",
|
|
9082
9226
|
name: "GLM-5",
|
|
@@ -9647,23 +9791,6 @@ export const MODELS = {
|
|
|
9647
9791
|
contextWindow: 256000,
|
|
9648
9792
|
maxTokens: 32000,
|
|
9649
9793
|
},
|
|
9650
|
-
"ling-3.0-tiny-free": {
|
|
9651
|
-
id: "ling-3.0-tiny-free",
|
|
9652
|
-
name: "Ling-3.0-tiny Free",
|
|
9653
|
-
api: "openai-completions",
|
|
9654
|
-
provider: "opencode",
|
|
9655
|
-
baseUrl: "https://opencode.ai/zen/v1",
|
|
9656
|
-
reasoning: true,
|
|
9657
|
-
input: ["text"],
|
|
9658
|
-
cost: {
|
|
9659
|
-
input: 0,
|
|
9660
|
-
output: 0,
|
|
9661
|
-
cacheRead: 0,
|
|
9662
|
-
cacheWrite: 0,
|
|
9663
|
-
},
|
|
9664
|
-
contextWindow: 262144,
|
|
9665
|
-
maxTokens: 32768,
|
|
9666
|
-
},
|
|
9667
9794
|
"mimo-v2.5-free": {
|
|
9668
9795
|
id: "mimo-v2.5-free",
|
|
9669
9796
|
name: "MiMo V2.5 Free",
|
|
@@ -11061,7 +11188,7 @@ export const MODELS = {
|
|
|
11061
11188
|
cacheRead: 0,
|
|
11062
11189
|
cacheWrite: 0,
|
|
11063
11190
|
},
|
|
11064
|
-
contextWindow:
|
|
11191
|
+
contextWindow: 64000,
|
|
11065
11192
|
maxTokens: 16000,
|
|
11066
11193
|
},
|
|
11067
11194
|
"deepseek/deepseek-r1-0528": {
|
|
@@ -11162,13 +11289,13 @@ export const MODELS = {
|
|
|
11162
11289
|
thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
|
|
11163
11290
|
input: ["text"],
|
|
11164
11291
|
cost: {
|
|
11165
|
-
input: 0.
|
|
11166
|
-
output: 0.
|
|
11167
|
-
cacheRead: 0.
|
|
11292
|
+
input: 0.14,
|
|
11293
|
+
output: 0.28,
|
|
11294
|
+
cacheRead: 0.028,
|
|
11168
11295
|
cacheWrite: 0,
|
|
11169
11296
|
},
|
|
11170
11297
|
contextWindow: 1048576,
|
|
11171
|
-
maxTokens:
|
|
11298
|
+
maxTokens: 393216,
|
|
11172
11299
|
},
|
|
11173
11300
|
"deepseek/deepseek-v4-pro": {
|
|
11174
11301
|
id: "deepseek/deepseek-v4-pro",
|
|
@@ -11574,10 +11701,10 @@ export const MODELS = {
|
|
|
11574
11701
|
reasoning: true,
|
|
11575
11702
|
input: ["text", "image"],
|
|
11576
11703
|
cost: {
|
|
11577
|
-
input:
|
|
11578
|
-
output:
|
|
11579
|
-
cacheRead: 0.
|
|
11580
|
-
cacheWrite: 0.
|
|
11704
|
+
input: 0.75,
|
|
11705
|
+
output: 3.75,
|
|
11706
|
+
cacheRead: 0.075,
|
|
11707
|
+
cacheWrite: 0.0416666666666667,
|
|
11581
11708
|
},
|
|
11582
11709
|
contextWindow: 1048576,
|
|
11583
11710
|
maxTokens: 65536,
|
|
@@ -11591,10 +11718,44 @@ export const MODELS = {
|
|
|
11591
11718
|
reasoning: true,
|
|
11592
11719
|
input: ["text", "image"],
|
|
11593
11720
|
cost: {
|
|
11594
|
-
input: 0.
|
|
11595
|
-
output:
|
|
11596
|
-
cacheRead: 0.
|
|
11597
|
-
cacheWrite: 0.
|
|
11721
|
+
input: 0.375,
|
|
11722
|
+
output: 1.875,
|
|
11723
|
+
cacheRead: 0.0375,
|
|
11724
|
+
cacheWrite: 0.0416666666666667,
|
|
11725
|
+
},
|
|
11726
|
+
contextWindow: 1048576,
|
|
11727
|
+
maxTokens: 65536,
|
|
11728
|
+
},
|
|
11729
|
+
"google/gemini-3.7-flash": {
|
|
11730
|
+
id: "google/gemini-3.7-flash",
|
|
11731
|
+
name: "Google: Gemini 3.7 Flash",
|
|
11732
|
+
api: "openai-completions",
|
|
11733
|
+
provider: "openrouter",
|
|
11734
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
11735
|
+
reasoning: true,
|
|
11736
|
+
input: ["text", "image"],
|
|
11737
|
+
cost: {
|
|
11738
|
+
input: 0.375,
|
|
11739
|
+
output: 1.875,
|
|
11740
|
+
cacheRead: 0.0375,
|
|
11741
|
+
cacheWrite: 0.0208333333333333,
|
|
11742
|
+
},
|
|
11743
|
+
contextWindow: 1048576,
|
|
11744
|
+
maxTokens: 65536,
|
|
11745
|
+
},
|
|
11746
|
+
"google/gemini-3.7-flash:batch": {
|
|
11747
|
+
id: "google/gemini-3.7-flash:batch",
|
|
11748
|
+
name: "Google: Gemini 3.7 Flash (batch)",
|
|
11749
|
+
api: "openai-completions",
|
|
11750
|
+
provider: "openrouter",
|
|
11751
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
11752
|
+
reasoning: true,
|
|
11753
|
+
input: ["text", "image"],
|
|
11754
|
+
cost: {
|
|
11755
|
+
input: 0.1875,
|
|
11756
|
+
output: 0.9375,
|
|
11757
|
+
cacheRead: 0.01875,
|
|
11758
|
+
cacheWrite: 0.0208333333333333,
|
|
11598
11759
|
},
|
|
11599
11760
|
contextWindow: 1048576,
|
|
11600
11761
|
maxTokens: 65536,
|
|
@@ -11787,23 +11948,6 @@ export const MODELS = {
|
|
|
11787
11948
|
contextWindow: 262144,
|
|
11788
11949
|
maxTokens: 32768,
|
|
11789
11950
|
},
|
|
11790
|
-
"inclusionai/ling-3.0-tiny:free": {
|
|
11791
|
-
id: "inclusionai/ling-3.0-tiny:free",
|
|
11792
|
-
name: "inclusionAI: Ling 3.0 Tiny (free)",
|
|
11793
|
-
api: "openai-completions",
|
|
11794
|
-
provider: "openrouter",
|
|
11795
|
-
baseUrl: "https://openrouter.ai/api/v1",
|
|
11796
|
-
reasoning: true,
|
|
11797
|
-
input: ["text"],
|
|
11798
|
-
cost: {
|
|
11799
|
-
input: 0,
|
|
11800
|
-
output: 0,
|
|
11801
|
-
cacheRead: 0,
|
|
11802
|
-
cacheWrite: 0,
|
|
11803
|
-
},
|
|
11804
|
-
contextWindow: 262144,
|
|
11805
|
-
maxTokens: 32768,
|
|
11806
|
-
},
|
|
11807
11951
|
"inclusionai/ring-2.6-1t": {
|
|
11808
11952
|
id: "inclusionai/ring-2.6-1t",
|
|
11809
11953
|
name: "inclusionAI: Ring-2.6-1T",
|
|
@@ -11887,7 +12031,7 @@ export const MODELS = {
|
|
|
11887
12031
|
cacheWrite: 0,
|
|
11888
12032
|
},
|
|
11889
12033
|
contextWindow: 128000,
|
|
11890
|
-
maxTokens:
|
|
12034
|
+
maxTokens: 8192,
|
|
11891
12035
|
},
|
|
11892
12036
|
"meituan/longcat-2.0": {
|
|
11893
12037
|
id: "meituan/longcat-2.0",
|
|
@@ -14837,13 +14981,13 @@ export const MODELS = {
|
|
|
14837
14981
|
reasoning: false,
|
|
14838
14982
|
input: ["text"],
|
|
14839
14983
|
cost: {
|
|
14840
|
-
input: 0.
|
|
14984
|
+
input: 0.09999999999999999,
|
|
14841
14985
|
output: 1.1,
|
|
14842
|
-
cacheRead: 0,
|
|
14986
|
+
cacheRead: 0.07,
|
|
14843
14987
|
cacheWrite: 0,
|
|
14844
14988
|
},
|
|
14845
14989
|
contextWindow: 262144,
|
|
14846
|
-
maxTokens:
|
|
14990
|
+
maxTokens: 262144,
|
|
14847
14991
|
},
|
|
14848
14992
|
"qwen/qwen3-next-80b-a3b-thinking": {
|
|
14849
14993
|
id: "qwen/qwen3-next-80b-a3b-thinking",
|
|
@@ -14860,7 +15004,7 @@ export const MODELS = {
|
|
|
14860
15004
|
cacheWrite: 0,
|
|
14861
15005
|
},
|
|
14862
15006
|
contextWindow: 262144,
|
|
14863
|
-
maxTokens:
|
|
15007
|
+
maxTokens: 32768,
|
|
14864
15008
|
},
|
|
14865
15009
|
"qwen/qwen3-vl-235b-a22b-instruct": {
|
|
14866
15010
|
id: "qwen/qwen3-vl-235b-a22b-instruct",
|
|
@@ -14905,13 +15049,13 @@ export const MODELS = {
|
|
|
14905
15049
|
reasoning: false,
|
|
14906
15050
|
input: ["text", "image"],
|
|
14907
15051
|
cost: {
|
|
14908
|
-
input: 0.
|
|
14909
|
-
output: 0.
|
|
15052
|
+
input: 0.13,
|
|
15053
|
+
output: 0.52,
|
|
14910
15054
|
cacheRead: 0,
|
|
14911
15055
|
cacheWrite: 0,
|
|
14912
15056
|
},
|
|
14913
15057
|
contextWindow: 262144,
|
|
14914
|
-
maxTokens:
|
|
15058
|
+
maxTokens: 32768,
|
|
14915
15059
|
},
|
|
14916
15060
|
"qwen/qwen3-vl-30b-a3b-thinking": {
|
|
14917
15061
|
id: "qwen/qwen3-vl-30b-a3b-thinking",
|
|
@@ -15041,13 +15185,13 @@ export const MODELS = {
|
|
|
15041
15185
|
reasoning: true,
|
|
15042
15186
|
input: ["text", "image"],
|
|
15043
15187
|
cost: {
|
|
15044
|
-
input: 0.
|
|
15045
|
-
output: 3,
|
|
15046
|
-
cacheRead: 0.
|
|
15188
|
+
input: 0.5,
|
|
15189
|
+
output: 3.5999999999999996,
|
|
15190
|
+
cacheRead: 0.3,
|
|
15047
15191
|
cacheWrite: 0,
|
|
15048
15192
|
},
|
|
15049
15193
|
contextWindow: 262144,
|
|
15050
|
-
maxTokens:
|
|
15194
|
+
maxTokens: 262144,
|
|
15051
15195
|
},
|
|
15052
15196
|
"qwen/qwen3.5-9b": {
|
|
15053
15197
|
id: "qwen/qwen3.5-9b",
|
|
@@ -15264,11 +15408,11 @@ export const MODELS = {
|
|
|
15264
15408
|
cost: {
|
|
15265
15409
|
input: 2,
|
|
15266
15410
|
output: 6,
|
|
15267
|
-
cacheRead: 0.
|
|
15411
|
+
cacheRead: 0.25,
|
|
15268
15412
|
cacheWrite: 0,
|
|
15269
15413
|
},
|
|
15270
|
-
contextWindow:
|
|
15271
|
-
maxTokens:
|
|
15414
|
+
contextWindow: 1010000,
|
|
15415
|
+
maxTokens: 262144,
|
|
15272
15416
|
},
|
|
15273
15417
|
"qwen/qwen3.8-max": {
|
|
15274
15418
|
id: "qwen/qwen3.8-max",
|
|
@@ -15840,13 +15984,13 @@ export const MODELS = {
|
|
|
15840
15984
|
reasoning: true,
|
|
15841
15985
|
input: ["text"],
|
|
15842
15986
|
cost: {
|
|
15843
|
-
input: 0.
|
|
15844
|
-
output:
|
|
15845
|
-
cacheRead: 0.
|
|
15987
|
+
input: 0.63,
|
|
15988
|
+
output: 1.9800000000000002,
|
|
15989
|
+
cacheRead: 0.0945,
|
|
15846
15990
|
cacheWrite: 0,
|
|
15847
15991
|
},
|
|
15848
15992
|
contextWindow: 1048576,
|
|
15849
|
-
maxTokens:
|
|
15993
|
+
maxTokens: 4096,
|
|
15850
15994
|
},
|
|
15851
15995
|
"z-ai/glm-5.2:batch": {
|
|
15852
15996
|
id: "z-ai/glm-5.2:batch",
|
|
@@ -15982,10 +16126,10 @@ export const MODELS = {
|
|
|
15982
16126
|
reasoning: true,
|
|
15983
16127
|
input: ["text", "image"],
|
|
15984
16128
|
cost: {
|
|
15985
|
-
input:
|
|
15986
|
-
output:
|
|
15987
|
-
cacheRead: 0.
|
|
15988
|
-
cacheWrite: 0.
|
|
16129
|
+
input: 0.375,
|
|
16130
|
+
output: 1.875,
|
|
16131
|
+
cacheRead: 0.0375,
|
|
16132
|
+
cacheWrite: 0.0208333333333333,
|
|
15989
16133
|
},
|
|
15990
16134
|
contextWindow: 1048576,
|
|
15991
16135
|
maxTokens: 65536,
|
|
@@ -17465,6 +17609,23 @@ export const MODELS = {
|
|
|
17465
17609
|
contextWindow: 1000000,
|
|
17466
17610
|
maxTokens: 64000,
|
|
17467
17611
|
},
|
|
17612
|
+
"alibaba/qwen3.8-2.4t-a95b": {
|
|
17613
|
+
id: "alibaba/qwen3.8-2.4t-a95b",
|
|
17614
|
+
name: "Qwen3.8 2.4T A95B",
|
|
17615
|
+
api: "anthropic-messages",
|
|
17616
|
+
provider: "vercel-ai-gateway",
|
|
17617
|
+
baseUrl: "https://ai-gateway.vercel.sh",
|
|
17618
|
+
reasoning: true,
|
|
17619
|
+
input: ["text"],
|
|
17620
|
+
cost: {
|
|
17621
|
+
input: 2,
|
|
17622
|
+
output: 6,
|
|
17623
|
+
cacheRead: 0.25,
|
|
17624
|
+
cacheWrite: 0,
|
|
17625
|
+
},
|
|
17626
|
+
contextWindow: 262144,
|
|
17627
|
+
maxTokens: 131072,
|
|
17628
|
+
},
|
|
17468
17629
|
"alibaba/qwen3.8-max": {
|
|
17469
17630
|
id: "alibaba/qwen3.8-max",
|
|
17470
17631
|
name: "Qwen 3.8 Max",
|
|
@@ -18231,6 +18392,23 @@ export const MODELS = {
|
|
|
18231
18392
|
contextWindow: 1000000,
|
|
18232
18393
|
maxTokens: 64000,
|
|
18233
18394
|
},
|
|
18395
|
+
"google/gemini-3.7-flash": {
|
|
18396
|
+
id: "google/gemini-3.7-flash",
|
|
18397
|
+
name: "Gemini 3.7 Flash",
|
|
18398
|
+
api: "anthropic-messages",
|
|
18399
|
+
provider: "vercel-ai-gateway",
|
|
18400
|
+
baseUrl: "https://ai-gateway.vercel.sh",
|
|
18401
|
+
reasoning: true,
|
|
18402
|
+
input: ["text", "image"],
|
|
18403
|
+
cost: {
|
|
18404
|
+
input: 0.75,
|
|
18405
|
+
output: 3.75,
|
|
18406
|
+
cacheRead: 0.075,
|
|
18407
|
+
cacheWrite: 0,
|
|
18408
|
+
},
|
|
18409
|
+
contextWindow: 1000000,
|
|
18410
|
+
maxTokens: 65536,
|
|
18411
|
+
},
|
|
18234
18412
|
"google/gemma-4-26b-a4b-it": {
|
|
18235
18413
|
id: "google/gemma-4-26b-a4b-it",
|
|
18236
18414
|
name: "Google Gemma 4 26B A4B",
|
|
@@ -18262,7 +18440,7 @@ export const MODELS = {
|
|
|
18262
18440
|
cacheRead: 0,
|
|
18263
18441
|
cacheWrite: 0,
|
|
18264
18442
|
},
|
|
18265
|
-
contextWindow:
|
|
18443
|
+
contextWindow: 262144,
|
|
18266
18444
|
maxTokens: 131072,
|
|
18267
18445
|
},
|
|
18268
18446
|
"inception/mercury-2": {
|
|
@@ -18316,23 +18494,6 @@ export const MODELS = {
|
|
|
18316
18494
|
contextWindow: 256000,
|
|
18317
18495
|
maxTokens: 32000,
|
|
18318
18496
|
},
|
|
18319
|
-
"inclusionai/ling-3.0-tiny-free": {
|
|
18320
|
-
id: "inclusionai/ling-3.0-tiny-free",
|
|
18321
|
-
name: "Ling 3.0 Tiny (Free)",
|
|
18322
|
-
api: "anthropic-messages",
|
|
18323
|
-
provider: "vercel-ai-gateway",
|
|
18324
|
-
baseUrl: "https://ai-gateway.vercel.sh",
|
|
18325
|
-
reasoning: true,
|
|
18326
|
-
input: ["text"],
|
|
18327
|
-
cost: {
|
|
18328
|
-
input: 0,
|
|
18329
|
-
output: 0,
|
|
18330
|
-
cacheRead: 0,
|
|
18331
|
-
cacheWrite: 0,
|
|
18332
|
-
},
|
|
18333
|
-
contextWindow: 256000,
|
|
18334
|
-
maxTokens: 32000,
|
|
18335
|
-
},
|
|
18336
18497
|
"interfaze/interfaze-beta": {
|
|
18337
18498
|
id: "interfaze/interfaze-beta",
|
|
18338
18499
|
name: "Interfaze Beta",
|
|
@@ -20473,91 +20634,6 @@ export const MODELS = {
|
|
|
20473
20634
|
},
|
|
20474
20635
|
},
|
|
20475
20636
|
"xai": {
|
|
20476
|
-
"grok-3": {
|
|
20477
|
-
id: "grok-3",
|
|
20478
|
-
name: "Grok 3",
|
|
20479
|
-
api: "openai-completions",
|
|
20480
|
-
provider: "xai",
|
|
20481
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20482
|
-
reasoning: false,
|
|
20483
|
-
input: ["text"],
|
|
20484
|
-
cost: {
|
|
20485
|
-
input: 3,
|
|
20486
|
-
output: 15,
|
|
20487
|
-
cacheRead: 0.75,
|
|
20488
|
-
cacheWrite: 0,
|
|
20489
|
-
},
|
|
20490
|
-
contextWindow: 131072,
|
|
20491
|
-
maxTokens: 8192,
|
|
20492
|
-
},
|
|
20493
|
-
"grok-3-fast": {
|
|
20494
|
-
id: "grok-3-fast",
|
|
20495
|
-
name: "Grok 3 Fast",
|
|
20496
|
-
api: "openai-completions",
|
|
20497
|
-
provider: "xai",
|
|
20498
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20499
|
-
reasoning: false,
|
|
20500
|
-
input: ["text"],
|
|
20501
|
-
cost: {
|
|
20502
|
-
input: 5,
|
|
20503
|
-
output: 25,
|
|
20504
|
-
cacheRead: 1.25,
|
|
20505
|
-
cacheWrite: 0,
|
|
20506
|
-
},
|
|
20507
|
-
contextWindow: 131072,
|
|
20508
|
-
maxTokens: 8192,
|
|
20509
|
-
},
|
|
20510
|
-
"grok-4.20-0309-non-reasoning": {
|
|
20511
|
-
id: "grok-4.20-0309-non-reasoning",
|
|
20512
|
-
name: "Grok 4.20 (Non-Reasoning)",
|
|
20513
|
-
api: "openai-completions",
|
|
20514
|
-
provider: "xai",
|
|
20515
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20516
|
-
reasoning: false,
|
|
20517
|
-
input: ["text", "image"],
|
|
20518
|
-
cost: {
|
|
20519
|
-
input: 1.25,
|
|
20520
|
-
output: 2.5,
|
|
20521
|
-
cacheRead: 0.2,
|
|
20522
|
-
cacheWrite: 0,
|
|
20523
|
-
},
|
|
20524
|
-
contextWindow: 1000000,
|
|
20525
|
-
maxTokens: 30000,
|
|
20526
|
-
},
|
|
20527
|
-
"grok-4.20-0309-reasoning": {
|
|
20528
|
-
id: "grok-4.20-0309-reasoning",
|
|
20529
|
-
name: "Grok 4.20 (Reasoning)",
|
|
20530
|
-
api: "openai-completions",
|
|
20531
|
-
provider: "xai",
|
|
20532
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20533
|
-
reasoning: true,
|
|
20534
|
-
input: ["text", "image"],
|
|
20535
|
-
cost: {
|
|
20536
|
-
input: 1.25,
|
|
20537
|
-
output: 2.5,
|
|
20538
|
-
cacheRead: 0.2,
|
|
20539
|
-
cacheWrite: 0,
|
|
20540
|
-
},
|
|
20541
|
-
contextWindow: 1000000,
|
|
20542
|
-
maxTokens: 30000,
|
|
20543
|
-
},
|
|
20544
|
-
"grok-4.3": {
|
|
20545
|
-
id: "grok-4.3",
|
|
20546
|
-
name: "Grok 4.3",
|
|
20547
|
-
api: "openai-completions",
|
|
20548
|
-
provider: "xai",
|
|
20549
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20550
|
-
reasoning: true,
|
|
20551
|
-
input: ["text", "image"],
|
|
20552
|
-
cost: {
|
|
20553
|
-
input: 1.25,
|
|
20554
|
-
output: 2.5,
|
|
20555
|
-
cacheRead: 0.2,
|
|
20556
|
-
cacheWrite: 0,
|
|
20557
|
-
},
|
|
20558
|
-
contextWindow: 1000000,
|
|
20559
|
-
maxTokens: 30000,
|
|
20560
|
-
},
|
|
20561
20637
|
"grok-4.5": {
|
|
20562
20638
|
id: "grok-4.5",
|
|
20563
20639
|
name: "Grok 4.5",
|
|
@@ -20566,6 +20642,7 @@ export const MODELS = {
|
|
|
20566
20642
|
baseUrl: "https://api.x.ai/v1",
|
|
20567
20643
|
compat: { "supportsLongCacheRetention": false },
|
|
20568
20644
|
reasoning: true,
|
|
20645
|
+
defaultThinkingLevel: "high",
|
|
20569
20646
|
thinkingLevelMap: { "off": null, "minimal": null },
|
|
20570
20647
|
input: ["text", "image"],
|
|
20571
20648
|
cost: {
|
|
@@ -20580,10 +20657,13 @@ export const MODELS = {
|
|
|
20580
20657
|
"grok-4.6": {
|
|
20581
20658
|
id: "grok-4.6",
|
|
20582
20659
|
name: "Grok 4.6",
|
|
20583
|
-
api: "openai-
|
|
20660
|
+
api: "openai-responses",
|
|
20584
20661
|
provider: "xai",
|
|
20585
20662
|
baseUrl: "https://api.x.ai/v1",
|
|
20663
|
+
compat: { "supportsLongCacheRetention": false },
|
|
20586
20664
|
reasoning: true,
|
|
20665
|
+
defaultThinkingLevel: "high",
|
|
20666
|
+
thinkingLevelMap: { "off": null, "minimal": null, "xhigh": "xhigh" },
|
|
20587
20667
|
input: ["text", "image"],
|
|
20588
20668
|
cost: {
|
|
20589
20669
|
input: 2,
|
|
@@ -20594,40 +20674,6 @@ export const MODELS = {
|
|
|
20594
20674
|
contextWindow: 500000,
|
|
20595
20675
|
maxTokens: 500000,
|
|
20596
20676
|
},
|
|
20597
|
-
"grok-build-0.1": {
|
|
20598
|
-
id: "grok-build-0.1",
|
|
20599
|
-
name: "Grok Build 0.1",
|
|
20600
|
-
api: "openai-completions",
|
|
20601
|
-
provider: "xai",
|
|
20602
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20603
|
-
reasoning: true,
|
|
20604
|
-
input: ["text", "image"],
|
|
20605
|
-
cost: {
|
|
20606
|
-
input: 1,
|
|
20607
|
-
output: 2,
|
|
20608
|
-
cacheRead: 0.2,
|
|
20609
|
-
cacheWrite: 0,
|
|
20610
|
-
},
|
|
20611
|
-
contextWindow: 256000,
|
|
20612
|
-
maxTokens: 256000,
|
|
20613
|
-
},
|
|
20614
|
-
"grok-code-fast-1": {
|
|
20615
|
-
id: "grok-code-fast-1",
|
|
20616
|
-
name: "Grok Code Fast 1",
|
|
20617
|
-
api: "openai-completions",
|
|
20618
|
-
provider: "xai",
|
|
20619
|
-
baseUrl: "https://api.x.ai/v1",
|
|
20620
|
-
reasoning: false,
|
|
20621
|
-
input: ["text"],
|
|
20622
|
-
cost: {
|
|
20623
|
-
input: 0.2,
|
|
20624
|
-
output: 1.5,
|
|
20625
|
-
cacheRead: 0.02,
|
|
20626
|
-
cacheWrite: 0,
|
|
20627
|
-
},
|
|
20628
|
-
contextWindow: 32768,
|
|
20629
|
-
maxTokens: 8192,
|
|
20630
|
-
},
|
|
20631
20677
|
},
|
|
20632
20678
|
"xiaomi": {
|
|
20633
20679
|
"mimo-v2-flash": {
|