@caupulican/pi-ai 0.90.6 → 0.90.11

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (41) hide show
  1. package/README.md +1 -1
  2. package/dist/image-models.generated.d.ts +30 -0
  3. package/dist/image-models.generated.d.ts.map +1 -1
  4. package/dist/image-models.generated.js +30 -0
  5. package/dist/image-models.generated.js.map +1 -1
  6. package/dist/models.generated.d.ts +262 -162
  7. package/dist/models.generated.d.ts.map +1 -1
  8. package/dist/models.generated.js +276 -213
  9. package/dist/models.generated.js.map +1 -1
  10. package/dist/providers/openai-responses.js +6 -0
  11. package/dist/providers/openai-responses.js.map +1 -1
  12. package/dist/utils/oauth/anthropic.d.ts.map +1 -1
  13. package/dist/utils/oauth/anthropic.js +1 -0
  14. package/dist/utils/oauth/anthropic.js.map +1 -1
  15. package/dist/utils/oauth/device-code.d.ts +3 -0
  16. package/dist/utils/oauth/device-code.d.ts.map +1 -1
  17. package/dist/utils/oauth/device-code.js +15 -2
  18. package/dist/utils/oauth/device-code.js.map +1 -1
  19. package/dist/utils/oauth/github-copilot.d.ts.map +1 -1
  20. package/dist/utils/oauth/github-copilot.js +1 -0
  21. package/dist/utils/oauth/github-copilot.js.map +1 -1
  22. package/dist/utils/oauth/kimi-coding.d.ts.map +1 -1
  23. package/dist/utils/oauth/kimi-coding.js +2 -0
  24. package/dist/utils/oauth/kimi-coding.js.map +1 -1
  25. package/dist/utils/oauth/openai-codex.d.ts.map +1 -1
  26. package/dist/utils/oauth/openai-codex.js +1 -0
  27. package/dist/utils/oauth/openai-codex.js.map +1 -1
  28. package/dist/utils/oauth/types.d.ts +4 -0
  29. package/dist/utils/oauth/types.d.ts.map +1 -1
  30. package/dist/utils/oauth/types.js.map +1 -1
  31. package/dist/utils/oauth/xai.d.ts.map +1 -1
  32. package/dist/utils/oauth/xai.js +15 -3
  33. package/dist/utils/oauth/xai.js.map +1 -1
  34. package/dist/utils/tool-repair/registry.d.ts +1 -1
  35. package/dist/utils/tool-repair/registry.d.ts.map +1 -1
  36. package/dist/utils/tool-repair/registry.js +14 -1
  37. package/dist/utils/tool-repair/registry.js.map +1 -1
  38. package/dist/utils/tool-repair/repairer.d.ts.map +1 -1
  39. package/dist/utils/tool-repair/repairer.js +62 -9
  40. package/dist/utils/tool-repair/repairer.js.map +1 -1
  41. package/package.json +1 -1
@@ -4148,6 +4148,24 @@ export const MODELS = {
4148
4148
  contextWindow: 131072,
4149
4149
  maxTokens: 32768,
4150
4150
  },
4151
+ "accounts/fireworks/models/inkling": {
4152
+ id: "accounts/fireworks/models/inkling",
4153
+ name: "Inkling",
4154
+ api: "anthropic-messages",
4155
+ provider: "fireworks",
4156
+ baseUrl: "https://api.fireworks.ai/inference",
4157
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4158
+ reasoning: true,
4159
+ input: ["text", "image"],
4160
+ cost: {
4161
+ input: 1,
4162
+ output: 4.05,
4163
+ cacheRead: 0.17,
4164
+ cacheWrite: 0,
4165
+ },
4166
+ contextWindow: 1048576,
4167
+ maxTokens: 1048576,
4168
+ },
4151
4169
  "accounts/fireworks/models/kimi-k2p6": {
4152
4170
  id: "accounts/fireworks/models/kimi-k2p6",
4153
4171
  name: "Kimi K2.6",
@@ -4238,6 +4256,60 @@ export const MODELS = {
4238
4256
  contextWindow: 512000,
4239
4257
  maxTokens: 512000,
4240
4258
  },
4259
+ "accounts/fireworks/models/muse-glimmer-30b": {
4260
+ id: "accounts/fireworks/models/muse-glimmer-30b",
4261
+ name: "Muse Glimmer 30B",
4262
+ api: "anthropic-messages",
4263
+ provider: "fireworks",
4264
+ baseUrl: "https://api.fireworks.ai/inference",
4265
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4266
+ reasoning: true,
4267
+ input: ["text", "image"],
4268
+ cost: {
4269
+ input: 0.35,
4270
+ output: 1.5,
4271
+ cacheRead: 0.04,
4272
+ cacheWrite: 0,
4273
+ },
4274
+ contextWindow: 131072,
4275
+ maxTokens: 131072,
4276
+ },
4277
+ "accounts/fireworks/models/nemotron-3-ultra-nvfp4": {
4278
+ id: "accounts/fireworks/models/nemotron-3-ultra-nvfp4",
4279
+ name: "Nemotron 3 Ultra 550B A55B",
4280
+ api: "anthropic-messages",
4281
+ provider: "fireworks",
4282
+ baseUrl: "https://api.fireworks.ai/inference",
4283
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4284
+ reasoning: true,
4285
+ input: ["text"],
4286
+ cost: {
4287
+ input: 0.6,
4288
+ output: 2.4,
4289
+ cacheRead: 0.119,
4290
+ cacheWrite: 0,
4291
+ },
4292
+ contextWindow: 262144,
4293
+ maxTokens: 128000,
4294
+ },
4295
+ "accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b": {
4296
+ id: "accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b",
4297
+ name: "Nemotron 3.5 Lightning 30B A3B",
4298
+ api: "anthropic-messages",
4299
+ provider: "fireworks",
4300
+ baseUrl: "https://api.fireworks.ai/inference",
4301
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4302
+ reasoning: true,
4303
+ input: ["text"],
4304
+ cost: {
4305
+ input: 0.05,
4306
+ output: 0.2,
4307
+ cacheRead: 0.01,
4308
+ cacheWrite: 0,
4309
+ },
4310
+ contextWindow: 262144,
4311
+ maxTokens: 262144,
4312
+ },
4241
4313
  "accounts/fireworks/models/qwen3p7-plus": {
4242
4314
  id: "accounts/fireworks/models/qwen3p7-plus",
4243
4315
  name: "Qwen 3.7 Plus",
@@ -4256,6 +4328,24 @@ export const MODELS = {
4256
4328
  contextWindow: 262144,
4257
4329
  maxTokens: 65536,
4258
4330
  },
4331
+ "accounts/fireworks/models/qwen3p8-max": {
4332
+ id: "accounts/fireworks/models/qwen3p8-max",
4333
+ name: "Qwen3.8 Max",
4334
+ api: "anthropic-messages",
4335
+ provider: "fireworks",
4336
+ baseUrl: "https://api.fireworks.ai/inference",
4337
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4338
+ reasoning: true,
4339
+ input: ["text"],
4340
+ cost: {
4341
+ input: 2,
4342
+ output: 6,
4343
+ cacheRead: 0.25,
4344
+ cacheWrite: 0,
4345
+ },
4346
+ contextWindow: 262144,
4347
+ maxTokens: 131072,
4348
+ },
4259
4349
  "accounts/fireworks/routers/glm-5p2-fast": {
4260
4350
  id: "accounts/fireworks/routers/glm-5p2-fast",
4261
4351
  name: "GLM 5.2 Fast",
@@ -5284,6 +5374,24 @@ export const MODELS = {
5284
5374
  contextWindow: 1048576,
5285
5375
  maxTokens: 65536,
5286
5376
  },
5377
+ "gemini-3.7-flash": {
5378
+ id: "gemini-3.7-flash",
5379
+ name: "Gemini 3.7 Flash",
5380
+ api: "google-generative-ai",
5381
+ provider: "google",
5382
+ baseUrl: "https://generativelanguage.googleapis.com/v1beta",
5383
+ reasoning: true,
5384
+ thinkingLevelMap: { "off": null },
5385
+ input: ["text", "image"],
5386
+ cost: {
5387
+ input: 0.75,
5388
+ output: 3.75,
5389
+ cacheRead: 0.075,
5390
+ cacheWrite: 0,
5391
+ },
5392
+ contextWindow: 1048576,
5393
+ maxTokens: 65536,
5394
+ },
5287
5395
  "gemini-flash-latest": {
5288
5396
  id: "gemini-flash-latest",
5289
5397
  name: "Gemini Flash Latest",
@@ -6191,6 +6299,24 @@ export const MODELS = {
6191
6299
  contextWindow: 64000,
6192
6300
  maxTokens: 8192,
6193
6301
  },
6302
+ "deepseek-ai/DeepSeek-V3-0324": {
6303
+ id: "deepseek-ai/DeepSeek-V3-0324",
6304
+ name: "DeepSeek V3 0324",
6305
+ api: "openai-completions",
6306
+ provider: "huggingface",
6307
+ baseUrl: "https://router.huggingface.co/v1",
6308
+ compat: { "supportsDeveloperRole": false },
6309
+ reasoning: false,
6310
+ input: ["text"],
6311
+ cost: {
6312
+ input: 0.27,
6313
+ output: 1.12,
6314
+ cacheRead: 0,
6315
+ cacheWrite: 0,
6316
+ },
6317
+ contextWindow: 163840,
6318
+ maxTokens: 163840,
6319
+ },
6194
6320
  "deepseek-ai/DeepSeek-V3.1": {
6195
6321
  id: "deepseek-ai/DeepSeek-V3.1",
6196
6322
  name: "DeepSeek-V3.1",
@@ -9077,6 +9203,24 @@ export const MODELS = {
9077
9203
  contextWindow: 1048576,
9078
9204
  maxTokens: 65536,
9079
9205
  },
9206
+ "gemini-3.7-flash": {
9207
+ id: "gemini-3.7-flash",
9208
+ name: "Gemini 3.7 Flash",
9209
+ api: "google-generative-ai",
9210
+ provider: "opencode",
9211
+ baseUrl: "https://opencode.ai/zen/v1",
9212
+ reasoning: true,
9213
+ thinkingLevelMap: { "off": null },
9214
+ input: ["text", "image"],
9215
+ cost: {
9216
+ input: 1.5,
9217
+ output: 7.5,
9218
+ cacheRead: 0.15,
9219
+ cacheWrite: 0,
9220
+ },
9221
+ contextWindow: 1048576,
9222
+ maxTokens: 65536,
9223
+ },
9080
9224
  "glm-5": {
9081
9225
  id: "glm-5",
9082
9226
  name: "GLM-5",
@@ -9647,23 +9791,6 @@ export const MODELS = {
9647
9791
  contextWindow: 256000,
9648
9792
  maxTokens: 32000,
9649
9793
  },
9650
- "ling-3.0-tiny-free": {
9651
- id: "ling-3.0-tiny-free",
9652
- name: "Ling-3.0-tiny Free",
9653
- api: "openai-completions",
9654
- provider: "opencode",
9655
- baseUrl: "https://opencode.ai/zen/v1",
9656
- reasoning: true,
9657
- input: ["text"],
9658
- cost: {
9659
- input: 0,
9660
- output: 0,
9661
- cacheRead: 0,
9662
- cacheWrite: 0,
9663
- },
9664
- contextWindow: 262144,
9665
- maxTokens: 32768,
9666
- },
9667
9794
  "mimo-v2.5-free": {
9668
9795
  id: "mimo-v2.5-free",
9669
9796
  name: "MiMo V2.5 Free",
@@ -11061,7 +11188,7 @@ export const MODELS = {
11061
11188
  cacheRead: 0,
11062
11189
  cacheWrite: 0,
11063
11190
  },
11064
- contextWindow: 163840,
11191
+ contextWindow: 64000,
11065
11192
  maxTokens: 16000,
11066
11193
  },
11067
11194
  "deepseek/deepseek-r1-0528": {
@@ -11162,13 +11289,13 @@ export const MODELS = {
11162
11289
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
11163
11290
  input: ["text"],
11164
11291
  cost: {
11165
- input: 0.08,
11166
- output: 0.18,
11167
- cacheRead: 0.016,
11292
+ input: 0.14,
11293
+ output: 0.28,
11294
+ cacheRead: 0.028,
11168
11295
  cacheWrite: 0,
11169
11296
  },
11170
11297
  contextWindow: 1048576,
11171
- maxTokens: 384000,
11298
+ maxTokens: 393216,
11172
11299
  },
11173
11300
  "deepseek/deepseek-v4-pro": {
11174
11301
  id: "deepseek/deepseek-v4-pro",
@@ -11574,10 +11701,10 @@ export const MODELS = {
11574
11701
  reasoning: true,
11575
11702
  input: ["text", "image"],
11576
11703
  cost: {
11577
- input: 1.5,
11578
- output: 7.5,
11579
- cacheRead: 0.15,
11580
- cacheWrite: 0.0833333333333333,
11704
+ input: 0.75,
11705
+ output: 3.75,
11706
+ cacheRead: 0.075,
11707
+ cacheWrite: 0.0416666666666667,
11581
11708
  },
11582
11709
  contextWindow: 1048576,
11583
11710
  maxTokens: 65536,
@@ -11591,10 +11718,44 @@ export const MODELS = {
11591
11718
  reasoning: true,
11592
11719
  input: ["text", "image"],
11593
11720
  cost: {
11594
- input: 0.75,
11595
- output: 3.75,
11596
- cacheRead: 0.075,
11597
- cacheWrite: 0.0833333333333333,
11721
+ input: 0.375,
11722
+ output: 1.875,
11723
+ cacheRead: 0.0375,
11724
+ cacheWrite: 0.0416666666666667,
11725
+ },
11726
+ contextWindow: 1048576,
11727
+ maxTokens: 65536,
11728
+ },
11729
+ "google/gemini-3.7-flash": {
11730
+ id: "google/gemini-3.7-flash",
11731
+ name: "Google: Gemini 3.7 Flash",
11732
+ api: "openai-completions",
11733
+ provider: "openrouter",
11734
+ baseUrl: "https://openrouter.ai/api/v1",
11735
+ reasoning: true,
11736
+ input: ["text", "image"],
11737
+ cost: {
11738
+ input: 0.375,
11739
+ output: 1.875,
11740
+ cacheRead: 0.0375,
11741
+ cacheWrite: 0.0208333333333333,
11742
+ },
11743
+ contextWindow: 1048576,
11744
+ maxTokens: 65536,
11745
+ },
11746
+ "google/gemini-3.7-flash:batch": {
11747
+ id: "google/gemini-3.7-flash:batch",
11748
+ name: "Google: Gemini 3.7 Flash (batch)",
11749
+ api: "openai-completions",
11750
+ provider: "openrouter",
11751
+ baseUrl: "https://openrouter.ai/api/v1",
11752
+ reasoning: true,
11753
+ input: ["text", "image"],
11754
+ cost: {
11755
+ input: 0.1875,
11756
+ output: 0.9375,
11757
+ cacheRead: 0.01875,
11758
+ cacheWrite: 0.0208333333333333,
11598
11759
  },
11599
11760
  contextWindow: 1048576,
11600
11761
  maxTokens: 65536,
@@ -11787,23 +11948,6 @@ export const MODELS = {
11787
11948
  contextWindow: 262144,
11788
11949
  maxTokens: 32768,
11789
11950
  },
11790
- "inclusionai/ling-3.0-tiny:free": {
11791
- id: "inclusionai/ling-3.0-tiny:free",
11792
- name: "inclusionAI: Ling 3.0 Tiny (free)",
11793
- api: "openai-completions",
11794
- provider: "openrouter",
11795
- baseUrl: "https://openrouter.ai/api/v1",
11796
- reasoning: true,
11797
- input: ["text"],
11798
- cost: {
11799
- input: 0,
11800
- output: 0,
11801
- cacheRead: 0,
11802
- cacheWrite: 0,
11803
- },
11804
- contextWindow: 262144,
11805
- maxTokens: 32768,
11806
- },
11807
11951
  "inclusionai/ring-2.6-1t": {
11808
11952
  id: "inclusionai/ring-2.6-1t",
11809
11953
  name: "inclusionAI: Ring-2.6-1T",
@@ -11887,7 +12031,7 @@ export const MODELS = {
11887
12031
  cacheWrite: 0,
11888
12032
  },
11889
12033
  contextWindow: 128000,
11890
- maxTokens: 32768,
12034
+ maxTokens: 8192,
11891
12035
  },
11892
12036
  "meituan/longcat-2.0": {
11893
12037
  id: "meituan/longcat-2.0",
@@ -12511,9 +12655,9 @@ export const MODELS = {
12511
12655
  reasoning: true,
12512
12656
  input: ["text", "image"],
12513
12657
  cost: {
12514
- input: 0.5795,
12515
- output: 2.44,
12516
- cacheRead: 0.0976,
12658
+ input: 0.95,
12659
+ output: 4,
12660
+ cacheRead: 0.16,
12517
12661
  cacheWrite: 0,
12518
12662
  },
12519
12663
  contextWindow: 262144,
@@ -12740,6 +12884,23 @@ export const MODELS = {
12740
12884
  contextWindow: 1000000,
12741
12885
  maxTokens: 65536,
12742
12886
  },
12887
+ "nvidia/nemotron-3.5-lightning": {
12888
+ id: "nvidia/nemotron-3.5-lightning",
12889
+ name: "NVIDIA: Nemotron 3.5 Lightning",
12890
+ api: "openai-completions",
12891
+ provider: "openrouter",
12892
+ baseUrl: "https://openrouter.ai/api/v1",
12893
+ reasoning: true,
12894
+ input: ["text"],
12895
+ cost: {
12896
+ input: 0.09999999999999999,
12897
+ output: 0.25,
12898
+ cacheRead: 0.049999999999999996,
12899
+ cacheWrite: 0,
12900
+ },
12901
+ contextWindow: 1000000,
12902
+ maxTokens: 262144,
12903
+ },
12743
12904
  "nvidia/nemotron-3.5-lightning:free": {
12744
12905
  id: "nvidia/nemotron-3.5-lightning:free",
12745
12906
  name: "NVIDIA: Nemotron 3.5 Lightning (free)",
@@ -14820,13 +14981,13 @@ export const MODELS = {
14820
14981
  reasoning: false,
14821
14982
  input: ["text"],
14822
14983
  cost: {
14823
- input: 0.09,
14984
+ input: 0.09999999999999999,
14824
14985
  output: 1.1,
14825
- cacheRead: 0,
14986
+ cacheRead: 0.07,
14826
14987
  cacheWrite: 0,
14827
14988
  },
14828
14989
  contextWindow: 262144,
14829
- maxTokens: 16384,
14990
+ maxTokens: 262144,
14830
14991
  },
14831
14992
  "qwen/qwen3-next-80b-a3b-thinking": {
14832
14993
  id: "qwen/qwen3-next-80b-a3b-thinking",
@@ -14843,7 +15004,7 @@ export const MODELS = {
14843
15004
  cacheWrite: 0,
14844
15005
  },
14845
15006
  contextWindow: 262144,
14846
- maxTokens: 262144,
15007
+ maxTokens: 32768,
14847
15008
  },
14848
15009
  "qwen/qwen3-vl-235b-a22b-instruct": {
14849
15010
  id: "qwen/qwen3-vl-235b-a22b-instruct",
@@ -14888,13 +15049,13 @@ export const MODELS = {
14888
15049
  reasoning: false,
14889
15050
  input: ["text", "image"],
14890
15051
  cost: {
14891
- input: 0.15,
14892
- output: 0.6,
15052
+ input: 0.13,
15053
+ output: 0.52,
14893
15054
  cacheRead: 0,
14894
15055
  cacheWrite: 0,
14895
15056
  },
14896
15057
  contextWindow: 262144,
14897
- maxTokens: 16384,
15058
+ maxTokens: 32768,
14898
15059
  },
14899
15060
  "qwen/qwen3-vl-30b-a3b-thinking": {
14900
15061
  id: "qwen/qwen3-vl-30b-a3b-thinking",
@@ -15007,9 +15168,9 @@ export const MODELS = {
15007
15168
  reasoning: true,
15008
15169
  input: ["text", "image"],
15009
15170
  cost: {
15010
- input: 0.14,
15011
- output: 1,
15012
- cacheRead: 0,
15171
+ input: 0.25,
15172
+ output: 1.25,
15173
+ cacheRead: 0.25,
15013
15174
  cacheWrite: 0,
15014
15175
  },
15015
15176
  contextWindow: 262144,
@@ -15247,11 +15408,11 @@ export const MODELS = {
15247
15408
  cost: {
15248
15409
  input: 2,
15249
15410
  output: 6,
15250
- cacheRead: 0.19999999999999998,
15411
+ cacheRead: 0.25,
15251
15412
  cacheWrite: 0,
15252
15413
  },
15253
- contextWindow: 262144,
15254
- maxTokens: 52429,
15414
+ contextWindow: 1010000,
15415
+ maxTokens: 262144,
15255
15416
  },
15256
15417
  "qwen/qwen3.8-max": {
15257
15418
  id: "qwen/qwen3.8-max",
@@ -15704,9 +15865,9 @@ export const MODELS = {
15704
15865
  reasoning: true,
15705
15866
  input: ["text"],
15706
15867
  cost: {
15707
- input: 0.55,
15708
- output: 2.2,
15709
- cacheRead: 0.11,
15868
+ input: 0.5,
15869
+ output: 2,
15870
+ cacheRead: 0.09999999999999999,
15710
15871
  cacheWrite: 0,
15711
15872
  },
15712
15873
  contextWindow: 204800,
@@ -15823,13 +15984,13 @@ export const MODELS = {
15823
15984
  reasoning: true,
15824
15985
  input: ["text"],
15825
15986
  cost: {
15826
- input: 0.49,
15827
- output: 1.54,
15828
- cacheRead: 0.091,
15987
+ input: 0.63,
15988
+ output: 1.9800000000000002,
15989
+ cacheRead: 0.0945,
15829
15990
  cacheWrite: 0,
15830
15991
  },
15831
15992
  contextWindow: 1048576,
15832
- maxTokens: 131072,
15993
+ maxTokens: 4096,
15833
15994
  },
15834
15995
  "z-ai/glm-5.2:batch": {
15835
15996
  id: "z-ai/glm-5.2:batch",
@@ -15965,10 +16126,10 @@ export const MODELS = {
15965
16126
  reasoning: true,
15966
16127
  input: ["text", "image"],
15967
16128
  cost: {
15968
- input: 1.5,
15969
- output: 7.5,
15970
- cacheRead: 0.15,
15971
- cacheWrite: 0.0833333333333333,
16129
+ input: 0.375,
16130
+ output: 1.875,
16131
+ cacheRead: 0.0375,
16132
+ cacheWrite: 0.0208333333333333,
15972
16133
  },
15973
16134
  contextWindow: 1048576,
15974
16135
  maxTokens: 65536,
@@ -17448,6 +17609,23 @@ export const MODELS = {
17448
17609
  contextWindow: 1000000,
17449
17610
  maxTokens: 64000,
17450
17611
  },
17612
+ "alibaba/qwen3.8-2.4t-a95b": {
17613
+ id: "alibaba/qwen3.8-2.4t-a95b",
17614
+ name: "Qwen3.8 2.4T A95B",
17615
+ api: "anthropic-messages",
17616
+ provider: "vercel-ai-gateway",
17617
+ baseUrl: "https://ai-gateway.vercel.sh",
17618
+ reasoning: true,
17619
+ input: ["text"],
17620
+ cost: {
17621
+ input: 2,
17622
+ output: 6,
17623
+ cacheRead: 0.25,
17624
+ cacheWrite: 0,
17625
+ },
17626
+ contextWindow: 262144,
17627
+ maxTokens: 131072,
17628
+ },
17451
17629
  "alibaba/qwen3.8-max": {
17452
17630
  id: "alibaba/qwen3.8-max",
17453
17631
  name: "Qwen 3.8 Max",
@@ -18214,6 +18392,23 @@ export const MODELS = {
18214
18392
  contextWindow: 1000000,
18215
18393
  maxTokens: 64000,
18216
18394
  },
18395
+ "google/gemini-3.7-flash": {
18396
+ id: "google/gemini-3.7-flash",
18397
+ name: "Gemini 3.7 Flash",
18398
+ api: "anthropic-messages",
18399
+ provider: "vercel-ai-gateway",
18400
+ baseUrl: "https://ai-gateway.vercel.sh",
18401
+ reasoning: true,
18402
+ input: ["text", "image"],
18403
+ cost: {
18404
+ input: 0.75,
18405
+ output: 3.75,
18406
+ cacheRead: 0.075,
18407
+ cacheWrite: 0,
18408
+ },
18409
+ contextWindow: 1000000,
18410
+ maxTokens: 65536,
18411
+ },
18217
18412
  "google/gemma-4-26b-a4b-it": {
18218
18413
  id: "google/gemma-4-26b-a4b-it",
18219
18414
  name: "Google Gemma 4 26B A4B",
@@ -18245,7 +18440,7 @@ export const MODELS = {
18245
18440
  cacheRead: 0,
18246
18441
  cacheWrite: 0,
18247
18442
  },
18248
- contextWindow: 256000,
18443
+ contextWindow: 262144,
18249
18444
  maxTokens: 131072,
18250
18445
  },
18251
18446
  "inception/mercury-2": {
@@ -18299,23 +18494,6 @@ export const MODELS = {
18299
18494
  contextWindow: 256000,
18300
18495
  maxTokens: 32000,
18301
18496
  },
18302
- "inclusionai/ling-3.0-tiny-free": {
18303
- id: "inclusionai/ling-3.0-tiny-free",
18304
- name: "Ling 3.0 Tiny (Free)",
18305
- api: "anthropic-messages",
18306
- provider: "vercel-ai-gateway",
18307
- baseUrl: "https://ai-gateway.vercel.sh",
18308
- reasoning: true,
18309
- input: ["text"],
18310
- cost: {
18311
- input: 0,
18312
- output: 0,
18313
- cacheRead: 0,
18314
- cacheWrite: 0,
18315
- },
18316
- contextWindow: 256000,
18317
- maxTokens: 32000,
18318
- },
18319
18497
  "interfaze/interfaze-beta": {
18320
18498
  id: "interfaze/interfaze-beta",
18321
18499
  name: "Interfaze Beta",
@@ -20456,91 +20634,6 @@ export const MODELS = {
20456
20634
  },
20457
20635
  },
20458
20636
  "xai": {
20459
- "grok-3": {
20460
- id: "grok-3",
20461
- name: "Grok 3",
20462
- api: "openai-completions",
20463
- provider: "xai",
20464
- baseUrl: "https://api.x.ai/v1",
20465
- reasoning: false,
20466
- input: ["text"],
20467
- cost: {
20468
- input: 3,
20469
- output: 15,
20470
- cacheRead: 0.75,
20471
- cacheWrite: 0,
20472
- },
20473
- contextWindow: 131072,
20474
- maxTokens: 8192,
20475
- },
20476
- "grok-3-fast": {
20477
- id: "grok-3-fast",
20478
- name: "Grok 3 Fast",
20479
- api: "openai-completions",
20480
- provider: "xai",
20481
- baseUrl: "https://api.x.ai/v1",
20482
- reasoning: false,
20483
- input: ["text"],
20484
- cost: {
20485
- input: 5,
20486
- output: 25,
20487
- cacheRead: 1.25,
20488
- cacheWrite: 0,
20489
- },
20490
- contextWindow: 131072,
20491
- maxTokens: 8192,
20492
- },
20493
- "grok-4.20-0309-non-reasoning": {
20494
- id: "grok-4.20-0309-non-reasoning",
20495
- name: "Grok 4.20 (Non-Reasoning)",
20496
- api: "openai-completions",
20497
- provider: "xai",
20498
- baseUrl: "https://api.x.ai/v1",
20499
- reasoning: false,
20500
- input: ["text", "image"],
20501
- cost: {
20502
- input: 1.25,
20503
- output: 2.5,
20504
- cacheRead: 0.2,
20505
- cacheWrite: 0,
20506
- },
20507
- contextWindow: 1000000,
20508
- maxTokens: 30000,
20509
- },
20510
- "grok-4.20-0309-reasoning": {
20511
- id: "grok-4.20-0309-reasoning",
20512
- name: "Grok 4.20 (Reasoning)",
20513
- api: "openai-completions",
20514
- provider: "xai",
20515
- baseUrl: "https://api.x.ai/v1",
20516
- reasoning: true,
20517
- input: ["text", "image"],
20518
- cost: {
20519
- input: 1.25,
20520
- output: 2.5,
20521
- cacheRead: 0.2,
20522
- cacheWrite: 0,
20523
- },
20524
- contextWindow: 1000000,
20525
- maxTokens: 30000,
20526
- },
20527
- "grok-4.3": {
20528
- id: "grok-4.3",
20529
- name: "Grok 4.3",
20530
- api: "openai-completions",
20531
- provider: "xai",
20532
- baseUrl: "https://api.x.ai/v1",
20533
- reasoning: true,
20534
- input: ["text", "image"],
20535
- cost: {
20536
- input: 1.25,
20537
- output: 2.5,
20538
- cacheRead: 0.2,
20539
- cacheWrite: 0,
20540
- },
20541
- contextWindow: 1000000,
20542
- maxTokens: 30000,
20543
- },
20544
20637
  "grok-4.5": {
20545
20638
  id: "grok-4.5",
20546
20639
  name: "Grok 4.5",
@@ -20549,6 +20642,7 @@ export const MODELS = {
20549
20642
  baseUrl: "https://api.x.ai/v1",
20550
20643
  compat: { "supportsLongCacheRetention": false },
20551
20644
  reasoning: true,
20645
+ defaultThinkingLevel: "high",
20552
20646
  thinkingLevelMap: { "off": null, "minimal": null },
20553
20647
  input: ["text", "image"],
20554
20648
  cost: {
@@ -20563,10 +20657,13 @@ export const MODELS = {
20563
20657
  "grok-4.6": {
20564
20658
  id: "grok-4.6",
20565
20659
  name: "Grok 4.6",
20566
- api: "openai-completions",
20660
+ api: "openai-responses",
20567
20661
  provider: "xai",
20568
20662
  baseUrl: "https://api.x.ai/v1",
20663
+ compat: { "supportsLongCacheRetention": false },
20569
20664
  reasoning: true,
20665
+ defaultThinkingLevel: "high",
20666
+ thinkingLevelMap: { "off": null, "minimal": null, "xhigh": "xhigh" },
20570
20667
  input: ["text", "image"],
20571
20668
  cost: {
20572
20669
  input: 2,
@@ -20577,40 +20674,6 @@ export const MODELS = {
20577
20674
  contextWindow: 500000,
20578
20675
  maxTokens: 500000,
20579
20676
  },
20580
- "grok-build-0.1": {
20581
- id: "grok-build-0.1",
20582
- name: "Grok Build 0.1",
20583
- api: "openai-completions",
20584
- provider: "xai",
20585
- baseUrl: "https://api.x.ai/v1",
20586
- reasoning: true,
20587
- input: ["text", "image"],
20588
- cost: {
20589
- input: 1,
20590
- output: 2,
20591
- cacheRead: 0.2,
20592
- cacheWrite: 0,
20593
- },
20594
- contextWindow: 256000,
20595
- maxTokens: 256000,
20596
- },
20597
- "grok-code-fast-1": {
20598
- id: "grok-code-fast-1",
20599
- name: "Grok Code Fast 1",
20600
- api: "openai-completions",
20601
- provider: "xai",
20602
- baseUrl: "https://api.x.ai/v1",
20603
- reasoning: false,
20604
- input: ["text"],
20605
- cost: {
20606
- input: 0.2,
20607
- output: 1.5,
20608
- cacheRead: 0.02,
20609
- cacheWrite: 0,
20610
- },
20611
- contextWindow: 32768,
20612
- maxTokens: 8192,
20613
- },
20614
20677
  },
20615
20678
  "xiaomi": {
20616
20679
  "mimo-v2-flash": {