@caupulican/pi-ai 0.90.7 → 0.90.12

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (41) hide show
  1. package/README.md +1 -1
  2. package/dist/image-models.generated.d.ts +15 -0
  3. package/dist/image-models.generated.d.ts.map +1 -1
  4. package/dist/image-models.generated.js +15 -0
  5. package/dist/image-models.generated.js.map +1 -1
  6. package/dist/models.generated.d.ts +245 -162
  7. package/dist/models.generated.d.ts.map +1 -1
  8. package/dist/models.generated.js +254 -208
  9. package/dist/models.generated.js.map +1 -1
  10. package/dist/providers/openai-responses.js +6 -0
  11. package/dist/providers/openai-responses.js.map +1 -1
  12. package/dist/utils/oauth/anthropic.d.ts.map +1 -1
  13. package/dist/utils/oauth/anthropic.js +1 -0
  14. package/dist/utils/oauth/anthropic.js.map +1 -1
  15. package/dist/utils/oauth/device-code.d.ts +3 -0
  16. package/dist/utils/oauth/device-code.d.ts.map +1 -1
  17. package/dist/utils/oauth/device-code.js +15 -2
  18. package/dist/utils/oauth/device-code.js.map +1 -1
  19. package/dist/utils/oauth/github-copilot.d.ts.map +1 -1
  20. package/dist/utils/oauth/github-copilot.js +1 -0
  21. package/dist/utils/oauth/github-copilot.js.map +1 -1
  22. package/dist/utils/oauth/kimi-coding.d.ts.map +1 -1
  23. package/dist/utils/oauth/kimi-coding.js +2 -0
  24. package/dist/utils/oauth/kimi-coding.js.map +1 -1
  25. package/dist/utils/oauth/openai-codex.d.ts.map +1 -1
  26. package/dist/utils/oauth/openai-codex.js +1 -0
  27. package/dist/utils/oauth/openai-codex.js.map +1 -1
  28. package/dist/utils/oauth/types.d.ts +4 -0
  29. package/dist/utils/oauth/types.d.ts.map +1 -1
  30. package/dist/utils/oauth/types.js.map +1 -1
  31. package/dist/utils/oauth/xai.d.ts.map +1 -1
  32. package/dist/utils/oauth/xai.js +15 -3
  33. package/dist/utils/oauth/xai.js.map +1 -1
  34. package/dist/utils/tool-repair/registry.d.ts +1 -1
  35. package/dist/utils/tool-repair/registry.d.ts.map +1 -1
  36. package/dist/utils/tool-repair/registry.js +14 -1
  37. package/dist/utils/tool-repair/registry.js.map +1 -1
  38. package/dist/utils/tool-repair/repairer.d.ts.map +1 -1
  39. package/dist/utils/tool-repair/repairer.js +62 -9
  40. package/dist/utils/tool-repair/repairer.js.map +1 -1
  41. package/package.json +1 -1
@@ -4148,6 +4148,24 @@ export const MODELS = {
4148
4148
  contextWindow: 131072,
4149
4149
  maxTokens: 32768,
4150
4150
  },
4151
+ "accounts/fireworks/models/inkling": {
4152
+ id: "accounts/fireworks/models/inkling",
4153
+ name: "Inkling",
4154
+ api: "anthropic-messages",
4155
+ provider: "fireworks",
4156
+ baseUrl: "https://api.fireworks.ai/inference",
4157
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4158
+ reasoning: true,
4159
+ input: ["text", "image"],
4160
+ cost: {
4161
+ input: 1,
4162
+ output: 4.05,
4163
+ cacheRead: 0.17,
4164
+ cacheWrite: 0,
4165
+ },
4166
+ contextWindow: 1048576,
4167
+ maxTokens: 1048576,
4168
+ },
4151
4169
  "accounts/fireworks/models/kimi-k2p6": {
4152
4170
  id: "accounts/fireworks/models/kimi-k2p6",
4153
4171
  name: "Kimi K2.6",
@@ -4238,6 +4256,60 @@ export const MODELS = {
4238
4256
  contextWindow: 512000,
4239
4257
  maxTokens: 512000,
4240
4258
  },
4259
+ "accounts/fireworks/models/muse-glimmer-30b": {
4260
+ id: "accounts/fireworks/models/muse-glimmer-30b",
4261
+ name: "Muse Glimmer 30B",
4262
+ api: "anthropic-messages",
4263
+ provider: "fireworks",
4264
+ baseUrl: "https://api.fireworks.ai/inference",
4265
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4266
+ reasoning: true,
4267
+ input: ["text", "image"],
4268
+ cost: {
4269
+ input: 0.35,
4270
+ output: 1.5,
4271
+ cacheRead: 0.04,
4272
+ cacheWrite: 0,
4273
+ },
4274
+ contextWindow: 131072,
4275
+ maxTokens: 131072,
4276
+ },
4277
+ "accounts/fireworks/models/nemotron-3-ultra-nvfp4": {
4278
+ id: "accounts/fireworks/models/nemotron-3-ultra-nvfp4",
4279
+ name: "Nemotron 3 Ultra 550B A55B",
4280
+ api: "anthropic-messages",
4281
+ provider: "fireworks",
4282
+ baseUrl: "https://api.fireworks.ai/inference",
4283
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4284
+ reasoning: true,
4285
+ input: ["text"],
4286
+ cost: {
4287
+ input: 0.6,
4288
+ output: 2.4,
4289
+ cacheRead: 0.119,
4290
+ cacheWrite: 0,
4291
+ },
4292
+ contextWindow: 262144,
4293
+ maxTokens: 128000,
4294
+ },
4295
+ "accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b": {
4296
+ id: "accounts/fireworks/models/nemotron-lightning-3p5-30b-a3b",
4297
+ name: "Nemotron 3.5 Lightning 30B A3B",
4298
+ api: "anthropic-messages",
4299
+ provider: "fireworks",
4300
+ baseUrl: "https://api.fireworks.ai/inference",
4301
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4302
+ reasoning: true,
4303
+ input: ["text"],
4304
+ cost: {
4305
+ input: 0.05,
4306
+ output: 0.2,
4307
+ cacheRead: 0.01,
4308
+ cacheWrite: 0,
4309
+ },
4310
+ contextWindow: 262144,
4311
+ maxTokens: 262144,
4312
+ },
4241
4313
  "accounts/fireworks/models/qwen3p7-plus": {
4242
4314
  id: "accounts/fireworks/models/qwen3p7-plus",
4243
4315
  name: "Qwen 3.7 Plus",
@@ -4256,6 +4328,24 @@ export const MODELS = {
4256
4328
  contextWindow: 262144,
4257
4329
  maxTokens: 65536,
4258
4330
  },
4331
+ "accounts/fireworks/models/qwen3p8-max": {
4332
+ id: "accounts/fireworks/models/qwen3p8-max",
4333
+ name: "Qwen3.8 Max",
4334
+ api: "anthropic-messages",
4335
+ provider: "fireworks",
4336
+ baseUrl: "https://api.fireworks.ai/inference",
4337
+ compat: { "sendSessionAffinityHeaders": true, "supportsEagerToolInputStreaming": false, "supportsCacheControlOnTools": false, "supportsLongCacheRetention": false },
4338
+ reasoning: true,
4339
+ input: ["text"],
4340
+ cost: {
4341
+ input: 2,
4342
+ output: 6,
4343
+ cacheRead: 0.25,
4344
+ cacheWrite: 0,
4345
+ },
4346
+ contextWindow: 262144,
4347
+ maxTokens: 131072,
4348
+ },
4259
4349
  "accounts/fireworks/routers/glm-5p2-fast": {
4260
4350
  id: "accounts/fireworks/routers/glm-5p2-fast",
4261
4351
  name: "GLM 5.2 Fast",
@@ -5284,6 +5374,24 @@ export const MODELS = {
5284
5374
  contextWindow: 1048576,
5285
5375
  maxTokens: 65536,
5286
5376
  },
5377
+ "gemini-3.7-flash": {
5378
+ id: "gemini-3.7-flash",
5379
+ name: "Gemini 3.7 Flash",
5380
+ api: "google-generative-ai",
5381
+ provider: "google",
5382
+ baseUrl: "https://generativelanguage.googleapis.com/v1beta",
5383
+ reasoning: true,
5384
+ thinkingLevelMap: { "off": null },
5385
+ input: ["text", "image"],
5386
+ cost: {
5387
+ input: 0.75,
5388
+ output: 3.75,
5389
+ cacheRead: 0.075,
5390
+ cacheWrite: 0,
5391
+ },
5392
+ contextWindow: 1048576,
5393
+ maxTokens: 65536,
5394
+ },
5287
5395
  "gemini-flash-latest": {
5288
5396
  id: "gemini-flash-latest",
5289
5397
  name: "Gemini Flash Latest",
@@ -6191,6 +6299,24 @@ export const MODELS = {
6191
6299
  contextWindow: 64000,
6192
6300
  maxTokens: 8192,
6193
6301
  },
6302
+ "deepseek-ai/DeepSeek-V3-0324": {
6303
+ id: "deepseek-ai/DeepSeek-V3-0324",
6304
+ name: "DeepSeek V3 0324",
6305
+ api: "openai-completions",
6306
+ provider: "huggingface",
6307
+ baseUrl: "https://router.huggingface.co/v1",
6308
+ compat: { "supportsDeveloperRole": false },
6309
+ reasoning: false,
6310
+ input: ["text"],
6311
+ cost: {
6312
+ input: 0.27,
6313
+ output: 1.12,
6314
+ cacheRead: 0,
6315
+ cacheWrite: 0,
6316
+ },
6317
+ contextWindow: 163840,
6318
+ maxTokens: 163840,
6319
+ },
6194
6320
  "deepseek-ai/DeepSeek-V3.1": {
6195
6321
  id: "deepseek-ai/DeepSeek-V3.1",
6196
6322
  name: "DeepSeek-V3.1",
@@ -9077,6 +9203,24 @@ export const MODELS = {
9077
9203
  contextWindow: 1048576,
9078
9204
  maxTokens: 65536,
9079
9205
  },
9206
+ "gemini-3.7-flash": {
9207
+ id: "gemini-3.7-flash",
9208
+ name: "Gemini 3.7 Flash",
9209
+ api: "google-generative-ai",
9210
+ provider: "opencode",
9211
+ baseUrl: "https://opencode.ai/zen/v1",
9212
+ reasoning: true,
9213
+ thinkingLevelMap: { "off": null },
9214
+ input: ["text", "image"],
9215
+ cost: {
9216
+ input: 1.5,
9217
+ output: 7.5,
9218
+ cacheRead: 0.15,
9219
+ cacheWrite: 0,
9220
+ },
9221
+ contextWindow: 1048576,
9222
+ maxTokens: 65536,
9223
+ },
9080
9224
  "glm-5": {
9081
9225
  id: "glm-5",
9082
9226
  name: "GLM-5",
@@ -9647,23 +9791,6 @@ export const MODELS = {
9647
9791
  contextWindow: 256000,
9648
9792
  maxTokens: 32000,
9649
9793
  },
9650
- "ling-3.0-tiny-free": {
9651
- id: "ling-3.0-tiny-free",
9652
- name: "Ling-3.0-tiny Free",
9653
- api: "openai-completions",
9654
- provider: "opencode",
9655
- baseUrl: "https://opencode.ai/zen/v1",
9656
- reasoning: true,
9657
- input: ["text"],
9658
- cost: {
9659
- input: 0,
9660
- output: 0,
9661
- cacheRead: 0,
9662
- cacheWrite: 0,
9663
- },
9664
- contextWindow: 262144,
9665
- maxTokens: 32768,
9666
- },
9667
9794
  "mimo-v2.5-free": {
9668
9795
  id: "mimo-v2.5-free",
9669
9796
  name: "MiMo V2.5 Free",
@@ -11061,7 +11188,7 @@ export const MODELS = {
11061
11188
  cacheRead: 0,
11062
11189
  cacheWrite: 0,
11063
11190
  },
11064
- contextWindow: 163840,
11191
+ contextWindow: 64000,
11065
11192
  maxTokens: 16000,
11066
11193
  },
11067
11194
  "deepseek/deepseek-r1-0528": {
@@ -11162,13 +11289,13 @@ export const MODELS = {
11162
11289
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
11163
11290
  input: ["text"],
11164
11291
  cost: {
11165
- input: 0.08,
11166
- output: 0.18,
11167
- cacheRead: 0.016,
11292
+ input: 0.14,
11293
+ output: 0.28,
11294
+ cacheRead: 0.028,
11168
11295
  cacheWrite: 0,
11169
11296
  },
11170
11297
  contextWindow: 1048576,
11171
- maxTokens: 384000,
11298
+ maxTokens: 393216,
11172
11299
  },
11173
11300
  "deepseek/deepseek-v4-pro": {
11174
11301
  id: "deepseek/deepseek-v4-pro",
@@ -11574,10 +11701,10 @@ export const MODELS = {
11574
11701
  reasoning: true,
11575
11702
  input: ["text", "image"],
11576
11703
  cost: {
11577
- input: 1.5,
11578
- output: 7.5,
11579
- cacheRead: 0.15,
11580
- cacheWrite: 0.0833333333333333,
11704
+ input: 0.75,
11705
+ output: 3.75,
11706
+ cacheRead: 0.075,
11707
+ cacheWrite: 0.0416666666666667,
11581
11708
  },
11582
11709
  contextWindow: 1048576,
11583
11710
  maxTokens: 65536,
@@ -11591,10 +11718,44 @@ export const MODELS = {
11591
11718
  reasoning: true,
11592
11719
  input: ["text", "image"],
11593
11720
  cost: {
11594
- input: 0.75,
11595
- output: 3.75,
11596
- cacheRead: 0.075,
11597
- cacheWrite: 0.0833333333333333,
11721
+ input: 0.375,
11722
+ output: 1.875,
11723
+ cacheRead: 0.0375,
11724
+ cacheWrite: 0.0416666666666667,
11725
+ },
11726
+ contextWindow: 1048576,
11727
+ maxTokens: 65536,
11728
+ },
11729
+ "google/gemini-3.7-flash": {
11730
+ id: "google/gemini-3.7-flash",
11731
+ name: "Google: Gemini 3.7 Flash",
11732
+ api: "openai-completions",
11733
+ provider: "openrouter",
11734
+ baseUrl: "https://openrouter.ai/api/v1",
11735
+ reasoning: true,
11736
+ input: ["text", "image"],
11737
+ cost: {
11738
+ input: 0.375,
11739
+ output: 1.875,
11740
+ cacheRead: 0.0375,
11741
+ cacheWrite: 0.0208333333333333,
11742
+ },
11743
+ contextWindow: 1048576,
11744
+ maxTokens: 65536,
11745
+ },
11746
+ "google/gemini-3.7-flash:batch": {
11747
+ id: "google/gemini-3.7-flash:batch",
11748
+ name: "Google: Gemini 3.7 Flash (batch)",
11749
+ api: "openai-completions",
11750
+ provider: "openrouter",
11751
+ baseUrl: "https://openrouter.ai/api/v1",
11752
+ reasoning: true,
11753
+ input: ["text", "image"],
11754
+ cost: {
11755
+ input: 0.1875,
11756
+ output: 0.9375,
11757
+ cacheRead: 0.01875,
11758
+ cacheWrite: 0.0208333333333333,
11598
11759
  },
11599
11760
  contextWindow: 1048576,
11600
11761
  maxTokens: 65536,
@@ -11787,23 +11948,6 @@ export const MODELS = {
11787
11948
  contextWindow: 262144,
11788
11949
  maxTokens: 32768,
11789
11950
  },
11790
- "inclusionai/ling-3.0-tiny:free": {
11791
- id: "inclusionai/ling-3.0-tiny:free",
11792
- name: "inclusionAI: Ling 3.0 Tiny (free)",
11793
- api: "openai-completions",
11794
- provider: "openrouter",
11795
- baseUrl: "https://openrouter.ai/api/v1",
11796
- reasoning: true,
11797
- input: ["text"],
11798
- cost: {
11799
- input: 0,
11800
- output: 0,
11801
- cacheRead: 0,
11802
- cacheWrite: 0,
11803
- },
11804
- contextWindow: 262144,
11805
- maxTokens: 32768,
11806
- },
11807
11951
  "inclusionai/ring-2.6-1t": {
11808
11952
  id: "inclusionai/ring-2.6-1t",
11809
11953
  name: "inclusionAI: Ring-2.6-1T",
@@ -11887,7 +12031,7 @@ export const MODELS = {
11887
12031
  cacheWrite: 0,
11888
12032
  },
11889
12033
  contextWindow: 128000,
11890
- maxTokens: 32768,
12034
+ maxTokens: 8192,
11891
12035
  },
11892
12036
  "meituan/longcat-2.0": {
11893
12037
  id: "meituan/longcat-2.0",
@@ -14837,13 +14981,13 @@ export const MODELS = {
14837
14981
  reasoning: false,
14838
14982
  input: ["text"],
14839
14983
  cost: {
14840
- input: 0.09,
14984
+ input: 0.09999999999999999,
14841
14985
  output: 1.1,
14842
- cacheRead: 0,
14986
+ cacheRead: 0.07,
14843
14987
  cacheWrite: 0,
14844
14988
  },
14845
14989
  contextWindow: 262144,
14846
- maxTokens: 16384,
14990
+ maxTokens: 262144,
14847
14991
  },
14848
14992
  "qwen/qwen3-next-80b-a3b-thinking": {
14849
14993
  id: "qwen/qwen3-next-80b-a3b-thinking",
@@ -14860,7 +15004,7 @@ export const MODELS = {
14860
15004
  cacheWrite: 0,
14861
15005
  },
14862
15006
  contextWindow: 262144,
14863
- maxTokens: 262144,
15007
+ maxTokens: 32768,
14864
15008
  },
14865
15009
  "qwen/qwen3-vl-235b-a22b-instruct": {
14866
15010
  id: "qwen/qwen3-vl-235b-a22b-instruct",
@@ -14905,13 +15049,13 @@ export const MODELS = {
14905
15049
  reasoning: false,
14906
15050
  input: ["text", "image"],
14907
15051
  cost: {
14908
- input: 0.15,
14909
- output: 0.6,
15052
+ input: 0.13,
15053
+ output: 0.52,
14910
15054
  cacheRead: 0,
14911
15055
  cacheWrite: 0,
14912
15056
  },
14913
15057
  contextWindow: 262144,
14914
- maxTokens: 16384,
15058
+ maxTokens: 32768,
14915
15059
  },
14916
15060
  "qwen/qwen3-vl-30b-a3b-thinking": {
14917
15061
  id: "qwen/qwen3-vl-30b-a3b-thinking",
@@ -15041,13 +15185,13 @@ export const MODELS = {
15041
15185
  reasoning: true,
15042
15186
  input: ["text", "image"],
15043
15187
  cost: {
15044
- input: 0.44999999999999996,
15045
- output: 3,
15046
- cacheRead: 0.045,
15188
+ input: 0.5,
15189
+ output: 3.5999999999999996,
15190
+ cacheRead: 0.3,
15047
15191
  cacheWrite: 0,
15048
15192
  },
15049
15193
  contextWindow: 262144,
15050
- maxTokens: 65536,
15194
+ maxTokens: 262144,
15051
15195
  },
15052
15196
  "qwen/qwen3.5-9b": {
15053
15197
  id: "qwen/qwen3.5-9b",
@@ -15264,11 +15408,11 @@ export const MODELS = {
15264
15408
  cost: {
15265
15409
  input: 2,
15266
15410
  output: 6,
15267
- cacheRead: 0.19999999999999998,
15411
+ cacheRead: 0.25,
15268
15412
  cacheWrite: 0,
15269
15413
  },
15270
- contextWindow: 1000000,
15271
- maxTokens: 52429,
15414
+ contextWindow: 1010000,
15415
+ maxTokens: 262144,
15272
15416
  },
15273
15417
  "qwen/qwen3.8-max": {
15274
15418
  id: "qwen/qwen3.8-max",
@@ -15840,13 +15984,13 @@ export const MODELS = {
15840
15984
  reasoning: true,
15841
15985
  input: ["text"],
15842
15986
  cost: {
15843
- input: 0.5,
15844
- output: 3.15,
15845
- cacheRead: 0.09999999999999999,
15987
+ input: 0.63,
15988
+ output: 1.9800000000000002,
15989
+ cacheRead: 0.0945,
15846
15990
  cacheWrite: 0,
15847
15991
  },
15848
15992
  contextWindow: 1048576,
15849
- maxTokens: 131072,
15993
+ maxTokens: 4096,
15850
15994
  },
15851
15995
  "z-ai/glm-5.2:batch": {
15852
15996
  id: "z-ai/glm-5.2:batch",
@@ -15982,10 +16126,10 @@ export const MODELS = {
15982
16126
  reasoning: true,
15983
16127
  input: ["text", "image"],
15984
16128
  cost: {
15985
- input: 1.5,
15986
- output: 7.5,
15987
- cacheRead: 0.15,
15988
- cacheWrite: 0.0833333333333333,
16129
+ input: 0.375,
16130
+ output: 1.875,
16131
+ cacheRead: 0.0375,
16132
+ cacheWrite: 0.0208333333333333,
15989
16133
  },
15990
16134
  contextWindow: 1048576,
15991
16135
  maxTokens: 65536,
@@ -17465,6 +17609,23 @@ export const MODELS = {
17465
17609
  contextWindow: 1000000,
17466
17610
  maxTokens: 64000,
17467
17611
  },
17612
+ "alibaba/qwen3.8-2.4t-a95b": {
17613
+ id: "alibaba/qwen3.8-2.4t-a95b",
17614
+ name: "Qwen3.8 2.4T A95B",
17615
+ api: "anthropic-messages",
17616
+ provider: "vercel-ai-gateway",
17617
+ baseUrl: "https://ai-gateway.vercel.sh",
17618
+ reasoning: true,
17619
+ input: ["text"],
17620
+ cost: {
17621
+ input: 2,
17622
+ output: 6,
17623
+ cacheRead: 0.25,
17624
+ cacheWrite: 0,
17625
+ },
17626
+ contextWindow: 262144,
17627
+ maxTokens: 131072,
17628
+ },
17468
17629
  "alibaba/qwen3.8-max": {
17469
17630
  id: "alibaba/qwen3.8-max",
17470
17631
  name: "Qwen 3.8 Max",
@@ -18231,6 +18392,23 @@ export const MODELS = {
18231
18392
  contextWindow: 1000000,
18232
18393
  maxTokens: 64000,
18233
18394
  },
18395
+ "google/gemini-3.7-flash": {
18396
+ id: "google/gemini-3.7-flash",
18397
+ name: "Gemini 3.7 Flash",
18398
+ api: "anthropic-messages",
18399
+ provider: "vercel-ai-gateway",
18400
+ baseUrl: "https://ai-gateway.vercel.sh",
18401
+ reasoning: true,
18402
+ input: ["text", "image"],
18403
+ cost: {
18404
+ input: 0.75,
18405
+ output: 3.75,
18406
+ cacheRead: 0.075,
18407
+ cacheWrite: 0,
18408
+ },
18409
+ contextWindow: 1000000,
18410
+ maxTokens: 65536,
18411
+ },
18234
18412
  "google/gemma-4-26b-a4b-it": {
18235
18413
  id: "google/gemma-4-26b-a4b-it",
18236
18414
  name: "Google Gemma 4 26B A4B",
@@ -18262,7 +18440,7 @@ export const MODELS = {
18262
18440
  cacheRead: 0,
18263
18441
  cacheWrite: 0,
18264
18442
  },
18265
- contextWindow: 256000,
18443
+ contextWindow: 262144,
18266
18444
  maxTokens: 131072,
18267
18445
  },
18268
18446
  "inception/mercury-2": {
@@ -18316,23 +18494,6 @@ export const MODELS = {
18316
18494
  contextWindow: 256000,
18317
18495
  maxTokens: 32000,
18318
18496
  },
18319
- "inclusionai/ling-3.0-tiny-free": {
18320
- id: "inclusionai/ling-3.0-tiny-free",
18321
- name: "Ling 3.0 Tiny (Free)",
18322
- api: "anthropic-messages",
18323
- provider: "vercel-ai-gateway",
18324
- baseUrl: "https://ai-gateway.vercel.sh",
18325
- reasoning: true,
18326
- input: ["text"],
18327
- cost: {
18328
- input: 0,
18329
- output: 0,
18330
- cacheRead: 0,
18331
- cacheWrite: 0,
18332
- },
18333
- contextWindow: 256000,
18334
- maxTokens: 32000,
18335
- },
18336
18497
  "interfaze/interfaze-beta": {
18337
18498
  id: "interfaze/interfaze-beta",
18338
18499
  name: "Interfaze Beta",
@@ -20473,91 +20634,6 @@ export const MODELS = {
20473
20634
  },
20474
20635
  },
20475
20636
  "xai": {
20476
- "grok-3": {
20477
- id: "grok-3",
20478
- name: "Grok 3",
20479
- api: "openai-completions",
20480
- provider: "xai",
20481
- baseUrl: "https://api.x.ai/v1",
20482
- reasoning: false,
20483
- input: ["text"],
20484
- cost: {
20485
- input: 3,
20486
- output: 15,
20487
- cacheRead: 0.75,
20488
- cacheWrite: 0,
20489
- },
20490
- contextWindow: 131072,
20491
- maxTokens: 8192,
20492
- },
20493
- "grok-3-fast": {
20494
- id: "grok-3-fast",
20495
- name: "Grok 3 Fast",
20496
- api: "openai-completions",
20497
- provider: "xai",
20498
- baseUrl: "https://api.x.ai/v1",
20499
- reasoning: false,
20500
- input: ["text"],
20501
- cost: {
20502
- input: 5,
20503
- output: 25,
20504
- cacheRead: 1.25,
20505
- cacheWrite: 0,
20506
- },
20507
- contextWindow: 131072,
20508
- maxTokens: 8192,
20509
- },
20510
- "grok-4.20-0309-non-reasoning": {
20511
- id: "grok-4.20-0309-non-reasoning",
20512
- name: "Grok 4.20 (Non-Reasoning)",
20513
- api: "openai-completions",
20514
- provider: "xai",
20515
- baseUrl: "https://api.x.ai/v1",
20516
- reasoning: false,
20517
- input: ["text", "image"],
20518
- cost: {
20519
- input: 1.25,
20520
- output: 2.5,
20521
- cacheRead: 0.2,
20522
- cacheWrite: 0,
20523
- },
20524
- contextWindow: 1000000,
20525
- maxTokens: 30000,
20526
- },
20527
- "grok-4.20-0309-reasoning": {
20528
- id: "grok-4.20-0309-reasoning",
20529
- name: "Grok 4.20 (Reasoning)",
20530
- api: "openai-completions",
20531
- provider: "xai",
20532
- baseUrl: "https://api.x.ai/v1",
20533
- reasoning: true,
20534
- input: ["text", "image"],
20535
- cost: {
20536
- input: 1.25,
20537
- output: 2.5,
20538
- cacheRead: 0.2,
20539
- cacheWrite: 0,
20540
- },
20541
- contextWindow: 1000000,
20542
- maxTokens: 30000,
20543
- },
20544
- "grok-4.3": {
20545
- id: "grok-4.3",
20546
- name: "Grok 4.3",
20547
- api: "openai-completions",
20548
- provider: "xai",
20549
- baseUrl: "https://api.x.ai/v1",
20550
- reasoning: true,
20551
- input: ["text", "image"],
20552
- cost: {
20553
- input: 1.25,
20554
- output: 2.5,
20555
- cacheRead: 0.2,
20556
- cacheWrite: 0,
20557
- },
20558
- contextWindow: 1000000,
20559
- maxTokens: 30000,
20560
- },
20561
20637
  "grok-4.5": {
20562
20638
  id: "grok-4.5",
20563
20639
  name: "Grok 4.5",
@@ -20566,6 +20642,7 @@ export const MODELS = {
20566
20642
  baseUrl: "https://api.x.ai/v1",
20567
20643
  compat: { "supportsLongCacheRetention": false },
20568
20644
  reasoning: true,
20645
+ defaultThinkingLevel: "high",
20569
20646
  thinkingLevelMap: { "off": null, "minimal": null },
20570
20647
  input: ["text", "image"],
20571
20648
  cost: {
@@ -20580,10 +20657,13 @@ export const MODELS = {
20580
20657
  "grok-4.6": {
20581
20658
  id: "grok-4.6",
20582
20659
  name: "Grok 4.6",
20583
- api: "openai-completions",
20660
+ api: "openai-responses",
20584
20661
  provider: "xai",
20585
20662
  baseUrl: "https://api.x.ai/v1",
20663
+ compat: { "supportsLongCacheRetention": false },
20586
20664
  reasoning: true,
20665
+ defaultThinkingLevel: "high",
20666
+ thinkingLevelMap: { "off": null, "minimal": null, "xhigh": "xhigh" },
20587
20667
  input: ["text", "image"],
20588
20668
  cost: {
20589
20669
  input: 2,
@@ -20594,40 +20674,6 @@ export const MODELS = {
20594
20674
  contextWindow: 500000,
20595
20675
  maxTokens: 500000,
20596
20676
  },
20597
- "grok-build-0.1": {
20598
- id: "grok-build-0.1",
20599
- name: "Grok Build 0.1",
20600
- api: "openai-completions",
20601
- provider: "xai",
20602
- baseUrl: "https://api.x.ai/v1",
20603
- reasoning: true,
20604
- input: ["text", "image"],
20605
- cost: {
20606
- input: 1,
20607
- output: 2,
20608
- cacheRead: 0.2,
20609
- cacheWrite: 0,
20610
- },
20611
- contextWindow: 256000,
20612
- maxTokens: 256000,
20613
- },
20614
- "grok-code-fast-1": {
20615
- id: "grok-code-fast-1",
20616
- name: "Grok Code Fast 1",
20617
- api: "openai-completions",
20618
- provider: "xai",
20619
- baseUrl: "https://api.x.ai/v1",
20620
- reasoning: false,
20621
- input: ["text"],
20622
- cost: {
20623
- input: 0.2,
20624
- output: 1.5,
20625
- cacheRead: 0.02,
20626
- cacheWrite: 0,
20627
- },
20628
- contextWindow: 32768,
20629
- maxTokens: 8192,
20630
- },
20631
20677
  },
20632
20678
  "xiaomi": {
20633
20679
  "mimo-v2-flash": {