@caupulican/pi-ai 0.81.38 → 0.81.39

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (32) hide show
  1. package/dist/image-models.generated.d.ts +60 -0
  2. package/dist/image-models.generated.d.ts.map +1 -1
  3. package/dist/image-models.generated.js +60 -0
  4. package/dist/image-models.generated.js.map +1 -1
  5. package/dist/models.generated.d.ts +163 -213
  6. package/dist/models.generated.d.ts.map +1 -1
  7. package/dist/models.generated.js +257 -279
  8. package/dist/models.generated.js.map +1 -1
  9. package/dist/providers/openai-completions.d.ts +1 -1
  10. package/dist/providers/openai-completions.d.ts.map +1 -1
  11. package/dist/providers/openai-completions.js +68 -19
  12. package/dist/providers/openai-completions.js.map +1 -1
  13. package/dist/stream.d.ts.map +1 -1
  14. package/dist/stream.js +145 -2
  15. package/dist/stream.js.map +1 -1
  16. package/dist/utils/tool-repair/repairer.d.ts +13 -1
  17. package/dist/utils/tool-repair/repairer.d.ts.map +1 -1
  18. package/dist/utils/tool-repair/repairer.js +44 -28
  19. package/dist/utils/tool-repair/repairer.js.map +1 -1
  20. package/dist/utils/tool-repair/text-protocol-history.d.ts +17 -0
  21. package/dist/utils/tool-repair/text-protocol-history.d.ts.map +1 -0
  22. package/dist/utils/tool-repair/text-protocol-history.js +22 -0
  23. package/dist/utils/tool-repair/text-protocol-history.js.map +1 -0
  24. package/dist/utils/tool-repair/text-protocol.d.ts +7 -1
  25. package/dist/utils/tool-repair/text-protocol.d.ts.map +1 -1
  26. package/dist/utils/tool-repair/text-protocol.js +44 -18
  27. package/dist/utils/tool-repair/text-protocol.js.map +1 -1
  28. package/dist/utils/validation.d.ts +9 -0
  29. package/dist/utils/validation.d.ts.map +1 -1
  30. package/dist/utils/validation.js +8 -1
  31. package/dist/utils/validation.js.map +1 -1
  32. package/package.json +1 -1
@@ -4806,7 +4806,7 @@ export const MODELS = {
4806
4806
  input: 1,
4807
4807
  output: 6,
4808
4808
  cacheRead: 0.1,
4809
- cacheWrite: 0,
4809
+ cacheWrite: 1.25,
4810
4810
  },
4811
4811
  contextWindow: 1050000,
4812
4812
  maxTokens: 128000,
@@ -4825,7 +4825,7 @@ export const MODELS = {
4825
4825
  input: 5,
4826
4826
  output: 30,
4827
4827
  cacheRead: 0.5,
4828
- cacheWrite: 0,
4828
+ cacheWrite: 6.25,
4829
4829
  },
4830
4830
  contextWindow: 1050000,
4831
4831
  maxTokens: 128000,
@@ -4844,7 +4844,7 @@ export const MODELS = {
4844
4844
  input: 2.5,
4845
4845
  output: 15,
4846
4846
  cacheRead: 0.25,
4847
- cacheWrite: 0,
4847
+ cacheWrite: 3.125,
4848
4848
  },
4849
4849
  contextWindow: 1050000,
4850
4850
  maxTokens: 128000,
@@ -6405,9 +6405,9 @@ export const MODELS = {
6405
6405
  },
6406
6406
  },
6407
6407
  "kimi-coding": {
6408
- "k2p7": {
6409
- id: "k2p7",
6410
- name: "Kimi K2.7 Code",
6408
+ "k3": {
6409
+ id: "k3",
6410
+ name: "Kimi K3",
6411
6411
  api: "anthropic-messages",
6412
6412
  provider: "kimi-coding",
6413
6413
  baseUrl: "https://api.kimi.com/coding",
@@ -6420,12 +6420,12 @@ export const MODELS = {
6420
6420
  cacheRead: 0,
6421
6421
  cacheWrite: 0,
6422
6422
  },
6423
- contextWindow: 262144,
6424
- maxTokens: 32768,
6423
+ contextWindow: 1048576,
6424
+ maxTokens: 131072,
6425
6425
  },
6426
6426
  "kimi-for-coding": {
6427
6427
  id: "kimi-for-coding",
6428
- name: "Kimi For Coding",
6428
+ name: "Kimi K2.7 Code",
6429
6429
  api: "anthropic-messages",
6430
6430
  provider: "kimi-coding",
6431
6431
  baseUrl: "https://api.kimi.com/coding",
@@ -6459,24 +6459,6 @@ export const MODELS = {
6459
6459
  contextWindow: 262144,
6460
6460
  maxTokens: 32768,
6461
6461
  },
6462
- "kimi-k2-thinking": {
6463
- id: "kimi-k2-thinking",
6464
- name: "Kimi K2 Thinking",
6465
- api: "anthropic-messages",
6466
- provider: "kimi-coding",
6467
- baseUrl: "https://api.kimi.com/coding",
6468
- headers: { "User-Agent": "KimiCLI/1.5" },
6469
- reasoning: true,
6470
- input: ["text"],
6471
- cost: {
6472
- input: 0,
6473
- output: 0,
6474
- cacheRead: 0,
6475
- cacheWrite: 0,
6476
- },
6477
- contextWindow: 262144,
6478
- maxTokens: 32768,
6479
- },
6480
6462
  },
6481
6463
  "llama-cpp": {
6482
6464
  "local": {
@@ -7245,6 +7227,24 @@ export const MODELS = {
7245
7227
  contextWindow: 262144,
7246
7228
  maxTokens: 262144,
7247
7229
  },
7230
+ "kimi-k3": {
7231
+ id: "kimi-k3",
7232
+ name: "Kimi K3",
7233
+ api: "openai-completions",
7234
+ provider: "moonshotai",
7235
+ baseUrl: "https://api.moonshot.ai/v1",
7236
+ compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false },
7237
+ reasoning: true,
7238
+ input: ["text", "image"],
7239
+ cost: {
7240
+ input: 3,
7241
+ output: 15,
7242
+ cacheRead: 0.3,
7243
+ cacheWrite: 0,
7244
+ },
7245
+ contextWindow: 1048576,
7246
+ maxTokens: 131072,
7247
+ },
7248
7248
  },
7249
7249
  "moonshotai-cn": {
7250
7250
  "kimi-k2-0711-preview": {
@@ -7409,6 +7409,24 @@ export const MODELS = {
7409
7409
  contextWindow: 262144,
7410
7410
  maxTokens: 262144,
7411
7411
  },
7412
+ "kimi-k3": {
7413
+ id: "kimi-k3",
7414
+ name: "Kimi K3",
7415
+ api: "openai-completions",
7416
+ provider: "moonshotai-cn",
7417
+ baseUrl: "https://api.moonshot.cn/v1",
7418
+ compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false },
7419
+ reasoning: true,
7420
+ input: ["text", "image"],
7421
+ cost: {
7422
+ input: 3,
7423
+ output: 15,
7424
+ cacheRead: 0.3,
7425
+ cacheWrite: 0,
7426
+ },
7427
+ contextWindow: 1048576,
7428
+ maxTokens: 131072,
7429
+ },
7412
7430
  },
7413
7431
  "openai": {
7414
7432
  "gpt-4": {
@@ -9136,7 +9154,7 @@ export const MODELS = {
9136
9154
  "grok-4.5": {
9137
9155
  id: "grok-4.5",
9138
9156
  name: "Grok 4.5",
9139
- api: "openai-completions",
9157
+ api: "openai-responses",
9140
9158
  provider: "opencode",
9141
9159
  baseUrl: "https://opencode.ai/zen/v1",
9142
9160
  reasoning: true,
@@ -9406,9 +9424,9 @@ export const MODELS = {
9406
9424
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "max" },
9407
9425
  input: ["text"],
9408
9426
  cost: {
9409
- input: 1.74,
9410
- output: 3.48,
9411
- cacheRead: 0.0145,
9427
+ input: 0.435,
9428
+ output: 0.87,
9429
+ cacheRead: 0.003625,
9412
9430
  cacheWrite: 0,
9413
9431
  },
9414
9432
  contextWindow: 1000000,
@@ -9448,6 +9466,23 @@ export const MODELS = {
9448
9466
  contextWindow: 1000000,
9449
9467
  maxTokens: 131072,
9450
9468
  },
9469
+ "grok-4.5": {
9470
+ id: "grok-4.5",
9471
+ name: "Grok 4.5",
9472
+ api: "openai-responses",
9473
+ provider: "opencode-go",
9474
+ baseUrl: "https://opencode.ai/zen/go/v1",
9475
+ reasoning: true,
9476
+ input: ["text", "image"],
9477
+ cost: {
9478
+ input: 2,
9479
+ output: 6,
9480
+ cacheRead: 0.5,
9481
+ cacheWrite: 0,
9482
+ },
9483
+ contextWindow: 500000,
9484
+ maxTokens: 500000,
9485
+ },
9451
9486
  "kimi-k2.6": {
9452
9487
  id: "kimi-k2.6",
9453
9488
  name: "Kimi K2.6",
@@ -9484,6 +9519,23 @@ export const MODELS = {
9484
9519
  contextWindow: 262144,
9485
9520
  maxTokens: 262144,
9486
9521
  },
9522
+ "kimi-k3": {
9523
+ id: "kimi-k3",
9524
+ name: "Kimi K3 (2x usage)",
9525
+ api: "openai-completions",
9526
+ provider: "opencode-go",
9527
+ baseUrl: "https://opencode.ai/zen/go/v1",
9528
+ reasoning: true,
9529
+ input: ["text", "image"],
9530
+ cost: {
9531
+ input: 3,
9532
+ output: 15,
9533
+ cacheRead: 0.3,
9534
+ cacheWrite: 0,
9535
+ },
9536
+ contextWindow: 1048576,
9537
+ maxTokens: 131072,
9538
+ },
9487
9539
  "mimo-v2.5": {
9488
9540
  id: "mimo-v2.5",
9489
9541
  name: "MiMo V2.5",
@@ -9510,9 +9562,9 @@ export const MODELS = {
9510
9562
  reasoning: true,
9511
9563
  input: ["text"],
9512
9564
  cost: {
9513
- input: 1.74,
9514
- output: 3.48,
9515
- cacheRead: 0.0145,
9565
+ input: 0.435,
9566
+ output: 0.87,
9567
+ cacheRead: 0.003625,
9516
9568
  cacheWrite: 0,
9517
9569
  },
9518
9570
  contextWindow: 1048576,
@@ -10338,7 +10390,7 @@ export const MODELS = {
10338
10390
  cost: {
10339
10391
  input: 0.098,
10340
10392
  output: 0.196,
10341
- cacheRead: 0.02,
10393
+ cacheRead: 0.0196,
10342
10394
  cacheWrite: 0,
10343
10395
  },
10344
10396
  contextWindow: 1048576,
@@ -10593,13 +10645,13 @@ export const MODELS = {
10593
10645
  reasoning: false,
10594
10646
  input: ["text", "image"],
10595
10647
  cost: {
10596
- input: 0.08,
10597
- output: 0.44999999999999996,
10598
- cacheRead: 0.04,
10648
+ input: 0.09999999999999999,
10649
+ output: 0.3,
10650
+ cacheRead: 0,
10599
10651
  cacheWrite: 0,
10600
10652
  },
10601
10653
  contextWindow: 131072,
10602
- maxTokens: 131072,
10654
+ maxTokens: 4096,
10603
10655
  },
10604
10656
  "google/gemma-4-26b-a4b-it": {
10605
10657
  id: "google/gemma-4-26b-a4b-it",
@@ -10610,13 +10662,13 @@ export const MODELS = {
10610
10662
  reasoning: true,
10611
10663
  input: ["text", "image"],
10612
10664
  cost: {
10613
- input: 0.09999999999999999,
10614
- output: 0.3,
10665
+ input: 0.07,
10666
+ output: 0.33999999999999997,
10615
10667
  cacheRead: 0,
10616
10668
  cacheWrite: 0,
10617
10669
  },
10618
10670
  contextWindow: 262144,
10619
- maxTokens: 256000,
10671
+ maxTokens: 16384,
10620
10672
  },
10621
10673
  "google/gemma-4-26b-a4b-it:free": {
10622
10674
  id: "google/gemma-4-26b-a4b-it:free",
@@ -10644,13 +10696,13 @@ export const MODELS = {
10644
10696
  reasoning: true,
10645
10697
  input: ["text", "image"],
10646
10698
  cost: {
10647
- input: 0.22,
10648
- output: 0.55,
10649
- cacheRead: 0.12,
10699
+ input: 0.12,
10700
+ output: 0.37,
10701
+ cacheRead: 0,
10650
10702
  cacheWrite: 0,
10651
10703
  },
10652
10704
  contextWindow: 262144,
10653
- maxTokens: 262144,
10705
+ maxTokens: 16384,
10654
10706
  },
10655
10707
  "google/gemma-4-31b-it:free": {
10656
10708
  id: "google/gemma-4-31b-it:free",
@@ -10806,6 +10858,23 @@ export const MODELS = {
10806
10858
  contextWindow: 256000,
10807
10859
  maxTokens: 80000,
10808
10860
  },
10861
+ "meituan/longcat-2.0": {
10862
+ id: "meituan/longcat-2.0",
10863
+ name: "Meituan: LongCat 2.0",
10864
+ api: "openai-completions",
10865
+ provider: "openrouter",
10866
+ baseUrl: "https://openrouter.ai/api/v1",
10867
+ reasoning: true,
10868
+ input: ["text"],
10869
+ cost: {
10870
+ input: 0.3,
10871
+ output: 1.2,
10872
+ cacheRead: 0.006,
10873
+ cacheWrite: 0,
10874
+ },
10875
+ contextWindow: 1048756,
10876
+ maxTokens: 262144,
10877
+ },
10809
10878
  "meta-llama/llama-3.1-70b-instruct": {
10810
10879
  id: "meta-llama/llama-3.1-70b-instruct",
10811
10880
  name: "Meta: Llama 3.1 70B Instruct",
@@ -10857,23 +10926,6 @@ export const MODELS = {
10857
10926
  contextWindow: 131072,
10858
10927
  maxTokens: 128000,
10859
10928
  },
10860
- "meta-llama/llama-3.3-70b-instruct:free": {
10861
- id: "meta-llama/llama-3.3-70b-instruct:free",
10862
- name: "Meta: Llama 3.3 70B Instruct (free)",
10863
- api: "openai-completions",
10864
- provider: "openrouter",
10865
- baseUrl: "https://openrouter.ai/api/v1",
10866
- reasoning: false,
10867
- input: ["text"],
10868
- cost: {
10869
- input: 0,
10870
- output: 0,
10871
- cacheRead: 0,
10872
- cacheWrite: 0,
10873
- },
10874
- contextWindow: 131072,
10875
- maxTokens: 4096,
10876
- },
10877
10929
  "meta-llama/llama-4-maverick": {
10878
10930
  id: "meta-llama/llama-4-maverick",
10879
10931
  name: "Meta: Llama 4 Maverick",
@@ -10908,6 +10960,23 @@ export const MODELS = {
10908
10960
  contextWindow: 10000000,
10909
10961
  maxTokens: 16384,
10910
10962
  },
10963
+ "meta/muse-spark-1.1": {
10964
+ id: "meta/muse-spark-1.1",
10965
+ name: "Meta: Muse Spark 1.1",
10966
+ api: "openai-completions",
10967
+ provider: "openrouter",
10968
+ baseUrl: "https://openrouter.ai/api/v1",
10969
+ reasoning: true,
10970
+ input: ["text", "image"],
10971
+ cost: {
10972
+ input: 1.25,
10973
+ output: 4.25,
10974
+ cacheRead: 0.15,
10975
+ cacheWrite: 0,
10976
+ },
10977
+ contextWindow: 1048576,
10978
+ maxTokens: 4096,
10979
+ },
10911
10980
  "minimax/minimax-m1": {
10912
10981
  id: "minimax/minimax-m1",
10913
10982
  name: "MiniMax: MiniMax M1",
@@ -10985,9 +11054,9 @@ export const MODELS = {
10985
11054
  reasoning: true,
10986
11055
  input: ["text"],
10987
11056
  cost: {
10988
- input: 0.3,
10989
- output: 1.2,
10990
- cacheRead: 0.06,
11057
+ input: 0.25,
11058
+ output: 1,
11059
+ cacheRead: 0.049999999999999996,
10991
11060
  cacheWrite: 0,
10992
11061
  },
10993
11062
  contextWindow: 204800,
@@ -11206,8 +11275,8 @@ export const MODELS = {
11206
11275
  reasoning: false,
11207
11276
  input: ["text"],
11208
11277
  cost: {
11209
- input: 0.02,
11210
- output: 0.04,
11278
+ input: 0.019000000000000003,
11279
+ output: 0.03,
11211
11280
  cacheRead: 0,
11212
11281
  cacheWrite: 0,
11213
11282
  },
@@ -11344,11 +11413,11 @@ export const MODELS = {
11344
11413
  cost: {
11345
11414
  input: 0.6,
11346
11415
  output: 2.5,
11347
- cacheRead: 0,
11416
+ cacheRead: 0.15,
11348
11417
  cacheWrite: 0,
11349
11418
  },
11350
11419
  contextWindow: 262144,
11351
- maxTokens: 262144,
11420
+ maxTokens: 100352,
11352
11421
  },
11353
11422
  "moonshotai/kimi-k2.5": {
11354
11423
  id: "moonshotai/kimi-k2.5",
@@ -11377,13 +11446,13 @@ export const MODELS = {
11377
11446
  reasoning: true,
11378
11447
  input: ["text", "image"],
11379
11448
  cost: {
11380
- input: 0.95,
11381
- output: 4,
11382
- cacheRead: 0.16,
11449
+ input: 0.684,
11450
+ output: 3.42,
11451
+ cacheRead: 0.144,
11383
11452
  cacheWrite: 0,
11384
11453
  },
11385
11454
  contextWindow: 262144,
11386
- maxTokens: 4096,
11455
+ maxTokens: 262144,
11387
11456
  },
11388
11457
  "moonshotai/kimi-k2.7-code": {
11389
11458
  id: "moonshotai/kimi-k2.7-code",
@@ -11394,14 +11463,31 @@ export const MODELS = {
11394
11463
  reasoning: true,
11395
11464
  input: ["text", "image"],
11396
11465
  cost: {
11397
- input: 0.719,
11398
- output: 3.49,
11399
- cacheRead: 0.149,
11466
+ input: 0.82,
11467
+ output: 3.75,
11468
+ cacheRead: 0.16,
11400
11469
  cacheWrite: 0,
11401
11470
  },
11402
11471
  contextWindow: 262144,
11403
11472
  maxTokens: 262144,
11404
11473
  },
11474
+ "moonshotai/kimi-k3": {
11475
+ id: "moonshotai/kimi-k3",
11476
+ name: "MoonshotAI: Kimi K3",
11477
+ api: "openai-completions",
11478
+ provider: "openrouter",
11479
+ baseUrl: "https://openrouter.ai/api/v1",
11480
+ reasoning: true,
11481
+ input: ["text", "image"],
11482
+ cost: {
11483
+ input: 3,
11484
+ output: 15,
11485
+ cacheRead: 0.3,
11486
+ cacheWrite: 0,
11487
+ },
11488
+ contextWindow: 1048576,
11489
+ maxTokens: 4096,
11490
+ },
11405
11491
  "nex-agi/nex-n2-mini": {
11406
11492
  id: "nex-agi/nex-n2-mini",
11407
11493
  name: "Nex AGI: Nex-N2-Mini",
@@ -11436,23 +11522,6 @@ export const MODELS = {
11436
11522
  contextWindow: 262144,
11437
11523
  maxTokens: 262144,
11438
11524
  },
11439
- "nvidia/llama-3.3-nemotron-super-49b-v1.5": {
11440
- id: "nvidia/llama-3.3-nemotron-super-49b-v1.5",
11441
- name: "NVIDIA: Llama 3.3 Nemotron Super 49B V1.5",
11442
- api: "openai-completions",
11443
- provider: "openrouter",
11444
- baseUrl: "https://openrouter.ai/api/v1",
11445
- reasoning: true,
11446
- input: ["text"],
11447
- cost: {
11448
- input: 0.39999999999999997,
11449
- output: 0.39999999999999997,
11450
- cacheRead: 0,
11451
- cacheWrite: 0,
11452
- },
11453
- contextWindow: 131072,
11454
- maxTokens: 16384,
11455
- },
11456
11525
  "nvidia/nemotron-3-nano-30b-a3b": {
11457
11526
  id: "nvidia/nemotron-3-nano-30b-a3b",
11458
11527
  name: "NVIDIA: Nemotron 3 Nano 30B A3B",
@@ -11513,13 +11582,13 @@ export const MODELS = {
11513
11582
  reasoning: true,
11514
11583
  input: ["text"],
11515
11584
  cost: {
11516
- input: 0.21,
11517
- output: 0.45499999999999996,
11518
- cacheRead: 0.06,
11585
+ input: 0.08499999999999999,
11586
+ output: 0.39999999999999997,
11587
+ cacheRead: 0,
11519
11588
  cacheWrite: 0,
11520
11589
  },
11521
11590
  contextWindow: 1000000,
11522
- maxTokens: 4096,
11591
+ maxTokens: 16384,
11523
11592
  },
11524
11593
  "nvidia/nemotron-3-super-120b-a12b:free": {
11525
11594
  id: "nvidia/nemotron-3-super-120b-a12b:free",
@@ -12644,6 +12713,23 @@ export const MODELS = {
12644
12713
  contextWindow: 2000000,
12645
12714
  maxTokens: 4096,
12646
12715
  },
12716
+ "openrouter/auto-beta": {
12717
+ id: "openrouter/auto-beta",
12718
+ name: "Auto Router (Beta)",
12719
+ api: "openai-completions",
12720
+ provider: "openrouter",
12721
+ baseUrl: "https://openrouter.ai/api/v1",
12722
+ reasoning: true,
12723
+ input: ["text", "image"],
12724
+ cost: {
12725
+ input: -1000000,
12726
+ output: -1000000,
12727
+ cacheRead: 0,
12728
+ cacheWrite: 0,
12729
+ },
12730
+ contextWindow: 2000000,
12731
+ maxTokens: 4096,
12732
+ },
12647
12733
  "openrouter/free": {
12648
12734
  id: "openrouter/free",
12649
12735
  name: "Free Models Router",
@@ -12840,13 +12926,13 @@ export const MODELS = {
12840
12926
  reasoning: true,
12841
12927
  input: ["text"],
12842
12928
  cost: {
12843
- input: 0.09999999999999999,
12929
+ input: 0.12,
12844
12930
  output: 0.24,
12845
12931
  cacheRead: 0,
12846
12932
  cacheWrite: 0,
12847
12933
  },
12848
12934
  contextWindow: 131702,
12849
- maxTokens: 40960,
12935
+ maxTokens: 16384,
12850
12936
  },
12851
12937
  "qwen/qwen3-235b-a22b": {
12852
12938
  id: "qwen/qwen3-235b-a22b",
@@ -12891,13 +12977,13 @@ export const MODELS = {
12891
12977
  reasoning: true,
12892
12978
  input: ["text"],
12893
12979
  cost: {
12894
- input: 0.14950000000000002,
12895
- output: 1.495,
12980
+ input: 0.3,
12981
+ output: 3,
12896
12982
  cacheRead: 0,
12897
12983
  cacheWrite: 0,
12898
12984
  },
12899
12985
  contextWindow: 262144,
12900
- maxTokens: 4096,
12986
+ maxTokens: 32768,
12901
12987
  },
12902
12988
  "qwen/qwen3-30b-a3b": {
12903
12989
  id: "qwen/qwen3-30b-a3b",
@@ -12908,13 +12994,13 @@ export const MODELS = {
12908
12994
  reasoning: true,
12909
12995
  input: ["text"],
12910
12996
  cost: {
12911
- input: 0.12,
12912
- output: 0.5,
12997
+ input: 0.13,
12998
+ output: 0.52,
12913
12999
  cacheRead: 0,
12914
13000
  cacheWrite: 0,
12915
13001
  },
12916
13002
  contextWindow: 131072,
12917
- maxTokens: 16384,
13003
+ maxTokens: 8192,
12918
13004
  },
12919
13005
  "qwen/qwen3-30b-a3b-instruct-2507": {
12920
13006
  id: "qwen/qwen3-30b-a3b-instruct-2507",
@@ -13069,23 +13155,6 @@ export const MODELS = {
13069
13155
  contextWindow: 1000000,
13070
13156
  maxTokens: 65536,
13071
13157
  },
13072
- "qwen/qwen3-coder:free": {
13073
- id: "qwen/qwen3-coder:free",
13074
- name: "Qwen: Qwen3 Coder 480B A35B (free)",
13075
- api: "openai-completions",
13076
- provider: "openrouter",
13077
- baseUrl: "https://openrouter.ai/api/v1",
13078
- reasoning: false,
13079
- input: ["text"],
13080
- cost: {
13081
- input: 0,
13082
- output: 0,
13083
- cacheRead: 0,
13084
- cacheWrite: 0,
13085
- },
13086
- contextWindow: 1048576,
13087
- maxTokens: 262000,
13088
- },
13089
13158
  "qwen/qwen3-max": {
13090
13159
  id: "qwen/qwen3-max",
13091
13160
  name: "Qwen: Qwen3 Max",
@@ -13137,23 +13206,6 @@ export const MODELS = {
13137
13206
  contextWindow: 262144,
13138
13207
  maxTokens: 262144,
13139
13208
  },
13140
- "qwen/qwen3-next-80b-a3b-instruct:free": {
13141
- id: "qwen/qwen3-next-80b-a3b-instruct:free",
13142
- name: "Qwen: Qwen3 Next 80B A3B Instruct (free)",
13143
- api: "openai-completions",
13144
- provider: "openrouter",
13145
- baseUrl: "https://openrouter.ai/api/v1",
13146
- reasoning: false,
13147
- input: ["text"],
13148
- cost: {
13149
- input: 0,
13150
- output: 0,
13151
- cacheRead: 0,
13152
- cacheWrite: 0,
13153
- },
13154
- contextWindow: 262144,
13155
- maxTokens: 4096,
13156
- },
13157
13209
  "qwen/qwen3-next-80b-a3b-thinking": {
13158
13210
  id: "qwen/qwen3-next-80b-a3b-thinking",
13159
13211
  name: "Qwen: Qwen3 Next 80B A3B Thinking",
@@ -13316,13 +13368,13 @@ export const MODELS = {
13316
13368
  reasoning: true,
13317
13369
  input: ["text", "image"],
13318
13370
  cost: {
13319
- input: 0.195,
13320
- output: 1.56,
13371
+ input: 0.26,
13372
+ output: 2.6,
13321
13373
  cacheRead: 0,
13322
13374
  cacheWrite: 0,
13323
13375
  },
13324
13376
  contextWindow: 262144,
13325
- maxTokens: 65536,
13377
+ maxTokens: 81920,
13326
13378
  },
13327
13379
  "qwen/qwen3.5-35b-a3b": {
13328
13380
  id: "qwen/qwen3.5-35b-a3b",
@@ -13350,9 +13402,9 @@ export const MODELS = {
13350
13402
  reasoning: true,
13351
13403
  input: ["text", "image"],
13352
13404
  cost: {
13353
- input: 0.44999999999999996,
13354
- output: 3,
13355
- cacheRead: 0.22499999999999998,
13405
+ input: 0.39,
13406
+ output: 2.34,
13407
+ cacheRead: 0,
13356
13408
  cacheWrite: 0,
13357
13409
  },
13358
13410
  contextWindow: 262144,
@@ -13715,6 +13767,23 @@ export const MODELS = {
13715
13767
  contextWindow: 32768,
13716
13768
  maxTokens: 32768,
13717
13769
  },
13770
+ "thinkingmachines/inkling": {
13771
+ id: "thinkingmachines/inkling",
13772
+ name: "Thinking Machines: Inkling",
13773
+ api: "openai-completions",
13774
+ provider: "openrouter",
13775
+ baseUrl: "https://openrouter.ai/api/v1",
13776
+ reasoning: true,
13777
+ input: ["text", "image"],
13778
+ cost: {
13779
+ input: 1,
13780
+ output: 4.05,
13781
+ cacheRead: 0.16999999999999998,
13782
+ cacheWrite: 0,
13783
+ },
13784
+ contextWindow: 1048576,
13785
+ maxTokens: 4096,
13786
+ },
13718
13787
  "upstage/solar-pro-3": {
13719
13788
  id: "upstage/solar-pro-3",
13720
13789
  name: "Upstage: Solar Pro 3",
@@ -13777,7 +13846,7 @@ export const MODELS = {
13777
13846
  cost: {
13778
13847
  input: 2,
13779
13848
  output: 6,
13780
- cacheRead: 0.5,
13849
+ cacheRead: 0.3,
13781
13850
  cacheWrite: 0,
13782
13851
  },
13783
13852
  contextWindow: 500000,
@@ -13945,13 +14014,13 @@ export const MODELS = {
13945
14014
  reasoning: true,
13946
14015
  input: ["text"],
13947
14016
  cost: {
13948
- input: 0.06,
14017
+ input: 0.060500000000000005,
13949
14018
  output: 0.39999999999999997,
13950
- cacheRead: 0.01,
14019
+ cacheRead: 0,
13951
14020
  cacheWrite: 0,
13952
14021
  },
13953
- contextWindow: 202752,
13954
- maxTokens: 16384,
14022
+ contextWindow: 200000,
14023
+ maxTokens: 131072,
13955
14024
  },
13956
14025
  "z-ai/glm-5": {
13957
14026
  id: "z-ai/glm-5",
@@ -13967,8 +14036,8 @@ export const MODELS = {
13967
14036
  cacheRead: 0.119,
13968
14037
  cacheWrite: 0,
13969
14038
  },
13970
- contextWindow: 202752,
13971
- maxTokens: 202752,
14039
+ contextWindow: 204800,
14040
+ maxTokens: 131072,
13972
14041
  },
13973
14042
  "z-ai/glm-5-turbo": {
13974
14043
  id: "z-ai/glm-5-turbo",
@@ -14013,9 +14082,9 @@ export const MODELS = {
14013
14082
  reasoning: true,
14014
14083
  input: ["text"],
14015
14084
  cost: {
14016
- input: 0.9478,
14017
- output: 2.9788,
14018
- cacheRead: 0.17602,
14085
+ input: 0.9786,
14086
+ output: 3.0755999999999997,
14087
+ cacheRead: 0.18174,
14019
14088
  cacheWrite: 0,
14020
14089
  },
14021
14090
  contextWindow: 1048576,
@@ -14149,13 +14218,13 @@ export const MODELS = {
14149
14218
  reasoning: true,
14150
14219
  input: ["text", "image"],
14151
14220
  cost: {
14152
- input: 0.66,
14153
- output: 3.41,
14154
- cacheRead: 0.15,
14221
+ input: 3,
14222
+ output: 15,
14223
+ cacheRead: 0.3,
14155
14224
  cacheWrite: 0,
14156
14225
  },
14157
- contextWindow: 262144,
14158
- maxTokens: 262144,
14226
+ contextWindow: 1048576,
14227
+ maxTokens: 4096,
14159
14228
  },
14160
14229
  "~openai/gpt-latest": {
14161
14230
  id: "~openai/gpt-latest",
@@ -14202,7 +14271,7 @@ export const MODELS = {
14202
14271
  cost: {
14203
14272
  input: 2,
14204
14273
  output: 6,
14205
- cacheRead: 0.5,
14274
+ cacheRead: 0.3,
14206
14275
  cacheWrite: 0,
14207
14276
  },
14208
14277
  contextWindow: 500000,
@@ -14266,43 +14335,6 @@ export const MODELS = {
14266
14335
  contextWindow: 32768,
14267
14336
  maxTokens: 32768,
14268
14337
  },
14269
- "Qwen/Qwen3-235B-A22B-Instruct-2507-tput": {
14270
- id: "Qwen/Qwen3-235B-A22B-Instruct-2507-tput",
14271
- name: "Qwen3 235B A22B Instruct 2507 FP8",
14272
- api: "openai-completions",
14273
- provider: "together",
14274
- baseUrl: "https://api.together.ai/v1",
14275
- compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false },
14276
- reasoning: false,
14277
- input: ["text"],
14278
- cost: {
14279
- input: 0.2,
14280
- output: 0.6,
14281
- cacheRead: 0,
14282
- cacheWrite: 0,
14283
- },
14284
- contextWindow: 262144,
14285
- maxTokens: 262144,
14286
- },
14287
- "Qwen/Qwen3.5-397B-A17B": {
14288
- id: "Qwen/Qwen3.5-397B-A17B",
14289
- name: "Qwen3.5 397B A17B",
14290
- api: "openai-completions",
14291
- provider: "together",
14292
- baseUrl: "https://api.together.ai/v1",
14293
- compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "thinkingFormat": "together" },
14294
- reasoning: true,
14295
- thinkingLevelMap: { "minimal": null, "low": null, "medium": null },
14296
- input: ["text", "image"],
14297
- cost: {
14298
- input: 0.6,
14299
- output: 3.6,
14300
- cacheRead: 0,
14301
- cacheWrite: 0,
14302
- },
14303
- contextWindow: 262144,
14304
- maxTokens: 130000,
14305
- },
14306
14338
  "Qwen/Qwen3.5-9B": {
14307
14339
  id: "Qwen/Qwen3.5-9B",
14308
14340
  name: "Qwen3.5 9B",
@@ -14378,24 +14410,6 @@ export const MODELS = {
14378
14410
  contextWindow: 512000,
14379
14411
  maxTokens: 384000,
14380
14412
  },
14381
- "essentialai/Rnj-1-Instruct": {
14382
- id: "essentialai/Rnj-1-Instruct",
14383
- name: "Rnj-1 Instruct",
14384
- api: "openai-completions",
14385
- provider: "together",
14386
- baseUrl: "https://api.together.ai/v1",
14387
- compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false },
14388
- reasoning: false,
14389
- input: ["text"],
14390
- cost: {
14391
- input: 0.15,
14392
- output: 0.15,
14393
- cacheRead: 0,
14394
- cacheWrite: 0,
14395
- },
14396
- contextWindow: 32768,
14397
- maxTokens: 32768,
14398
- },
14399
14413
  "google/gemma-4-31B-it": {
14400
14414
  id: "google/gemma-4-31B-it",
14401
14415
  name: "Gemma 4 31B Instruct",
@@ -14528,42 +14542,23 @@ export const MODELS = {
14528
14542
  contextWindow: 131072,
14529
14543
  maxTokens: 131072,
14530
14544
  },
14531
- "zai-org/GLM-5": {
14532
- id: "zai-org/GLM-5",
14533
- name: "GLM-5",
14545
+ "thinkingmachines/Inkling": {
14546
+ id: "thinkingmachines/Inkling",
14547
+ name: "Inkling",
14534
14548
  api: "openai-completions",
14535
14549
  provider: "together",
14536
14550
  baseUrl: "https://api.together.ai/v1",
14537
14551
  compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "thinkingFormat": "together" },
14538
14552
  reasoning: true,
14539
14553
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null },
14540
- input: ["text"],
14554
+ input: ["text", "image"],
14541
14555
  cost: {
14542
14556
  input: 1,
14543
- output: 3.2,
14544
- cacheRead: 0,
14545
- cacheWrite: 0,
14546
- },
14547
- contextWindow: 202752,
14548
- maxTokens: 131072,
14549
- },
14550
- "zai-org/GLM-5.1": {
14551
- id: "zai-org/GLM-5.1",
14552
- name: "GLM-5.1",
14553
- api: "openai-completions",
14554
- provider: "together",
14555
- baseUrl: "https://api.together.ai/v1",
14556
- compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "thinkingFormat": "together" },
14557
- reasoning: true,
14558
- thinkingLevelMap: { "minimal": null, "low": null, "medium": null },
14559
- input: ["text"],
14560
- cost: {
14561
- input: 1.4,
14562
- output: 4.4,
14563
- cacheRead: 0,
14557
+ output: 4.05,
14558
+ cacheRead: 0.17,
14564
14559
  cacheWrite: 0,
14565
14560
  },
14566
- contextWindow: 202752,
14561
+ contextWindow: 524288,
14567
14562
  maxTokens: 131072,
14568
14563
  },
14569
14564
  "zai-org/GLM-5.2": {
@@ -15890,40 +15885,6 @@ export const MODELS = {
15890
15885
  contextWindow: 128000,
15891
15886
  maxTokens: 8192,
15892
15887
  },
15893
- "meta/llama-3.2-11b": {
15894
- id: "meta/llama-3.2-11b",
15895
- name: "Llama 3.2 11B Vision Instruct",
15896
- api: "anthropic-messages",
15897
- provider: "vercel-ai-gateway",
15898
- baseUrl: "https://ai-gateway.vercel.sh",
15899
- reasoning: false,
15900
- input: ["text", "image"],
15901
- cost: {
15902
- input: 0.16,
15903
- output: 0.16,
15904
- cacheRead: 0,
15905
- cacheWrite: 0,
15906
- },
15907
- contextWindow: 128000,
15908
- maxTokens: 8192,
15909
- },
15910
- "meta/llama-3.2-90b": {
15911
- id: "meta/llama-3.2-90b",
15912
- name: "Llama 3.2 90B Vision Instruct",
15913
- api: "anthropic-messages",
15914
- provider: "vercel-ai-gateway",
15915
- baseUrl: "https://ai-gateway.vercel.sh",
15916
- reasoning: false,
15917
- input: ["text", "image"],
15918
- cost: {
15919
- input: 0.72,
15920
- output: 0.72,
15921
- cacheRead: 0,
15922
- cacheWrite: 0,
15923
- },
15924
- contextWindow: 128000,
15925
- maxTokens: 8192,
15926
- },
15927
15888
  "meta/llama-3.3-70b": {
15928
15889
  id: "meta/llama-3.3-70b",
15929
15890
  name: "Llama 3.3 70B Instruct",
@@ -16468,6 +16429,23 @@ export const MODELS = {
16468
16429
  contextWindow: 262144,
16469
16430
  maxTokens: 32768,
16470
16431
  },
16432
+ "moonshotai/kimi-k3": {
16433
+ id: "moonshotai/kimi-k3",
16434
+ name: "Kimi K3",
16435
+ api: "anthropic-messages",
16436
+ provider: "vercel-ai-gateway",
16437
+ baseUrl: "https://ai-gateway.vercel.sh",
16438
+ reasoning: true,
16439
+ input: ["text", "image"],
16440
+ cost: {
16441
+ input: 3,
16442
+ output: 15,
16443
+ cacheRead: 0.3,
16444
+ cacheWrite: 0,
16445
+ },
16446
+ contextWindow: 1000000,
16447
+ maxTokens: 131072,
16448
+ },
16471
16449
  "nvidia/nemotron-3-nano-30b-a3b": {
16472
16450
  id: "nvidia/nemotron-3-nano-30b-a3b",
16473
16451
  name: "Nemotron 3 Nano 30B A3B",
@@ -17514,7 +17492,7 @@ export const MODELS = {
17514
17492
  cost: {
17515
17493
  input: 2,
17516
17494
  output: 6,
17517
- cacheRead: 0.5,
17495
+ cacheRead: 0.3,
17518
17496
  cacheWrite: 0,
17519
17497
  },
17520
17498
  contextWindow: 500000,
@@ -17924,7 +17902,7 @@ export const MODELS = {
17924
17902
  cost: {
17925
17903
  input: 2,
17926
17904
  output: 6,
17927
- cacheRead: 0.5,
17905
+ cacheRead: 0.3,
17928
17906
  cacheWrite: 0,
17929
17907
  },
17930
17908
  contextWindow: 500000,