omk-ai 0.96.2 → 0.97.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -797,7 +797,7 @@ export const MODELS = {
797
797
  cacheRead: 0.022,
798
798
  cacheWrite: 0.275,
799
799
  },
800
- contextWindow: 1050000,
800
+ contextWindow: 1000000,
801
801
  maxTokens: 128000,
802
802
  },
803
803
  "global.openai.gpt-5.6-sol": {
@@ -815,7 +815,7 @@ export const MODELS = {
815
815
  cacheRead: 0.55,
816
816
  cacheWrite: 6.875,
817
817
  },
818
- contextWindow: 1050000,
818
+ contextWindow: 1000000,
819
819
  maxTokens: 128000,
820
820
  },
821
821
  "global.openai.gpt-5.6-terra": {
@@ -833,7 +833,7 @@ export const MODELS = {
833
833
  cacheRead: 0.22,
834
834
  cacheWrite: 2.75,
835
835
  },
836
- contextWindow: 1050000,
836
+ contextWindow: 1000000,
837
837
  maxTokens: 128000,
838
838
  },
839
839
  "google.gemma-3-27b-it": {
@@ -1434,7 +1434,7 @@ export const MODELS = {
1434
1434
  cacheRead: 0.022,
1435
1435
  cacheWrite: 0.275,
1436
1436
  },
1437
- contextWindow: 1050000,
1437
+ contextWindow: 1000000,
1438
1438
  maxTokens: 128000,
1439
1439
  },
1440
1440
  "openai.gpt-5.6-sol": {
@@ -1452,7 +1452,7 @@ export const MODELS = {
1452
1452
  cacheRead: 0.55,
1453
1453
  cacheWrite: 6.875,
1454
1454
  },
1455
- contextWindow: 1050000,
1455
+ contextWindow: 1000000,
1456
1456
  maxTokens: 128000,
1457
1457
  },
1458
1458
  "openai.gpt-5.6-terra": {
@@ -1470,7 +1470,7 @@ export const MODELS = {
1470
1470
  cacheRead: 0.22,
1471
1471
  cacheWrite: 2.75,
1472
1472
  },
1473
- contextWindow: 1050000,
1473
+ contextWindow: 1000000,
1474
1474
  maxTokens: 128000,
1475
1475
  },
1476
1476
  "openai.gpt-oss-120b": {
@@ -2894,7 +2894,7 @@ export const MODELS = {
2894
2894
  cacheRead: 0.5,
2895
2895
  cacheWrite: 6.25,
2896
2896
  },
2897
- contextWindow: 1050000,
2897
+ contextWindow: 1000000,
2898
2898
  maxTokens: 128000,
2899
2899
  },
2900
2900
  "gpt-5.6-luna": {
@@ -2912,7 +2912,7 @@ export const MODELS = {
2912
2912
  cacheRead: 0.02,
2913
2913
  cacheWrite: 0.25,
2914
2914
  },
2915
- contextWindow: 1050000,
2915
+ contextWindow: 1000000,
2916
2916
  maxTokens: 128000,
2917
2917
  },
2918
2918
  "gpt-5.6-sol": {
@@ -2930,7 +2930,7 @@ export const MODELS = {
2930
2930
  cacheRead: 0.5,
2931
2931
  cacheWrite: 6.25,
2932
2932
  },
2933
- contextWindow: 1050000,
2933
+ contextWindow: 1000000,
2934
2934
  maxTokens: 128000,
2935
2935
  },
2936
2936
  "gpt-5.6-terra": {
@@ -2948,7 +2948,7 @@ export const MODELS = {
2948
2948
  cacheRead: 0.2,
2949
2949
  cacheWrite: 2.5,
2950
2950
  },
2951
- contextWindow: 1050000,
2951
+ contextWindow: 1000000,
2952
2952
  maxTokens: 128000,
2953
2953
  },
2954
2954
  "gpt-realtime-2.1": {
@@ -3727,7 +3727,7 @@ export const MODELS = {
3727
3727
  cacheRead: 0.5,
3728
3728
  cacheWrite: 6.25,
3729
3729
  },
3730
- contextWindow: 1050000,
3730
+ contextWindow: 1000000,
3731
3731
  maxTokens: 128000,
3732
3732
  },
3733
3733
  "gpt-5.6-luna": {
@@ -3745,7 +3745,7 @@ export const MODELS = {
3745
3745
  cacheRead: 0.02,
3746
3746
  cacheWrite: 0.25,
3747
3747
  },
3748
- contextWindow: 1050000,
3748
+ contextWindow: 1000000,
3749
3749
  maxTokens: 128000,
3750
3750
  },
3751
3751
  "gpt-5.6-sol": {
@@ -3763,7 +3763,7 @@ export const MODELS = {
3763
3763
  cacheRead: 0.5,
3764
3764
  cacheWrite: 6.25,
3765
3765
  },
3766
- contextWindow: 1050000,
3766
+ contextWindow: 1000000,
3767
3767
  maxTokens: 128000,
3768
3768
  },
3769
3769
  "gpt-5.6-terra": {
@@ -3781,7 +3781,7 @@ export const MODELS = {
3781
3781
  cacheRead: 0.2,
3782
3782
  cacheWrite: 2.5,
3783
3783
  },
3784
- contextWindow: 1050000,
3784
+ contextWindow: 1000000,
3785
3785
  maxTokens: 128000,
3786
3786
  },
3787
3787
  "o1": {
@@ -5354,7 +5354,7 @@ export const MODELS = {
5354
5354
  cacheRead: 0.02,
5355
5355
  cacheWrite: 0,
5356
5356
  },
5357
- contextWindow: 1050000,
5357
+ contextWindow: 1000000,
5358
5358
  maxTokens: 128000,
5359
5359
  },
5360
5360
  "gpt-5.6-sol": {
@@ -5368,12 +5368,12 @@ export const MODELS = {
5368
5368
  thinkingLevelMap: { "off": null, "minimal": "low", "xhigh": "xhigh", "max": "max" },
5369
5369
  input: ["text", "image"],
5370
5370
  cost: {
5371
- input: 5,
5372
- output: 30,
5373
- cacheRead: 0.5,
5374
- cacheWrite: 6.25,
5371
+ input: 2.5,
5372
+ output: 15,
5373
+ cacheRead: 0.25,
5374
+ cacheWrite: 3.125,
5375
5375
  },
5376
- contextWindow: 1050000,
5376
+ contextWindow: 1000000,
5377
5377
  maxTokens: 128000,
5378
5378
  },
5379
5379
  "gpt-5.6-terra": {
@@ -5392,7 +5392,7 @@ export const MODELS = {
5392
5392
  cacheRead: 0.2,
5393
5393
  cacheWrite: 0,
5394
5394
  },
5395
- contextWindow: 1050000,
5395
+ contextWindow: 1000000,
5396
5396
  maxTokens: 128000,
5397
5397
  },
5398
5398
  "grok-4.5": {
@@ -8550,6 +8550,26 @@ export const MODELS = {
8550
8550
  },
8551
8551
  },
8552
8552
  "nvidia": {
8553
+ "deepseek-ai/deepseek-v4-flash-0731": {
8554
+ id: "deepseek-ai/deepseek-v4-flash-0731",
8555
+ name: "DeepSeek V4 Flash 0731",
8556
+ api: "openai-completions",
8557
+ provider: "nvidia",
8558
+ baseUrl: "https://integrate.api.nvidia.com/v1",
8559
+ headers: { "NVCF-POLL-SECONDS": "3600" },
8560
+ compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "requiresReasoningContentOnAssistantMessages": true, "thinkingFormat": "deepseek" },
8561
+ reasoning: true,
8562
+ thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "max" },
8563
+ input: ["text"],
8564
+ cost: {
8565
+ input: 0,
8566
+ output: 0,
8567
+ cacheRead: 0,
8568
+ cacheWrite: 0,
8569
+ },
8570
+ contextWindow: 1000000,
8571
+ maxTokens: 384000,
8572
+ },
8553
8573
  "google/gemma-3-12b-it": {
8554
8574
  id: "google/gemma-3-12b-it",
8555
8575
  name: "Gemma 3 12B IT",
@@ -8759,6 +8779,25 @@ export const MODELS = {
8759
8779
  contextWindow: 262144,
8760
8780
  maxTokens: 262144,
8761
8781
  },
8782
+ "moonshotai/kimi-k3": {
8783
+ id: "moonshotai/kimi-k3",
8784
+ name: "Kimi K3",
8785
+ api: "openai-completions",
8786
+ provider: "nvidia",
8787
+ baseUrl: "https://integrate.api.nvidia.com/v1",
8788
+ headers: { "NVCF-POLL-SECONDS": "3600" },
8789
+ compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false },
8790
+ reasoning: true,
8791
+ input: ["text", "image"],
8792
+ cost: {
8793
+ input: 0,
8794
+ output: 0,
8795
+ cacheRead: 0,
8796
+ cacheWrite: 0,
8797
+ },
8798
+ contextWindow: 1048576,
8799
+ maxTokens: 131072,
8800
+ },
8762
8801
  "nvidia/cosmos-reason2-8b": {
8763
8802
  id: "nvidia/cosmos-reason2-8b",
8764
8803
  name: "Cosmos Reason2 8B",
@@ -9667,7 +9706,7 @@ export const MODELS = {
9667
9706
  cacheRead: 0.5,
9668
9707
  cacheWrite: 6.25,
9669
9708
  },
9670
- contextWindow: 1050000,
9709
+ contextWindow: 1000000,
9671
9710
  maxTokens: 128000,
9672
9711
  },
9673
9712
  "gpt-5.6-luna": {
@@ -9685,7 +9724,7 @@ export const MODELS = {
9685
9724
  cacheRead: 0.02,
9686
9725
  cacheWrite: 0.25,
9687
9726
  },
9688
- contextWindow: 1050000,
9727
+ contextWindow: 1000000,
9689
9728
  maxTokens: 128000,
9690
9729
  },
9691
9730
  "gpt-5.6-sol": {
@@ -9703,7 +9742,7 @@ export const MODELS = {
9703
9742
  cacheRead: 0.5,
9704
9743
  cacheWrite: 6.25,
9705
9744
  },
9706
- contextWindow: 1050000,
9745
+ contextWindow: 1000000,
9707
9746
  maxTokens: 128000,
9708
9747
  },
9709
9748
  "gpt-5.6-terra": {
@@ -9721,7 +9760,7 @@ export const MODELS = {
9721
9760
  cacheRead: 0.2,
9722
9761
  cacheWrite: 2.5,
9723
9762
  },
9724
- contextWindow: 1050000,
9763
+ contextWindow: 1000000,
9725
9764
  maxTokens: 128000,
9726
9765
  },
9727
9766
  "gpt-realtime-2.1": {
@@ -9932,7 +9971,7 @@ export const MODELS = {
9932
9971
  cacheRead: 0.1,
9933
9972
  cacheWrite: 0,
9934
9973
  },
9935
- contextWindow: 372000,
9974
+ contextWindow: 1000000,
9936
9975
  maxTokens: 128000,
9937
9976
  },
9938
9977
  "gpt-5.6-moa": {
@@ -9950,7 +9989,7 @@ export const MODELS = {
9950
9989
  cacheRead: 0,
9951
9990
  cacheWrite: 0,
9952
9991
  },
9953
- contextWindow: 300000,
9992
+ contextWindow: 1000000,
9954
9993
  maxTokens: 128000,
9955
9994
  },
9956
9995
  "gpt-5.6-sol": {
@@ -9986,7 +10025,7 @@ export const MODELS = {
9986
10025
  cacheRead: 0.25,
9987
10026
  cacheWrite: 0,
9988
10027
  },
9989
- contextWindow: 372000,
10028
+ contextWindow: 1000000,
9990
10029
  maxTokens: 128000,
9991
10030
  },
9992
10031
  },
@@ -10705,7 +10744,7 @@ export const MODELS = {
10705
10744
  cacheRead: 0.02,
10706
10745
  cacheWrite: 0.25,
10707
10746
  },
10708
- contextWindow: 1050000,
10747
+ contextWindow: 1000000,
10709
10748
  maxTokens: 128000,
10710
10749
  },
10711
10750
  "gpt-5.6-sol": {
@@ -10718,12 +10757,12 @@ export const MODELS = {
10718
10757
  thinkingLevelMap: { "off": null, "xhigh": "xhigh", "max": "max" },
10719
10758
  input: ["text", "image"],
10720
10759
  cost: {
10721
- input: 2.5,
10722
- output: 15,
10723
- cacheRead: 0.25,
10724
- cacheWrite: 3.125,
10760
+ input: 2,
10761
+ output: 10,
10762
+ cacheRead: 0.2,
10763
+ cacheWrite: 2.5,
10725
10764
  },
10726
- contextWindow: 1050000,
10765
+ contextWindow: 1000000,
10727
10766
  maxTokens: 128000,
10728
10767
  },
10729
10768
  "gpt-5.6-terra": {
@@ -10741,7 +10780,7 @@ export const MODELS = {
10741
10780
  cacheRead: 0.25,
10742
10781
  cacheWrite: 3.125,
10743
10782
  },
10744
- contextWindow: 1050000,
10783
+ contextWindow: 1000000,
10745
10784
  maxTokens: 128000,
10746
10785
  },
10747
10786
  "grok-4.5": {
@@ -11093,6 +11132,25 @@ export const MODELS = {
11093
11132
  contextWindow: 1000000,
11094
11133
  maxTokens: 384000,
11095
11134
  },
11135
+ "deepseek-v4-flash-vision-exp": {
11136
+ id: "deepseek-v4-flash-vision-exp",
11137
+ name: "DeepSeek V4 Flash Vision Exp",
11138
+ api: "openai-completions",
11139
+ provider: "opencode-go",
11140
+ baseUrl: "https://opencode.ai/zen/go/v1",
11141
+ compat: { "requiresReasoningContentOnAssistantMessages": true, "thinkingFormat": "deepseek" },
11142
+ reasoning: true,
11143
+ thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "max" },
11144
+ input: ["text", "image"],
11145
+ cost: {
11146
+ input: 0.22,
11147
+ output: 0.66,
11148
+ cacheRead: 0.007,
11149
+ cacheWrite: 0,
11150
+ },
11151
+ contextWindow: 1000000,
11152
+ maxTokens: 384000,
11153
+ },
11096
11154
  "deepseek-v4-pro": {
11097
11155
  id: "deepseek-v4-pro",
11098
11156
  name: "DeepSeek V4 Pro (New)",
@@ -11180,7 +11238,7 @@ export const MODELS = {
11180
11238
  cacheRead: 0.02,
11181
11239
  cacheWrite: 0.25,
11182
11240
  },
11183
- contextWindow: 1050000,
11241
+ contextWindow: 1000000,
11184
11242
  maxTokens: 128000,
11185
11243
  },
11186
11244
  "grok-4.5": {
@@ -11270,6 +11328,23 @@ export const MODELS = {
11270
11328
  contextWindow: 1048576,
11271
11329
  maxTokens: 131072,
11272
11330
  },
11331
+ "longcat-2.0": {
11332
+ id: "longcat-2.0",
11333
+ name: "LongCat-2.0",
11334
+ api: "openai-completions",
11335
+ provider: "opencode-go",
11336
+ baseUrl: "https://opencode.ai/zen/go/v1",
11337
+ reasoning: true,
11338
+ input: ["text"],
11339
+ cost: {
11340
+ input: 0.3,
11341
+ output: 1.2,
11342
+ cacheRead: 0.006,
11343
+ cacheWrite: 0,
11344
+ },
11345
+ contextWindow: 1000000,
11346
+ maxTokens: 131072,
11347
+ },
11273
11348
  "mimo-v2.5": {
11274
11349
  id: "mimo-v2.5",
11275
11350
  name: "MiMo V2.5",
@@ -12316,13 +12391,13 @@ export const MODELS = {
12316
12391
  reasoning: true,
12317
12392
  input: ["text"],
12318
12393
  cost: {
12319
- input: 0.25,
12320
- output: 0.95,
12321
- cacheRead: 0.13,
12394
+ input: 0.55,
12395
+ output: 1.6500000000000001,
12396
+ cacheRead: 0.55,
12322
12397
  cacheWrite: 0,
12323
12398
  },
12324
12399
  contextWindow: 163840,
12325
- maxTokens: 32768,
12400
+ maxTokens: 161000,
12326
12401
  },
12327
12402
  "deepseek/deepseek-r1": {
12328
12403
  id: "deepseek/deepseek-r1",
@@ -12384,13 +12459,13 @@ export const MODELS = {
12384
12459
  reasoning: true,
12385
12460
  input: ["text"],
12386
12461
  cost: {
12387
- input: 0.26899999999999996,
12388
- output: 0.39999999999999997,
12389
- cacheRead: 0.13449999999999998,
12462
+ input: 0.26,
12463
+ output: 0.38,
12464
+ cacheRead: 0.13,
12390
12465
  cacheWrite: 0,
12391
12466
  },
12392
12467
  contextWindow: 163840,
12393
- maxTokens: 65536,
12468
+ maxTokens: 163840,
12394
12469
  },
12395
12470
  "deepseek/deepseek-v3.2-exp": {
12396
12471
  id: "deepseek/deepseek-v3.2-exp",
@@ -12420,9 +12495,9 @@ export const MODELS = {
12420
12495
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
12421
12496
  input: ["text"],
12422
12497
  cost: {
12423
- input: 0.08259999999999999,
12424
- output: 0.16519999999999999,
12425
- cacheRead: 0.01652,
12498
+ input: 0.056,
12499
+ output: 0.112,
12500
+ cacheRead: 0.0112,
12426
12501
  cacheWrite: 0,
12427
12502
  },
12428
12503
  contextWindow: 1048576,
@@ -12439,12 +12514,31 @@ export const MODELS = {
12439
12514
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
12440
12515
  input: ["text"],
12441
12516
  cost: {
12442
- input: 0.08,
12443
- output: 0.18,
12444
- cacheRead: 0.016,
12517
+ input: 0.0658,
12518
+ output: 0.1316,
12519
+ cacheRead: 0.01316,
12445
12520
  cacheWrite: 0,
12446
12521
  },
12447
12522
  contextWindow: 1310720,
12523
+ maxTokens: 131072,
12524
+ },
12525
+ "deepseek/deepseek-v4-flash-vision-exp": {
12526
+ id: "deepseek/deepseek-v4-flash-vision-exp",
12527
+ name: "DeepSeek: DeepSeek V4 Flash Vision Exp",
12528
+ api: "openai-completions",
12529
+ provider: "openrouter",
12530
+ baseUrl: "https://openrouter.ai/api/v1",
12531
+ compat: { "requiresReasoningContentOnAssistantMessages": true },
12532
+ reasoning: true,
12533
+ thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
12534
+ input: ["text", "image"],
12535
+ cost: {
12536
+ input: 0.22,
12537
+ output: 0.66,
12538
+ cacheRead: 0.007,
12539
+ cacheWrite: 0,
12540
+ },
12541
+ contextWindow: 1048576,
12448
12542
  maxTokens: 384000,
12449
12543
  },
12450
12544
  "deepseek/deepseek-v4-pro": {
@@ -12458,13 +12552,13 @@ export const MODELS = {
12458
12552
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
12459
12553
  input: ["text"],
12460
12554
  cost: {
12461
- input: 1.5999999999999999,
12462
- output: 3.1999999999999997,
12463
- cacheRead: 0.135,
12555
+ input: 0.522,
12556
+ output: 1.044,
12557
+ cacheRead: 0.0435,
12464
12558
  cacheWrite: 0,
12465
12559
  },
12466
12560
  contextWindow: 1048576,
12467
- maxTokens: 393216,
12561
+ maxTokens: 384000,
12468
12562
  },
12469
12563
  "deepseek/deepseek-v4-pro-0813": {
12470
12564
  id: "deepseek/deepseek-v4-pro-0813",
@@ -12477,9 +12571,9 @@ export const MODELS = {
12477
12571
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
12478
12572
  input: ["text"],
12479
12573
  cost: {
12480
- input: 1.1880000000000002,
12481
- output: 3.564,
12482
- cacheRead: 0.039599999999999996,
12574
+ input: 1.122,
12575
+ output: 3.366,
12576
+ cacheRead: 0.037399999999999996,
12483
12577
  cacheWrite: 0,
12484
12578
  },
12485
12579
  contextWindow: 1048576,
@@ -13064,40 +13158,6 @@ export const MODELS = {
13064
13158
  contextWindow: 128000,
13065
13159
  maxTokens: 50000,
13066
13160
  },
13067
- "inclusionai/ling-2.6-1t": {
13068
- id: "inclusionai/ling-2.6-1t",
13069
- name: "inclusionAI: Ling-2.6-1T",
13070
- api: "openai-completions",
13071
- provider: "openrouter",
13072
- baseUrl: "https://openrouter.ai/api/v1",
13073
- reasoning: false,
13074
- input: ["text"],
13075
- cost: {
13076
- input: 0.075,
13077
- output: 0.625,
13078
- cacheRead: 0.015,
13079
- cacheWrite: 0,
13080
- },
13081
- contextWindow: 262144,
13082
- maxTokens: 32768,
13083
- },
13084
- "inclusionai/ling-2.6-flash": {
13085
- id: "inclusionai/ling-2.6-flash",
13086
- name: "inclusionAI: Ling-2.6-flash",
13087
- api: "openai-completions",
13088
- provider: "openrouter",
13089
- baseUrl: "https://openrouter.ai/api/v1",
13090
- reasoning: false,
13091
- input: ["text"],
13092
- cost: {
13093
- input: 0.01,
13094
- output: 0.03,
13095
- cacheRead: 0.002,
13096
- cacheWrite: 0,
13097
- },
13098
- contextWindow: 262144,
13099
- maxTokens: 32768,
13100
- },
13101
13161
  "inclusionai/ling-3.0-flash": {
13102
13162
  id: "inclusionai/ling-3.0-flash",
13103
13163
  name: "Ling-3.0-flash",
@@ -13115,23 +13175,6 @@ export const MODELS = {
13115
13175
  contextWindow: 262144,
13116
13176
  maxTokens: 32768,
13117
13177
  },
13118
- "inclusionai/ring-2.6-1t": {
13119
- id: "inclusionai/ring-2.6-1t",
13120
- name: "inclusionAI: Ring-2.6-1T",
13121
- api: "openai-completions",
13122
- provider: "openrouter",
13123
- baseUrl: "https://openrouter.ai/api/v1",
13124
- reasoning: true,
13125
- input: ["text"],
13126
- cost: {
13127
- input: 0.075,
13128
- output: 0.625,
13129
- cacheRead: 0.015,
13130
- cacheWrite: 0,
13131
- },
13132
- contextWindow: 262144,
13133
- maxTokens: 65536,
13134
- },
13135
13178
  "kwaipilot/kat-coder-air-v2.5": {
13136
13179
  id: "kwaipilot/kat-coder-air-v2.5",
13137
13180
  name: "Kwaipilot: KAT-Coder-Air V2.5",
@@ -13311,13 +13354,13 @@ export const MODELS = {
13311
13354
  reasoning: true,
13312
13355
  input: ["text", "image"],
13313
13356
  cost: {
13314
- input: 0.3,
13315
- output: 1.1,
13357
+ input: 0.35,
13358
+ output: 1.5,
13316
13359
  cacheRead: 0.04,
13317
13360
  cacheWrite: 0,
13318
13361
  },
13319
13362
  contextWindow: 131072,
13320
- maxTokens: 131072,
13363
+ maxTokens: 4096,
13321
13364
  },
13322
13365
  "meta/muse-spark-1.1": {
13323
13366
  id: "meta/muse-spark-1.1",
@@ -13353,6 +13396,23 @@ export const MODELS = {
13353
13396
  contextWindow: 1048576,
13354
13397
  maxTokens: 4096,
13355
13398
  },
13399
+ "meta/muse-spark-1.2-contributor": {
13400
+ id: "meta/muse-spark-1.2-contributor",
13401
+ name: "Meta: Muse Spark 1.2 Contributor",
13402
+ api: "openai-completions",
13403
+ provider: "openrouter",
13404
+ baseUrl: "https://openrouter.ai/api/v1",
13405
+ reasoning: true,
13406
+ input: ["text", "image"],
13407
+ cost: {
13408
+ input: 0.09999999999999999,
13409
+ output: 0.19999999999999998,
13410
+ cacheRead: 0.002,
13411
+ cacheWrite: 0,
13412
+ },
13413
+ contextWindow: 1048576,
13414
+ maxTokens: 4096,
13415
+ },
13356
13416
  "minimax/minimax-m1": {
13357
13417
  id: "minimax/minimax-m1",
13358
13418
  name: "MiniMax: MiniMax M1",
@@ -13414,12 +13474,12 @@ export const MODELS = {
13414
13474
  input: ["text"],
13415
13475
  cost: {
13416
13476
  input: 0.27,
13417
- output: 0.95,
13418
- cacheRead: 0.03,
13477
+ output: 1.08,
13478
+ cacheRead: 0.027,
13419
13479
  cacheWrite: 0,
13420
13480
  },
13421
13481
  contextWindow: 204800,
13422
- maxTokens: 32768,
13482
+ maxTokens: 128000,
13423
13483
  },
13424
13484
  "minimax/minimax-m2.7": {
13425
13485
  id: "minimax/minimax-m2.7",
@@ -13436,7 +13496,7 @@ export const MODELS = {
13436
13496
  cacheWrite: 0,
13437
13497
  },
13438
13498
  contextWindow: 204800,
13439
- maxTokens: 131072,
13499
+ maxTokens: 4096,
13440
13500
  },
13441
13501
  "minimax/minimax-m3": {
13442
13502
  id: "minimax/minimax-m3",
@@ -13489,6 +13549,23 @@ export const MODELS = {
13489
13549
  contextWindow: 256000,
13490
13550
  maxTokens: 4096,
13491
13551
  },
13552
+ "mistralai/devstral-2512": {
13553
+ id: "mistralai/devstral-2512",
13554
+ name: "Mistral: Devstral 2 2512",
13555
+ api: "openai-completions",
13556
+ provider: "openrouter",
13557
+ baseUrl: "https://openrouter.ai/api/v1",
13558
+ reasoning: false,
13559
+ input: ["text"],
13560
+ cost: {
13561
+ input: 0.44,
13562
+ output: 2.2,
13563
+ cacheRead: 0.044,
13564
+ cacheWrite: 0,
13565
+ },
13566
+ contextWindow: 262144,
13567
+ maxTokens: 4096,
13568
+ },
13492
13569
  "mistralai/ministral-14b-2512": {
13493
13570
  id: "mistralai/ministral-14b-2512",
13494
13571
  name: "Mistral: Ministral 3 14B 2512",
@@ -13702,12 +13779,12 @@ export const MODELS = {
13702
13779
  reasoning: false,
13703
13780
  input: ["text", "image"],
13704
13781
  cost: {
13705
- input: 0.09375,
13706
- output: 0.25,
13707
- cacheRead: 0,
13782
+ input: 0.075,
13783
+ output: 0.19999999999999998,
13784
+ cacheRead: 0,
13708
13785
  cacheWrite: 0,
13709
13786
  },
13710
- contextWindow: 256000,
13787
+ contextWindow: 131072,
13711
13788
  maxTokens: 16384,
13712
13789
  },
13713
13790
  "mistralai/mixtral-8x22b-instruct": {
@@ -13828,7 +13905,7 @@ export const MODELS = {
13828
13905
  cacheWrite: 0,
13829
13906
  },
13830
13907
  contextWindow: 262144,
13831
- maxTokens: 262144,
13908
+ maxTokens: 4096,
13832
13909
  },
13833
13910
  "moonshotai/kimi-k2.7-code": {
13834
13911
  id: "moonshotai/kimi-k2.7-code",
@@ -13841,7 +13918,7 @@ export const MODELS = {
13841
13918
  cost: {
13842
13919
  input: 0.67,
13843
13920
  output: 3.4,
13844
- cacheRead: 0.16999999999999998,
13921
+ cacheRead: 0.19,
13845
13922
  cacheWrite: 0,
13846
13923
  },
13847
13924
  contextWindow: 262144,
@@ -13879,7 +13956,7 @@ export const MODELS = {
13879
13956
  cacheWrite: 0,
13880
13957
  },
13881
13958
  contextWindow: 1048576,
13882
- maxTokens: 4096,
13959
+ maxTokens: 1048576,
13883
13960
  },
13884
13961
  "nex-agi/nex-n2-mini": {
13885
13962
  id: "nex-agi/nex-n2-mini",
@@ -13932,23 +14009,6 @@ export const MODELS = {
13932
14009
  contextWindow: 262144,
13933
14010
  maxTokens: 262144,
13934
14011
  },
13935
- "nvidia/nemotron-3-nano-30b-a3b:free": {
13936
- id: "nvidia/nemotron-3-nano-30b-a3b:free",
13937
- name: "NVIDIA: Nemotron 3 Nano 30B A3B (free)",
13938
- api: "openai-completions",
13939
- provider: "openrouter",
13940
- baseUrl: "https://openrouter.ai/api/v1",
13941
- reasoning: true,
13942
- input: ["text"],
13943
- cost: {
13944
- input: 0,
13945
- output: 0,
13946
- cacheRead: 0,
13947
- cacheWrite: 0,
13948
- },
13949
- contextWindow: 256000,
13950
- maxTokens: 4096,
13951
- },
13952
14012
  "nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free": {
13953
14013
  id: "nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free",
13954
14014
  name: "NVIDIA: Nemotron 3 Nano Omni (free)",
@@ -14085,40 +14145,6 @@ export const MODELS = {
14085
14145
  contextWindow: 1000000,
14086
14146
  maxTokens: 65536,
14087
14147
  },
14088
- "nvidia/nemotron-nano-12b-v2-vl:free": {
14089
- id: "nvidia/nemotron-nano-12b-v2-vl:free",
14090
- name: "NVIDIA: Nemotron Nano 12B 2 VL (free)",
14091
- api: "openai-completions",
14092
- provider: "openrouter",
14093
- baseUrl: "https://openrouter.ai/api/v1",
14094
- reasoning: true,
14095
- input: ["text", "image"],
14096
- cost: {
14097
- input: 0,
14098
- output: 0,
14099
- cacheRead: 0,
14100
- cacheWrite: 0,
14101
- },
14102
- contextWindow: 128000,
14103
- maxTokens: 128000,
14104
- },
14105
- "nvidia/nemotron-nano-9b-v2:free": {
14106
- id: "nvidia/nemotron-nano-9b-v2:free",
14107
- name: "NVIDIA: Nemotron Nano 9B V2 (free)",
14108
- api: "openai-completions",
14109
- provider: "openrouter",
14110
- baseUrl: "https://openrouter.ai/api/v1",
14111
- reasoning: true,
14112
- input: ["text"],
14113
- cost: {
14114
- input: 0,
14115
- output: 0,
14116
- cacheRead: 0,
14117
- cacheWrite: 0,
14118
- },
14119
- contextWindow: 128000,
14120
- maxTokens: 4096,
14121
- },
14122
14148
  "openai/gpt-3.5-turbo": {
14123
14149
  id: "openai/gpt-3.5-turbo",
14124
14150
  name: "OpenAI: GPT-3.5 Turbo",
@@ -15071,7 +15097,7 @@ export const MODELS = {
15071
15097
  cacheRead: 0.02,
15072
15098
  cacheWrite: 0.25,
15073
15099
  },
15074
- contextWindow: 1050000,
15100
+ contextWindow: 1000000,
15075
15101
  maxTokens: 128000,
15076
15102
  },
15077
15103
  "openai/gpt-5.6-luna-pro": {
@@ -15089,7 +15115,7 @@ export const MODELS = {
15089
15115
  cacheRead: 0.02,
15090
15116
  cacheWrite: 0.25,
15091
15117
  },
15092
- contextWindow: 1050000,
15118
+ contextWindow: 1000000,
15093
15119
  maxTokens: 128000,
15094
15120
  },
15095
15121
  "openai/gpt-5.6-luna-pro:batch": {
@@ -15107,7 +15133,7 @@ export const MODELS = {
15107
15133
  cacheRead: 0.01,
15108
15134
  cacheWrite: 0,
15109
15135
  },
15110
- contextWindow: 1050000,
15136
+ contextWindow: 1000000,
15111
15137
  maxTokens: 128000,
15112
15138
  },
15113
15139
  "openai/gpt-5.6-luna:batch": {
@@ -15125,7 +15151,7 @@ export const MODELS = {
15125
15151
  cacheRead: 0.01,
15126
15152
  cacheWrite: 0,
15127
15153
  },
15128
- contextWindow: 1050000,
15154
+ contextWindow: 1000000,
15129
15155
  maxTokens: 128000,
15130
15156
  },
15131
15157
  "openai/gpt-5.6-sol": {
@@ -15138,12 +15164,12 @@ export const MODELS = {
15138
15164
  thinkingLevelMap: { "xhigh": "xhigh", "max": "max" },
15139
15165
  input: ["text", "image"],
15140
15166
  cost: {
15141
- input: 2.5,
15142
- output: 15,
15143
- cacheRead: 0.25,
15144
- cacheWrite: 3.125,
15167
+ input: 2,
15168
+ output: 10,
15169
+ cacheRead: 0.19999999999999998,
15170
+ cacheWrite: 2.5,
15145
15171
  },
15146
- contextWindow: 1050000,
15172
+ contextWindow: 1000000,
15147
15173
  maxTokens: 128000,
15148
15174
  },
15149
15175
  "openai/gpt-5.6-sol-pro": {
@@ -15156,12 +15182,12 @@ export const MODELS = {
15156
15182
  thinkingLevelMap: { "xhigh": "xhigh", "max": "max" },
15157
15183
  input: ["text", "image"],
15158
15184
  cost: {
15159
- input: 2.5,
15160
- output: 15,
15161
- cacheRead: 0.25,
15162
- cacheWrite: 3.125,
15185
+ input: 2,
15186
+ output: 10,
15187
+ cacheRead: 0.19999999999999998,
15188
+ cacheWrite: 2.5,
15163
15189
  },
15164
- contextWindow: 1050000,
15190
+ contextWindow: 1000000,
15165
15191
  maxTokens: 128000,
15166
15192
  },
15167
15193
  "openai/gpt-5.6-sol-pro:batch": {
@@ -15174,12 +15200,12 @@ export const MODELS = {
15174
15200
  thinkingLevelMap: { "xhigh": "xhigh", "max": "max" },
15175
15201
  input: ["text", "image"],
15176
15202
  cost: {
15177
- input: 1.25,
15178
- output: 7.5,
15179
- cacheRead: 0.125,
15180
- cacheWrite: 0,
15203
+ input: 1,
15204
+ output: 5,
15205
+ cacheRead: 0.09999999999999999,
15206
+ cacheWrite: 1.25,
15181
15207
  },
15182
- contextWindow: 1050000,
15208
+ contextWindow: 1000000,
15183
15209
  maxTokens: 128000,
15184
15210
  },
15185
15211
  "openai/gpt-5.6-sol:batch": {
@@ -15192,12 +15218,12 @@ export const MODELS = {
15192
15218
  thinkingLevelMap: { "xhigh": "xhigh", "max": "max" },
15193
15219
  input: ["text", "image"],
15194
15220
  cost: {
15195
- input: 1.25,
15196
- output: 7.5,
15197
- cacheRead: 0.125,
15198
- cacheWrite: 0,
15221
+ input: 1,
15222
+ output: 5,
15223
+ cacheRead: 0.09999999999999999,
15224
+ cacheWrite: 1.25,
15199
15225
  },
15200
- contextWindow: 1050000,
15226
+ contextWindow: 1000000,
15201
15227
  maxTokens: 128000,
15202
15228
  },
15203
15229
  "openai/gpt-5.6-terra": {
@@ -15215,7 +15241,7 @@ export const MODELS = {
15215
15241
  cacheRead: 0.19999999999999998,
15216
15242
  cacheWrite: 2.5,
15217
15243
  },
15218
- contextWindow: 1050000,
15244
+ contextWindow: 1000000,
15219
15245
  maxTokens: 128000,
15220
15246
  },
15221
15247
  "openai/gpt-5.6-terra-pro": {
@@ -15233,7 +15259,7 @@ export const MODELS = {
15233
15259
  cacheRead: 0.19999999999999998,
15234
15260
  cacheWrite: 2.5,
15235
15261
  },
15236
- contextWindow: 1050000,
15262
+ contextWindow: 1000000,
15237
15263
  maxTokens: 128000,
15238
15264
  },
15239
15265
  "openai/gpt-5.6-terra-pro:batch": {
@@ -15251,7 +15277,7 @@ export const MODELS = {
15251
15277
  cacheRead: 0.09999999999999999,
15252
15278
  cacheWrite: 0,
15253
15279
  },
15254
- contextWindow: 1050000,
15280
+ contextWindow: 1000000,
15255
15281
  maxTokens: 128000,
15256
15282
  },
15257
15283
  "openai/gpt-5.6-terra:batch": {
@@ -15269,7 +15295,7 @@ export const MODELS = {
15269
15295
  cacheRead: 0.09999999999999999,
15270
15296
  cacheWrite: 0,
15271
15297
  },
15272
- contextWindow: 1050000,
15298
+ contextWindow: 1000000,
15273
15299
  maxTokens: 128000,
15274
15300
  },
15275
15301
  "openai/gpt-5:batch": {
@@ -15349,9 +15375,9 @@ export const MODELS = {
15349
15375
  reasoning: true,
15350
15376
  input: ["text"],
15351
15377
  cost: {
15352
- input: 0.03,
15378
+ input: 0.037,
15353
15379
  output: 0.16999999999999998,
15354
- cacheRead: 0.03,
15380
+ cacheRead: 0,
15355
15381
  cacheWrite: 0,
15356
15382
  },
15357
15383
  contextWindow: 131072,
@@ -15374,23 +15400,6 @@ export const MODELS = {
15374
15400
  contextWindow: 131072,
15375
15401
  maxTokens: 131072,
15376
15402
  },
15377
- "openai/gpt-oss-20b:free": {
15378
- id: "openai/gpt-oss-20b:free",
15379
- name: "OpenAI: gpt-oss-20b (free)",
15380
- api: "openai-completions",
15381
- provider: "openrouter",
15382
- baseUrl: "https://openrouter.ai/api/v1",
15383
- reasoning: true,
15384
- input: ["text"],
15385
- cost: {
15386
- input: 0,
15387
- output: 0,
15388
- cacheRead: 0,
15389
- cacheWrite: 0,
15390
- },
15391
- contextWindow: 131072,
15392
- maxTokens: 32768,
15393
- },
15394
15403
  "openai/gpt-oss-safeguard-20b": {
15395
15404
  id: "openai/gpt-oss-safeguard-20b",
15396
15405
  name: "OpenAI: gpt-oss-safeguard-20b",
@@ -15919,7 +15928,7 @@ export const MODELS = {
15919
15928
  cacheRead: 0,
15920
15929
  cacheWrite: 0,
15921
15930
  },
15922
- contextWindow: 262144,
15931
+ contextWindow: 131072,
15923
15932
  maxTokens: 4096,
15924
15933
  },
15925
15934
  "qwen/qwen3-30b-a3b": {
@@ -15932,13 +15941,13 @@ export const MODELS = {
15932
15941
  thinkingLevelMap: { "xhigh": "high", "max": "max" },
15933
15942
  input: ["text"],
15934
15943
  cost: {
15935
- input: 0.13,
15936
- output: 0.52,
15944
+ input: 0.12,
15945
+ output: 0.5,
15937
15946
  cacheRead: 0,
15938
15947
  cacheWrite: 0,
15939
15948
  },
15940
15949
  contextWindow: 131072,
15941
- maxTokens: 8192,
15950
+ maxTokens: 16384,
15942
15951
  },
15943
15952
  "qwen/qwen3-30b-a3b-instruct-2507": {
15944
15953
  id: "qwen/qwen3-30b-a3b-instruct-2507",
@@ -16140,13 +16149,13 @@ export const MODELS = {
16140
16149
  reasoning: false,
16141
16150
  input: ["text"],
16142
16151
  cost: {
16143
- input: 0.09,
16152
+ input: 0.09999999999999999,
16144
16153
  output: 1.1,
16145
- cacheRead: 0,
16154
+ cacheRead: 0.07,
16146
16155
  cacheWrite: 0,
16147
16156
  },
16148
16157
  contextWindow: 262144,
16149
- maxTokens: 16384,
16158
+ maxTokens: 262144,
16150
16159
  },
16151
16160
  "qwen/qwen3-next-80b-a3b-thinking": {
16152
16161
  id: "qwen/qwen3-next-80b-a3b-thinking",
@@ -16304,7 +16313,7 @@ export const MODELS = {
16304
16313
  cacheWrite: 0,
16305
16314
  },
16306
16315
  contextWindow: 262144,
16307
- maxTokens: 65536,
16316
+ maxTokens: 262144,
16308
16317
  },
16309
16318
  "qwen/qwen3.5-27b": {
16310
16319
  id: "qwen/qwen3.5-27b",
@@ -16352,13 +16361,13 @@ export const MODELS = {
16352
16361
  thinkingLevelMap: { "xhigh": "high", "max": "max" },
16353
16362
  input: ["text", "image"],
16354
16363
  cost: {
16355
- input: 0.39,
16356
- output: 2.34,
16357
- cacheRead: 0,
16364
+ input: 0.5,
16365
+ output: 3.5999999999999996,
16366
+ cacheRead: 0.3,
16358
16367
  cacheWrite: 0,
16359
16368
  },
16360
16369
  contextWindow: 262144,
16361
- maxTokens: 65536,
16370
+ maxTokens: 262144,
16362
16371
  },
16363
16372
  "qwen/qwen3.5-9b": {
16364
16373
  id: "qwen/qwen3.5-9b",
@@ -16442,13 +16451,13 @@ export const MODELS = {
16442
16451
  thinkingLevelMap: { "xhigh": "high", "max": "max" },
16443
16452
  input: ["text", "image"],
16444
16453
  cost: {
16445
- input: 0.6,
16446
- output: 3.5999999999999996,
16447
- cacheRead: 0.12,
16454
+ input: 0.32,
16455
+ output: 3.1999999999999997,
16456
+ cacheRead: 0,
16448
16457
  cacheWrite: 0,
16449
16458
  },
16450
16459
  contextWindow: 262144,
16451
- maxTokens: 262144,
16460
+ maxTokens: 81920,
16452
16461
  },
16453
16462
  "qwen/qwen3.6-35b-a3b": {
16454
16463
  id: "qwen/qwen3.6-35b-a3b",
@@ -16592,7 +16601,7 @@ export const MODELS = {
16592
16601
  cacheWrite: 0,
16593
16602
  },
16594
16603
  contextWindow: 1048576,
16595
- maxTokens: 262144,
16604
+ maxTokens: 131072,
16596
16605
  },
16597
16606
  "qwen/qwen3.8-27b": {
16598
16607
  id: "qwen/qwen3.8-27b",
@@ -16604,8 +16613,8 @@ export const MODELS = {
16604
16613
  thinkingLevelMap: { "xhigh": "high", "max": "max" },
16605
16614
  input: ["text", "image"],
16606
16615
  cost: {
16607
- input: 0.44999999999999996,
16608
- output: 3.1999999999999997,
16616
+ input: 0.39999999999999997,
16617
+ output: 3,
16609
16618
  cacheRead: 0.049999999999999996,
16610
16619
  cacheWrite: 0,
16611
16620
  },
@@ -16816,7 +16825,7 @@ export const MODELS = {
16816
16825
  cacheWrite: 0,
16817
16826
  },
16818
16827
  contextWindow: 1024000,
16819
- maxTokens: 1024000,
16828
+ maxTokens: 32768,
16820
16829
  },
16821
16830
  "thinkingmachines/inkling": {
16822
16831
  id: "thinkingmachines/inkling",
@@ -16827,13 +16836,13 @@ export const MODELS = {
16827
16836
  reasoning: true,
16828
16837
  input: ["text", "image"],
16829
16838
  cost: {
16830
- input: 0.95,
16839
+ input: 1,
16831
16840
  output: 4.05,
16832
- cacheRead: 0.16,
16841
+ cacheRead: 0.16999999999999998,
16833
16842
  cacheWrite: 0,
16834
16843
  },
16835
16844
  contextWindow: 1048576,
16836
- maxTokens: 262144,
16845
+ maxTokens: 4096,
16837
16846
  },
16838
16847
  "thinkingmachines/inkling-small": {
16839
16848
  id: "thinkingmachines/inkling-small",
@@ -16852,6 +16861,23 @@ export const MODELS = {
16852
16861
  contextWindow: 1048576,
16853
16862
  maxTokens: 262144,
16854
16863
  },
16864
+ "thinkingmachines/inkling-small:free": {
16865
+ id: "thinkingmachines/inkling-small:free",
16866
+ name: "Thinking Machines: Inkling Small (free)",
16867
+ api: "openai-completions",
16868
+ provider: "openrouter",
16869
+ baseUrl: "https://openrouter.ai/api/v1",
16870
+ reasoning: true,
16871
+ input: ["text", "image"],
16872
+ cost: {
16873
+ input: 0,
16874
+ output: 0,
16875
+ cacheRead: 0,
16876
+ cacheWrite: 0,
16877
+ },
16878
+ contextWindow: 262144,
16879
+ maxTokens: 262144,
16880
+ },
16855
16881
  "thinkingmachines/inkling:batch": {
16856
16882
  id: "thinkingmachines/inkling:batch",
16857
16883
  name: "Thinking Machines: Inkling (batch)",
@@ -16869,6 +16895,23 @@ export const MODELS = {
16869
16895
  contextWindow: 524288,
16870
16896
  maxTokens: 4096,
16871
16897
  },
16898
+ "thinkingmachines/inkling:free": {
16899
+ id: "thinkingmachines/inkling:free",
16900
+ name: "Thinking Machines: Inkling (free)",
16901
+ api: "openai-completions",
16902
+ provider: "openrouter",
16903
+ baseUrl: "https://openrouter.ai/api/v1",
16904
+ reasoning: true,
16905
+ input: ["text", "image"],
16906
+ cost: {
16907
+ input: 0,
16908
+ output: 0,
16909
+ cacheRead: 0,
16910
+ cacheWrite: 0,
16911
+ },
16912
+ contextWindow: 262144,
16913
+ maxTokens: 262144,
16914
+ },
16872
16915
  "upstage/solar-pro-3": {
16873
16916
  id: "upstage/solar-pro-3",
16874
16917
  name: "Upstage: Solar Pro 3",
@@ -17361,9 +17404,9 @@ export const MODELS = {
17361
17404
  thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "xhigh" },
17362
17405
  input: ["text"],
17363
17406
  cost: {
17364
- input: 0.065,
17365
- output: 0.18,
17366
- cacheRead: 0.02,
17407
+ input: 0.04,
17408
+ output: 0.08,
17409
+ cacheRead: 0.008,
17367
17410
  cacheWrite: 0,
17368
17411
  },
17369
17412
  contextWindow: 1310720,
@@ -17429,10 +17472,10 @@ export const MODELS = {
17429
17472
  reasoning: true,
17430
17473
  input: ["text", "image"],
17431
17474
  cost: {
17432
- input: 2.5,
17433
- output: 15,
17434
- cacheRead: 0.25,
17435
- cacheWrite: 3.125,
17475
+ input: 2,
17476
+ output: 10,
17477
+ cacheRead: 0.19999999999999998,
17478
+ cacheWrite: 2.5,
17436
17479
  },
17437
17480
  contextWindow: 1050000,
17438
17481
  maxTokens: 128000,
@@ -18893,9 +18936,26 @@ export const MODELS = {
18893
18936
  reasoning: true,
18894
18937
  input: ["text"],
18895
18938
  cost: {
18896
- input: 0.13,
18897
- output: 0.26,
18898
- cacheRead: 0.028,
18939
+ input: 0.07600000000000001,
18940
+ output: 0.153,
18941
+ cacheRead: 0.014,
18942
+ cacheWrite: 0,
18943
+ },
18944
+ contextWindow: 1000000,
18945
+ maxTokens: 384000,
18946
+ },
18947
+ "deepseek/deepseek-v4-flash-vision-exp": {
18948
+ id: "deepseek/deepseek-v4-flash-vision-exp",
18949
+ name: "DeepSeek V4 Flash Vision Exp",
18950
+ api: "anthropic-messages",
18951
+ provider: "vercel-ai-gateway",
18952
+ baseUrl: "https://ai-gateway.vercel.sh",
18953
+ reasoning: true,
18954
+ input: ["text", "image"],
18955
+ cost: {
18956
+ input: 0.22,
18957
+ output: 0.66,
18958
+ cacheRead: 0.007,
18899
18959
  cacheWrite: 0,
18900
18960
  },
18901
18961
  contextWindow: 1000000,
@@ -19998,13 +20058,30 @@ export const MODELS = {
19998
20058
  reasoning: true,
19999
20059
  input: ["text"],
20000
20060
  cost: {
20001
- input: 0.049999999999999996,
20002
- output: 0.19999999999999998,
20003
- cacheRead: 0.01,
20061
+ input: 0,
20062
+ output: 0,
20063
+ cacheRead: 0,
20004
20064
  cacheWrite: 0,
20005
20065
  },
20006
- contextWindow: 262144,
20007
- maxTokens: 131072,
20066
+ contextWindow: 1000000,
20067
+ maxTokens: 32768,
20068
+ },
20069
+ "nvidia/nemotron-3.5-lightning-free": {
20070
+ id: "nvidia/nemotron-3.5-lightning-free",
20071
+ name: "Nemotron 3.5 Lightning 30B (Free)",
20072
+ api: "anthropic-messages",
20073
+ provider: "vercel-ai-gateway",
20074
+ baseUrl: "https://ai-gateway.vercel.sh",
20075
+ reasoning: true,
20076
+ input: ["text"],
20077
+ cost: {
20078
+ input: 0,
20079
+ output: 0,
20080
+ cacheRead: 0,
20081
+ cacheWrite: 0,
20082
+ },
20083
+ contextWindow: 1000000,
20084
+ maxTokens: 32768,
20008
20085
  },
20009
20086
  "nvidia/nemotron-nano-12b-v2-vl": {
20010
20087
  id: "nvidia/nemotron-nano-12b-v2-vl",
@@ -20733,7 +20810,7 @@ export const MODELS = {
20733
20810
  cacheRead: 0.02,
20734
20811
  cacheWrite: 0.25,
20735
20812
  },
20736
- contextWindow: 1050000,
20813
+ contextWindow: 1000000,
20737
20814
  maxTokens: 128000,
20738
20815
  },
20739
20816
  "openai/gpt-5.6-luna-fast": {
@@ -20751,7 +20828,7 @@ export const MODELS = {
20751
20828
  cacheRead: 0.04,
20752
20829
  cacheWrite: 0.25,
20753
20830
  },
20754
- contextWindow: 1050000,
20831
+ contextWindow: 1000000,
20755
20832
  maxTokens: 128000,
20756
20833
  },
20757
20834
  "openai/gpt-5.6-sol": {
@@ -20764,12 +20841,12 @@ export const MODELS = {
20764
20841
  thinkingLevelMap: { "xhigh": "xhigh", "max": "max" },
20765
20842
  input: ["text", "image"],
20766
20843
  cost: {
20767
- input: 2.5,
20768
- output: 15,
20769
- cacheRead: 0.25,
20770
- cacheWrite: 3.125,
20844
+ input: 2,
20845
+ output: 10,
20846
+ cacheRead: 0.19999999999999998,
20847
+ cacheWrite: 2.5,
20771
20848
  },
20772
- contextWindow: 1050000,
20849
+ contextWindow: 1000000,
20773
20850
  maxTokens: 128000,
20774
20851
  },
20775
20852
  "openai/gpt-5.6-sol-fast": {
@@ -20782,12 +20859,12 @@ export const MODELS = {
20782
20859
  thinkingLevelMap: { "xhigh": "xhigh", "max": "max" },
20783
20860
  input: ["text", "image"],
20784
20861
  cost: {
20785
- input: 5,
20786
- output: 30,
20787
- cacheRead: 0.5,
20788
- cacheWrite: 3.125,
20862
+ input: 4,
20863
+ output: 20,
20864
+ cacheRead: 0.39999999999999997,
20865
+ cacheWrite: 2.5,
20789
20866
  },
20790
- contextWindow: 1050000,
20867
+ contextWindow: 1000000,
20791
20868
  maxTokens: 128000,
20792
20869
  },
20793
20870
  "openai/gpt-5.6-terra": {
@@ -20805,7 +20882,7 @@ export const MODELS = {
20805
20882
  cacheRead: 0.19999999999999998,
20806
20883
  cacheWrite: 2.5,
20807
20884
  },
20808
- contextWindow: 1050000,
20885
+ contextWindow: 1000000,
20809
20886
  maxTokens: 128000,
20810
20887
  },
20811
20888
  "openai/gpt-5.6-terra-fast": {
@@ -20823,7 +20900,7 @@ export const MODELS = {
20823
20900
  cacheRead: 0.39999999999999997,
20824
20901
  cacheWrite: 2.5,
20825
20902
  },
20826
- contextWindow: 1050000,
20903
+ contextWindow: 1000000,
20827
20904
  maxTokens: 128000,
20828
20905
  },
20829
20906
  "openai/gpt-oss-120b": {