@caupulican/pi-ai 0.81.38 → 0.81.39
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/image-models.generated.d.ts +60 -0
- package/dist/image-models.generated.d.ts.map +1 -1
- package/dist/image-models.generated.js +60 -0
- package/dist/image-models.generated.js.map +1 -1
- package/dist/models.generated.d.ts +163 -213
- package/dist/models.generated.d.ts.map +1 -1
- package/dist/models.generated.js +257 -279
- package/dist/models.generated.js.map +1 -1
- package/dist/providers/openai-completions.d.ts +1 -1
- package/dist/providers/openai-completions.d.ts.map +1 -1
- package/dist/providers/openai-completions.js +68 -19
- package/dist/providers/openai-completions.js.map +1 -1
- package/dist/stream.d.ts.map +1 -1
- package/dist/stream.js +145 -2
- package/dist/stream.js.map +1 -1
- package/dist/utils/tool-repair/repairer.d.ts +13 -1
- package/dist/utils/tool-repair/repairer.d.ts.map +1 -1
- package/dist/utils/tool-repair/repairer.js +44 -28
- package/dist/utils/tool-repair/repairer.js.map +1 -1
- package/dist/utils/tool-repair/text-protocol-history.d.ts +17 -0
- package/dist/utils/tool-repair/text-protocol-history.d.ts.map +1 -0
- package/dist/utils/tool-repair/text-protocol-history.js +22 -0
- package/dist/utils/tool-repair/text-protocol-history.js.map +1 -0
- package/dist/utils/tool-repair/text-protocol.d.ts +7 -1
- package/dist/utils/tool-repair/text-protocol.d.ts.map +1 -1
- package/dist/utils/tool-repair/text-protocol.js +44 -18
- package/dist/utils/tool-repair/text-protocol.js.map +1 -1
- package/dist/utils/validation.d.ts +9 -0
- package/dist/utils/validation.d.ts.map +1 -1
- package/dist/utils/validation.js +8 -1
- package/dist/utils/validation.js.map +1 -1
- package/package.json +1 -1
package/dist/models.generated.js
CHANGED
|
@@ -4806,7 +4806,7 @@ export const MODELS = {
|
|
|
4806
4806
|
input: 1,
|
|
4807
4807
|
output: 6,
|
|
4808
4808
|
cacheRead: 0.1,
|
|
4809
|
-
cacheWrite:
|
|
4809
|
+
cacheWrite: 1.25,
|
|
4810
4810
|
},
|
|
4811
4811
|
contextWindow: 1050000,
|
|
4812
4812
|
maxTokens: 128000,
|
|
@@ -4825,7 +4825,7 @@ export const MODELS = {
|
|
|
4825
4825
|
input: 5,
|
|
4826
4826
|
output: 30,
|
|
4827
4827
|
cacheRead: 0.5,
|
|
4828
|
-
cacheWrite:
|
|
4828
|
+
cacheWrite: 6.25,
|
|
4829
4829
|
},
|
|
4830
4830
|
contextWindow: 1050000,
|
|
4831
4831
|
maxTokens: 128000,
|
|
@@ -4844,7 +4844,7 @@ export const MODELS = {
|
|
|
4844
4844
|
input: 2.5,
|
|
4845
4845
|
output: 15,
|
|
4846
4846
|
cacheRead: 0.25,
|
|
4847
|
-
cacheWrite:
|
|
4847
|
+
cacheWrite: 3.125,
|
|
4848
4848
|
},
|
|
4849
4849
|
contextWindow: 1050000,
|
|
4850
4850
|
maxTokens: 128000,
|
|
@@ -6405,9 +6405,9 @@ export const MODELS = {
|
|
|
6405
6405
|
},
|
|
6406
6406
|
},
|
|
6407
6407
|
"kimi-coding": {
|
|
6408
|
-
"
|
|
6409
|
-
id: "
|
|
6410
|
-
name: "Kimi
|
|
6408
|
+
"k3": {
|
|
6409
|
+
id: "k3",
|
|
6410
|
+
name: "Kimi K3",
|
|
6411
6411
|
api: "anthropic-messages",
|
|
6412
6412
|
provider: "kimi-coding",
|
|
6413
6413
|
baseUrl: "https://api.kimi.com/coding",
|
|
@@ -6420,12 +6420,12 @@ export const MODELS = {
|
|
|
6420
6420
|
cacheRead: 0,
|
|
6421
6421
|
cacheWrite: 0,
|
|
6422
6422
|
},
|
|
6423
|
-
contextWindow:
|
|
6424
|
-
maxTokens:
|
|
6423
|
+
contextWindow: 1048576,
|
|
6424
|
+
maxTokens: 131072,
|
|
6425
6425
|
},
|
|
6426
6426
|
"kimi-for-coding": {
|
|
6427
6427
|
id: "kimi-for-coding",
|
|
6428
|
-
name: "Kimi
|
|
6428
|
+
name: "Kimi K2.7 Code",
|
|
6429
6429
|
api: "anthropic-messages",
|
|
6430
6430
|
provider: "kimi-coding",
|
|
6431
6431
|
baseUrl: "https://api.kimi.com/coding",
|
|
@@ -6459,24 +6459,6 @@ export const MODELS = {
|
|
|
6459
6459
|
contextWindow: 262144,
|
|
6460
6460
|
maxTokens: 32768,
|
|
6461
6461
|
},
|
|
6462
|
-
"kimi-k2-thinking": {
|
|
6463
|
-
id: "kimi-k2-thinking",
|
|
6464
|
-
name: "Kimi K2 Thinking",
|
|
6465
|
-
api: "anthropic-messages",
|
|
6466
|
-
provider: "kimi-coding",
|
|
6467
|
-
baseUrl: "https://api.kimi.com/coding",
|
|
6468
|
-
headers: { "User-Agent": "KimiCLI/1.5" },
|
|
6469
|
-
reasoning: true,
|
|
6470
|
-
input: ["text"],
|
|
6471
|
-
cost: {
|
|
6472
|
-
input: 0,
|
|
6473
|
-
output: 0,
|
|
6474
|
-
cacheRead: 0,
|
|
6475
|
-
cacheWrite: 0,
|
|
6476
|
-
},
|
|
6477
|
-
contextWindow: 262144,
|
|
6478
|
-
maxTokens: 32768,
|
|
6479
|
-
},
|
|
6480
6462
|
},
|
|
6481
6463
|
"llama-cpp": {
|
|
6482
6464
|
"local": {
|
|
@@ -7245,6 +7227,24 @@ export const MODELS = {
|
|
|
7245
7227
|
contextWindow: 262144,
|
|
7246
7228
|
maxTokens: 262144,
|
|
7247
7229
|
},
|
|
7230
|
+
"kimi-k3": {
|
|
7231
|
+
id: "kimi-k3",
|
|
7232
|
+
name: "Kimi K3",
|
|
7233
|
+
api: "openai-completions",
|
|
7234
|
+
provider: "moonshotai",
|
|
7235
|
+
baseUrl: "https://api.moonshot.ai/v1",
|
|
7236
|
+
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false },
|
|
7237
|
+
reasoning: true,
|
|
7238
|
+
input: ["text", "image"],
|
|
7239
|
+
cost: {
|
|
7240
|
+
input: 3,
|
|
7241
|
+
output: 15,
|
|
7242
|
+
cacheRead: 0.3,
|
|
7243
|
+
cacheWrite: 0,
|
|
7244
|
+
},
|
|
7245
|
+
contextWindow: 1048576,
|
|
7246
|
+
maxTokens: 131072,
|
|
7247
|
+
},
|
|
7248
7248
|
},
|
|
7249
7249
|
"moonshotai-cn": {
|
|
7250
7250
|
"kimi-k2-0711-preview": {
|
|
@@ -7409,6 +7409,24 @@ export const MODELS = {
|
|
|
7409
7409
|
contextWindow: 262144,
|
|
7410
7410
|
maxTokens: 262144,
|
|
7411
7411
|
},
|
|
7412
|
+
"kimi-k3": {
|
|
7413
|
+
id: "kimi-k3",
|
|
7414
|
+
name: "Kimi K3",
|
|
7415
|
+
api: "openai-completions",
|
|
7416
|
+
provider: "moonshotai-cn",
|
|
7417
|
+
baseUrl: "https://api.moonshot.cn/v1",
|
|
7418
|
+
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false },
|
|
7419
|
+
reasoning: true,
|
|
7420
|
+
input: ["text", "image"],
|
|
7421
|
+
cost: {
|
|
7422
|
+
input: 3,
|
|
7423
|
+
output: 15,
|
|
7424
|
+
cacheRead: 0.3,
|
|
7425
|
+
cacheWrite: 0,
|
|
7426
|
+
},
|
|
7427
|
+
contextWindow: 1048576,
|
|
7428
|
+
maxTokens: 131072,
|
|
7429
|
+
},
|
|
7412
7430
|
},
|
|
7413
7431
|
"openai": {
|
|
7414
7432
|
"gpt-4": {
|
|
@@ -9136,7 +9154,7 @@ export const MODELS = {
|
|
|
9136
9154
|
"grok-4.5": {
|
|
9137
9155
|
id: "grok-4.5",
|
|
9138
9156
|
name: "Grok 4.5",
|
|
9139
|
-
api: "openai-
|
|
9157
|
+
api: "openai-responses",
|
|
9140
9158
|
provider: "opencode",
|
|
9141
9159
|
baseUrl: "https://opencode.ai/zen/v1",
|
|
9142
9160
|
reasoning: true,
|
|
@@ -9406,9 +9424,9 @@ export const MODELS = {
|
|
|
9406
9424
|
thinkingLevelMap: { "minimal": null, "low": null, "medium": null, "high": "high", "xhigh": "max" },
|
|
9407
9425
|
input: ["text"],
|
|
9408
9426
|
cost: {
|
|
9409
|
-
input:
|
|
9410
|
-
output:
|
|
9411
|
-
cacheRead: 0.
|
|
9427
|
+
input: 0.435,
|
|
9428
|
+
output: 0.87,
|
|
9429
|
+
cacheRead: 0.003625,
|
|
9412
9430
|
cacheWrite: 0,
|
|
9413
9431
|
},
|
|
9414
9432
|
contextWindow: 1000000,
|
|
@@ -9448,6 +9466,23 @@ export const MODELS = {
|
|
|
9448
9466
|
contextWindow: 1000000,
|
|
9449
9467
|
maxTokens: 131072,
|
|
9450
9468
|
},
|
|
9469
|
+
"grok-4.5": {
|
|
9470
|
+
id: "grok-4.5",
|
|
9471
|
+
name: "Grok 4.5",
|
|
9472
|
+
api: "openai-responses",
|
|
9473
|
+
provider: "opencode-go",
|
|
9474
|
+
baseUrl: "https://opencode.ai/zen/go/v1",
|
|
9475
|
+
reasoning: true,
|
|
9476
|
+
input: ["text", "image"],
|
|
9477
|
+
cost: {
|
|
9478
|
+
input: 2,
|
|
9479
|
+
output: 6,
|
|
9480
|
+
cacheRead: 0.5,
|
|
9481
|
+
cacheWrite: 0,
|
|
9482
|
+
},
|
|
9483
|
+
contextWindow: 500000,
|
|
9484
|
+
maxTokens: 500000,
|
|
9485
|
+
},
|
|
9451
9486
|
"kimi-k2.6": {
|
|
9452
9487
|
id: "kimi-k2.6",
|
|
9453
9488
|
name: "Kimi K2.6",
|
|
@@ -9484,6 +9519,23 @@ export const MODELS = {
|
|
|
9484
9519
|
contextWindow: 262144,
|
|
9485
9520
|
maxTokens: 262144,
|
|
9486
9521
|
},
|
|
9522
|
+
"kimi-k3": {
|
|
9523
|
+
id: "kimi-k3",
|
|
9524
|
+
name: "Kimi K3 (2x usage)",
|
|
9525
|
+
api: "openai-completions",
|
|
9526
|
+
provider: "opencode-go",
|
|
9527
|
+
baseUrl: "https://opencode.ai/zen/go/v1",
|
|
9528
|
+
reasoning: true,
|
|
9529
|
+
input: ["text", "image"],
|
|
9530
|
+
cost: {
|
|
9531
|
+
input: 3,
|
|
9532
|
+
output: 15,
|
|
9533
|
+
cacheRead: 0.3,
|
|
9534
|
+
cacheWrite: 0,
|
|
9535
|
+
},
|
|
9536
|
+
contextWindow: 1048576,
|
|
9537
|
+
maxTokens: 131072,
|
|
9538
|
+
},
|
|
9487
9539
|
"mimo-v2.5": {
|
|
9488
9540
|
id: "mimo-v2.5",
|
|
9489
9541
|
name: "MiMo V2.5",
|
|
@@ -9510,9 +9562,9 @@ export const MODELS = {
|
|
|
9510
9562
|
reasoning: true,
|
|
9511
9563
|
input: ["text"],
|
|
9512
9564
|
cost: {
|
|
9513
|
-
input:
|
|
9514
|
-
output:
|
|
9515
|
-
cacheRead: 0.
|
|
9565
|
+
input: 0.435,
|
|
9566
|
+
output: 0.87,
|
|
9567
|
+
cacheRead: 0.003625,
|
|
9516
9568
|
cacheWrite: 0,
|
|
9517
9569
|
},
|
|
9518
9570
|
contextWindow: 1048576,
|
|
@@ -10338,7 +10390,7 @@ export const MODELS = {
|
|
|
10338
10390
|
cost: {
|
|
10339
10391
|
input: 0.098,
|
|
10340
10392
|
output: 0.196,
|
|
10341
|
-
cacheRead: 0.
|
|
10393
|
+
cacheRead: 0.0196,
|
|
10342
10394
|
cacheWrite: 0,
|
|
10343
10395
|
},
|
|
10344
10396
|
contextWindow: 1048576,
|
|
@@ -10593,13 +10645,13 @@ export const MODELS = {
|
|
|
10593
10645
|
reasoning: false,
|
|
10594
10646
|
input: ["text", "image"],
|
|
10595
10647
|
cost: {
|
|
10596
|
-
input: 0.
|
|
10597
|
-
output: 0.
|
|
10598
|
-
cacheRead: 0
|
|
10648
|
+
input: 0.09999999999999999,
|
|
10649
|
+
output: 0.3,
|
|
10650
|
+
cacheRead: 0,
|
|
10599
10651
|
cacheWrite: 0,
|
|
10600
10652
|
},
|
|
10601
10653
|
contextWindow: 131072,
|
|
10602
|
-
maxTokens:
|
|
10654
|
+
maxTokens: 4096,
|
|
10603
10655
|
},
|
|
10604
10656
|
"google/gemma-4-26b-a4b-it": {
|
|
10605
10657
|
id: "google/gemma-4-26b-a4b-it",
|
|
@@ -10610,13 +10662,13 @@ export const MODELS = {
|
|
|
10610
10662
|
reasoning: true,
|
|
10611
10663
|
input: ["text", "image"],
|
|
10612
10664
|
cost: {
|
|
10613
|
-
input: 0.
|
|
10614
|
-
output: 0.
|
|
10665
|
+
input: 0.07,
|
|
10666
|
+
output: 0.33999999999999997,
|
|
10615
10667
|
cacheRead: 0,
|
|
10616
10668
|
cacheWrite: 0,
|
|
10617
10669
|
},
|
|
10618
10670
|
contextWindow: 262144,
|
|
10619
|
-
maxTokens:
|
|
10671
|
+
maxTokens: 16384,
|
|
10620
10672
|
},
|
|
10621
10673
|
"google/gemma-4-26b-a4b-it:free": {
|
|
10622
10674
|
id: "google/gemma-4-26b-a4b-it:free",
|
|
@@ -10644,13 +10696,13 @@ export const MODELS = {
|
|
|
10644
10696
|
reasoning: true,
|
|
10645
10697
|
input: ["text", "image"],
|
|
10646
10698
|
cost: {
|
|
10647
|
-
input: 0.
|
|
10648
|
-
output: 0.
|
|
10649
|
-
cacheRead: 0
|
|
10699
|
+
input: 0.12,
|
|
10700
|
+
output: 0.37,
|
|
10701
|
+
cacheRead: 0,
|
|
10650
10702
|
cacheWrite: 0,
|
|
10651
10703
|
},
|
|
10652
10704
|
contextWindow: 262144,
|
|
10653
|
-
maxTokens:
|
|
10705
|
+
maxTokens: 16384,
|
|
10654
10706
|
},
|
|
10655
10707
|
"google/gemma-4-31b-it:free": {
|
|
10656
10708
|
id: "google/gemma-4-31b-it:free",
|
|
@@ -10806,6 +10858,23 @@ export const MODELS = {
|
|
|
10806
10858
|
contextWindow: 256000,
|
|
10807
10859
|
maxTokens: 80000,
|
|
10808
10860
|
},
|
|
10861
|
+
"meituan/longcat-2.0": {
|
|
10862
|
+
id: "meituan/longcat-2.0",
|
|
10863
|
+
name: "Meituan: LongCat 2.0",
|
|
10864
|
+
api: "openai-completions",
|
|
10865
|
+
provider: "openrouter",
|
|
10866
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
10867
|
+
reasoning: true,
|
|
10868
|
+
input: ["text"],
|
|
10869
|
+
cost: {
|
|
10870
|
+
input: 0.3,
|
|
10871
|
+
output: 1.2,
|
|
10872
|
+
cacheRead: 0.006,
|
|
10873
|
+
cacheWrite: 0,
|
|
10874
|
+
},
|
|
10875
|
+
contextWindow: 1048756,
|
|
10876
|
+
maxTokens: 262144,
|
|
10877
|
+
},
|
|
10809
10878
|
"meta-llama/llama-3.1-70b-instruct": {
|
|
10810
10879
|
id: "meta-llama/llama-3.1-70b-instruct",
|
|
10811
10880
|
name: "Meta: Llama 3.1 70B Instruct",
|
|
@@ -10857,23 +10926,6 @@ export const MODELS = {
|
|
|
10857
10926
|
contextWindow: 131072,
|
|
10858
10927
|
maxTokens: 128000,
|
|
10859
10928
|
},
|
|
10860
|
-
"meta-llama/llama-3.3-70b-instruct:free": {
|
|
10861
|
-
id: "meta-llama/llama-3.3-70b-instruct:free",
|
|
10862
|
-
name: "Meta: Llama 3.3 70B Instruct (free)",
|
|
10863
|
-
api: "openai-completions",
|
|
10864
|
-
provider: "openrouter",
|
|
10865
|
-
baseUrl: "https://openrouter.ai/api/v1",
|
|
10866
|
-
reasoning: false,
|
|
10867
|
-
input: ["text"],
|
|
10868
|
-
cost: {
|
|
10869
|
-
input: 0,
|
|
10870
|
-
output: 0,
|
|
10871
|
-
cacheRead: 0,
|
|
10872
|
-
cacheWrite: 0,
|
|
10873
|
-
},
|
|
10874
|
-
contextWindow: 131072,
|
|
10875
|
-
maxTokens: 4096,
|
|
10876
|
-
},
|
|
10877
10929
|
"meta-llama/llama-4-maverick": {
|
|
10878
10930
|
id: "meta-llama/llama-4-maverick",
|
|
10879
10931
|
name: "Meta: Llama 4 Maverick",
|
|
@@ -10908,6 +10960,23 @@ export const MODELS = {
|
|
|
10908
10960
|
contextWindow: 10000000,
|
|
10909
10961
|
maxTokens: 16384,
|
|
10910
10962
|
},
|
|
10963
|
+
"meta/muse-spark-1.1": {
|
|
10964
|
+
id: "meta/muse-spark-1.1",
|
|
10965
|
+
name: "Meta: Muse Spark 1.1",
|
|
10966
|
+
api: "openai-completions",
|
|
10967
|
+
provider: "openrouter",
|
|
10968
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
10969
|
+
reasoning: true,
|
|
10970
|
+
input: ["text", "image"],
|
|
10971
|
+
cost: {
|
|
10972
|
+
input: 1.25,
|
|
10973
|
+
output: 4.25,
|
|
10974
|
+
cacheRead: 0.15,
|
|
10975
|
+
cacheWrite: 0,
|
|
10976
|
+
},
|
|
10977
|
+
contextWindow: 1048576,
|
|
10978
|
+
maxTokens: 4096,
|
|
10979
|
+
},
|
|
10911
10980
|
"minimax/minimax-m1": {
|
|
10912
10981
|
id: "minimax/minimax-m1",
|
|
10913
10982
|
name: "MiniMax: MiniMax M1",
|
|
@@ -10985,9 +11054,9 @@ export const MODELS = {
|
|
|
10985
11054
|
reasoning: true,
|
|
10986
11055
|
input: ["text"],
|
|
10987
11056
|
cost: {
|
|
10988
|
-
input: 0.
|
|
10989
|
-
output: 1
|
|
10990
|
-
cacheRead: 0.
|
|
11057
|
+
input: 0.25,
|
|
11058
|
+
output: 1,
|
|
11059
|
+
cacheRead: 0.049999999999999996,
|
|
10991
11060
|
cacheWrite: 0,
|
|
10992
11061
|
},
|
|
10993
11062
|
contextWindow: 204800,
|
|
@@ -11206,8 +11275,8 @@ export const MODELS = {
|
|
|
11206
11275
|
reasoning: false,
|
|
11207
11276
|
input: ["text"],
|
|
11208
11277
|
cost: {
|
|
11209
|
-
input: 0.
|
|
11210
|
-
output: 0.
|
|
11278
|
+
input: 0.019000000000000003,
|
|
11279
|
+
output: 0.03,
|
|
11211
11280
|
cacheRead: 0,
|
|
11212
11281
|
cacheWrite: 0,
|
|
11213
11282
|
},
|
|
@@ -11344,11 +11413,11 @@ export const MODELS = {
|
|
|
11344
11413
|
cost: {
|
|
11345
11414
|
input: 0.6,
|
|
11346
11415
|
output: 2.5,
|
|
11347
|
-
cacheRead: 0,
|
|
11416
|
+
cacheRead: 0.15,
|
|
11348
11417
|
cacheWrite: 0,
|
|
11349
11418
|
},
|
|
11350
11419
|
contextWindow: 262144,
|
|
11351
|
-
maxTokens:
|
|
11420
|
+
maxTokens: 100352,
|
|
11352
11421
|
},
|
|
11353
11422
|
"moonshotai/kimi-k2.5": {
|
|
11354
11423
|
id: "moonshotai/kimi-k2.5",
|
|
@@ -11377,13 +11446,13 @@ export const MODELS = {
|
|
|
11377
11446
|
reasoning: true,
|
|
11378
11447
|
input: ["text", "image"],
|
|
11379
11448
|
cost: {
|
|
11380
|
-
input: 0.
|
|
11381
|
-
output:
|
|
11382
|
-
cacheRead: 0.
|
|
11449
|
+
input: 0.684,
|
|
11450
|
+
output: 3.42,
|
|
11451
|
+
cacheRead: 0.144,
|
|
11383
11452
|
cacheWrite: 0,
|
|
11384
11453
|
},
|
|
11385
11454
|
contextWindow: 262144,
|
|
11386
|
-
maxTokens:
|
|
11455
|
+
maxTokens: 262144,
|
|
11387
11456
|
},
|
|
11388
11457
|
"moonshotai/kimi-k2.7-code": {
|
|
11389
11458
|
id: "moonshotai/kimi-k2.7-code",
|
|
@@ -11394,14 +11463,31 @@ export const MODELS = {
|
|
|
11394
11463
|
reasoning: true,
|
|
11395
11464
|
input: ["text", "image"],
|
|
11396
11465
|
cost: {
|
|
11397
|
-
input: 0.
|
|
11398
|
-
output: 3.
|
|
11399
|
-
cacheRead: 0.
|
|
11466
|
+
input: 0.82,
|
|
11467
|
+
output: 3.75,
|
|
11468
|
+
cacheRead: 0.16,
|
|
11400
11469
|
cacheWrite: 0,
|
|
11401
11470
|
},
|
|
11402
11471
|
contextWindow: 262144,
|
|
11403
11472
|
maxTokens: 262144,
|
|
11404
11473
|
},
|
|
11474
|
+
"moonshotai/kimi-k3": {
|
|
11475
|
+
id: "moonshotai/kimi-k3",
|
|
11476
|
+
name: "MoonshotAI: Kimi K3",
|
|
11477
|
+
api: "openai-completions",
|
|
11478
|
+
provider: "openrouter",
|
|
11479
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
11480
|
+
reasoning: true,
|
|
11481
|
+
input: ["text", "image"],
|
|
11482
|
+
cost: {
|
|
11483
|
+
input: 3,
|
|
11484
|
+
output: 15,
|
|
11485
|
+
cacheRead: 0.3,
|
|
11486
|
+
cacheWrite: 0,
|
|
11487
|
+
},
|
|
11488
|
+
contextWindow: 1048576,
|
|
11489
|
+
maxTokens: 4096,
|
|
11490
|
+
},
|
|
11405
11491
|
"nex-agi/nex-n2-mini": {
|
|
11406
11492
|
id: "nex-agi/nex-n2-mini",
|
|
11407
11493
|
name: "Nex AGI: Nex-N2-Mini",
|
|
@@ -11436,23 +11522,6 @@ export const MODELS = {
|
|
|
11436
11522
|
contextWindow: 262144,
|
|
11437
11523
|
maxTokens: 262144,
|
|
11438
11524
|
},
|
|
11439
|
-
"nvidia/llama-3.3-nemotron-super-49b-v1.5": {
|
|
11440
|
-
id: "nvidia/llama-3.3-nemotron-super-49b-v1.5",
|
|
11441
|
-
name: "NVIDIA: Llama 3.3 Nemotron Super 49B V1.5",
|
|
11442
|
-
api: "openai-completions",
|
|
11443
|
-
provider: "openrouter",
|
|
11444
|
-
baseUrl: "https://openrouter.ai/api/v1",
|
|
11445
|
-
reasoning: true,
|
|
11446
|
-
input: ["text"],
|
|
11447
|
-
cost: {
|
|
11448
|
-
input: 0.39999999999999997,
|
|
11449
|
-
output: 0.39999999999999997,
|
|
11450
|
-
cacheRead: 0,
|
|
11451
|
-
cacheWrite: 0,
|
|
11452
|
-
},
|
|
11453
|
-
contextWindow: 131072,
|
|
11454
|
-
maxTokens: 16384,
|
|
11455
|
-
},
|
|
11456
11525
|
"nvidia/nemotron-3-nano-30b-a3b": {
|
|
11457
11526
|
id: "nvidia/nemotron-3-nano-30b-a3b",
|
|
11458
11527
|
name: "NVIDIA: Nemotron 3 Nano 30B A3B",
|
|
@@ -11513,13 +11582,13 @@ export const MODELS = {
|
|
|
11513
11582
|
reasoning: true,
|
|
11514
11583
|
input: ["text"],
|
|
11515
11584
|
cost: {
|
|
11516
|
-
input: 0.
|
|
11517
|
-
output: 0.
|
|
11518
|
-
cacheRead: 0
|
|
11585
|
+
input: 0.08499999999999999,
|
|
11586
|
+
output: 0.39999999999999997,
|
|
11587
|
+
cacheRead: 0,
|
|
11519
11588
|
cacheWrite: 0,
|
|
11520
11589
|
},
|
|
11521
11590
|
contextWindow: 1000000,
|
|
11522
|
-
maxTokens:
|
|
11591
|
+
maxTokens: 16384,
|
|
11523
11592
|
},
|
|
11524
11593
|
"nvidia/nemotron-3-super-120b-a12b:free": {
|
|
11525
11594
|
id: "nvidia/nemotron-3-super-120b-a12b:free",
|
|
@@ -12644,6 +12713,23 @@ export const MODELS = {
|
|
|
12644
12713
|
contextWindow: 2000000,
|
|
12645
12714
|
maxTokens: 4096,
|
|
12646
12715
|
},
|
|
12716
|
+
"openrouter/auto-beta": {
|
|
12717
|
+
id: "openrouter/auto-beta",
|
|
12718
|
+
name: "Auto Router (Beta)",
|
|
12719
|
+
api: "openai-completions",
|
|
12720
|
+
provider: "openrouter",
|
|
12721
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
12722
|
+
reasoning: true,
|
|
12723
|
+
input: ["text", "image"],
|
|
12724
|
+
cost: {
|
|
12725
|
+
input: -1000000,
|
|
12726
|
+
output: -1000000,
|
|
12727
|
+
cacheRead: 0,
|
|
12728
|
+
cacheWrite: 0,
|
|
12729
|
+
},
|
|
12730
|
+
contextWindow: 2000000,
|
|
12731
|
+
maxTokens: 4096,
|
|
12732
|
+
},
|
|
12647
12733
|
"openrouter/free": {
|
|
12648
12734
|
id: "openrouter/free",
|
|
12649
12735
|
name: "Free Models Router",
|
|
@@ -12840,13 +12926,13 @@ export const MODELS = {
|
|
|
12840
12926
|
reasoning: true,
|
|
12841
12927
|
input: ["text"],
|
|
12842
12928
|
cost: {
|
|
12843
|
-
input: 0.
|
|
12929
|
+
input: 0.12,
|
|
12844
12930
|
output: 0.24,
|
|
12845
12931
|
cacheRead: 0,
|
|
12846
12932
|
cacheWrite: 0,
|
|
12847
12933
|
},
|
|
12848
12934
|
contextWindow: 131702,
|
|
12849
|
-
maxTokens:
|
|
12935
|
+
maxTokens: 16384,
|
|
12850
12936
|
},
|
|
12851
12937
|
"qwen/qwen3-235b-a22b": {
|
|
12852
12938
|
id: "qwen/qwen3-235b-a22b",
|
|
@@ -12891,13 +12977,13 @@ export const MODELS = {
|
|
|
12891
12977
|
reasoning: true,
|
|
12892
12978
|
input: ["text"],
|
|
12893
12979
|
cost: {
|
|
12894
|
-
input: 0.
|
|
12895
|
-
output:
|
|
12980
|
+
input: 0.3,
|
|
12981
|
+
output: 3,
|
|
12896
12982
|
cacheRead: 0,
|
|
12897
12983
|
cacheWrite: 0,
|
|
12898
12984
|
},
|
|
12899
12985
|
contextWindow: 262144,
|
|
12900
|
-
maxTokens:
|
|
12986
|
+
maxTokens: 32768,
|
|
12901
12987
|
},
|
|
12902
12988
|
"qwen/qwen3-30b-a3b": {
|
|
12903
12989
|
id: "qwen/qwen3-30b-a3b",
|
|
@@ -12908,13 +12994,13 @@ export const MODELS = {
|
|
|
12908
12994
|
reasoning: true,
|
|
12909
12995
|
input: ["text"],
|
|
12910
12996
|
cost: {
|
|
12911
|
-
input: 0.
|
|
12912
|
-
output: 0.
|
|
12997
|
+
input: 0.13,
|
|
12998
|
+
output: 0.52,
|
|
12913
12999
|
cacheRead: 0,
|
|
12914
13000
|
cacheWrite: 0,
|
|
12915
13001
|
},
|
|
12916
13002
|
contextWindow: 131072,
|
|
12917
|
-
maxTokens:
|
|
13003
|
+
maxTokens: 8192,
|
|
12918
13004
|
},
|
|
12919
13005
|
"qwen/qwen3-30b-a3b-instruct-2507": {
|
|
12920
13006
|
id: "qwen/qwen3-30b-a3b-instruct-2507",
|
|
@@ -13069,23 +13155,6 @@ export const MODELS = {
|
|
|
13069
13155
|
contextWindow: 1000000,
|
|
13070
13156
|
maxTokens: 65536,
|
|
13071
13157
|
},
|
|
13072
|
-
"qwen/qwen3-coder:free": {
|
|
13073
|
-
id: "qwen/qwen3-coder:free",
|
|
13074
|
-
name: "Qwen: Qwen3 Coder 480B A35B (free)",
|
|
13075
|
-
api: "openai-completions",
|
|
13076
|
-
provider: "openrouter",
|
|
13077
|
-
baseUrl: "https://openrouter.ai/api/v1",
|
|
13078
|
-
reasoning: false,
|
|
13079
|
-
input: ["text"],
|
|
13080
|
-
cost: {
|
|
13081
|
-
input: 0,
|
|
13082
|
-
output: 0,
|
|
13083
|
-
cacheRead: 0,
|
|
13084
|
-
cacheWrite: 0,
|
|
13085
|
-
},
|
|
13086
|
-
contextWindow: 1048576,
|
|
13087
|
-
maxTokens: 262000,
|
|
13088
|
-
},
|
|
13089
13158
|
"qwen/qwen3-max": {
|
|
13090
13159
|
id: "qwen/qwen3-max",
|
|
13091
13160
|
name: "Qwen: Qwen3 Max",
|
|
@@ -13137,23 +13206,6 @@ export const MODELS = {
|
|
|
13137
13206
|
contextWindow: 262144,
|
|
13138
13207
|
maxTokens: 262144,
|
|
13139
13208
|
},
|
|
13140
|
-
"qwen/qwen3-next-80b-a3b-instruct:free": {
|
|
13141
|
-
id: "qwen/qwen3-next-80b-a3b-instruct:free",
|
|
13142
|
-
name: "Qwen: Qwen3 Next 80B A3B Instruct (free)",
|
|
13143
|
-
api: "openai-completions",
|
|
13144
|
-
provider: "openrouter",
|
|
13145
|
-
baseUrl: "https://openrouter.ai/api/v1",
|
|
13146
|
-
reasoning: false,
|
|
13147
|
-
input: ["text"],
|
|
13148
|
-
cost: {
|
|
13149
|
-
input: 0,
|
|
13150
|
-
output: 0,
|
|
13151
|
-
cacheRead: 0,
|
|
13152
|
-
cacheWrite: 0,
|
|
13153
|
-
},
|
|
13154
|
-
contextWindow: 262144,
|
|
13155
|
-
maxTokens: 4096,
|
|
13156
|
-
},
|
|
13157
13209
|
"qwen/qwen3-next-80b-a3b-thinking": {
|
|
13158
13210
|
id: "qwen/qwen3-next-80b-a3b-thinking",
|
|
13159
13211
|
name: "Qwen: Qwen3 Next 80B A3B Thinking",
|
|
@@ -13316,13 +13368,13 @@ export const MODELS = {
|
|
|
13316
13368
|
reasoning: true,
|
|
13317
13369
|
input: ["text", "image"],
|
|
13318
13370
|
cost: {
|
|
13319
|
-
input: 0.
|
|
13320
|
-
output:
|
|
13371
|
+
input: 0.26,
|
|
13372
|
+
output: 2.6,
|
|
13321
13373
|
cacheRead: 0,
|
|
13322
13374
|
cacheWrite: 0,
|
|
13323
13375
|
},
|
|
13324
13376
|
contextWindow: 262144,
|
|
13325
|
-
maxTokens:
|
|
13377
|
+
maxTokens: 81920,
|
|
13326
13378
|
},
|
|
13327
13379
|
"qwen/qwen3.5-35b-a3b": {
|
|
13328
13380
|
id: "qwen/qwen3.5-35b-a3b",
|
|
@@ -13350,9 +13402,9 @@ export const MODELS = {
|
|
|
13350
13402
|
reasoning: true,
|
|
13351
13403
|
input: ["text", "image"],
|
|
13352
13404
|
cost: {
|
|
13353
|
-
input: 0.
|
|
13354
|
-
output:
|
|
13355
|
-
cacheRead: 0
|
|
13405
|
+
input: 0.39,
|
|
13406
|
+
output: 2.34,
|
|
13407
|
+
cacheRead: 0,
|
|
13356
13408
|
cacheWrite: 0,
|
|
13357
13409
|
},
|
|
13358
13410
|
contextWindow: 262144,
|
|
@@ -13715,6 +13767,23 @@ export const MODELS = {
|
|
|
13715
13767
|
contextWindow: 32768,
|
|
13716
13768
|
maxTokens: 32768,
|
|
13717
13769
|
},
|
|
13770
|
+
"thinkingmachines/inkling": {
|
|
13771
|
+
id: "thinkingmachines/inkling",
|
|
13772
|
+
name: "Thinking Machines: Inkling",
|
|
13773
|
+
api: "openai-completions",
|
|
13774
|
+
provider: "openrouter",
|
|
13775
|
+
baseUrl: "https://openrouter.ai/api/v1",
|
|
13776
|
+
reasoning: true,
|
|
13777
|
+
input: ["text", "image"],
|
|
13778
|
+
cost: {
|
|
13779
|
+
input: 1,
|
|
13780
|
+
output: 4.05,
|
|
13781
|
+
cacheRead: 0.16999999999999998,
|
|
13782
|
+
cacheWrite: 0,
|
|
13783
|
+
},
|
|
13784
|
+
contextWindow: 1048576,
|
|
13785
|
+
maxTokens: 4096,
|
|
13786
|
+
},
|
|
13718
13787
|
"upstage/solar-pro-3": {
|
|
13719
13788
|
id: "upstage/solar-pro-3",
|
|
13720
13789
|
name: "Upstage: Solar Pro 3",
|
|
@@ -13777,7 +13846,7 @@ export const MODELS = {
|
|
|
13777
13846
|
cost: {
|
|
13778
13847
|
input: 2,
|
|
13779
13848
|
output: 6,
|
|
13780
|
-
cacheRead: 0.
|
|
13849
|
+
cacheRead: 0.3,
|
|
13781
13850
|
cacheWrite: 0,
|
|
13782
13851
|
},
|
|
13783
13852
|
contextWindow: 500000,
|
|
@@ -13945,13 +14014,13 @@ export const MODELS = {
|
|
|
13945
14014
|
reasoning: true,
|
|
13946
14015
|
input: ["text"],
|
|
13947
14016
|
cost: {
|
|
13948
|
-
input: 0.
|
|
14017
|
+
input: 0.060500000000000005,
|
|
13949
14018
|
output: 0.39999999999999997,
|
|
13950
|
-
cacheRead: 0
|
|
14019
|
+
cacheRead: 0,
|
|
13951
14020
|
cacheWrite: 0,
|
|
13952
14021
|
},
|
|
13953
|
-
contextWindow:
|
|
13954
|
-
maxTokens:
|
|
14022
|
+
contextWindow: 200000,
|
|
14023
|
+
maxTokens: 131072,
|
|
13955
14024
|
},
|
|
13956
14025
|
"z-ai/glm-5": {
|
|
13957
14026
|
id: "z-ai/glm-5",
|
|
@@ -13967,8 +14036,8 @@ export const MODELS = {
|
|
|
13967
14036
|
cacheRead: 0.119,
|
|
13968
14037
|
cacheWrite: 0,
|
|
13969
14038
|
},
|
|
13970
|
-
contextWindow:
|
|
13971
|
-
maxTokens:
|
|
14039
|
+
contextWindow: 204800,
|
|
14040
|
+
maxTokens: 131072,
|
|
13972
14041
|
},
|
|
13973
14042
|
"z-ai/glm-5-turbo": {
|
|
13974
14043
|
id: "z-ai/glm-5-turbo",
|
|
@@ -14013,9 +14082,9 @@ export const MODELS = {
|
|
|
14013
14082
|
reasoning: true,
|
|
14014
14083
|
input: ["text"],
|
|
14015
14084
|
cost: {
|
|
14016
|
-
input: 0.
|
|
14017
|
-
output:
|
|
14018
|
-
cacheRead: 0.
|
|
14085
|
+
input: 0.9786,
|
|
14086
|
+
output: 3.0755999999999997,
|
|
14087
|
+
cacheRead: 0.18174,
|
|
14019
14088
|
cacheWrite: 0,
|
|
14020
14089
|
},
|
|
14021
14090
|
contextWindow: 1048576,
|
|
@@ -14149,13 +14218,13 @@ export const MODELS = {
|
|
|
14149
14218
|
reasoning: true,
|
|
14150
14219
|
input: ["text", "image"],
|
|
14151
14220
|
cost: {
|
|
14152
|
-
input:
|
|
14153
|
-
output:
|
|
14154
|
-
cacheRead: 0.
|
|
14221
|
+
input: 3,
|
|
14222
|
+
output: 15,
|
|
14223
|
+
cacheRead: 0.3,
|
|
14155
14224
|
cacheWrite: 0,
|
|
14156
14225
|
},
|
|
14157
|
-
contextWindow:
|
|
14158
|
-
maxTokens:
|
|
14226
|
+
contextWindow: 1048576,
|
|
14227
|
+
maxTokens: 4096,
|
|
14159
14228
|
},
|
|
14160
14229
|
"~openai/gpt-latest": {
|
|
14161
14230
|
id: "~openai/gpt-latest",
|
|
@@ -14202,7 +14271,7 @@ export const MODELS = {
|
|
|
14202
14271
|
cost: {
|
|
14203
14272
|
input: 2,
|
|
14204
14273
|
output: 6,
|
|
14205
|
-
cacheRead: 0.
|
|
14274
|
+
cacheRead: 0.3,
|
|
14206
14275
|
cacheWrite: 0,
|
|
14207
14276
|
},
|
|
14208
14277
|
contextWindow: 500000,
|
|
@@ -14266,43 +14335,6 @@ export const MODELS = {
|
|
|
14266
14335
|
contextWindow: 32768,
|
|
14267
14336
|
maxTokens: 32768,
|
|
14268
14337
|
},
|
|
14269
|
-
"Qwen/Qwen3-235B-A22B-Instruct-2507-tput": {
|
|
14270
|
-
id: "Qwen/Qwen3-235B-A22B-Instruct-2507-tput",
|
|
14271
|
-
name: "Qwen3 235B A22B Instruct 2507 FP8",
|
|
14272
|
-
api: "openai-completions",
|
|
14273
|
-
provider: "together",
|
|
14274
|
-
baseUrl: "https://api.together.ai/v1",
|
|
14275
|
-
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false },
|
|
14276
|
-
reasoning: false,
|
|
14277
|
-
input: ["text"],
|
|
14278
|
-
cost: {
|
|
14279
|
-
input: 0.2,
|
|
14280
|
-
output: 0.6,
|
|
14281
|
-
cacheRead: 0,
|
|
14282
|
-
cacheWrite: 0,
|
|
14283
|
-
},
|
|
14284
|
-
contextWindow: 262144,
|
|
14285
|
-
maxTokens: 262144,
|
|
14286
|
-
},
|
|
14287
|
-
"Qwen/Qwen3.5-397B-A17B": {
|
|
14288
|
-
id: "Qwen/Qwen3.5-397B-A17B",
|
|
14289
|
-
name: "Qwen3.5 397B A17B",
|
|
14290
|
-
api: "openai-completions",
|
|
14291
|
-
provider: "together",
|
|
14292
|
-
baseUrl: "https://api.together.ai/v1",
|
|
14293
|
-
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "thinkingFormat": "together" },
|
|
14294
|
-
reasoning: true,
|
|
14295
|
-
thinkingLevelMap: { "minimal": null, "low": null, "medium": null },
|
|
14296
|
-
input: ["text", "image"],
|
|
14297
|
-
cost: {
|
|
14298
|
-
input: 0.6,
|
|
14299
|
-
output: 3.6,
|
|
14300
|
-
cacheRead: 0,
|
|
14301
|
-
cacheWrite: 0,
|
|
14302
|
-
},
|
|
14303
|
-
contextWindow: 262144,
|
|
14304
|
-
maxTokens: 130000,
|
|
14305
|
-
},
|
|
14306
14338
|
"Qwen/Qwen3.5-9B": {
|
|
14307
14339
|
id: "Qwen/Qwen3.5-9B",
|
|
14308
14340
|
name: "Qwen3.5 9B",
|
|
@@ -14378,24 +14410,6 @@ export const MODELS = {
|
|
|
14378
14410
|
contextWindow: 512000,
|
|
14379
14411
|
maxTokens: 384000,
|
|
14380
14412
|
},
|
|
14381
|
-
"essentialai/Rnj-1-Instruct": {
|
|
14382
|
-
id: "essentialai/Rnj-1-Instruct",
|
|
14383
|
-
name: "Rnj-1 Instruct",
|
|
14384
|
-
api: "openai-completions",
|
|
14385
|
-
provider: "together",
|
|
14386
|
-
baseUrl: "https://api.together.ai/v1",
|
|
14387
|
-
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false },
|
|
14388
|
-
reasoning: false,
|
|
14389
|
-
input: ["text"],
|
|
14390
|
-
cost: {
|
|
14391
|
-
input: 0.15,
|
|
14392
|
-
output: 0.15,
|
|
14393
|
-
cacheRead: 0,
|
|
14394
|
-
cacheWrite: 0,
|
|
14395
|
-
},
|
|
14396
|
-
contextWindow: 32768,
|
|
14397
|
-
maxTokens: 32768,
|
|
14398
|
-
},
|
|
14399
14413
|
"google/gemma-4-31B-it": {
|
|
14400
14414
|
id: "google/gemma-4-31B-it",
|
|
14401
14415
|
name: "Gemma 4 31B Instruct",
|
|
@@ -14528,42 +14542,23 @@ export const MODELS = {
|
|
|
14528
14542
|
contextWindow: 131072,
|
|
14529
14543
|
maxTokens: 131072,
|
|
14530
14544
|
},
|
|
14531
|
-
"
|
|
14532
|
-
id: "
|
|
14533
|
-
name: "
|
|
14545
|
+
"thinkingmachines/Inkling": {
|
|
14546
|
+
id: "thinkingmachines/Inkling",
|
|
14547
|
+
name: "Inkling",
|
|
14534
14548
|
api: "openai-completions",
|
|
14535
14549
|
provider: "together",
|
|
14536
14550
|
baseUrl: "https://api.together.ai/v1",
|
|
14537
14551
|
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "thinkingFormat": "together" },
|
|
14538
14552
|
reasoning: true,
|
|
14539
14553
|
thinkingLevelMap: { "minimal": null, "low": null, "medium": null },
|
|
14540
|
-
input: ["text"],
|
|
14554
|
+
input: ["text", "image"],
|
|
14541
14555
|
cost: {
|
|
14542
14556
|
input: 1,
|
|
14543
|
-
output:
|
|
14544
|
-
cacheRead: 0,
|
|
14545
|
-
cacheWrite: 0,
|
|
14546
|
-
},
|
|
14547
|
-
contextWindow: 202752,
|
|
14548
|
-
maxTokens: 131072,
|
|
14549
|
-
},
|
|
14550
|
-
"zai-org/GLM-5.1": {
|
|
14551
|
-
id: "zai-org/GLM-5.1",
|
|
14552
|
-
name: "GLM-5.1",
|
|
14553
|
-
api: "openai-completions",
|
|
14554
|
-
provider: "together",
|
|
14555
|
-
baseUrl: "https://api.together.ai/v1",
|
|
14556
|
-
compat: { "supportsStore": false, "supportsDeveloperRole": false, "supportsReasoningEffort": false, "maxTokensField": "max_tokens", "supportsStrictMode": false, "supportsLongCacheRetention": false, "thinkingFormat": "together" },
|
|
14557
|
-
reasoning: true,
|
|
14558
|
-
thinkingLevelMap: { "minimal": null, "low": null, "medium": null },
|
|
14559
|
-
input: ["text"],
|
|
14560
|
-
cost: {
|
|
14561
|
-
input: 1.4,
|
|
14562
|
-
output: 4.4,
|
|
14563
|
-
cacheRead: 0,
|
|
14557
|
+
output: 4.05,
|
|
14558
|
+
cacheRead: 0.17,
|
|
14564
14559
|
cacheWrite: 0,
|
|
14565
14560
|
},
|
|
14566
|
-
contextWindow:
|
|
14561
|
+
contextWindow: 524288,
|
|
14567
14562
|
maxTokens: 131072,
|
|
14568
14563
|
},
|
|
14569
14564
|
"zai-org/GLM-5.2": {
|
|
@@ -15890,40 +15885,6 @@ export const MODELS = {
|
|
|
15890
15885
|
contextWindow: 128000,
|
|
15891
15886
|
maxTokens: 8192,
|
|
15892
15887
|
},
|
|
15893
|
-
"meta/llama-3.2-11b": {
|
|
15894
|
-
id: "meta/llama-3.2-11b",
|
|
15895
|
-
name: "Llama 3.2 11B Vision Instruct",
|
|
15896
|
-
api: "anthropic-messages",
|
|
15897
|
-
provider: "vercel-ai-gateway",
|
|
15898
|
-
baseUrl: "https://ai-gateway.vercel.sh",
|
|
15899
|
-
reasoning: false,
|
|
15900
|
-
input: ["text", "image"],
|
|
15901
|
-
cost: {
|
|
15902
|
-
input: 0.16,
|
|
15903
|
-
output: 0.16,
|
|
15904
|
-
cacheRead: 0,
|
|
15905
|
-
cacheWrite: 0,
|
|
15906
|
-
},
|
|
15907
|
-
contextWindow: 128000,
|
|
15908
|
-
maxTokens: 8192,
|
|
15909
|
-
},
|
|
15910
|
-
"meta/llama-3.2-90b": {
|
|
15911
|
-
id: "meta/llama-3.2-90b",
|
|
15912
|
-
name: "Llama 3.2 90B Vision Instruct",
|
|
15913
|
-
api: "anthropic-messages",
|
|
15914
|
-
provider: "vercel-ai-gateway",
|
|
15915
|
-
baseUrl: "https://ai-gateway.vercel.sh",
|
|
15916
|
-
reasoning: false,
|
|
15917
|
-
input: ["text", "image"],
|
|
15918
|
-
cost: {
|
|
15919
|
-
input: 0.72,
|
|
15920
|
-
output: 0.72,
|
|
15921
|
-
cacheRead: 0,
|
|
15922
|
-
cacheWrite: 0,
|
|
15923
|
-
},
|
|
15924
|
-
contextWindow: 128000,
|
|
15925
|
-
maxTokens: 8192,
|
|
15926
|
-
},
|
|
15927
15888
|
"meta/llama-3.3-70b": {
|
|
15928
15889
|
id: "meta/llama-3.3-70b",
|
|
15929
15890
|
name: "Llama 3.3 70B Instruct",
|
|
@@ -16468,6 +16429,23 @@ export const MODELS = {
|
|
|
16468
16429
|
contextWindow: 262144,
|
|
16469
16430
|
maxTokens: 32768,
|
|
16470
16431
|
},
|
|
16432
|
+
"moonshotai/kimi-k3": {
|
|
16433
|
+
id: "moonshotai/kimi-k3",
|
|
16434
|
+
name: "Kimi K3",
|
|
16435
|
+
api: "anthropic-messages",
|
|
16436
|
+
provider: "vercel-ai-gateway",
|
|
16437
|
+
baseUrl: "https://ai-gateway.vercel.sh",
|
|
16438
|
+
reasoning: true,
|
|
16439
|
+
input: ["text", "image"],
|
|
16440
|
+
cost: {
|
|
16441
|
+
input: 3,
|
|
16442
|
+
output: 15,
|
|
16443
|
+
cacheRead: 0.3,
|
|
16444
|
+
cacheWrite: 0,
|
|
16445
|
+
},
|
|
16446
|
+
contextWindow: 1000000,
|
|
16447
|
+
maxTokens: 131072,
|
|
16448
|
+
},
|
|
16471
16449
|
"nvidia/nemotron-3-nano-30b-a3b": {
|
|
16472
16450
|
id: "nvidia/nemotron-3-nano-30b-a3b",
|
|
16473
16451
|
name: "Nemotron 3 Nano 30B A3B",
|
|
@@ -17514,7 +17492,7 @@ export const MODELS = {
|
|
|
17514
17492
|
cost: {
|
|
17515
17493
|
input: 2,
|
|
17516
17494
|
output: 6,
|
|
17517
|
-
cacheRead: 0.
|
|
17495
|
+
cacheRead: 0.3,
|
|
17518
17496
|
cacheWrite: 0,
|
|
17519
17497
|
},
|
|
17520
17498
|
contextWindow: 500000,
|
|
@@ -17924,7 +17902,7 @@ export const MODELS = {
|
|
|
17924
17902
|
cost: {
|
|
17925
17903
|
input: 2,
|
|
17926
17904
|
output: 6,
|
|
17927
|
-
cacheRead: 0.
|
|
17905
|
+
cacheRead: 0.3,
|
|
17928
17906
|
cacheWrite: 0,
|
|
17929
17907
|
},
|
|
17930
17908
|
contextWindow: 500000,
|