smoltalk 0.14.1 → 0.14.2
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/models.d.ts +57 -13
- package/dist/models.js +66 -15
- package/package.json +1 -1
package/dist/models.d.ts
CHANGED
|
@@ -176,6 +176,13 @@ export declare const textToSpeechModels: readonly [{
|
|
|
176
176
|
readonly perCharacterCost: 0.00004;
|
|
177
177
|
readonly maxInputChars: 200;
|
|
178
178
|
readonly formats: readonly ["wav"];
|
|
179
|
+
}, {
|
|
180
|
+
readonly type: "text-to-speech";
|
|
181
|
+
readonly modelName: "gemini-3.1-flash-tts-preview";
|
|
182
|
+
readonly provider: "google";
|
|
183
|
+
readonly inputTokenCost: 1;
|
|
184
|
+
readonly outputAudioTokenCost: 20;
|
|
185
|
+
readonly formats: readonly ["pcm", "wav"];
|
|
179
186
|
}, {
|
|
180
187
|
readonly type: "text-to-speech";
|
|
181
188
|
readonly modelName: "gemini-2.5-flash-preview-tts";
|
|
@@ -688,7 +695,7 @@ export declare const textModels: readonly [{
|
|
|
688
695
|
readonly inputTokenCost: 5;
|
|
689
696
|
readonly cachedInputTokenCost: 0.5;
|
|
690
697
|
readonly outputTokenCost: 22.5;
|
|
691
|
-
readonly thresholdTokens:
|
|
698
|
+
readonly thresholdTokens: 272000;
|
|
692
699
|
};
|
|
693
700
|
readonly reasoning: {
|
|
694
701
|
readonly levels: readonly ["none", "low", "medium", "high", "xhigh"];
|
|
@@ -778,7 +785,7 @@ export declare const textModels: readonly [{
|
|
|
778
785
|
readonly longContext: {
|
|
779
786
|
readonly inputTokenCost: 60;
|
|
780
787
|
readonly outputTokenCost: 270;
|
|
781
|
-
readonly thresholdTokens:
|
|
788
|
+
readonly thresholdTokens: 272000;
|
|
782
789
|
};
|
|
783
790
|
readonly reasoning: {
|
|
784
791
|
readonly levels: readonly ["medium", "high", "xhigh"];
|
|
@@ -813,7 +820,7 @@ export declare const textModels: readonly [{
|
|
|
813
820
|
readonly inputTokenCost: 10;
|
|
814
821
|
readonly cachedInputTokenCost: 1;
|
|
815
822
|
readonly outputTokenCost: 45;
|
|
816
|
-
readonly thresholdTokens:
|
|
823
|
+
readonly thresholdTokens: 272000;
|
|
817
824
|
};
|
|
818
825
|
readonly reasoning: {
|
|
819
826
|
readonly levels: readonly ["none", "low", "medium", "high", "xhigh"];
|
|
@@ -845,7 +852,7 @@ export declare const textModels: readonly [{
|
|
|
845
852
|
readonly longContext: {
|
|
846
853
|
readonly inputTokenCost: 60;
|
|
847
854
|
readonly outputTokenCost: 270;
|
|
848
|
-
readonly thresholdTokens:
|
|
855
|
+
readonly thresholdTokens: 272000;
|
|
849
856
|
};
|
|
850
857
|
readonly reasoning: {
|
|
851
858
|
readonly levels: readonly ["none", "low", "medium", "high", "xhigh"];
|
|
@@ -866,6 +873,41 @@ export declare const textModels: readonly [{
|
|
|
866
873
|
readonly structuredOutput: true;
|
|
867
874
|
readonly temperatureSupported: false;
|
|
868
875
|
readonly provider: "openai-responses";
|
|
876
|
+
}, {
|
|
877
|
+
readonly type: "text";
|
|
878
|
+
readonly modelName: "gpt-6-astra";
|
|
879
|
+
readonly description: "GPT-6 Astra is OpenAI's most capable model, built for the hardest end-to-end work: complex reasoning, coding, computer use, research, and document creation. 1M context window. Standard pricing for ≤272K input tokens; prompts above that are billed at 2x input/cache and 1.5x output for the whole request. Knowledge cutoff: April 2026.";
|
|
880
|
+
readonly maxInputTokens: 1050000;
|
|
881
|
+
readonly maxOutputTokens: 128000;
|
|
882
|
+
readonly inputTokenCost: 10;
|
|
883
|
+
readonly cachedInputTokenCost: 1;
|
|
884
|
+
readonly outputTokenCost: 50;
|
|
885
|
+
readonly outputTokensPerSecond: 69;
|
|
886
|
+
readonly longContext: {
|
|
887
|
+
readonly inputTokenCost: 20;
|
|
888
|
+
readonly cachedInputTokenCost: 2;
|
|
889
|
+
readonly outputTokenCost: 75;
|
|
890
|
+
readonly thresholdTokens: 272000;
|
|
891
|
+
};
|
|
892
|
+
readonly reasoning: {
|
|
893
|
+
readonly levels: readonly ["low", "medium", "high", "xhigh", "max"];
|
|
894
|
+
readonly defaultLevel: "medium";
|
|
895
|
+
readonly canDisable: false;
|
|
896
|
+
readonly outputsThinking: false;
|
|
897
|
+
readonly outputsSignatures: false;
|
|
898
|
+
};
|
|
899
|
+
readonly modalities: {
|
|
900
|
+
readonly input: readonly ["text", "image", "pdf"];
|
|
901
|
+
readonly output: readonly ["text"];
|
|
902
|
+
};
|
|
903
|
+
readonly knowledge: "2026-04-30";
|
|
904
|
+
readonly releaseDate: "2026-09-04";
|
|
905
|
+
readonly lastUpdated: "2026-09-04";
|
|
906
|
+
readonly family: "gpt";
|
|
907
|
+
readonly openWeights: false;
|
|
908
|
+
readonly structuredOutput: true;
|
|
909
|
+
readonly temperatureSupported: false;
|
|
910
|
+
readonly provider: "openai";
|
|
869
911
|
}, {
|
|
870
912
|
readonly type: "text";
|
|
871
913
|
readonly modelName: "gpt-5.6-sol";
|
|
@@ -875,12 +917,12 @@ export declare const textModels: readonly [{
|
|
|
875
917
|
readonly inputTokenCost: 4;
|
|
876
918
|
readonly cachedInputTokenCost: 0.4;
|
|
877
919
|
readonly outputTokenCost: 20;
|
|
878
|
-
readonly outputTokensPerSecond:
|
|
920
|
+
readonly outputTokensPerSecond: 77;
|
|
879
921
|
readonly longContext: {
|
|
880
922
|
readonly inputTokenCost: 8;
|
|
881
923
|
readonly cachedInputTokenCost: 0.8;
|
|
882
924
|
readonly outputTokenCost: 30;
|
|
883
|
-
readonly thresholdTokens:
|
|
925
|
+
readonly thresholdTokens: 272000;
|
|
884
926
|
};
|
|
885
927
|
readonly reasoning: {
|
|
886
928
|
readonly levels: readonly ["none", "low", "medium", "high", "xhigh", "max"];
|
|
@@ -910,12 +952,12 @@ export declare const textModels: readonly [{
|
|
|
910
952
|
readonly inputTokenCost: 2;
|
|
911
953
|
readonly cachedInputTokenCost: 0.2;
|
|
912
954
|
readonly outputTokenCost: 12;
|
|
913
|
-
readonly outputTokensPerSecond:
|
|
955
|
+
readonly outputTokensPerSecond: 106;
|
|
914
956
|
readonly longContext: {
|
|
915
957
|
readonly inputTokenCost: 4;
|
|
916
958
|
readonly cachedInputTokenCost: 0.4;
|
|
917
959
|
readonly outputTokenCost: 18;
|
|
918
|
-
readonly thresholdTokens:
|
|
960
|
+
readonly thresholdTokens: 272000;
|
|
919
961
|
};
|
|
920
962
|
readonly reasoning: {
|
|
921
963
|
readonly levels: readonly ["none", "low", "medium", "high", "xhigh", "max"];
|
|
@@ -945,12 +987,12 @@ export declare const textModels: readonly [{
|
|
|
945
987
|
readonly inputTokenCost: 0.2;
|
|
946
988
|
readonly cachedInputTokenCost: 0.02;
|
|
947
989
|
readonly outputTokenCost: 1.2;
|
|
948
|
-
readonly outputTokensPerSecond:
|
|
990
|
+
readonly outputTokensPerSecond: 165;
|
|
949
991
|
readonly longContext: {
|
|
950
992
|
readonly inputTokenCost: 0.4;
|
|
951
993
|
readonly cachedInputTokenCost: 0.04;
|
|
952
994
|
readonly outputTokenCost: 1.8;
|
|
953
|
-
readonly thresholdTokens:
|
|
995
|
+
readonly thresholdTokens: 272000;
|
|
954
996
|
};
|
|
955
997
|
readonly reasoning: {
|
|
956
998
|
readonly levels: readonly ["none", "low", "medium", "high", "xhigh", "max"];
|
|
@@ -980,7 +1022,7 @@ export declare const textModels: readonly [{
|
|
|
980
1022
|
readonly inputTokenCost: 2;
|
|
981
1023
|
readonly cachedInputTokenCost: 0.2;
|
|
982
1024
|
readonly outputTokenCost: 12;
|
|
983
|
-
readonly outputTokensPerSecond:
|
|
1025
|
+
readonly outputTokensPerSecond: 124;
|
|
984
1026
|
readonly longContext: {
|
|
985
1027
|
readonly inputTokenCost: 4;
|
|
986
1028
|
readonly cachedInputTokenCost: 0.4;
|
|
@@ -1043,7 +1085,7 @@ export declare const textModels: readonly [{
|
|
|
1043
1085
|
readonly inputTokenCost: 0.75;
|
|
1044
1086
|
readonly cachedInputTokenCost: 0.075;
|
|
1045
1087
|
readonly outputTokenCost: 3.75;
|
|
1046
|
-
readonly outputTokensPerSecond:
|
|
1088
|
+
readonly outputTokensPerSecond: 329;
|
|
1047
1089
|
readonly inputAudioTokenCost: 1.5;
|
|
1048
1090
|
readonly reasoning: {
|
|
1049
1091
|
readonly levels: readonly ["low", "medium", "high"];
|
|
@@ -1455,6 +1497,7 @@ export declare const textModels: readonly [{
|
|
|
1455
1497
|
readonly cachedInputTokenCost: 0.25;
|
|
1456
1498
|
readonly cacheCreationInputTokenCost: 12.5;
|
|
1457
1499
|
readonly outputTokenCost: 50;
|
|
1500
|
+
readonly outputTokensPerSecond: 69;
|
|
1458
1501
|
readonly reasoning: {
|
|
1459
1502
|
readonly thinkingStyle: "adaptive";
|
|
1460
1503
|
readonly levels: readonly ["low", "medium", "high", "xhigh", "max"];
|
|
@@ -1484,6 +1527,7 @@ export declare const textModels: readonly [{
|
|
|
1484
1527
|
readonly cachedInputTokenCost: 0.5;
|
|
1485
1528
|
readonly cacheCreationInputTokenCost: 6.25;
|
|
1486
1529
|
readonly outputTokenCost: 25;
|
|
1530
|
+
readonly outputTokensPerSecond: 59;
|
|
1487
1531
|
readonly reasoning: {
|
|
1488
1532
|
readonly thinkingStyle: "adaptive";
|
|
1489
1533
|
readonly levels: readonly ["low", "medium", "high", "xhigh", "max"];
|
|
@@ -1625,7 +1669,7 @@ export declare const textModels: readonly [{
|
|
|
1625
1669
|
readonly cachedInputTokenCost: 0.2;
|
|
1626
1670
|
readonly cacheCreationInputTokenCost: 2.5;
|
|
1627
1671
|
readonly outputTokenCost: 10;
|
|
1628
|
-
readonly outputTokensPerSecond:
|
|
1672
|
+
readonly outputTokensPerSecond: 81;
|
|
1629
1673
|
readonly reasoning: {
|
|
1630
1674
|
readonly thinkingStyle: "adaptive";
|
|
1631
1675
|
readonly levels: readonly ["low", "medium", "high", "xhigh", "max"];
|
package/dist/models.js
CHANGED
|
@@ -92,6 +92,14 @@ export const textToSpeechModels = [
|
|
|
92
92
|
},
|
|
93
93
|
// Gemini TTS is token-billed (text input + audio output). No maxInputChars:
|
|
94
94
|
// Gemini documents a 32k-token context, and characters are not a sound proxy.
|
|
95
|
+
{
|
|
96
|
+
type: "text-to-speech",
|
|
97
|
+
modelName: "gemini-3.1-flash-tts-preview",
|
|
98
|
+
provider: "google",
|
|
99
|
+
inputTokenCost: 1.0, // $/1M text-input tokens, verified 2026-09-21
|
|
100
|
+
outputAudioTokenCost: 20.0, // $/1M audio-output tokens
|
|
101
|
+
formats: ["pcm", "wav"],
|
|
102
|
+
},
|
|
95
103
|
{
|
|
96
104
|
type: "text-to-speech",
|
|
97
105
|
modelName: "gemini-2.5-flash-preview-tts",
|
|
@@ -626,7 +634,7 @@ export const textModels = [
|
|
|
626
634
|
inputTokenCost: 5,
|
|
627
635
|
cachedInputTokenCost: 0.5,
|
|
628
636
|
outputTokenCost: 22.5,
|
|
629
|
-
thresholdTokens:
|
|
637
|
+
thresholdTokens: 272000,
|
|
630
638
|
},
|
|
631
639
|
reasoning: {
|
|
632
640
|
levels: ["none", "low", "medium", "high", "xhigh"],
|
|
@@ -719,7 +727,7 @@ export const textModels = [
|
|
|
719
727
|
longContext: {
|
|
720
728
|
inputTokenCost: 60,
|
|
721
729
|
outputTokenCost: 270,
|
|
722
|
-
thresholdTokens:
|
|
730
|
+
thresholdTokens: 272000,
|
|
723
731
|
},
|
|
724
732
|
reasoning: {
|
|
725
733
|
levels: ["medium", "high", "xhigh"],
|
|
@@ -755,7 +763,7 @@ export const textModels = [
|
|
|
755
763
|
inputTokenCost: 10,
|
|
756
764
|
cachedInputTokenCost: 1,
|
|
757
765
|
outputTokenCost: 45,
|
|
758
|
-
thresholdTokens:
|
|
766
|
+
thresholdTokens: 272000,
|
|
759
767
|
},
|
|
760
768
|
reasoning: {
|
|
761
769
|
levels: ["none", "low", "medium", "high", "xhigh"],
|
|
@@ -788,7 +796,7 @@ export const textModels = [
|
|
|
788
796
|
longContext: {
|
|
789
797
|
inputTokenCost: 60,
|
|
790
798
|
outputTokenCost: 270,
|
|
791
|
-
thresholdTokens:
|
|
799
|
+
thresholdTokens: 272000,
|
|
792
800
|
},
|
|
793
801
|
reasoning: {
|
|
794
802
|
levels: ["none", "low", "medium", "high", "xhigh"],
|
|
@@ -810,6 +818,42 @@ export const textModels = [
|
|
|
810
818
|
temperatureSupported: false,
|
|
811
819
|
provider: "openai-responses",
|
|
812
820
|
},
|
|
821
|
+
{
|
|
822
|
+
type: "text",
|
|
823
|
+
modelName: "gpt-6-astra",
|
|
824
|
+
description: "GPT-6 Astra is OpenAI's most capable model, built for the hardest end-to-end work: complex reasoning, coding, computer use, research, and document creation. 1M context window. Standard pricing for ≤272K input tokens; prompts above that are billed at 2x input/cache and 1.5x output for the whole request. Knowledge cutoff: April 2026.",
|
|
825
|
+
maxInputTokens: 1050000,
|
|
826
|
+
maxOutputTokens: 128000,
|
|
827
|
+
inputTokenCost: 10,
|
|
828
|
+
cachedInputTokenCost: 1,
|
|
829
|
+
outputTokenCost: 50,
|
|
830
|
+
outputTokensPerSecond: 69,
|
|
831
|
+
longContext: {
|
|
832
|
+
inputTokenCost: 20,
|
|
833
|
+
cachedInputTokenCost: 2,
|
|
834
|
+
outputTokenCost: 75,
|
|
835
|
+
thresholdTokens: 272000,
|
|
836
|
+
},
|
|
837
|
+
reasoning: {
|
|
838
|
+
levels: ["low", "medium", "high", "xhigh", "max"],
|
|
839
|
+
defaultLevel: "medium",
|
|
840
|
+
canDisable: false,
|
|
841
|
+
outputsThinking: false,
|
|
842
|
+
outputsSignatures: false,
|
|
843
|
+
},
|
|
844
|
+
modalities: {
|
|
845
|
+
input: ["text", "image", "pdf"],
|
|
846
|
+
output: ["text"],
|
|
847
|
+
},
|
|
848
|
+
knowledge: "2026-04-30",
|
|
849
|
+
releaseDate: "2026-09-04",
|
|
850
|
+
lastUpdated: "2026-09-04",
|
|
851
|
+
family: "gpt",
|
|
852
|
+
openWeights: false,
|
|
853
|
+
structuredOutput: true,
|
|
854
|
+
temperatureSupported: false,
|
|
855
|
+
provider: "openai",
|
|
856
|
+
},
|
|
813
857
|
{
|
|
814
858
|
type: "text",
|
|
815
859
|
modelName: "gpt-5.6-sol",
|
|
@@ -819,12 +863,12 @@ export const textModels = [
|
|
|
819
863
|
inputTokenCost: 4,
|
|
820
864
|
cachedInputTokenCost: 0.4,
|
|
821
865
|
outputTokenCost: 20,
|
|
822
|
-
outputTokensPerSecond:
|
|
866
|
+
outputTokensPerSecond: 77,
|
|
823
867
|
longContext: {
|
|
824
868
|
inputTokenCost: 8,
|
|
825
869
|
cachedInputTokenCost: 0.8,
|
|
826
870
|
outputTokenCost: 30,
|
|
827
|
-
thresholdTokens:
|
|
871
|
+
thresholdTokens: 272000,
|
|
828
872
|
},
|
|
829
873
|
reasoning: {
|
|
830
874
|
levels: ["none", "low", "medium", "high", "xhigh", "max"],
|
|
@@ -855,12 +899,12 @@ export const textModels = [
|
|
|
855
899
|
inputTokenCost: 2,
|
|
856
900
|
cachedInputTokenCost: 0.2,
|
|
857
901
|
outputTokenCost: 12,
|
|
858
|
-
outputTokensPerSecond:
|
|
902
|
+
outputTokensPerSecond: 106,
|
|
859
903
|
longContext: {
|
|
860
904
|
inputTokenCost: 4,
|
|
861
905
|
cachedInputTokenCost: 0.4,
|
|
862
906
|
outputTokenCost: 18,
|
|
863
|
-
thresholdTokens:
|
|
907
|
+
thresholdTokens: 272000,
|
|
864
908
|
},
|
|
865
909
|
reasoning: {
|
|
866
910
|
levels: ["none", "low", "medium", "high", "xhigh", "max"],
|
|
@@ -891,12 +935,12 @@ export const textModels = [
|
|
|
891
935
|
inputTokenCost: 0.2,
|
|
892
936
|
cachedInputTokenCost: 0.02,
|
|
893
937
|
outputTokenCost: 1.2,
|
|
894
|
-
outputTokensPerSecond:
|
|
938
|
+
outputTokensPerSecond: 165,
|
|
895
939
|
longContext: {
|
|
896
940
|
inputTokenCost: 0.4,
|
|
897
941
|
cachedInputTokenCost: 0.04,
|
|
898
942
|
outputTokenCost: 1.8,
|
|
899
|
-
thresholdTokens:
|
|
943
|
+
thresholdTokens: 272000,
|
|
900
944
|
},
|
|
901
945
|
reasoning: {
|
|
902
946
|
levels: ["none", "low", "medium", "high", "xhigh", "max"],
|
|
@@ -927,7 +971,7 @@ export const textModels = [
|
|
|
927
971
|
inputTokenCost: 2,
|
|
928
972
|
cachedInputTokenCost: 0.2,
|
|
929
973
|
outputTokenCost: 12,
|
|
930
|
-
outputTokensPerSecond:
|
|
974
|
+
outputTokensPerSecond: 124,
|
|
931
975
|
longContext: {
|
|
932
976
|
inputTokenCost: 4,
|
|
933
977
|
cachedInputTokenCost: 0.4,
|
|
@@ -992,7 +1036,7 @@ export const textModels = [
|
|
|
992
1036
|
inputTokenCost: 0.75,
|
|
993
1037
|
cachedInputTokenCost: 0.075,
|
|
994
1038
|
outputTokenCost: 3.75,
|
|
995
|
-
outputTokensPerSecond:
|
|
1039
|
+
outputTokensPerSecond: 329,
|
|
996
1040
|
inputAudioTokenCost: 1.5,
|
|
997
1041
|
reasoning: {
|
|
998
1042
|
levels: ["low", "medium", "high"],
|
|
@@ -1424,6 +1468,7 @@ export const textModels = [
|
|
|
1424
1468
|
cachedInputTokenCost: 0.25,
|
|
1425
1469
|
cacheCreationInputTokenCost: 12.5,
|
|
1426
1470
|
outputTokenCost: 50,
|
|
1471
|
+
outputTokensPerSecond: 69,
|
|
1427
1472
|
reasoning: {
|
|
1428
1473
|
thinkingStyle: "adaptive",
|
|
1429
1474
|
levels: ["low", "medium", "high", "xhigh", "max"],
|
|
@@ -1454,6 +1499,7 @@ export const textModels = [
|
|
|
1454
1499
|
cachedInputTokenCost: 0.5,
|
|
1455
1500
|
cacheCreationInputTokenCost: 6.25,
|
|
1456
1501
|
outputTokenCost: 25,
|
|
1502
|
+
outputTokensPerSecond: 59,
|
|
1457
1503
|
reasoning: {
|
|
1458
1504
|
thinkingStyle: "adaptive",
|
|
1459
1505
|
levels: ["low", "medium", "high", "xhigh", "max"],
|
|
@@ -1600,7 +1646,7 @@ export const textModels = [
|
|
|
1600
1646
|
cachedInputTokenCost: 0.2,
|
|
1601
1647
|
cacheCreationInputTokenCost: 2.5,
|
|
1602
1648
|
outputTokenCost: 10,
|
|
1603
|
-
outputTokensPerSecond:
|
|
1649
|
+
outputTokensPerSecond: 81,
|
|
1604
1650
|
reasoning: {
|
|
1605
1651
|
thinkingStyle: "adaptive",
|
|
1606
1652
|
levels: ["low", "medium", "high", "xhigh", "max"],
|
|
@@ -1939,7 +1985,7 @@ export const hostedTools = [
|
|
|
1939
1985
|
category: "code_execution",
|
|
1940
1986
|
description: "Run code in a sandboxed container.",
|
|
1941
1987
|
providerToolId: "code_execution",
|
|
1942
|
-
pricing: { unit: "per_hour", amount: 0.05, freeAllowance: "
|
|
1988
|
+
pricing: { unit: "per_hour", amount: 0.05, freeAllowance: "1,550 container-hours/month", note: "Free when used with web_search or web_fetch. 5-minute minimum execution time per container." },
|
|
1943
1989
|
},
|
|
1944
1990
|
{
|
|
1945
1991
|
name: "web_search",
|
|
@@ -2014,7 +2060,12 @@ export const hostedTools = [
|
|
|
2014
2060
|
description: "Grounding with Google Maps (Gemini 3 only).",
|
|
2015
2061
|
providerToolId: "google_maps",
|
|
2016
2062
|
models: ["gemini-3-pro-preview", "gemini-3.1-pro-preview", "gemini-3-flash-preview", "gemini-3.5-flash", "gemini-3.6-flash", "gemini-3.7-flash", "gemini-3.8-flash", "gemini-3.1-flash-lite", "gemini-3.5-flash-lite"],
|
|
2017
|
-
pricing: {
|
|
2063
|
+
pricing: {
|
|
2064
|
+
unit: "per_call",
|
|
2065
|
+
amount: 0.014,
|
|
2066
|
+
freeAllowance: "5,000 grounded prompts/month (Gemini 3)",
|
|
2067
|
+
note: "$14 per 1,000 search queries on the Gemini 3 family.",
|
|
2068
|
+
},
|
|
2018
2069
|
},
|
|
2019
2070
|
{
|
|
2020
2071
|
name: "web_search",
|