@molecule/api-resource-ai-models 1.7.0 → 1.8.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +7 -5
- package/dist/models.d.ts +6 -4
- package/dist/models.d.ts.map +1 -1
- package/dist/models.js +75 -63
- package/package.json +1 -1
package/README.md
CHANGED
|
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
|
|
|
3
3
|
Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
|
|
4
4
|
Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
|
|
5
5
|
To change this document, edit the module-level JSDoc in src/index.ts.
|
|
6
|
-
Generated: 2026-09-
|
|
6
|
+
Generated: 2026-09-24T07:06:51.686Z
|
|
7
7
|
-->
|
|
8
8
|
|
|
9
9
|
# @molecule/api-resource-ai-models
|
|
@@ -893,10 +893,12 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
893
893
|
request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
|
|
894
894
|
output. The catalog's price fields are flat per-MTok rates with no
|
|
895
895
|
context-band dimension, so that band is NOT modeled, exactly as the
|
|
896
|
-
> 200K tiers above are not.
|
|
897
|
-
>
|
|
898
|
-
>
|
|
899
|
-
>
|
|
896
|
+
> 200K tiers above are not. Astra was withheld at the time: it could not serve a
|
|
897
|
+
> tool-carrying request on /v1/chat/completions.)
|
|
898
|
+
> (re-verified 2026-09-24 on the gpt-6-astra model page: ADDED gpt-6-astra
|
|
899
|
+
> ($10/$50, cached $1, cache writes $12.50; 1.05M ctx / 128K out; image
|
|
900
|
+
> input; knowledge cutoff Apr 30 2026; effort low|medium|high|xhigh|max),
|
|
901
|
+
> now that the bond calls /v1/responses.)
|
|
900
902
|
> (re-verified 2026-09-23: ADDED gpt-6-sol ($2/$10, cached $0.20, cache
|
|
901
903
|
writes $2.50) and gpt-6-luna ($0.10/$0.50, cached $0.01, cache writes
|
|
902
904
|
$0.125), released 2026-09-22 — both with the same unmodeled >272K band,
|
package/dist/models.d.ts
CHANGED
|
@@ -116,10 +116,12 @@ import type { ModelDefinition } from './types.js';
|
|
|
116
116
|
* request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
|
|
117
117
|
* output. The catalog's price fields are flat per-MTok rates with no
|
|
118
118
|
* context-band dimension, so that band is NOT modeled, exactly as the
|
|
119
|
-
* >200K tiers above are not.
|
|
120
|
-
*
|
|
121
|
-
*
|
|
122
|
-
*
|
|
119
|
+
* >200K tiers above are not. Astra was withheld at the time: it could not serve a
|
|
120
|
+
* tool-carrying request on /v1/chat/completions.)
|
|
121
|
+
* (re-verified 2026-09-24 on the gpt-6-astra model page: ADDED gpt-6-astra
|
|
122
|
+
* ($10/$50, cached $1, cache writes $12.50; 1.05M ctx / 128K out; image
|
|
123
|
+
* input; knowledge cutoff Apr 30 2026; effort low|medium|high|xhigh|max),
|
|
124
|
+
* now that the bond calls /v1/responses.)
|
|
123
125
|
* (re-verified 2026-09-23: ADDED gpt-6-sol ($2/$10, cached $0.20, cache
|
|
124
126
|
* writes $2.50) and gpt-6-luna ($0.10/$0.50, cached $0.01, cache writes
|
|
125
127
|
* $0.125), released 2026-09-22 — both with the same unmodeled >272K band,
|
package/dist/models.d.ts.map
CHANGED
|
@@ -1 +1 @@
|
|
|
1
|
-
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAwDjD
|
|
1
|
+
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAwDjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAoQG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAqkEnC,CAAA"}
|
package/dist/models.js
CHANGED
|
@@ -165,10 +165,12 @@ const CHINA_PUBLIC_HOLIDAYS = [
|
|
|
165
165
|
* request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
|
|
166
166
|
* output. The catalog's price fields are flat per-MTok rates with no
|
|
167
167
|
* context-band dimension, so that band is NOT modeled, exactly as the
|
|
168
|
-
* >200K tiers above are not.
|
|
169
|
-
*
|
|
170
|
-
*
|
|
171
|
-
*
|
|
168
|
+
* >200K tiers above are not. Astra was withheld at the time: it could not serve a
|
|
169
|
+
* tool-carrying request on /v1/chat/completions.)
|
|
170
|
+
* (re-verified 2026-09-24 on the gpt-6-astra model page: ADDED gpt-6-astra
|
|
171
|
+
* ($10/$50, cached $1, cache writes $12.50; 1.05M ctx / 128K out; image
|
|
172
|
+
* input; knowledge cutoff Apr 30 2026; effort low|medium|high|xhigh|max),
|
|
173
|
+
* now that the bond calls /v1/responses.)
|
|
172
174
|
* (re-verified 2026-09-23: ADDED gpt-6-sol ($2/$10, cached $0.20, cache
|
|
173
175
|
* writes $2.50) and gpt-6-luna ($0.10/$0.50, cached $0.01, cache writes
|
|
174
176
|
* $0.125), released 2026-09-22 — both with the same unmodeled >272K band,
|
|
@@ -746,32 +748,14 @@ export const MODELS = [
|
|
|
746
748
|
// exist upstream — not modeled (same as the Gemini/Grok >200K tiers), and
|
|
747
749
|
// neither are the Batch (0.5×) or Sol "Fast mode" (2×) cards.
|
|
748
750
|
//
|
|
749
|
-
//
|
|
750
|
-
//
|
|
751
|
-
//
|
|
752
|
-
//
|
|
753
|
-
//
|
|
754
|
-
//
|
|
755
|
-
//
|
|
756
|
-
//
|
|
757
|
-
// reasoning_effort to 'none'."
|
|
758
|
-
// - tools, field omitted → the same 400 (the model applies its own default)
|
|
759
|
-
// - tools + reasoning_effort 'none' → 400 "Unsupported value: … does not
|
|
760
|
-
// support 'none' with this model. Supported values are: 'low', 'medium',
|
|
761
|
-
// 'high', and 'xhigh'."
|
|
762
|
-
// So the gpt-5.6 workaround (`toolsRequireReasoningOff`, which pins an
|
|
763
|
-
// explicit 'none') does NOT carry over: OpenAI removed 'none' from this
|
|
764
|
-
// model's ladder, which closes the one door that made the 5.6 family usable.
|
|
765
|
-
// Nothing is wrong with the model or the account — no-tools chat and
|
|
766
|
-
// /v1/responses WITH tools both return 200. Both candidate entries were built
|
|
767
|
-
// and run through molecule-dev's verify:model-dispatch; both FAILED, so the
|
|
768
|
-
// entry was withheld rather than shipped broken (the glm-5.3-flash lesson:
|
|
769
|
-
// a catalog entry that cannot serve a Synthase-shaped turn breaks every turn
|
|
770
|
-
// on it). The freshness gate WILL keep listing it as a new-model candidate —
|
|
771
|
-
// that is correct; it becomes addable the day the bond speaks /v1/responses.
|
|
772
|
-
// Note also that the docs page advertises effort 'max', which
|
|
773
|
-
// /v1/chat/completions rejects for this model — verify the ladder against the
|
|
774
|
-
// endpoint, not the docs, when this is revisited.
|
|
751
|
+
// gpt-6-astra (OpenAI's flagship since 2026-09-04) was withheld until the
|
|
752
|
+
// bond moved to /v1/responses: on /v1/chat/completions it rejects every
|
|
753
|
+
// tool-carrying request (tools + any reasoning_effort → 400, and it has no
|
|
754
|
+
// 'none' effort to fall back on — probed live 2026-09-21). The bond calls
|
|
755
|
+
// /v1/responses on OpenAI's own endpoint since 2026-09-24, where tools +
|
|
756
|
+
// every effort low..max return 200 (probed live 2026-09-24), so astra needs
|
|
757
|
+
// no `toolsRequireReasoningOff` pin. The docs page lists chat/completions
|
|
758
|
+
// tool support too; the live endpoint disagrees — trust the endpoint.
|
|
775
759
|
//
|
|
776
760
|
// gpt-6-sol and gpt-6-luna (released 2026-09-22) are NOT astra's case: both
|
|
777
761
|
// still accept reasoning_effort 'none', and tools + 'none' returns 200 on
|
|
@@ -782,6 +766,37 @@ export const MODELS = [
|
|
|
782
766
|
// $2/$10 is priced at 5.6-terra's tier and undercuts 5.6-sol, so it supersedes
|
|
783
767
|
// both; luna supersedes 5.6-luna at half the price.
|
|
784
768
|
// ---------------------------------------------------------------------------
|
|
769
|
+
{
|
|
770
|
+
id: 'gpt-6-astra',
|
|
771
|
+
provider: 'openai',
|
|
772
|
+
label: 'GPT-6 Astra',
|
|
773
|
+
description: 'OpenAI flagship — the hardest reasoning & coding work',
|
|
774
|
+
// Documented as 1.05M; floored to 1M like the other OpenAI entries.
|
|
775
|
+
contextWindow: 1_000_000,
|
|
776
|
+
maxOutputTokens: 128_000,
|
|
777
|
+
supportsThinking: true,
|
|
778
|
+
thinkingBudgetTokens: 16_000,
|
|
779
|
+
thinkingConfigurable: true,
|
|
780
|
+
// No 'none' on this model. 'max' verified accepted with tools on
|
|
781
|
+
// /v1/responses (2026-09-24).
|
|
782
|
+
supportedEffortLevels: ['low', 'medium', 'high', 'xhigh', 'max'],
|
|
783
|
+
defaultEffortLevel: 'medium',
|
|
784
|
+
// temperature → 400 "Unsupported parameter: 'temperature' is not supported
|
|
785
|
+
// with this model" (probed live 2026-09-24).
|
|
786
|
+
rejectsTemperature: true,
|
|
787
|
+
supportsVision: true,
|
|
788
|
+
supportsPromptCaching: true,
|
|
789
|
+
supportsTools: true,
|
|
790
|
+
// NO webSearchToolType yet: /v1/responses serves `web_search` (verified
|
|
791
|
+
// 2026-09-24), but its per-call fee is not metered — add with metering.
|
|
792
|
+
codeExecutionToolType: 'code_interpreter',
|
|
793
|
+
// Standard tier; >272K band ($20/$75, cached $2, writes $25) not modeled.
|
|
794
|
+
inputPricePerMTok: 10,
|
|
795
|
+
outputPricePerMTok: 50,
|
|
796
|
+
cacheReadPricePerMTok: 1,
|
|
797
|
+
cacheWritePricePerMTok: 12.5,
|
|
798
|
+
knowledgeCutoff: '2026-04-30',
|
|
799
|
+
},
|
|
785
800
|
{
|
|
786
801
|
id: 'gpt-6-sol',
|
|
787
802
|
provider: 'openai',
|
|
@@ -799,10 +814,11 @@ export const MODELS = [
|
|
|
799
814
|
supportsPromptCaching: true,
|
|
800
815
|
supportsTools: true,
|
|
801
816
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
817
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
818
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
802
819
|
toolsRequireReasoningOff: true,
|
|
803
|
-
// NO webSearchToolType: see gpt-
|
|
804
|
-
//
|
|
805
|
-
// bond moves to /v1/responses.
|
|
820
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
821
|
+
// its per-call fee is not metered yet.
|
|
806
822
|
codeExecutionToolType: 'code_interpreter',
|
|
807
823
|
// Standard tier. A long-context band above 272K prompt tokens reprices the
|
|
808
824
|
// whole request ($4/$15, cached $0.40, cache writes $5) — not modeled, same
|
|
@@ -829,6 +845,8 @@ export const MODELS = [
|
|
|
829
845
|
supportsPromptCaching: true,
|
|
830
846
|
supportsTools: true,
|
|
831
847
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
848
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
849
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
832
850
|
toolsRequireReasoningOff: true,
|
|
833
851
|
// NO webSearchToolType: see gpt-5.6-sol.
|
|
834
852
|
codeExecutionToolType: 'code_interpreter',
|
|
@@ -858,12 +876,11 @@ export const MODELS = [
|
|
|
858
876
|
supportsPromptCaching: true,
|
|
859
877
|
supportsTools: true,
|
|
860
878
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
879
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
880
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
861
881
|
toolsRequireReasoningOff: true,
|
|
862
|
-
// NO webSearchToolType:
|
|
863
|
-
//
|
|
864
|
-
// bond deliberately forwards no server tools. Advertising one here surfaced
|
|
865
|
-
// web search in the system prompt while it could never work. Re-add when
|
|
866
|
-
// the bond moves to /v1/responses (verified 2026-08-28).
|
|
882
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
883
|
+
// its per-call fee is not metered yet.
|
|
867
884
|
codeExecutionToolType: 'code_interpreter',
|
|
868
885
|
// LIST price. OpenAI ran a >20% PROMO from 2026-08-22 ($4/$20, cache read
|
|
869
886
|
// $0.40, cache write $5) — "GPT-5.6 Sol's promotional pricing is available
|
|
@@ -899,12 +916,11 @@ export const MODELS = [
|
|
|
899
916
|
supportsPromptCaching: true,
|
|
900
917
|
supportsTools: true,
|
|
901
918
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
919
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
920
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
902
921
|
toolsRequireReasoningOff: true,
|
|
903
|
-
// NO webSearchToolType:
|
|
904
|
-
//
|
|
905
|
-
// bond deliberately forwards no server tools. Advertising one here surfaced
|
|
906
|
-
// web search in the system prompt while it could never work. Re-add when
|
|
907
|
-
// the bond moves to /v1/responses (verified 2026-08-28).
|
|
922
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
923
|
+
// its per-call fee is not metered yet.
|
|
908
924
|
codeExecutionToolType: 'code_interpreter',
|
|
909
925
|
// Repriced 2026-07-30 (20% cut from $2.50/$15).
|
|
910
926
|
inputPricePerMTok: 2,
|
|
@@ -935,12 +951,11 @@ export const MODELS = [
|
|
|
935
951
|
supportsPromptCaching: true,
|
|
936
952
|
supportsTools: true,
|
|
937
953
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
954
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
955
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
938
956
|
toolsRequireReasoningOff: true,
|
|
939
|
-
// NO webSearchToolType:
|
|
940
|
-
//
|
|
941
|
-
// bond deliberately forwards no server tools. Advertising one here surfaced
|
|
942
|
-
// web search in the system prompt while it could never work. Re-add when
|
|
943
|
-
// the bond moves to /v1/responses (verified 2026-08-28).
|
|
957
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
958
|
+
// its per-call fee is not metered yet.
|
|
944
959
|
codeExecutionToolType: 'code_interpreter',
|
|
945
960
|
// Repriced 2026-07-30 (80% cut from $1/$6).
|
|
946
961
|
inputPricePerMTok: 0.2,
|
|
@@ -977,12 +992,11 @@ export const MODELS = [
|
|
|
977
992
|
supportsPromptCaching: true,
|
|
978
993
|
supportsTools: true,
|
|
979
994
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
995
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
996
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
980
997
|
toolsRequireReasoningOff: true,
|
|
981
|
-
// NO webSearchToolType:
|
|
982
|
-
//
|
|
983
|
-
// bond deliberately forwards no server tools. Advertising one here surfaced
|
|
984
|
-
// web search in the system prompt while it could never work. Re-add when
|
|
985
|
-
// the bond moves to /v1/responses (verified 2026-08-28).
|
|
998
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
999
|
+
// its per-call fee is not metered yet.
|
|
986
1000
|
codeExecutionToolType: 'code_interpreter',
|
|
987
1001
|
inputPricePerMTok: 5,
|
|
988
1002
|
outputPricePerMTok: 30,
|
|
@@ -1012,12 +1026,11 @@ export const MODELS = [
|
|
|
1012
1026
|
supportsPromptCaching: true,
|
|
1013
1027
|
supportsTools: true,
|
|
1014
1028
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
1029
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
1030
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
1015
1031
|
toolsRequireReasoningOff: true,
|
|
1016
|
-
// NO webSearchToolType:
|
|
1017
|
-
//
|
|
1018
|
-
// bond deliberately forwards no server tools. Advertising one here surfaced
|
|
1019
|
-
// web search in the system prompt while it could never work. Re-add when
|
|
1020
|
-
// the bond moves to /v1/responses (verified 2026-08-28).
|
|
1032
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
1033
|
+
// its per-call fee is not metered yet.
|
|
1021
1034
|
codeExecutionToolType: 'code_interpreter',
|
|
1022
1035
|
inputPricePerMTok: 2.5,
|
|
1023
1036
|
outputPricePerMTok: 15,
|
|
@@ -1049,12 +1062,11 @@ export const MODELS = [
|
|
|
1049
1062
|
supportsPromptCaching: true,
|
|
1050
1063
|
supportsTools: true,
|
|
1051
1064
|
// Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
|
|
1065
|
+
// The bond now calls /v1/responses, where it is not (verified 2026-09-24);
|
|
1066
|
+
// lifting this pin changes reasoning quality/cost, so it waits for a model eval.
|
|
1052
1067
|
toolsRequireReasoningOff: true,
|
|
1053
|
-
// NO webSearchToolType:
|
|
1054
|
-
//
|
|
1055
|
-
// bond deliberately forwards no server tools. Advertising one here surfaced
|
|
1056
|
-
// web search in the system prompt while it could never work. Re-add when
|
|
1057
|
-
// the bond moves to /v1/responses (verified 2026-08-28).
|
|
1068
|
+
// NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
|
|
1069
|
+
// its per-call fee is not metered yet.
|
|
1058
1070
|
codeExecutionToolType: 'code_interpreter',
|
|
1059
1071
|
inputPricePerMTok: 0.75,
|
|
1060
1072
|
outputPricePerMTok: 4.5,
|
package/package.json
CHANGED