@molecule/api-resource-ai-models 1.7.0 → 1.8.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/README.md CHANGED
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
3
3
  Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
4
4
  Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
5
5
  To change this document, edit the module-level JSDoc in src/index.ts.
6
- Generated: 2026-09-23T14:26:31.162Z
6
+ Generated: 2026-09-24T07:06:51.686Z
7
7
  -->
8
8
 
9
9
  # @molecule/api-resource-ai-models
@@ -893,10 +893,12 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
893
893
  request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
894
894
  output. The catalog's price fields are flat per-MTok rates with no
895
895
  context-band dimension, so that band is NOT modeled, exactly as the
896
- > 200K tiers above are not. It is moot for now — astra is deliberately NOT
897
- > in the catalog because it cannot serve a tool-carrying request on the
898
- > bond's /v1/chat/completions endpoint; see the DO NOT ADD block above the
899
- > OpenAI entries for the live 400s and the condition that lifts it.)
896
+ > 200K tiers above are not. Astra was withheld at the time: it could not serve a
897
+ > tool-carrying request on /v1/chat/completions.)
898
+ > (re-verified 2026-09-24 on the gpt-6-astra model page: ADDED gpt-6-astra
899
+ > ($10/$50, cached $1, cache writes $12.50; 1.05M ctx / 128K out; image
900
+ > input; knowledge cutoff Apr 30 2026; effort low|medium|high|xhigh|max),
901
+ > now that the bond calls /v1/responses.)
900
902
  > (re-verified 2026-09-23: ADDED gpt-6-sol ($2/$10, cached $0.20, cache
901
903
  writes $2.50) and gpt-6-luna ($0.10/$0.50, cached $0.01, cache writes
902
904
  $0.125), released 2026-09-22 — both with the same unmodeled >272K band,
package/dist/models.d.ts CHANGED
@@ -116,10 +116,12 @@ import type { ModelDefinition } from './types.js';
116
116
  * request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
117
117
  * output. The catalog's price fields are flat per-MTok rates with no
118
118
  * context-band dimension, so that band is NOT modeled, exactly as the
119
- * >200K tiers above are not. It is moot for now — astra is deliberately NOT
120
- * in the catalog because it cannot serve a tool-carrying request on the
121
- * bond's /v1/chat/completions endpoint; see the DO NOT ADD block above the
122
- * OpenAI entries for the live 400s and the condition that lifts it.)
119
+ * >200K tiers above are not. Astra was withheld at the time: it could not serve a
120
+ * tool-carrying request on /v1/chat/completions.)
121
+ * (re-verified 2026-09-24 on the gpt-6-astra model page: ADDED gpt-6-astra
122
+ * ($10/$50, cached $1, cache writes $12.50; 1.05M ctx / 128K out; image
123
+ * input; knowledge cutoff Apr 30 2026; effort low|medium|high|xhigh|max),
124
+ * now that the bond calls /v1/responses.)
123
125
  * (re-verified 2026-09-23: ADDED gpt-6-sol ($2/$10, cached $0.20, cache
124
126
  * writes $2.50) and gpt-6-luna ($0.10/$0.50, cached $0.01, cache writes
125
127
  * $0.125), released 2026-09-22 — both with the same unmodeled >272K band,
@@ -1 +1 @@
1
- {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAwDjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAkQG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EA2jEnC,CAAA"}
1
+ {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAwDjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAoQG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAqkEnC,CAAA"}
package/dist/models.js CHANGED
@@ -165,10 +165,12 @@ const CHINA_PUBLIC_HOLIDAYS = [
165
165
  * request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
166
166
  * output. The catalog's price fields are flat per-MTok rates with no
167
167
  * context-band dimension, so that band is NOT modeled, exactly as the
168
- * >200K tiers above are not. It is moot for now — astra is deliberately NOT
169
- * in the catalog because it cannot serve a tool-carrying request on the
170
- * bond's /v1/chat/completions endpoint; see the DO NOT ADD block above the
171
- * OpenAI entries for the live 400s and the condition that lifts it.)
168
+ * >200K tiers above are not. Astra was withheld at the time: it could not serve a
169
+ * tool-carrying request on /v1/chat/completions.)
170
+ * (re-verified 2026-09-24 on the gpt-6-astra model page: ADDED gpt-6-astra
171
+ * ($10/$50, cached $1, cache writes $12.50; 1.05M ctx / 128K out; image
172
+ * input; knowledge cutoff Apr 30 2026; effort low|medium|high|xhigh|max),
173
+ * now that the bond calls /v1/responses.)
172
174
  * (re-verified 2026-09-23: ADDED gpt-6-sol ($2/$10, cached $0.20, cache
173
175
  * writes $2.50) and gpt-6-luna ($0.10/$0.50, cached $0.01, cache writes
174
176
  * $0.125), released 2026-09-22 — both with the same unmodeled >272K band,
@@ -746,32 +748,14 @@ export const MODELS = [
746
748
  // exist upstream — not modeled (same as the Gemini/Grok >200K tiers), and
747
749
  // neither are the Batch (0.5×) or Sol "Fast mode" (2×) cards.
748
750
  //
749
- // DO NOT ADD gpt-6-astra (OpenAI's flagship since 2026-09-04) UNTIL THE BOND
750
- // MOVES TO /v1/responses. It is deliberately absent, not overlooked. The
751
- // openai bond posts to /v1/chat/completions, and on that endpoint this model
752
- // rejects EVERY request Synthase can send, because Synthase always carries
753
- // function tools (probed live 2026-09-21 on our own key):
754
- // - tools + reasoning_effort low|medium|high|xhigh → 400 "Function tools
755
- // with reasoning_effort are not supported for gpt-6-astra in
756
- // /v1/chat/completions. To use function tools, use /v1/responses or set
757
- // reasoning_effort to 'none'."
758
- // - tools, field omitted → the same 400 (the model applies its own default)
759
- // - tools + reasoning_effort 'none' → 400 "Unsupported value: … does not
760
- // support 'none' with this model. Supported values are: 'low', 'medium',
761
- // 'high', and 'xhigh'."
762
- // So the gpt-5.6 workaround (`toolsRequireReasoningOff`, which pins an
763
- // explicit 'none') does NOT carry over: OpenAI removed 'none' from this
764
- // model's ladder, which closes the one door that made the 5.6 family usable.
765
- // Nothing is wrong with the model or the account — no-tools chat and
766
- // /v1/responses WITH tools both return 200. Both candidate entries were built
767
- // and run through molecule-dev's verify:model-dispatch; both FAILED, so the
768
- // entry was withheld rather than shipped broken (the glm-5.3-flash lesson:
769
- // a catalog entry that cannot serve a Synthase-shaped turn breaks every turn
770
- // on it). The freshness gate WILL keep listing it as a new-model candidate —
771
- // that is correct; it becomes addable the day the bond speaks /v1/responses.
772
- // Note also that the docs page advertises effort 'max', which
773
- // /v1/chat/completions rejects for this model — verify the ladder against the
774
- // endpoint, not the docs, when this is revisited.
751
+ // gpt-6-astra (OpenAI's flagship since 2026-09-04) was withheld until the
752
+ // bond moved to /v1/responses: on /v1/chat/completions it rejects every
753
+ // tool-carrying request (tools + any reasoning_effort → 400, and it has no
754
+ // 'none' effort to fall back on — probed live 2026-09-21). The bond calls
755
+ // /v1/responses on OpenAI's own endpoint since 2026-09-24, where tools +
756
+ // every effort low..max return 200 (probed live 2026-09-24), so astra needs
757
+ // no `toolsRequireReasoningOff` pin. The docs page lists chat/completions
758
+ // tool support too; the live endpoint disagrees — trust the endpoint.
775
759
  //
776
760
  // gpt-6-sol and gpt-6-luna (released 2026-09-22) are NOT astra's case: both
777
761
  // still accept reasoning_effort 'none', and tools + 'none' returns 200 on
@@ -782,6 +766,37 @@ export const MODELS = [
782
766
  // $2/$10 is priced at 5.6-terra's tier and undercuts 5.6-sol, so it supersedes
783
767
  // both; luna supersedes 5.6-luna at half the price.
784
768
  // ---------------------------------------------------------------------------
769
+ {
770
+ id: 'gpt-6-astra',
771
+ provider: 'openai',
772
+ label: 'GPT-6 Astra',
773
+ description: 'OpenAI flagship — the hardest reasoning & coding work',
774
+ // Documented as 1.05M; floored to 1M like the other OpenAI entries.
775
+ contextWindow: 1_000_000,
776
+ maxOutputTokens: 128_000,
777
+ supportsThinking: true,
778
+ thinkingBudgetTokens: 16_000,
779
+ thinkingConfigurable: true,
780
+ // No 'none' on this model. 'max' verified accepted with tools on
781
+ // /v1/responses (2026-09-24).
782
+ supportedEffortLevels: ['low', 'medium', 'high', 'xhigh', 'max'],
783
+ defaultEffortLevel: 'medium',
784
+ // temperature → 400 "Unsupported parameter: 'temperature' is not supported
785
+ // with this model" (probed live 2026-09-24).
786
+ rejectsTemperature: true,
787
+ supportsVision: true,
788
+ supportsPromptCaching: true,
789
+ supportsTools: true,
790
+ // NO webSearchToolType yet: /v1/responses serves `web_search` (verified
791
+ // 2026-09-24), but its per-call fee is not metered — add with metering.
792
+ codeExecutionToolType: 'code_interpreter',
793
+ // Standard tier; >272K band ($20/$75, cached $2, writes $25) not modeled.
794
+ inputPricePerMTok: 10,
795
+ outputPricePerMTok: 50,
796
+ cacheReadPricePerMTok: 1,
797
+ cacheWritePricePerMTok: 12.5,
798
+ knowledgeCutoff: '2026-04-30',
799
+ },
785
800
  {
786
801
  id: 'gpt-6-sol',
787
802
  provider: 'openai',
@@ -799,10 +814,11 @@ export const MODELS = [
799
814
  supportsPromptCaching: true,
800
815
  supportsTools: true,
801
816
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
817
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
818
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
802
819
  toolsRequireReasoningOff: true,
803
- // NO webSearchToolType: see gpt-5.6-sol — the bond calls
804
- // /v1/chat/completions, which has no web_search tool type. Re-add when the
805
- // bond moves to /v1/responses.
820
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
821
+ // its per-call fee is not metered yet.
806
822
  codeExecutionToolType: 'code_interpreter',
807
823
  // Standard tier. A long-context band above 272K prompt tokens reprices the
808
824
  // whole request ($4/$15, cached $0.40, cache writes $5) — not modeled, same
@@ -829,6 +845,8 @@ export const MODELS = [
829
845
  supportsPromptCaching: true,
830
846
  supportsTools: true,
831
847
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
848
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
849
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
832
850
  toolsRequireReasoningOff: true,
833
851
  // NO webSearchToolType: see gpt-5.6-sol.
834
852
  codeExecutionToolType: 'code_interpreter',
@@ -858,12 +876,11 @@ export const MODELS = [
858
876
  supportsPromptCaching: true,
859
877
  supportsTools: true,
860
878
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
879
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
880
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
861
881
  toolsRequireReasoningOff: true,
862
- // NO webSearchToolType: the OpenAI bond calls /v1/chat/completions, which
863
- // has no web_search tool type (it is a Responses-API construct), and the
864
- // bond deliberately forwards no server tools. Advertising one here surfaced
865
- // web search in the system prompt while it could never work. Re-add when
866
- // the bond moves to /v1/responses (verified 2026-08-28).
882
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
883
+ // its per-call fee is not metered yet.
867
884
  codeExecutionToolType: 'code_interpreter',
868
885
  // LIST price. OpenAI ran a >20% PROMO from 2026-08-22 ($4/$20, cache read
869
886
  // $0.40, cache write $5) — "GPT-5.6 Sol's promotional pricing is available
@@ -899,12 +916,11 @@ export const MODELS = [
899
916
  supportsPromptCaching: true,
900
917
  supportsTools: true,
901
918
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
919
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
920
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
902
921
  toolsRequireReasoningOff: true,
903
- // NO webSearchToolType: the OpenAI bond calls /v1/chat/completions, which
904
- // has no web_search tool type (it is a Responses-API construct), and the
905
- // bond deliberately forwards no server tools. Advertising one here surfaced
906
- // web search in the system prompt while it could never work. Re-add when
907
- // the bond moves to /v1/responses (verified 2026-08-28).
922
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
923
+ // its per-call fee is not metered yet.
908
924
  codeExecutionToolType: 'code_interpreter',
909
925
  // Repriced 2026-07-30 (20% cut from $2.50/$15).
910
926
  inputPricePerMTok: 2,
@@ -935,12 +951,11 @@ export const MODELS = [
935
951
  supportsPromptCaching: true,
936
952
  supportsTools: true,
937
953
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
954
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
955
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
938
956
  toolsRequireReasoningOff: true,
939
- // NO webSearchToolType: the OpenAI bond calls /v1/chat/completions, which
940
- // has no web_search tool type (it is a Responses-API construct), and the
941
- // bond deliberately forwards no server tools. Advertising one here surfaced
942
- // web search in the system prompt while it could never work. Re-add when
943
- // the bond moves to /v1/responses (verified 2026-08-28).
957
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
958
+ // its per-call fee is not metered yet.
944
959
  codeExecutionToolType: 'code_interpreter',
945
960
  // Repriced 2026-07-30 (80% cut from $1/$6).
946
961
  inputPricePerMTok: 0.2,
@@ -977,12 +992,11 @@ export const MODELS = [
977
992
  supportsPromptCaching: true,
978
993
  supportsTools: true,
979
994
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
995
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
996
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
980
997
  toolsRequireReasoningOff: true,
981
- // NO webSearchToolType: the OpenAI bond calls /v1/chat/completions, which
982
- // has no web_search tool type (it is a Responses-API construct), and the
983
- // bond deliberately forwards no server tools. Advertising one here surfaced
984
- // web search in the system prompt while it could never work. Re-add when
985
- // the bond moves to /v1/responses (verified 2026-08-28).
998
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
999
+ // its per-call fee is not metered yet.
986
1000
  codeExecutionToolType: 'code_interpreter',
987
1001
  inputPricePerMTok: 5,
988
1002
  outputPricePerMTok: 30,
@@ -1012,12 +1026,11 @@ export const MODELS = [
1012
1026
  supportsPromptCaching: true,
1013
1027
  supportsTools: true,
1014
1028
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
1029
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
1030
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
1015
1031
  toolsRequireReasoningOff: true,
1016
- // NO webSearchToolType: the OpenAI bond calls /v1/chat/completions, which
1017
- // has no web_search tool type (it is a Responses-API construct), and the
1018
- // bond deliberately forwards no server tools. Advertising one here surfaced
1019
- // web search in the system prompt while it could never work. Re-add when
1020
- // the bond moves to /v1/responses (verified 2026-08-28).
1032
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
1033
+ // its per-call fee is not metered yet.
1021
1034
  codeExecutionToolType: 'code_interpreter',
1022
1035
  inputPricePerMTok: 2.5,
1023
1036
  outputPricePerMTok: 15,
@@ -1049,12 +1062,11 @@ export const MODELS = [
1049
1062
  supportsPromptCaching: true,
1050
1063
  supportsTools: true,
1051
1064
  // Tools + ANY reasoning is a 400 on /v1/chat/completions for this family.
1065
+ // The bond now calls /v1/responses, where it is not (verified 2026-09-24);
1066
+ // lifting this pin changes reasoning quality/cost, so it waits for a model eval.
1052
1067
  toolsRequireReasoningOff: true,
1053
- // NO webSearchToolType: the OpenAI bond calls /v1/chat/completions, which
1054
- // has no web_search tool type (it is a Responses-API construct), and the
1055
- // bond deliberately forwards no server tools. Advertising one here surfaced
1056
- // web search in the system prompt while it could never work. Re-add when
1057
- // the bond moves to /v1/responses (verified 2026-08-28).
1068
+ // NO webSearchToolType: see gpt-6-astra — served on /v1/responses, but
1069
+ // its per-call fee is not metered yet.
1058
1070
  codeExecutionToolType: 'code_interpreter',
1059
1071
  inputPricePerMTok: 0.75,
1060
1072
  outputPricePerMTok: 4.5,
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@molecule/api-resource-ai-models",
3
- "version": "1.7.0",
3
+ "version": "1.8.0",
4
4
  "description": "AI model catalog — server-side source of truth plus an authentication-gated discovery endpoint",
5
5
  "type": "module",
6
6
  "main": "dist/index.js",