@prestyj/core 5.28.0 → 5.29.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -16,8 +16,8 @@ interface ModelInfo {
16
16
  /**
17
17
  * Vendor-declared default reasoning level (Codex models.json
18
18
  * `default_reasoning_level`). When present, fresh sessions start here rather
19
- * than at the ceiling: the deep-reasoning flagships (Astra ships "low",
20
- * GPT-6 Sol/Luna "medium") think dramatically longer per rung, so defaulting to
19
+ * than at the ceiling: the deep-reasoning models (Astra and GPT-6.1 Sol ship
20
+ * "low", GPT-6 Luna "medium") think dramatically longer per rung, so defaulting to
21
21
  * `maxThinkingLevel` made new sessions pathologically slow.
22
22
  */
23
23
  defaultThinkingLevel?: ThinkingLevel;
@@ -39,7 +39,7 @@ interface ModelInfo {
39
39
  /**
40
40
  * The top reasoning tier this model genuinely uses. Used when thinking is
41
41
  * enabled to pick the strongest setting per model:
42
- * - OpenAI GPT-6 Astra / Sol: `ultra` (Codex orchestration preset above `max`)
42
+ * - OpenAI GPT-6 Astra / GPT-6.1 Sol: `ultra` (Codex orchestration preset above `max`)
43
43
  * - OpenAI GPT-6 Luna: `max`
44
44
  * - OpenAI Pro/Codex/old: clamped to what the model accepts
45
45
  * - Claude Fable 5.1 / Fable 5 / Mythos 5, Opus 5.5 / Opus 5 and Sonnet 5.5: `max`
@@ -134,25 +134,17 @@ declare function getDefaultThinkingLevel(modelId: string, options?: {
134
134
  baseUrl?: string;
135
135
  }): ThinkingLevel;
136
136
  /**
137
- * Get the model to use for compaction summarization.
138
- * - Anthropic: always Sonnet 5.5
139
- * - OpenAI: cheapest (Codex Mini)
140
- * - Gemini: use the current model
141
- * - GLM: GLM-5.3-Flash (the registered low-cost sibling)
142
- * - Moonshot: use the current model (no cheap alternative registered)
143
- */
144
- declare function getSummaryModel(provider: Provider, currentModelId: string): ModelInfo;
145
- /**
146
- * Fastest/cheapest sibling within the SAME provider, for scout-style read-only
147
- * sub-agents (recon, research) where a low-latency model is enough and the
148
- * frontier model is wasted spend + latency.
137
+ * Get the model to use for compaction summarization — the ONLY place a
138
+ * cheaper model is used. Everything else (sub-agents, reviewers, judges)
139
+ * runs on the parent's model.
149
140
  *
150
- * Routes off each model's `costTier` — the single source of truth that already
151
- * travels with the registry entry — so a model rename/bump needs no change
152
- * here. Providers with no low-tier sibling (Moonshot, MiniMax, Sakana,
153
- * OpenRouter) gracefully keep the parent model, so there's never a
154
- * crash or a cross-provider jump to a login the user may not have.
141
+ * Picks the provider's first model tagged with its summary tier
142
+ * ({@link SUMMARY_COST_TIER}). When none is registered — the tier was removed,
143
+ * or the provider has no cheap sibling — it summarizes on the current model.
144
+ * Never throws. The caller still retries on the current model when the
145
+ * provider rejects the chosen one at runtime (retired upstream, or not on the
146
+ * user's plan), since the registry cannot know that.
155
147
  */
156
- declare function getFastModel(provider: Provider, currentModelId: string): ModelInfo;
148
+ declare function getSummaryModel(provider: Provider, currentModelId: string): ModelInfo;
157
149
 
158
- export { type ContextWindowOptions, DEFAULT_MAX_VIDEO_BYTES, MODELS, type ModelInfo, clearRuntimeModels, getAllModels, getAuthStorageKey, getAuthStorageKeys, getContextWindow, getDefaultModel, getDefaultThinkingLevel, getFastModel, getMaxThinkingLevel, getModel, getModelsForProvider, getSummaryModel, getToolResultCharLimit, getVideoByteLimit, registerRuntimeModels, usesOpenAICodexTransport };
150
+ export { type ContextWindowOptions, DEFAULT_MAX_VIDEO_BYTES, MODELS, type ModelInfo, clearRuntimeModels, getAllModels, getAuthStorageKey, getAuthStorageKeys, getContextWindow, getDefaultModel, getDefaultThinkingLevel, getMaxThinkingLevel, getModel, getModelsForProvider, getSummaryModel, getToolResultCharLimit, getVideoByteLimit, registerRuntimeModels, usesOpenAICodexTransport };
@@ -16,8 +16,8 @@ interface ModelInfo {
16
16
  /**
17
17
  * Vendor-declared default reasoning level (Codex models.json
18
18
  * `default_reasoning_level`). When present, fresh sessions start here rather
19
- * than at the ceiling: the deep-reasoning flagships (Astra ships "low",
20
- * GPT-6 Sol/Luna "medium") think dramatically longer per rung, so defaulting to
19
+ * than at the ceiling: the deep-reasoning models (Astra and GPT-6.1 Sol ship
20
+ * "low", GPT-6 Luna "medium") think dramatically longer per rung, so defaulting to
21
21
  * `maxThinkingLevel` made new sessions pathologically slow.
22
22
  */
23
23
  defaultThinkingLevel?: ThinkingLevel;
@@ -39,7 +39,7 @@ interface ModelInfo {
39
39
  /**
40
40
  * The top reasoning tier this model genuinely uses. Used when thinking is
41
41
  * enabled to pick the strongest setting per model:
42
- * - OpenAI GPT-6 Astra / Sol: `ultra` (Codex orchestration preset above `max`)
42
+ * - OpenAI GPT-6 Astra / GPT-6.1 Sol: `ultra` (Codex orchestration preset above `max`)
43
43
  * - OpenAI GPT-6 Luna: `max`
44
44
  * - OpenAI Pro/Codex/old: clamped to what the model accepts
45
45
  * - Claude Fable 5.1 / Fable 5 / Mythos 5, Opus 5.5 / Opus 5 and Sonnet 5.5: `max`
@@ -134,25 +134,17 @@ declare function getDefaultThinkingLevel(modelId: string, options?: {
134
134
  baseUrl?: string;
135
135
  }): ThinkingLevel;
136
136
  /**
137
- * Get the model to use for compaction summarization.
138
- * - Anthropic: always Sonnet 5.5
139
- * - OpenAI: cheapest (Codex Mini)
140
- * - Gemini: use the current model
141
- * - GLM: GLM-5.3-Flash (the registered low-cost sibling)
142
- * - Moonshot: use the current model (no cheap alternative registered)
143
- */
144
- declare function getSummaryModel(provider: Provider, currentModelId: string): ModelInfo;
145
- /**
146
- * Fastest/cheapest sibling within the SAME provider, for scout-style read-only
147
- * sub-agents (recon, research) where a low-latency model is enough and the
148
- * frontier model is wasted spend + latency.
137
+ * Get the model to use for compaction summarization — the ONLY place a
138
+ * cheaper model is used. Everything else (sub-agents, reviewers, judges)
139
+ * runs on the parent's model.
149
140
  *
150
- * Routes off each model's `costTier` — the single source of truth that already
151
- * travels with the registry entry — so a model rename/bump needs no change
152
- * here. Providers with no low-tier sibling (Moonshot, MiniMax, Sakana,
153
- * OpenRouter) gracefully keep the parent model, so there's never a
154
- * crash or a cross-provider jump to a login the user may not have.
141
+ * Picks the provider's first model tagged with its summary tier
142
+ * ({@link SUMMARY_COST_TIER}). When none is registered — the tier was removed,
143
+ * or the provider has no cheap sibling — it summarizes on the current model.
144
+ * Never throws. The caller still retries on the current model when the
145
+ * provider rejects the chosen one at runtime (retired upstream, or not on the
146
+ * user's plan), since the registry cannot know that.
155
147
  */
156
- declare function getFastModel(provider: Provider, currentModelId: string): ModelInfo;
148
+ declare function getSummaryModel(provider: Provider, currentModelId: string): ModelInfo;
157
149
 
158
- export { type ContextWindowOptions, DEFAULT_MAX_VIDEO_BYTES, MODELS, type ModelInfo, clearRuntimeModels, getAllModels, getAuthStorageKey, getAuthStorageKeys, getContextWindow, getDefaultModel, getDefaultThinkingLevel, getFastModel, getMaxThinkingLevel, getModel, getModelsForProvider, getSummaryModel, getToolResultCharLimit, getVideoByteLimit, registerRuntimeModels, usesOpenAICodexTransport };
150
+ export { type ContextWindowOptions, DEFAULT_MAX_VIDEO_BYTES, MODELS, type ModelInfo, clearRuntimeModels, getAllModels, getAuthStorageKey, getAuthStorageKeys, getContextWindow, getDefaultModel, getDefaultThinkingLevel, getMaxThinkingLevel, getModel, getModelsForProvider, getSummaryModel, getToolResultCharLimit, getVideoByteLimit, registerRuntimeModels, usesOpenAICodexTransport };
@@ -8,7 +8,6 @@ import {
8
8
  getContextWindow,
9
9
  getDefaultModel,
10
10
  getDefaultThinkingLevel,
11
- getFastModel,
12
11
  getMaxThinkingLevel,
13
12
  getModel,
14
13
  getModelsForProvider,
@@ -17,7 +16,7 @@ import {
17
16
  getVideoByteLimit,
18
17
  registerRuntimeModels,
19
18
  usesOpenAICodexTransport
20
- } from "./chunk-ZCXAIMZF.js";
19
+ } from "./chunk-N73LFQ3N.js";
21
20
  import "./chunk-CRU3SSNX.js";
22
21
  export {
23
22
  DEFAULT_MAX_VIDEO_BYTES,
@@ -29,7 +28,6 @@ export {
29
28
  getContextWindow,
30
29
  getDefaultModel,
31
30
  getDefaultThinkingLevel,
32
- getFastModel,
33
31
  getMaxThinkingLevel,
34
32
  getModel,
35
33
  getModelsForProvider,
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@prestyj/core",
3
- "version": "5.28.0",
3
+ "version": "5.29.0",
4
4
  "type": "module",
5
5
  "description": "Provider-agnostic, UI-free shared foundation: model registry, auth, paths, telegram, voice",
6
6
  "license": "MIT",
@@ -33,7 +33,7 @@
33
33
  "dist"
34
34
  ],
35
35
  "dependencies": {
36
- "@prestyj/ai": "5.28.0"
36
+ "@prestyj/ai": "5.29.0"
37
37
  },
38
38
  "optionalDependencies": {
39
39
  "@huggingface/transformers": "4.2.0",