rovecode 0.3.2
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/LICENSE +662 -0
- package/README.md +737 -0
- package/THIRD_PARTY_NOTICES.md +268 -0
- package/bin/rovecode.ts +21 -0
- package/package.json +56 -0
- package/src/acp/server.ts +374 -0
- package/src/cli/auth-login.ts +122 -0
- package/src/cli/connect.ts +244 -0
- package/src/cli/context-cmd.ts +199 -0
- package/src/cli/dispatch.ts +82 -0
- package/src/cli/doctor.ts +362 -0
- package/src/cli/export.ts +276 -0
- package/src/cli/help.ts +293 -0
- package/src/cli/is-tui-invocation.ts +8 -0
- package/src/cli/main.ts +583 -0
- package/src/cli/market-cmd.ts +658 -0
- package/src/cli/mcp-login.ts +141 -0
- package/src/cli/mcp-market-cmd.ts +302 -0
- package/src/cli/output.ts +382 -0
- package/src/cli/repl.ts +250 -0
- package/src/cli/repomap-root.ts +14 -0
- package/src/cli/resume.ts +57 -0
- package/src/cli/run-flags.ts +43 -0
- package/src/cli/run-limits.ts +78 -0
- package/src/cli/runtime.ts +931 -0
- package/src/cli/session-arg.ts +30 -0
- package/src/cli/sessions-cmd.ts +145 -0
- package/src/cli/setup.ts +153 -0
- package/src/cli/skills-cmd.ts +194 -0
- package/src/cli/start-chat.ts +65 -0
- package/src/cli/trust-cmd.ts +52 -0
- package/src/coding/bash.ts +148 -0
- package/src/coding/checkpoints.ts +327 -0
- package/src/coding/diff.ts +138 -0
- package/src/coding/files.ts +341 -0
- package/src/coding/hashline.ts +274 -0
- package/src/coding/lsp-gate.ts +254 -0
- package/src/coding/lsp-servers.ts +147 -0
- package/src/coding/lsp.ts +283 -0
- package/src/coding/repomap-cache.ts +99 -0
- package/src/coding/repomap-files.ts +192 -0
- package/src/coding/repomap.ts +481 -0
- package/src/core/agents.ts +255 -0
- package/src/core/compaction.ts +259 -0
- package/src/core/config.ts +289 -0
- package/src/core/context-report.ts +228 -0
- package/src/core/context.ts +60 -0
- package/src/core/count-remote.ts +107 -0
- package/src/core/execpolicy-rules.ts +196 -0
- package/src/core/execpolicy.ts +385 -0
- package/src/core/executor.ts +454 -0
- package/src/core/guardrails.ts +400 -0
- package/src/core/hooks.ts +411 -0
- package/src/core/images.ts +230 -0
- package/src/core/intro.ts +266 -0
- package/src/core/loop.ts +567 -0
- package/src/core/modes.ts +372 -0
- package/src/core/orchestrator.ts +245 -0
- package/src/core/proc-group.ts +48 -0
- package/src/core/project-trust.ts +98 -0
- package/src/core/reflection.ts +165 -0
- package/src/core/sandbox-config.ts +186 -0
- package/src/core/session-id.ts +24 -0
- package/src/core/session-images.ts +73 -0
- package/src/core/session-ops.ts +183 -0
- package/src/core/session-text.ts +29 -0
- package/src/core/session.ts +469 -0
- package/src/core/settings.ts +170 -0
- package/src/core/tasks.ts +646 -0
- package/src/core/token-scale.ts +108 -0
- package/src/core/tools.ts +309 -0
- package/src/core/trust.ts +104 -0
- package/src/core/types.ts +330 -0
- package/src/core/update-check.ts +171 -0
- package/src/core/usage.ts +204 -0
- package/src/core/validate.ts +121 -0
- package/src/core/verify-gate.ts +159 -0
- package/src/core/verify.ts +236 -0
- package/src/core/voice.ts +158 -0
- package/src/core/win-job.ts +183 -0
- package/src/core/workspace.ts +184 -0
- package/src/design/audit.ts +797 -0
- package/src/design/direction.ts +190 -0
- package/src/design/rules.ts +157 -0
- package/src/eval/bench.ts +150 -0
- package/src/eval/gauntlet-runner.ts +215 -0
- package/src/eval/gauntlet-support.ts +84 -0
- package/src/eval/gauntlet-wave3.ts +269 -0
- package/src/eval/gauntlet-wave4.ts +217 -0
- package/src/eval/gauntlet.ts +253 -0
- package/src/index.ts +17 -0
- package/src/lanes/agy.ts +95 -0
- package/src/lanes/approval.ts +24 -0
- package/src/lanes/claude.ts +129 -0
- package/src/lanes/codex.ts +127 -0
- package/src/lanes/events.ts +130 -0
- package/src/lanes/job.ts +142 -0
- package/src/lanes/opencode.ts +122 -0
- package/src/lanes/process.ts +184 -0
- package/src/lanes/progress.ts +183 -0
- package/src/lanes/registry.ts +178 -0
- package/src/lanes/runner.ts +124 -0
- package/src/lanes/types.ts +112 -0
- package/src/market/catalogs/mcp-docs.json +111 -0
- package/src/market/catalogs/plugins.json +111 -0
- package/src/market/catalogs/skills.json +478 -0
- package/src/market/clone.ts +72 -0
- package/src/market/context-cost.ts +121 -0
- package/src/market/digest.ts +106 -0
- package/src/market/index.ts +22 -0
- package/src/market/install.ts +578 -0
- package/src/market/manifest.ts +187 -0
- package/src/market/prereq.ts +145 -0
- package/src/market/registry.ts +363 -0
- package/src/market/resolve.ts +111 -0
- package/src/market/types.ts +236 -0
- package/src/market/validate.ts +227 -0
- package/src/mcp/client.ts +449 -0
- package/src/mcp/config.ts +252 -0
- package/src/mcp/local-package.ts +211 -0
- package/src/mcp/market-catalog.ts +84 -0
- package/src/mcp/market-install.ts +289 -0
- package/src/mcp/market.ts +362 -0
- package/src/mcp/oauth.ts +251 -0
- package/src/mcp/prompts-resources.ts +249 -0
- package/src/mcp/shared.ts +149 -0
- package/src/mcp/status.ts +67 -0
- package/src/mcp/tools.ts +275 -0
- package/src/mcp/transport.ts +122 -0
- package/src/mcp/trust.ts +25 -0
- package/src/memory/blocks.ts +278 -0
- package/src/memory/recall.ts +355 -0
- package/src/memory/scope.ts +182 -0
- package/src/memory/store.ts +105 -0
- package/src/memory/tools.ts +99 -0
- package/src/plugins/cli.ts +119 -0
- package/src/plugins/discover.ts +108 -0
- package/src/plugins/index.ts +50 -0
- package/src/plugins/install.ts +184 -0
- package/src/plugins/load.ts +124 -0
- package/src/plugins/manifest.ts +92 -0
- package/src/plugins/state.ts +83 -0
- package/src/providers/auth.ts +408 -0
- package/src/providers/cache.ts +223 -0
- package/src/providers/catalog-local.ts +160 -0
- package/src/providers/catalog.ts +421 -0
- package/src/providers/middleware-context.ts +86 -0
- package/src/providers/middleware.ts +373 -0
- package/src/providers/model-list.ts +23 -0
- package/src/providers/models-index.json +1 -0
- package/src/providers/oauth/common.ts +105 -0
- package/src/providers/oauth/device-code.ts +107 -0
- package/src/providers/oauth/github-copilot.ts +146 -0
- package/src/providers/oauth/loopback.ts +158 -0
- package/src/providers/oauth/openai.ts +163 -0
- package/src/providers/oauth/openrouter.ts +89 -0
- package/src/providers/oauth/pkce.ts +45 -0
- package/src/providers/oauth/registry.ts +39 -0
- package/src/providers/oauth/seam.ts +89 -0
- package/src/providers/profile-glm53.ts +111 -0
- package/src/providers/profile-sonnet5-persona.ts +65 -0
- package/src/providers/profile-sonnet5-voice.ts +23 -0
- package/src/providers/profiles.ts +156 -0
- package/src/providers/provider-config.ts +311 -0
- package/src/providers/registry.ts +333 -0
- package/src/providers/responses.ts +209 -0
- package/src/providers/retry.ts +234 -0
- package/src/providers/router.ts +294 -0
- package/src/providers/sse.ts +26 -0
- package/src/providers/stream-errors.ts +117 -0
- package/src/providers/stream.ts +566 -0
- package/src/providers/thinking.ts +189 -0
- package/src/providers/wire-messages.ts +129 -0
- package/src/providers/wire-responses.ts +79 -0
- package/src/providers/wire-select.ts +53 -0
- package/src/server/http.ts +291 -0
- package/src/server/openapi.ts +246 -0
- package/src/sextant/card-hits.ts +102 -0
- package/src/sextant/card-keys.ts +55 -0
- package/src/sextant/context-source.ts +157 -0
- package/src/sextant/crew-cards.ts +350 -0
- package/src/sextant/draw-agents.ts +273 -0
- package/src/sextant/draw-code.ts +388 -0
- package/src/sextant/draw-context.ts +222 -0
- package/src/sextant/draw-frame.ts +164 -0
- package/src/sextant/draw-market.ts +573 -0
- package/src/sextant/draw-messages.ts +386 -0
- package/src/sextant/draw-pet.ts +230 -0
- package/src/sextant/draw-plan.ts +187 -0
- package/src/sextant/draw-tabs.ts +85 -0
- package/src/sextant/draw-util.ts +65 -0
- package/src/sextant/draw-wizard.ts +378 -0
- package/src/sextant/engine.ts +230 -0
- package/src/sextant/frame-hits.ts +25 -0
- package/src/sextant/frame.ts +101 -0
- package/src/sextant/git-status.ts +197 -0
- package/src/sextant/grid.ts +59 -0
- package/src/sextant/input.ts +119 -0
- package/src/sextant/keys.ts +521 -0
- package/src/sextant/layout.ts +86 -0
- package/src/sextant/local-commands.ts +169 -0
- package/src/sextant/market-source.ts +287 -0
- package/src/sextant/mentions.ts +200 -0
- package/src/sextant/message-hits.ts +26 -0
- package/src/sextant/model.ts +387 -0
- package/src/sextant/overlays.ts +456 -0
- package/src/sextant/panel-hits.ts +38 -0
- package/src/sextant/pet.ts +399 -0
- package/src/sextant/screen.ts +324 -0
- package/src/sextant/scroll-hits.ts +66 -0
- package/src/sextant/scrollbar.ts +82 -0
- package/src/sextant/sextant-bridge.ts +174 -0
- package/src/sextant/sextant-cards.ts +142 -0
- package/src/sextant/sextant-diff-base.ts +63 -0
- package/src/sextant/sextant-files.ts +154 -0
- package/src/sextant/sextant-frame-loop.ts +335 -0
- package/src/sextant/sextant-renderer.ts +574 -0
- package/src/sextant/sextant-repo.ts +140 -0
- package/src/sextant/theme.ts +66 -0
- package/src/sextant/tool-rows.ts +189 -0
- package/src/sextant/types.ts +493 -0
- package/src/skills/index.ts +387 -0
- package/src/skills/pack.ts +220 -0
- package/src/skills/spec.ts +162 -0
- package/src/skills/tools.ts +69 -0
- package/src/skills/versioned.ts +227 -0
- package/src/telemetry/otel-export.ts +122 -0
- package/src/telemetry/otel-lanes.ts +89 -0
- package/src/telemetry/otel-logs.ts +131 -0
- package/src/telemetry/otel-metrics.ts +136 -0
- package/src/telemetry/otel.ts +397 -0
- package/src/telemetry/otlp.ts +76 -0
- package/src/tools/ask-user.ts +156 -0
- package/src/tools/bash-bg.ts +94 -0
- package/src/tools/bash-jobs.ts +237 -0
- package/src/tools/design.ts +151 -0
- package/src/tools/evalcell.ts +338 -0
- package/src/tools/html-text.ts +139 -0
- package/src/tools/provider.ts +149 -0
- package/src/tools/task.ts +250 -0
- package/src/tools/todo.ts +320 -0
- package/src/tools/webfetch.ts +332 -0
- package/src/tools/websearch.ts +359 -0
- package/src/tui/agents-cmd.ts +41 -0
- package/src/tui/app.ts +749 -0
- package/src/tui/attach.ts +127 -0
- package/src/tui/boot-notes.ts +41 -0
- package/src/tui/builtin-prompts.ts +59 -0
- package/src/tui/checkpoints-cmd.ts +70 -0
- package/src/tui/clipboard-image.ts +81 -0
- package/src/tui/clipboard.ts +78 -0
- package/src/tui/commands.ts +283 -0
- package/src/tui/config-view.ts +53 -0
- package/src/tui/context-cmds.ts +282 -0
- package/src/tui/cost.ts +108 -0
- package/src/tui/crash-guard.ts +173 -0
- package/src/tui/focus-terminal.ts +34 -0
- package/src/tui/git-cmds.ts +273 -0
- package/src/tui/git-plain.ts +58 -0
- package/src/tui/info-cmd.ts +150 -0
- package/src/tui/input-plain.ts +76 -0
- package/src/tui/mcp-cmd.ts +128 -0
- package/src/tui/memory-note.ts +77 -0
- package/src/tui/modes-cmd.ts +45 -0
- package/src/tui/notify-seq.ts +100 -0
- package/src/tui/notify.ts +318 -0
- package/src/tui/overlays.ts +97 -0
- package/src/tui/pi-renderer.ts +428 -0
- package/src/tui/providers-cmd.ts +377 -0
- package/src/tui/reasoning-view.ts +56 -0
- package/src/tui/renderer.ts +128 -0
- package/src/tui/replay-marker.ts +29 -0
- package/src/tui/session-cmd.ts +148 -0
- package/src/tui/session-manage.ts +95 -0
- package/src/tui/sextant-attach.ts +102 -0
- package/src/tui/sextant-io.ts +202 -0
- package/src/tui/sextant-smoke.ts +110 -0
- package/src/tui/shell-cmd.ts +158 -0
- package/src/tui/smoke.ts +72 -0
- package/src/tui/staged-terminal.ts +50 -0
- package/src/tui/startup.ts +12 -0
- package/src/tui/theme.ts +59 -0
- package/src/tui/todo-label.ts +7 -0
- package/src/tui/trust-card.ts +107 -0
- package/src/tui/tui-commands.ts +87 -0
- package/tsconfig.json +30 -0
- package/vendor/pi-tui/LICENSE +21 -0
- package/vendor/pi-tui/PATCHES.md +12 -0
- package/vendor/pi-tui/PROVENANCE.md +12 -0
- package/vendor/pi-tui/README.upstream.md +854 -0
- package/vendor/pi-tui/native/win32/prebuilds/win32-arm64/win32-console-mode.node +0 -0
- package/vendor/pi-tui/native/win32/prebuilds/win32-x64/win32-console-mode.node +0 -0
- package/vendor/pi-tui/src/alt-screen-search.ts +158 -0
- package/vendor/pi-tui/src/autocomplete.ts +827 -0
- package/vendor/pi-tui/src/components/alt-screen-flash.ts +52 -0
- package/vendor/pi-tui/src/components/box.ts +138 -0
- package/vendor/pi-tui/src/components/cancellable-loader.ts +41 -0
- package/vendor/pi-tui/src/components/editor.ts +2364 -0
- package/vendor/pi-tui/src/components/h-stack.ts +45 -0
- package/vendor/pi-tui/src/components/image.ts +128 -0
- package/vendor/pi-tui/src/components/input.ts +448 -0
- package/vendor/pi-tui/src/components/loader.ts +93 -0
- package/vendor/pi-tui/src/components/markdown.ts +1016 -0
- package/vendor/pi-tui/src/components/scroll-view.ts +217 -0
- package/vendor/pi-tui/src/components/select-list.ts +230 -0
- package/vendor/pi-tui/src/components/settings-list.ts +277 -0
- package/vendor/pi-tui/src/components/spacer.ts +29 -0
- package/vendor/pi-tui/src/components/stack.ts +155 -0
- package/vendor/pi-tui/src/components/text.ts +108 -0
- package/vendor/pi-tui/src/components/truncated-text.ts +66 -0
- package/vendor/pi-tui/src/components/v-stack.ts +34 -0
- package/vendor/pi-tui/src/editor-component.ts +75 -0
- package/vendor/pi-tui/src/fuzzy.ts +138 -0
- package/vendor/pi-tui/src/index.ts +149 -0
- package/vendor/pi-tui/src/keybindings.ts +321 -0
- package/vendor/pi-tui/src/keys.ts +1402 -0
- package/vendor/pi-tui/src/kill-ring.ts +47 -0
- package/vendor/pi-tui/src/latex.ts +1381 -0
- package/vendor/pi-tui/src/layout-node.ts +52 -0
- package/vendor/pi-tui/src/layout.ts +411 -0
- package/vendor/pi-tui/src/native-modifiers.ts +60 -0
- package/vendor/pi-tui/src/native-module-path.ts +32 -0
- package/vendor/pi-tui/src/stdin-buffer.ts +445 -0
- package/vendor/pi-tui/src/terminal-colors.ts +74 -0
- package/vendor/pi-tui/src/terminal-image.ts +701 -0
- package/vendor/pi-tui/src/terminal.ts +554 -0
- package/vendor/pi-tui/src/tui-alt-screen.ts +1379 -0
- package/vendor/pi-tui/src/tui-main-screen.ts +655 -0
- package/vendor/pi-tui/src/tui.ts +1264 -0
- package/vendor/pi-tui/src/undo-stack.ts +29 -0
- package/vendor/pi-tui/src/utils.ts +1327 -0
- package/vendor/pi-tui/src/word-navigation.ts +118 -0
- package/vendor/pi-tui/test/test-themes.ts +39 -0
- package/vendor/pi-tui/test/virtual-terminal.ts +219 -0
|
@@ -0,0 +1,234 @@
|
|
|
1
|
+
/** Same-model retry with exponential backoff (port #23). A StreamFn wrapper — NO second agent
|
|
2
|
+
* loop (ADR-003): bounded re-invocations of the SAME (model, messages, options) inside one
|
|
3
|
+
* stream invocation, the way the router makes one bounded pass over chain candidates. Wired
|
|
4
|
+
* INSIDE the router (`router.wrap(withRetry(stream))`, cli/runtime.ts) so retries exhaust on
|
|
5
|
+
* candidate N before the chain advances to N+1, and the router's notes stay one-per-ADVANCE.
|
|
6
|
+
*
|
|
7
|
+
* Source: gemini-cli (Apache-2.0 @0bd1d43 — research/source_snapshots/google-gemini-gemini-cli,
|
|
8
|
+
* packages/core/src/utils):
|
|
9
|
+
* - retryWithBackoff loop shape: attempt counter vs maxAttempts, delay doubling up to
|
|
10
|
+
* maxDelayMs (retry.ts:296-310, :494-501, :517-524); defaults 10 attempts / 5s initial /
|
|
11
|
+
* 30s max (retry.ts:20, :42-47).
|
|
12
|
+
* - Retryable set: 429 and 5xx; "Explicitly do not retry 400" (retry.ts:193-199); transport
|
|
13
|
+
* failures without a status are retryable (retry.ts:49-62, :174-189). Shared with the router
|
|
14
|
+
* through classifyStreamError (router.ts) — one classifier, two consumers (gemini-cli's 499
|
|
15
|
+
* is excluded there, documented deviation).
|
|
16
|
+
* - Aborts pass through untouched (retry.ts:337-340); the backoff sleep itself is abortable
|
|
17
|
+
* (delay.ts:22-48: timeout + abort listener, both cleared on settle).
|
|
18
|
+
* - A server-suggested delay is a FLOOR on the next sleep (retry.ts:472-476
|
|
19
|
+
* `Math.max(currentDelay, retryDelayMs)`); a suggestion beyond the cap is terminal — no wait,
|
|
20
|
+
* immediate fallback (googleQuotaErrors.ts:120 MAX_RETRYABLE_DELAY_SECONDS, :286-289). Here
|
|
21
|
+
* the suggestion is the HTTP Retry-After header (RFC 9110 §10.2.3: delay-seconds or
|
|
22
|
+
* HTTP-date), OpenAI's retry-after-ms, or Anthropic's ratelimit-reset timestamps, all recorded
|
|
23
|
+
* by stream-errors.ts; the cap is the per-invocation total budget and the run's own deadline.
|
|
24
|
+
* - Deviations: FULL jitter — delay = U[0,1) × min(max, base·2ⁿ) (AWS "Exponential Backoff And
|
|
25
|
+
* Jitter") instead of gemini-cli's ±30% around the current delay (retry.ts:494-495) or +20%
|
|
26
|
+
* over the server floor (:478): fewer synchronized retries, and the schedule is pinnable with
|
|
27
|
+
* an injected random. A total wall-clock cap per invocation (gemini-cli has none): with a
|
|
28
|
+
* fallback chain every candidate pays the retry budget, so it is kept short. No content-based
|
|
29
|
+
* retry (:314-328), no quota classification / fallback dialog (:342-420) — the router owns
|
|
30
|
+
* model fallback. A thrown inner stream is folded into an error turn (never-throw) but NOT
|
|
31
|
+
* retried: a seam-contract violation is a harness bug, not a provider outcome.
|
|
32
|
+
* - Per-invocation state only: attempt counter and deadline live inside one generator run;
|
|
33
|
+
* nothing is remembered across calls (unlike the router's sticky switch).
|
|
34
|
+
* - IDEMPOTENCY (2026-09-04): a retry is taken only while NOTHING has been streamed to the consumer.
|
|
35
|
+
* Once a text or reasoning delta has gone out, a failure ends the turn — with the partial text kept as
|
|
36
|
+
* the turn's parts and an error that says so — because a re-drive would print a second answer under
|
|
37
|
+
* the first (the router applies the same rule to its chain advance). Deltas from a failed attempt that
|
|
38
|
+
* streamed nothing cannot exist, so the "pass through live" rule and this one never meet.
|
|
39
|
+
*
|
|
40
|
+
* Env (retryOptionsFromEnv, read once by cli/runtime.ts):
|
|
41
|
+
* - ROVECODE_RETRY_MAX retries after the first attempt; 0 disables. Default 3 (→ 4 attempts;
|
|
42
|
+
* upstream 10 attempts — a fallback chain multiplies attempts per candidate).
|
|
43
|
+
* - ROVECODE_RETRY_BASE_MS cap of the first backoff in ms. Default 1000 (upstream 5000; full jitter
|
|
44
|
+
* halves the expected wait — ~0.5 s, 1 s, 2 s, 4 s expected).
|
|
45
|
+
* Fixed: max backoff 20 s, total budget 60 s per invocation, and never past StreamOptions.deadlineAt
|
|
46
|
+
* (the run's --max-seconds clock). */
|
|
47
|
+
|
|
48
|
+
import type { AssistantTurn, ModelRef, StreamEvent, StreamFn } from "../core/types.ts";
|
|
49
|
+
import { classifyStreamError } from "./router.ts";
|
|
50
|
+
import { httpErrorMeta, providerMessage } from "./stream-errors.ts";
|
|
51
|
+
|
|
52
|
+
export const DEFAULT_MAX_RETRIES = 3;
|
|
53
|
+
export const DEFAULT_BASE_MS = 1_000;
|
|
54
|
+
export const DEFAULT_MAX_DELAY_MS = 20_000;
|
|
55
|
+
export const DEFAULT_TOTAL_MS = 60_000;
|
|
56
|
+
|
|
57
|
+
export interface RetryNote {
|
|
58
|
+
model: ModelRef;
|
|
59
|
+
/** Attempts made so far (the one that just failed); the retry about to happen is attempt+1. */
|
|
60
|
+
attempt: number;
|
|
61
|
+
/** attempts the policy allows in total (maxRetries + 1) — for "(2/4)" */
|
|
62
|
+
maxAttempts: number;
|
|
63
|
+
delayMs: number;
|
|
64
|
+
/** The failed turn's error text, e.g. "HTTP 429: ...". */
|
|
65
|
+
reason: string;
|
|
66
|
+
/** the HTTP status when the failure was a response; undefined for a transport failure */
|
|
67
|
+
status?: number;
|
|
68
|
+
/** Parsed server wait hint (Retry-After / retry-after-ms / anthropic-ratelimit-*-reset) when the error turn carried one. */
|
|
69
|
+
retryAfterMs?: number;
|
|
70
|
+
}
|
|
71
|
+
|
|
72
|
+
/** why the wrapper stopped retrying a retryable failure */
|
|
73
|
+
export type GiveUpWhy = "attempts" | "budget" | "deadline";
|
|
74
|
+
export interface GiveUpNote extends Omit<RetryNote, "delayMs"> { why: GiveUpWhy; delayMs?: number }
|
|
75
|
+
|
|
76
|
+
export interface RetryOptions {
|
|
77
|
+
/** Retries after the first attempt (ROVECODE_RETRY_MAX). 0 = never retry. */
|
|
78
|
+
maxRetries?: number;
|
|
79
|
+
/** Cap of the first backoff, doubling per attempt (ROVECODE_RETRY_BASE_MS). */
|
|
80
|
+
baseMs?: number;
|
|
81
|
+
/** Ceiling for the doubling backoff term. */
|
|
82
|
+
maxDelayMs?: number;
|
|
83
|
+
/** Wall-clock budget per invocation incl. sleeps; a wait that would end past it is not taken. */
|
|
84
|
+
totalMs?: number;
|
|
85
|
+
/** Retry visibility — a note-style callback, not a StreamEvent (grammar is shared/untouchable). */
|
|
86
|
+
onRetry?: (note: RetryNote) => void;
|
|
87
|
+
/** the last word: a retryable failure that will not be retried (attempts, budget or the run's deadline) */
|
|
88
|
+
onGiveUp?: (note: GiveUpNote) => void;
|
|
89
|
+
/** Test seams: abortable sleep, jitter source, clock (epoch ms — also anchors HTTP-date). */
|
|
90
|
+
sleep?: (ms: number, signal?: AbortSignal) => Promise<void>;
|
|
91
|
+
random?: () => number;
|
|
92
|
+
now?: () => number;
|
|
93
|
+
}
|
|
94
|
+
|
|
95
|
+
/** Abortable sleep (delay.ts:22-48 shape) on a REF'D setTimeout — Bun unrefs AbortSignal.timeout
|
|
96
|
+
* timers, so a sleep built on one could not hold the process/test runner — plus an abort
|
|
97
|
+
* listener; whichever settles first clears the other. Resolves (never rejects) on abort: the
|
|
98
|
+
* caller re-checks signal.aborted. */
|
|
99
|
+
export function sleepMs(ms: number, signal?: AbortSignal): Promise<void> {
|
|
100
|
+
return new Promise<void>((resolve) => {
|
|
101
|
+
if (signal?.aborted) { resolve(); return; }
|
|
102
|
+
const settle = (): void => { clearTimeout(timer); signal?.removeEventListener("abort", settle); resolve(); };
|
|
103
|
+
const timer = setTimeout(settle, ms);
|
|
104
|
+
signal?.addEventListener("abort", settle, { once: true });
|
|
105
|
+
});
|
|
106
|
+
}
|
|
107
|
+
|
|
108
|
+
/** RFC 9110 §10.2.3 Retry-After: delay-seconds (non-negative; a fractional value is tolerated)
|
|
109
|
+
* or an HTTP-date, converted to ms from `now` (a past date → 0). Unparseable → undefined. */
|
|
110
|
+
export function parseRetryAfter(value: string | undefined, now: number): number | undefined {
|
|
111
|
+
const v = value?.trim();
|
|
112
|
+
if (!v) return undefined;
|
|
113
|
+
if (/^\d+(?:\.\d+)?$/.test(v)) return Math.round(Number(v) * 1000);
|
|
114
|
+
const at = Date.parse(v);
|
|
115
|
+
return Number.isNaN(at) ? undefined : Math.max(0, at - now);
|
|
116
|
+
}
|
|
117
|
+
|
|
118
|
+
/** the server's wait hint for a failed turn, whichever header it used: the largest of Retry-After,
|
|
119
|
+
* retry-after-ms and the Anthropic ratelimit reset timestamps (RFC 3339) — the longest wait is the
|
|
120
|
+
* one that will actually clear the limit */
|
|
121
|
+
export function serverWaitMs(turn: AssistantTurn, now: number): number | undefined {
|
|
122
|
+
const meta = httpErrorMeta(turn);
|
|
123
|
+
if (!meta) return undefined;
|
|
124
|
+
const hints: number[] = [];
|
|
125
|
+
const ra = parseRetryAfter(meta.retryAfter, now);
|
|
126
|
+
if (ra !== undefined) hints.push(ra);
|
|
127
|
+
if (meta.retryAfterMs !== undefined && /^\d+(?:\.\d+)?$/.test(meta.retryAfterMs.trim())) hints.push(Math.round(Number(meta.retryAfterMs)));
|
|
128
|
+
if (meta.resetAt !== undefined) { const at = Date.parse(meta.resetAt); if (!Number.isNaN(at)) hints.push(Math.max(0, at - now)); }
|
|
129
|
+
return hints.length ? Math.max(...hints) : undefined;
|
|
130
|
+
}
|
|
131
|
+
|
|
132
|
+
/** Env knobs (header). Blank/invalid/out-of-range values fall back to the defaults;
|
|
133
|
+
* ROVECODE_RETRY_MAX=0 is honored (retry off). */
|
|
134
|
+
export function retryOptionsFromEnv(env: Record<string, string | undefined> = process.env): RetryOptions {
|
|
135
|
+
return { maxRetries: envInt(env.ROVECODE_RETRY_MAX, DEFAULT_MAX_RETRIES, 0), baseMs: envInt(env.ROVECODE_RETRY_BASE_MS, DEFAULT_BASE_MS, 1) };
|
|
136
|
+
}
|
|
137
|
+
|
|
138
|
+
function envInt(raw: string | undefined, dflt: number, min: number): number {
|
|
139
|
+
if (raw === undefined || raw.trim() === "") return dflt;
|
|
140
|
+
const n = Number(raw);
|
|
141
|
+
return Number.isFinite(n) && n >= min ? Math.floor(n) : dflt;
|
|
142
|
+
}
|
|
143
|
+
|
|
144
|
+
const errorTurn = (error: string): AssistantTurn => ({ parts: [], stopReason: "error", usage: { input: 0, output: 0 }, error });
|
|
145
|
+
|
|
146
|
+
// ---------- the human's words ----------
|
|
147
|
+
|
|
148
|
+
/** one noun phrase per failure class — what the notice and the give-up line lead with */
|
|
149
|
+
export function failureWord(status: number | undefined, error: string): string {
|
|
150
|
+
if (status === 429) return "rate limited";
|
|
151
|
+
if (status === 529 || status === 503) return "overloaded";
|
|
152
|
+
if (status !== undefined && status >= 500) return `server error (HTTP ${status})`;
|
|
153
|
+
if (status !== undefined) return `HTTP ${status}`;
|
|
154
|
+
if (/^no response from /.test(error)) return "no response";
|
|
155
|
+
return "connection failed";
|
|
156
|
+
}
|
|
157
|
+
|
|
158
|
+
const fmtSeconds = (ms: number): string => { const s = ms / 1000; return `${s >= 10 ? Math.round(s) : Math.round(s * 10) / 10} s`; };
|
|
159
|
+
|
|
160
|
+
/** "anthropic: overloaded — retrying in 4 s (2/4)" */
|
|
161
|
+
export function describeRetry(n: RetryNote): string {
|
|
162
|
+
return `${n.model.provider}: ${failureWord(n.status, n.reason)} — retrying in ${fmtSeconds(n.delayMs)} (${n.attempt + 1}/${n.maxAttempts})`;
|
|
163
|
+
}
|
|
164
|
+
|
|
165
|
+
/** "anthropic: overloaded (HTTP 529) — gave up after 4 attempts: Overloaded" */
|
|
166
|
+
export function describeGiveUp(n: GiveUpNote): string {
|
|
167
|
+
const why = n.why === "attempts" ? `gave up after ${n.attempt} attempt${n.attempt === 1 ? "" : "s"}`
|
|
168
|
+
: n.why === "deadline" ? `not retried: the run's time limit is closer than the ${fmtSeconds(n.delayMs ?? 0)} wait`
|
|
169
|
+
: `not retried: the ${fmtSeconds(n.delayMs ?? 0)} wait would pass the retry budget`;
|
|
170
|
+
const status = n.status !== undefined && !/HTTP/.test(failureWord(n.status, n.reason)) ? ` (HTTP ${n.status})` : "";
|
|
171
|
+
const msg = providerMessage(n.reason);
|
|
172
|
+
return `${n.model.provider}: ${failureWord(n.status, n.reason)}${status} — ${why}${msg ? `: ${msg}` : ""}`;
|
|
173
|
+
}
|
|
174
|
+
|
|
175
|
+
/** Wrap a StreamFn: a retryable failed turn (429 / 5xx / transport, never 400, never abort, never after
|
|
176
|
+
* a delta went out) is re-driven against the SAME model with the same args after an abortable
|
|
177
|
+
* full-jitter backoff, until it succeeds, a non-retryable outcome lands, or a cap (attempts / total
|
|
178
|
+
* budget / the run's deadline) stops it — then the LAST turn is yielded untouched so the router sees
|
|
179
|
+
* the genuine provider error. Never throws (ADR-003). */
|
|
180
|
+
export function withRetry(inner: StreamFn, opts: RetryOptions = {}): StreamFn {
|
|
181
|
+
const maxRetries = opts.maxRetries ?? DEFAULT_MAX_RETRIES;
|
|
182
|
+
const baseMs = opts.baseMs ?? DEFAULT_BASE_MS;
|
|
183
|
+
const maxDelayMs = opts.maxDelayMs ?? DEFAULT_MAX_DELAY_MS;
|
|
184
|
+
const totalMs = opts.totalMs ?? DEFAULT_TOTAL_MS;
|
|
185
|
+
const sleep = opts.sleep ?? sleepMs;
|
|
186
|
+
const random = opts.random ?? Math.random;
|
|
187
|
+
const now = opts.now ?? Date.now;
|
|
188
|
+
return async function* (model, messages, options): AsyncGenerator<StreamEvent> {
|
|
189
|
+
const startedAt = now();
|
|
190
|
+
for (let attempt = 1; ; attempt++) {
|
|
191
|
+
let turn: AssistantTurn | null = null;
|
|
192
|
+
let streamed = ""; // every text delta this attempt let through — the answer the consumer has already seen
|
|
193
|
+
try {
|
|
194
|
+
for await (const ev of inner(model, messages, options)) {
|
|
195
|
+
if (ev.type === "turn") turn = ev.turn;
|
|
196
|
+
else { if (ev.type === "text_delta") streamed += ev.text; else if (ev.type === "reasoning_delta") streamed ||= " "; yield ev; } // deltas pass through live (header)
|
|
197
|
+
}
|
|
198
|
+
} catch (e) {
|
|
199
|
+
yield { type: "turn", turn: errorTurn(e instanceof Error ? e.message : String(e)) }; // folded, not retried (header)
|
|
200
|
+
return;
|
|
201
|
+
}
|
|
202
|
+
if (turn === null) { yield { type: "turn", turn: errorTurn("stream ended without a terminal turn") }; return; }
|
|
203
|
+
const meta = httpErrorMeta(turn);
|
|
204
|
+
const status = meta?.status ?? classifyStreamError(turn.error).status;
|
|
205
|
+
// retry.ts:337-340: aborts never retry; 400/4xx-non-429 and non-error turns stand as they are
|
|
206
|
+
if (turn.stopReason !== "error" || options?.signal?.aborted === true || !classifyStreamError(turn.error).retryable) {
|
|
207
|
+
yield { type: "turn", turn };
|
|
208
|
+
return;
|
|
209
|
+
}
|
|
210
|
+
if (streamed.length > 0) {
|
|
211
|
+
// idempotency (header): part of the answer is on the screen — end the turn, keep that text as the
|
|
212
|
+
// turn's parts (the loop stores it and the summary shows it above the error), say why no retry
|
|
213
|
+
const kept = streamed.trim().length > 0 && turn.parts.length === 0 ? [{ kind: "text" as const, text: streamed }] : turn.parts;
|
|
214
|
+
yield { type: "turn", turn: { ...turn, parts: kept, error: `${turn.error ?? "provider stream failed"} — the connection dropped after part of the answer had arrived; not retried, a retry would repeat it` } };
|
|
215
|
+
return;
|
|
216
|
+
}
|
|
217
|
+
const note = { model, attempt, maxAttempts: maxRetries + 1, reason: turn.error ?? "error", ...(status !== undefined ? { status } : {}) };
|
|
218
|
+
// retries off (ROVECODE_RETRY_MAX=0): nothing was ever going to be retried, so there is no giving up to announce
|
|
219
|
+
if (attempt > maxRetries) { if (maxRetries > 0) opts.onGiveUp?.({ ...note, why: "attempts" }); yield { type: "turn", turn }; return; }
|
|
220
|
+
const retryAfterMs = serverWaitMs(turn, now());
|
|
221
|
+
const cap = Math.min(maxDelayMs, baseMs * 2 ** (attempt - 1));
|
|
222
|
+
const delayMs = Math.max(Math.round(random() * cap), retryAfterMs ?? 0); // full jitter, server floor (retry.ts:476)
|
|
223
|
+
const withHint = retryAfterMs !== undefined ? { retryAfterMs } : {};
|
|
224
|
+
// budget: a wait that would end past the deadline is not taken — the failure surfaces now
|
|
225
|
+
// and the router may advance at once (googleQuotaErrors.ts:286-289 shape)
|
|
226
|
+
if (now() - startedAt + delayMs > totalMs) { opts.onGiveUp?.({ ...note, ...withHint, delayMs, why: "budget" }); yield { type: "turn", turn }; return; }
|
|
227
|
+
// the run's own clock (--max-seconds): a wait that ends past it would only be cut off at the turn boundary
|
|
228
|
+
if (options?.deadlineAt !== undefined && now() + delayMs > options.deadlineAt) { opts.onGiveUp?.({ ...note, ...withHint, delayMs, why: "deadline" }); yield { type: "turn", turn }; return; }
|
|
229
|
+
opts.onRetry?.({ ...note, delayMs, ...withHint });
|
|
230
|
+
await sleep(delayMs, options?.signal);
|
|
231
|
+
if (options?.signal?.aborted) { yield { type: "turn", turn }; return; } // abort landed during backoff: last turn stands
|
|
232
|
+
}
|
|
233
|
+
};
|
|
234
|
+
}
|
|
@@ -0,0 +1,294 @@
|
|
|
1
|
+
/** Role router + provider fallback chains (port #14). ModelRef resolver + StreamFn wrapper —
|
|
2
|
+
* NO second agent loop (ADR-003): the wrapper makes one bounded pass over a chain's remaining
|
|
3
|
+
* candidates inside a single stream invocation; same-model retry/backoff is the port-#23
|
|
4
|
+
* withRetry wrapper, which sits INSIDE this wrap (retries exhaust before the chain advances).
|
|
5
|
+
*
|
|
6
|
+
* Role table (oh-my-pi, MIT — research/source_snapshots/can1357-oh-my-pi):
|
|
7
|
+
* - Role names ported as the documented SUBSET default/smol/plan/commit/task of OMP's
|
|
8
|
+
* ModelRole union (packages/coding-agent/src/config/model-roles.ts:22-32; full set adds
|
|
9
|
+
* slow/vision/designer/tiny/advisor). Config shape mirrors OMP's `modelRoles` record
|
|
10
|
+
* (config/settings-schema.ts:668, :6315 — Record<role, selector>).
|
|
11
|
+
* - Missing/unknown role → the "default" role (config/model-resolver.ts:1017
|
|
12
|
+
* DEFAULT_MODEL_ROLE; configured-value-else-defaults resolution :1155-1189).
|
|
13
|
+
* - Precedence: explicit request selector > configured role chain > default chain, per
|
|
14
|
+
* OMP's resolveEffectiveAgentModelSelection (model-resolver.ts:1229-1261
|
|
15
|
+
* requestModel > settingsOverride > agentModel > default).
|
|
16
|
+
* - Selector grammar: "provider/model" split on the FIRST slash so model ids keep their own
|
|
17
|
+
* slashes (model-resolver.ts:214-216); comma-separated ordered candidate lists
|
|
18
|
+
* (model-resolver.ts:1051-1055 normalizeModelPatternList); per-role ordered chains as in
|
|
19
|
+
* OMP's priority.json:2-23 (rolePriorityDefaults, model-resolver.ts:1110-1113).
|
|
20
|
+
* - NOT ported: OMP's fuzzy/glob matching, thinking-level suffixes, custom role aliases and
|
|
21
|
+
* alias cycle guard (rovecode chains are flat ModelRefs — no aliases, so no cycles).
|
|
22
|
+
*
|
|
23
|
+
* Fallback chains (gemini-cli, Apache-2.0 @0bd1d43 — research/source_snapshots/
|
|
24
|
+
* google-gemini-gemini-cli, packages/core/src):
|
|
25
|
+
* - Ordered chain, "the first model in the chain is the primary model"
|
|
26
|
+
* (availability/modelPolicy.ts:52-56 ModelPolicyChain).
|
|
27
|
+
* - Trigger conditions ported from utils/retry.ts isRetryableError (:170-209): advance on
|
|
28
|
+
* HTTP 429 or 5xx; NEVER on 400 (:193-194 "Explicitly do not retry 400"); non-HTTP
|
|
29
|
+
* transport failures (fetch failed / network codes / incomplete stream JSON) are
|
|
30
|
+
* retryable (:49-62, :122-123, :141-147, :180-189). gemini-cli additionally retries 499;
|
|
31
|
+
* the bar pins 429/5xx, so 499 is deliberately NOT retryable here (documented deviation).
|
|
32
|
+
* - Aborts never advance the chain (retry.ts:337-339 rethrows AbortError untouched).
|
|
33
|
+
* - On failure the handler picks the FIRST AVAILABLE later candidate and never falls back
|
|
34
|
+
* to the failed model itself (fallback/handler.ts:59-76); the switch is STICKY for the
|
|
35
|
+
* session via activateFallbackMode (handler.ts:163-169). Sticky here = per wrapped
|
|
36
|
+
* StreamFn, reset when a chain exhausts (rovecode simplification: gemini-cli instead tracks
|
|
37
|
+
* per-model health and marks models healthy again on success, retry.ts:330-334).
|
|
38
|
+
* - gemini-cli retries the new model immediately (retry.ts:404, :459 `attempt = 0; continue`)
|
|
39
|
+
* inside its retryWithBackoff loop; rovecode tries each candidate ONCE per invocation — no
|
|
40
|
+
* delays, no attempt reset (ADR-003: no second loop). Status is read from the seam's
|
|
41
|
+
* error TEXT ("HTTP <status>: <body>" — built by src/providers/stream-errors.ts httpErrorTurn), the same
|
|
42
|
+
* message-sniffing fallback gemini-cli itself uses (retry.ts:553-558).
|
|
43
|
+
* - Mid-stream failure: gemini-cli re-streams and signals the consumer with a RETRY event
|
|
44
|
+
* (core/geminiChat.ts:655-679). Rovecode's StreamEvent grammar (core/types.ts:47-50) has no
|
|
45
|
+
* such variant and is shared/untouchable, so forwarded text_delta events from a failed
|
|
46
|
+
* attempt simply stand; canonical content is the terminal turn only (core/loop.ts
|
|
47
|
+
* collectTurn:219-226), so the final message is never corrupted. The "note" on each
|
|
48
|
+
* advance is therefore an onNote CALLBACK, not a StreamEvent.
|
|
49
|
+
* - Exhausted chain → terminal turn with stopReason "error"; this wrapper NEVER throws
|
|
50
|
+
* (ADR-003 seam contract), even when the wrapped stream does.
|
|
51
|
+
* - A SINGLE-candidate chain has nothing to advance to: a retryable failure is yielded
|
|
52
|
+
* UNTOUCHED — no exhausted-rewrite, no note — so plain provider errors survive verbatim
|
|
53
|
+
* (R2 #14 LOW/MED-4).
|
|
54
|
+
* - A requested model outside every configured chain is a passthrough singleton, unless
|
|
55
|
+
* `looseFallback` is set (runtime sets it when chains are explicit user config): then the
|
|
56
|
+
* request is PREPENDED to the default chain — gemini-cli's shape, where "the first model
|
|
57
|
+
* in the chain is the primary model" is whatever was requested and configured fallbacks
|
|
58
|
+
* follow (R2 #14 MED-3).
|
|
59
|
+
* - The candidate that produced each terminal turn is recorded per turn object (servedBy)
|
|
60
|
+
* so the loop stamps Message.origin with the model that SERVED, not the one it asked
|
|
61
|
+
* for (R2 #14 HIGH-2). */
|
|
62
|
+
|
|
63
|
+
import type { AssistantTurn, ModelRef, StreamFn, StreamEvent, TokenUsage } from "../core/types.ts";
|
|
64
|
+
|
|
65
|
+
// ---------- roles (OMP subset, documented above) ----------
|
|
66
|
+
|
|
67
|
+
export type ModelRole = "default" | "smol" | "plan" | "commit" | "task";
|
|
68
|
+
|
|
69
|
+
export const MODEL_ROLES: readonly ModelRole[] = ["default", "smol", "plan", "commit", "task"];
|
|
70
|
+
|
|
71
|
+
export interface RouterNote {
|
|
72
|
+
/** Chain identity: role name when the model belongs to a configured role chain,
|
|
73
|
+
* else "provider/model" of the loose (passthrough) model. */
|
|
74
|
+
chain: string;
|
|
75
|
+
from: ModelRef;
|
|
76
|
+
/** Next candidate, or null when this failure exhausted the chain. */
|
|
77
|
+
to: ModelRef | null;
|
|
78
|
+
/** The failed turn's error text, e.g. "HTTP 429: ...". */
|
|
79
|
+
reason: string;
|
|
80
|
+
}
|
|
81
|
+
|
|
82
|
+
export interface RouterConfig {
|
|
83
|
+
/** Role → primary ModelRef or ordered fallback chain (first = primary,
|
|
84
|
+
* gemini-cli modelPolicy.ts:52-56). `default` is required; other roles optional. */
|
|
85
|
+
roles: Partial<Record<ModelRole, ModelRef | readonly ModelRef[]>> & { default: ModelRef | readonly ModelRef[] };
|
|
86
|
+
/** Keep fallback switches for later calls on the same wrapped stream (gemini-cli
|
|
87
|
+
* activateFallbackMode, handler.ts:163-169). Default true; exhaustion resets. */
|
|
88
|
+
sticky?: boolean;
|
|
89
|
+
/** Treat the default chain as the fallback pool for models outside EVERY configured
|
|
90
|
+
* chain: the requested model is prepended as the primary (header: MED-3). Default
|
|
91
|
+
* false — a synthesized single-model default (no explicit chain config) must not
|
|
92
|
+
* drag loose models onto a placeholder ref. */
|
|
93
|
+
looseFallback?: boolean;
|
|
94
|
+
/** Advance notification — see header for why this is a callback, not a StreamEvent. */
|
|
95
|
+
onNote?: (note: RouterNote) => void;
|
|
96
|
+
}
|
|
97
|
+
|
|
98
|
+
export interface Router {
|
|
99
|
+
/** Full ordered chain for a role; unknown/unconfigured roles get the default chain. */
|
|
100
|
+
chain(role: string): readonly ModelRef[];
|
|
101
|
+
/** Primary ModelRef for a role. `explicit` (request selector or ModelRef) wins over the
|
|
102
|
+
* role table — OMP request>config precedence (model-resolver.ts:1229-1261). */
|
|
103
|
+
resolve(role: string, explicit?: string | ModelRef): ModelRef;
|
|
104
|
+
/** Wrap a StreamFn with chain-advance-on-failure. Non-turn events pass through live. */
|
|
105
|
+
wrap(stream: StreamFn): StreamFn;
|
|
106
|
+
}
|
|
107
|
+
|
|
108
|
+
// ---------- selector parsing (OMP model-resolver.ts:214-216, :1051-1055) ----------
|
|
109
|
+
|
|
110
|
+
/** "provider/model" split on the FIRST slash (model ids keep their own slashes, e.g.
|
|
111
|
+
* "kaesra/zai-org/glm-5.3-flash"); no slash → model on `defaultProvider`. */
|
|
112
|
+
export function parseModelRef(selector: string, defaultProvider: string): ModelRef {
|
|
113
|
+
const s = selector.trim();
|
|
114
|
+
const slash = s.indexOf("/");
|
|
115
|
+
if (slash <= 0) return { provider: defaultProvider, model: s };
|
|
116
|
+
return { provider: s.slice(0, slash), model: s.slice(slash + 1) };
|
|
117
|
+
}
|
|
118
|
+
|
|
119
|
+
/** Comma-separated ordered candidate list → ModelRef chain. */
|
|
120
|
+
export function parseModelChain(selectors: string, defaultProvider: string): ModelRef[] {
|
|
121
|
+
return selectors
|
|
122
|
+
.split(",")
|
|
123
|
+
.map((s) => s.trim())
|
|
124
|
+
.filter(Boolean)
|
|
125
|
+
.map((s) => parseModelRef(s, defaultProvider));
|
|
126
|
+
}
|
|
127
|
+
|
|
128
|
+
const ROLE_ENV: Record<ModelRole, string> = {
|
|
129
|
+
default: "ROVECODE_MODEL_DEFAULT",
|
|
130
|
+
smol: "ROVECODE_MODEL_SMOL",
|
|
131
|
+
plan: "ROVECODE_MODEL_PLAN",
|
|
132
|
+
commit: "ROVECODE_MODEL_COMMIT",
|
|
133
|
+
task: "ROVECODE_MODEL_TASK",
|
|
134
|
+
};
|
|
135
|
+
|
|
136
|
+
/** Role table from env (ROVECODE_MODEL_DEFAULT/SMOL/PLAN/COMMIT/TASK, each a comma-separated
|
|
137
|
+
* "provider/model" chain). Unset default → [fallback] (the provider-derived model). */
|
|
138
|
+
export function roleTableFromEnv(
|
|
139
|
+
fallback: ModelRef,
|
|
140
|
+
env: Record<string, string | undefined> = process.env,
|
|
141
|
+
): RouterConfig["roles"] {
|
|
142
|
+
const roles: Partial<Record<ModelRole, ModelRef[]>> = {};
|
|
143
|
+
for (const role of MODEL_ROLES) {
|
|
144
|
+
const raw = env[ROLE_ENV[role]];
|
|
145
|
+
if (raw && raw.trim().length > 0) {
|
|
146
|
+
const chain = parseModelChain(raw, fallback.provider);
|
|
147
|
+
if (chain.length > 0) roles[role] = chain;
|
|
148
|
+
}
|
|
149
|
+
}
|
|
150
|
+
return { ...roles, default: roles.default ?? [fallback] };
|
|
151
|
+
}
|
|
152
|
+
|
|
153
|
+
// ---------- failure classification (gemini-cli retry.ts:170-209; rovecode stream.ts:78,109,171) ----------
|
|
154
|
+
|
|
155
|
+
export interface StreamErrorClass {
|
|
156
|
+
/** HTTP status parsed from the seam's "HTTP <status>: ..." error text, if present. */
|
|
157
|
+
status?: number;
|
|
158
|
+
/** Advance-the-chain eligible: 429, 5xx, or a non-HTTP transport failure. */
|
|
159
|
+
retryable: boolean;
|
|
160
|
+
}
|
|
161
|
+
|
|
162
|
+
export function classifyStreamError(error: string | undefined): StreamErrorClass {
|
|
163
|
+
const m = /^HTTP (\d{3})\b/.exec(error ?? "");
|
|
164
|
+
if (m) {
|
|
165
|
+
const status = Number(m[1]);
|
|
166
|
+
return { status, retryable: status === 429 || (status >= 500 && status < 600) };
|
|
167
|
+
}
|
|
168
|
+
// `config: …` (providers/registry.ts dispatcher: unknown provider, missing key) — a configuration
|
|
169
|
+
// mistake neither a retry nor the next chain candidate can fix; the human can. Never retryable.
|
|
170
|
+
if ((error ?? "").startsWith("config: ")) return { retryable: false };
|
|
171
|
+
// No HTTP prefix → the fetch itself failed (network/SSL/parse) — retryable per
|
|
172
|
+
// gemini-cli retry.ts:49-62,122-123,141-147,180-189.
|
|
173
|
+
return { retryable: true };
|
|
174
|
+
}
|
|
175
|
+
|
|
176
|
+
// ---------- served-model tagging (R2 #14 HIGH-2) ----------
|
|
177
|
+
|
|
178
|
+
/** ModelRef that actually produced a terminal turn, keyed on the turn object itself — a
|
|
179
|
+
* WeakMap side-channel because the StreamEvent grammar (core/types.ts:49-52) is shared/
|
|
180
|
+
* untouchable (no new event variant, no new turn field). Per-invocation by construction:
|
|
181
|
+
* each terminal turn is a distinct object. Unwrapped streams never mark their turns, so
|
|
182
|
+
* consumers fall back to the model they asked for. */
|
|
183
|
+
const SERVED = new WeakMap<AssistantTurn, ModelRef>();
|
|
184
|
+
|
|
185
|
+
/** The chain candidate that served `turn`, when the router produced it. loop.ts stamps
|
|
186
|
+
* Message.origin with this so /cost prices the model that ANSWERED after a fallback. */
|
|
187
|
+
export function servedBy(turn: AssistantTurn): ModelRef | undefined {
|
|
188
|
+
return SERVED.get(turn);
|
|
189
|
+
}
|
|
190
|
+
|
|
191
|
+
// ---------- router ----------
|
|
192
|
+
|
|
193
|
+
const sameRef = (a: ModelRef, b: ModelRef): boolean => a.provider === b.provider && a.model === b.model;
|
|
194
|
+
const refKey = (m: ModelRef): string => `${m.provider}/${m.model}`;
|
|
195
|
+
const zeroUsage = (): TokenUsage => ({ input: 0, output: 0 });
|
|
196
|
+
const errorTurn = (error: string): AssistantTurn => ({ parts: [], stopReason: "error", usage: zeroUsage(), error });
|
|
197
|
+
|
|
198
|
+
export function createRouter(config: RouterConfig): Router {
|
|
199
|
+
const table = new Map<ModelRole, readonly ModelRef[]>();
|
|
200
|
+
for (const role of MODEL_ROLES) {
|
|
201
|
+
const entry = config.roles[role];
|
|
202
|
+
if (!entry) continue;
|
|
203
|
+
const chain: readonly ModelRef[] = Array.isArray(entry) ? entry : [entry];
|
|
204
|
+
if (chain.length > 0) table.set(role, chain);
|
|
205
|
+
}
|
|
206
|
+
const defaults = table.get("default");
|
|
207
|
+
// Factory-time misconfiguration (empty default chain) is a programmer error and may
|
|
208
|
+
// throw; the ADR-003 never-throw contract applies to the stream path below.
|
|
209
|
+
if (!defaults) throw new Error("router: roles.default must contain at least one ModelRef");
|
|
210
|
+
|
|
211
|
+
const chainFor = (role: string): readonly ModelRef[] => table.get(role as ModelRole) ?? defaults;
|
|
212
|
+
|
|
213
|
+
/** Chain containing `model`: prefer the role whose chain HEAD is the model (that is what
|
|
214
|
+
* resolve() hands the loop), else the first role chain containing it anywhere, else —
|
|
215
|
+
* under `looseFallback` — the model PREPENDED to the default chain (header: MED-3; the
|
|
216
|
+
* request stays the primary, configured models become its fallbacks), else a singleton
|
|
217
|
+
* passthrough chain. Scan order = MODEL_ROLES order (deterministic). */
|
|
218
|
+
const locate = (model: ModelRef): { key: string; chain: readonly ModelRef[]; index: number } => {
|
|
219
|
+
let containing: { key: string; chain: readonly ModelRef[]; index: number } | null = null;
|
|
220
|
+
for (const role of MODEL_ROLES) {
|
|
221
|
+
const chain = table.get(role);
|
|
222
|
+
if (!chain) continue;
|
|
223
|
+
const idx = chain.findIndex((c) => sameRef(c, model));
|
|
224
|
+
if (idx === 0) return { key: role, chain, index: 0 };
|
|
225
|
+
if (idx > 0 && containing === null) containing = { key: role, chain, index: idx };
|
|
226
|
+
}
|
|
227
|
+
if (containing) return containing;
|
|
228
|
+
// model ∉ any chain (the scan above covered `default`), so the prepend never duplicates
|
|
229
|
+
const chain = config.looseFallback === true ? [model, ...defaults] : [model];
|
|
230
|
+
return { key: refKey(model), chain, index: 0 };
|
|
231
|
+
};
|
|
232
|
+
|
|
233
|
+
return {
|
|
234
|
+
chain: chainFor,
|
|
235
|
+
|
|
236
|
+
resolve(role: string, explicit?: string | ModelRef): ModelRef {
|
|
237
|
+
const chain = chainFor(role);
|
|
238
|
+
const head = chain[0]!; // chains in `table` are never empty
|
|
239
|
+
if (explicit !== undefined) {
|
|
240
|
+
return typeof explicit === "string" ? parseModelRef(explicit, head.provider) : explicit;
|
|
241
|
+
}
|
|
242
|
+
return head;
|
|
243
|
+
},
|
|
244
|
+
|
|
245
|
+
wrap(stream: StreamFn): StreamFn {
|
|
246
|
+
const sticky = config.sticky ?? true;
|
|
247
|
+
const survivors = new Map<string, number>(); // chain key → sticky start index
|
|
248
|
+
return async function* (model, messages, options): AsyncGenerator<StreamEvent> {
|
|
249
|
+
const { key, chain, index } = locate(model);
|
|
250
|
+
const start = sticky
|
|
251
|
+
? Math.min(Math.max(index, survivors.get(key) ?? 0), chain.length - 1)
|
|
252
|
+
: index;
|
|
253
|
+
let last: AssistantTurn | null = null;
|
|
254
|
+
for (let i = start; i < chain.length; i++) {
|
|
255
|
+
const candidate = chain[i]!;
|
|
256
|
+
let turn: AssistantTurn | null = null;
|
|
257
|
+
let streamed = false; // any delta reached the consumer from THIS candidate
|
|
258
|
+
try {
|
|
259
|
+
for await (const ev of stream(candidate, messages, options)) {
|
|
260
|
+
if (ev.type === "turn") turn = ev.turn;
|
|
261
|
+
else { if (ev.type === "text_delta" || ev.type === "reasoning_delta") streamed = true; yield ev; } // deltas pass through live (header: mid-stream failure note)
|
|
262
|
+
}
|
|
263
|
+
} catch (e) {
|
|
264
|
+
// Defensive: the seam contract says streams never throw; if one does, keep the
|
|
265
|
+
// never-throw guarantee here by folding it into an error turn.
|
|
266
|
+
turn = errorTurn(e instanceof Error ? e.message : String(e));
|
|
267
|
+
}
|
|
268
|
+
const t: AssistantTurn = turn ?? errorTurn("stream ended without a terminal turn");
|
|
269
|
+
const aborted = options?.signal?.aborted === true; // retry.ts:337-339: aborts never advance
|
|
270
|
+
// chain.length === 1: nothing to advance to — surface the provider error untouched,
|
|
271
|
+
// no exhausted-rewrite, no note (header: LOW/MED-4).
|
|
272
|
+
// streamed: part of an answer already reached the screen — a re-drive on the next candidate would
|
|
273
|
+
// print a second answer under the first; the failure stands (retry.ts applies the same rule)
|
|
274
|
+
if (t.stopReason !== "error" || aborted || streamed || !classifyStreamError(t.error).retryable || chain.length === 1) {
|
|
275
|
+
SERVED.set(t, candidate); // header: HIGH-2 — this candidate produced the turn
|
|
276
|
+
yield { type: "turn", turn: t }; // success or non-retryable: NO advance
|
|
277
|
+
return;
|
|
278
|
+
}
|
|
279
|
+
last = t;
|
|
280
|
+
const next = i + 1 < chain.length ? chain[i + 1]! : null;
|
|
281
|
+
if (sticky && next !== null) survivors.set(key, i + 1); // handler.ts:163-169 sticky switch
|
|
282
|
+
config.onNote?.({ chain: key, from: candidate, to: next, reason: t.error ?? "error" });
|
|
283
|
+
}
|
|
284
|
+
if (sticky) survivors.delete(key); // exhausted: reset so recovered models get retried
|
|
285
|
+
const tried = chain.length - start;
|
|
286
|
+
const exhausted = errorTurn(
|
|
287
|
+
`model chain '${key}' exhausted (${tried} candidate${tried === 1 ? "" : "s"} failed); last: ${last?.error ?? "unknown error"}`,
|
|
288
|
+
);
|
|
289
|
+
SERVED.set(exhausted, chain[chain.length - 1]!); // last candidate attempted (status honesty)
|
|
290
|
+
yield { type: "turn", turn: exhausted };
|
|
291
|
+
};
|
|
292
|
+
},
|
|
293
|
+
};
|
|
294
|
+
}
|
|
@@ -0,0 +1,26 @@
|
|
|
1
|
+
/** SSE line reader shared by the streaming adapters (moved VERBATIM out of stream.ts — aion port #75 — so a
|
|
2
|
+
* second wire (the Codex Responses wire, not in this build yet) can import it without importing stream.ts).
|
|
3
|
+
* Yields every `data:` payload except the `[DONE]` sentinel; `event:` lines, comments and blank lines
|
|
4
|
+
* are skipped — the OpenAI Responses stream names its event in the payload's `type` field, so the
|
|
5
|
+
* `event:` line carries nothing the consumer needs (codex codex-api/src/sse/responses.rs:843-845 fixtures
|
|
6
|
+
* pair `event:` + `data:` lines the same way). */
|
|
7
|
+
|
|
8
|
+
export async function* sseLines(body: ReadableStream<Uint8Array>): AsyncGenerator<string> {
|
|
9
|
+
const reader = body.getReader();
|
|
10
|
+
const dec = new TextDecoder();
|
|
11
|
+
let buf = "";
|
|
12
|
+
while (true) {
|
|
13
|
+
const { done, value } = await reader.read();
|
|
14
|
+
if (done) break;
|
|
15
|
+
buf += dec.decode(value, { stream: true });
|
|
16
|
+
const lines = buf.split("\n");
|
|
17
|
+
buf = lines.pop() ?? "";
|
|
18
|
+
for (const l of lines) {
|
|
19
|
+
const t = l.trim();
|
|
20
|
+
if (t.startsWith("data:")) {
|
|
21
|
+
const payload = t.slice(5).trim();
|
|
22
|
+
if (payload && payload !== "[DONE]") yield payload;
|
|
23
|
+
}
|
|
24
|
+
}
|
|
25
|
+
}
|
|
26
|
+
}
|