@sayknow-cli/ai 0.4.2 → 0.4.3

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (186) hide show
  1. package/dist/types/api-registry.d.ts +30 -0
  2. package/dist/types/auth-broker/client.d.ts +67 -0
  3. package/dist/types/auth-broker/index.d.ts +5 -0
  4. package/dist/types/auth-broker/refresher.d.ts +25 -0
  5. package/dist/types/auth-broker/remote-store.d.ts +99 -0
  6. package/dist/types/auth-broker/server.d.ts +32 -0
  7. package/dist/types/auth-broker/types.d.ts +110 -0
  8. package/dist/types/auth-broker/wire-schemas.d.ts +443 -0
  9. package/dist/types/auth-gateway/http.d.ts +40 -0
  10. package/dist/types/auth-gateway/index.d.ts +3 -0
  11. package/dist/types/auth-gateway/server.d.ts +36 -0
  12. package/dist/types/auth-gateway/types.d.ts +115 -0
  13. package/dist/types/auth-storage.d.ts +699 -0
  14. package/dist/types/cli.d.ts +2 -0
  15. package/dist/types/context-cap-policy.d.ts +10 -0
  16. package/dist/types/index.d.ts +53 -0
  17. package/dist/types/model-cache.d.ts +19 -0
  18. package/dist/types/model-manager.d.ts +64 -0
  19. package/dist/types/model-retirements.d.ts +6 -0
  20. package/dist/types/model-thinking.d.ts +74 -0
  21. package/dist/types/models.d.ts +21 -0
  22. package/dist/types/provider-details.d.ts +24 -0
  23. package/dist/types/provider-models/bundled-references.d.ts +4 -0
  24. package/dist/types/provider-models/descriptors.d.ts +48 -0
  25. package/dist/types/provider-models/google.d.ts +20 -0
  26. package/dist/types/provider-models/index.d.ts +5 -0
  27. package/dist/types/provider-models/ollama.d.ts +7 -0
  28. package/dist/types/provider-models/openai-compat.d.ts +256 -0
  29. package/dist/types/provider-models/special.d.ts +19 -0
  30. package/dist/types/providers/amazon-bedrock.d.ts +60 -0
  31. package/dist/types/providers/anthropic-messages-server-schema.d.ts +450 -0
  32. package/dist/types/providers/anthropic-messages-server.d.ts +17 -0
  33. package/dist/types/providers/anthropic.d.ts +209 -0
  34. package/dist/types/providers/aws-credential-config.d.ts +19 -0
  35. package/dist/types/providers/aws-credentials.d.ts +43 -0
  36. package/dist/types/providers/aws-eventstream.d.ts +38 -0
  37. package/dist/types/providers/aws-sigv4.d.ts +55 -0
  38. package/dist/types/providers/azure-openai-responses.d.ts +15 -0
  39. package/dist/types/providers/composer-discipline.d.ts +26 -0
  40. package/dist/types/providers/cursor/client-version.d.ts +10 -0
  41. package/dist/types/providers/cursor/gen/agent_pb.d.ts +13022 -0
  42. package/dist/types/providers/cursor.d.ts +45 -0
  43. package/dist/types/providers/error-message.d.ts +27 -0
  44. package/dist/types/providers/github-copilot-headers.d.ts +40 -0
  45. package/dist/types/providers/gitlab-duo.d.ts +27 -0
  46. package/dist/types/providers/google-auth.d.ts +24 -0
  47. package/dist/types/providers/google-gemini-cli.d.ts +75 -0
  48. package/dist/types/providers/google-gemini-headers.d.ts +43 -0
  49. package/dist/types/providers/google-shared.d.ts +179 -0
  50. package/dist/types/providers/google-types.d.ts +138 -0
  51. package/dist/types/providers/google-vertex.d.ts +7 -0
  52. package/dist/types/providers/google.d.ts +4 -0
  53. package/dist/types/providers/grammar.d.ts +1 -0
  54. package/dist/types/providers/kimi.d.ts +27 -0
  55. package/dist/types/providers/mock.d.ts +177 -0
  56. package/dist/types/providers/ollama.d.ts +41 -0
  57. package/dist/types/providers/openai-anthropic-shim.d.ts +31 -0
  58. package/dist/types/providers/openai-bounded-rate-limits.d.ts +3 -0
  59. package/dist/types/providers/openai-chat-server-schema.d.ts +815 -0
  60. package/dist/types/providers/openai-chat-server.d.ts +16 -0
  61. package/dist/types/providers/openai-codex/constants.d.ts +26 -0
  62. package/dist/types/providers/openai-codex/request-transformer.d.ts +50 -0
  63. package/dist/types/providers/openai-codex/response-handler.d.ts +18 -0
  64. package/dist/types/providers/openai-codex-responses.d.ts +68 -0
  65. package/dist/types/providers/openai-completions-compat.d.ts +27 -0
  66. package/dist/types/providers/openai-completions.d.ts +33 -0
  67. package/dist/types/providers/openai-request-transform.d.ts +4 -0
  68. package/dist/types/providers/openai-responses-server-schema.d.ts +392 -0
  69. package/dist/types/providers/openai-responses-server.d.ts +17 -0
  70. package/dist/types/providers/openai-responses-shared.d.ts +104 -0
  71. package/dist/types/providers/openai-responses.d.ts +32 -0
  72. package/dist/types/providers/pi-native-client.d.ts +13 -0
  73. package/dist/types/providers/pi-native-server.d.ts +60 -0
  74. package/dist/types/providers/register-builtins.d.ts +31 -0
  75. package/dist/types/providers/synthetic.d.ts +26 -0
  76. package/dist/types/providers/transform-messages.d.ts +14 -0
  77. package/dist/types/providers/vision-guard.d.ts +8 -0
  78. package/dist/types/rate-limit-utils.d.ts +19 -0
  79. package/dist/types/stream.d.ts +43 -0
  80. package/dist/types/types.d.ts +882 -0
  81. package/dist/types/usage/claude.d.ts +3 -0
  82. package/dist/types/usage/gemini.d.ts +2 -0
  83. package/dist/types/usage/github-copilot.d.ts +7 -0
  84. package/dist/types/usage/google-antigravity.d.ts +2 -0
  85. package/dist/types/usage/grok-cli.d.ts +10 -0
  86. package/dist/types/usage/kimi.d.ts +2 -0
  87. package/dist/types/usage/minimax-code.d.ts +2 -0
  88. package/dist/types/usage/openai-codex.d.ts +3 -0
  89. package/dist/types/usage/shared.d.ts +1 -0
  90. package/dist/types/usage/zai.d.ts +2 -0
  91. package/dist/types/usage.d.ts +258 -0
  92. package/dist/types/utils/abort.d.ts +19 -0
  93. package/dist/types/utils/anthropic-auth.d.ts +31 -0
  94. package/dist/types/utils/discovery/antigravity.d.ts +67 -0
  95. package/dist/types/utils/discovery/codex.d.ts +38 -0
  96. package/dist/types/utils/discovery/cursor.d.ts +23 -0
  97. package/dist/types/utils/discovery/gemini.d.ts +25 -0
  98. package/dist/types/utils/discovery/index.d.ts +4 -0
  99. package/dist/types/utils/discovery/openai-compatible.d.ts +74 -0
  100. package/dist/types/utils/event-stream.d.ts +39 -0
  101. package/dist/types/utils/fallback-transport.d.ts +66 -0
  102. package/dist/types/utils/fireworks-model-id.d.ts +10 -0
  103. package/dist/types/utils/foundry.d.ts +1 -0
  104. package/dist/types/utils/h2-fetch.d.ts +22 -0
  105. package/dist/types/utils/http-inspector.d.ts +35 -0
  106. package/dist/types/utils/idle-iterator.d.ts +67 -0
  107. package/dist/types/utils/json-parse.d.ts +18 -0
  108. package/dist/types/utils/oauth/alibaba-coding-plan.d.ts +18 -0
  109. package/dist/types/utils/oauth/anthropic.d.ts +22 -0
  110. package/dist/types/utils/oauth/api-key-login.d.ts +35 -0
  111. package/dist/types/utils/oauth/api-key-validation.d.ts +27 -0
  112. package/dist/types/utils/oauth/callback-server.d.ts +60 -0
  113. package/dist/types/utils/oauth/cerebras.d.ts +1 -0
  114. package/dist/types/utils/oauth/cloudflare-ai-gateway.d.ts +18 -0
  115. package/dist/types/utils/oauth/cursor.d.ts +15 -0
  116. package/dist/types/utils/oauth/deepinfra.d.ts +1 -0
  117. package/dist/types/utils/oauth/deepseek.d.ts +10 -0
  118. package/dist/types/utils/oauth/firepass.d.ts +1 -0
  119. package/dist/types/utils/oauth/fireworks.d.ts +1 -0
  120. package/dist/types/utils/oauth/fugu.d.ts +1 -0
  121. package/dist/types/utils/oauth/github-copilot.d.ts +38 -0
  122. package/dist/types/utils/oauth/gitlab-duo.d.ts +3 -0
  123. package/dist/types/utils/oauth/glm-zcode.d.ts +71 -0
  124. package/dist/types/utils/oauth/google-antigravity.d.ts +11 -0
  125. package/dist/types/utils/oauth/google-gemini-cli.d.ts +10 -0
  126. package/dist/types/utils/oauth/google-oauth-shared.d.ts +28 -0
  127. package/dist/types/utils/oauth/huggingface.d.ts +19 -0
  128. package/dist/types/utils/oauth/index.d.ts +39 -0
  129. package/dist/types/utils/oauth/kagi.d.ts +17 -0
  130. package/dist/types/utils/oauth/kilo.d.ts +5 -0
  131. package/dist/types/utils/oauth/kimi.d.ts +21 -0
  132. package/dist/types/utils/oauth/litellm.d.ts +18 -0
  133. package/dist/types/utils/oauth/lm-studio.d.ts +17 -0
  134. package/dist/types/utils/oauth/minimax-code.d.ts +28 -0
  135. package/dist/types/utils/oauth/moonshot.d.ts +1 -0
  136. package/dist/types/utils/oauth/nanogpt.d.ts +1 -0
  137. package/dist/types/utils/oauth/nvidia.d.ts +18 -0
  138. package/dist/types/utils/oauth/ollama-cloud.d.ts +2 -0
  139. package/dist/types/utils/oauth/ollama.d.ts +18 -0
  140. package/dist/types/utils/oauth/openai-codex.d.ts +21 -0
  141. package/dist/types/utils/oauth/opencode.d.ts +18 -0
  142. package/dist/types/utils/oauth/parallel.d.ts +17 -0
  143. package/dist/types/utils/oauth/perplexity.d.ts +9 -0
  144. package/dist/types/utils/oauth/pkce.d.ts +8 -0
  145. package/dist/types/utils/oauth/qianfan.d.ts +17 -0
  146. package/dist/types/utils/oauth/qwen-portal.d.ts +19 -0
  147. package/dist/types/utils/oauth/synthetic.d.ts +1 -0
  148. package/dist/types/utils/oauth/tavily.d.ts +17 -0
  149. package/dist/types/utils/oauth/together.d.ts +1 -0
  150. package/dist/types/utils/oauth/types.d.ts +45 -0
  151. package/dist/types/utils/oauth/venice.d.ts +18 -0
  152. package/dist/types/utils/oauth/vercel-ai-gateway.d.ts +18 -0
  153. package/dist/types/utils/oauth/vllm.d.ts +16 -0
  154. package/dist/types/utils/oauth/xai.d.ts +30 -0
  155. package/dist/types/utils/oauth/xiaomi.d.ts +25 -0
  156. package/dist/types/utils/oauth/zai.d.ts +18 -0
  157. package/dist/types/utils/oauth/zenmux.d.ts +1 -0
  158. package/dist/types/utils/overflow.d.ts +14 -0
  159. package/dist/types/utils/parse-bind.d.ts +23 -0
  160. package/dist/types/utils/provider-response.d.ts +3 -0
  161. package/dist/types/utils/retry-after.d.ts +3 -0
  162. package/dist/types/utils/retry-budget.d.ts +1 -0
  163. package/dist/types/utils/retry.d.ts +27 -0
  164. package/dist/types/utils/schema/adapt.d.ts +24 -0
  165. package/dist/types/utils/schema/compatibility.d.ts +30 -0
  166. package/dist/types/utils/schema/dereference.d.ts +11 -0
  167. package/dist/types/utils/schema/draft.d.ts +10 -0
  168. package/dist/types/utils/schema/equality.d.ts +4 -0
  169. package/dist/types/utils/schema/fields.d.ts +49 -0
  170. package/dist/types/utils/schema/index.d.ts +14 -0
  171. package/dist/types/utils/schema/json-schema-validator.d.ts +12 -0
  172. package/dist/types/utils/schema/meta-validator.d.ts +2 -0
  173. package/dist/types/utils/schema/normalize.d.ts +93 -0
  174. package/dist/types/utils/schema/root-combinator.d.ts +12 -0
  175. package/dist/types/utils/schema/spill.d.ts +8 -0
  176. package/dist/types/utils/schema/stamps.d.ts +25 -0
  177. package/dist/types/utils/schema/types.d.ts +4 -0
  178. package/dist/types/utils/schema/wire.d.ts +54 -0
  179. package/dist/types/utils/schema/zod-decontaminate.d.ts +31 -0
  180. package/dist/types/utils/sse-debug.d.ts +10 -0
  181. package/dist/types/utils/tool-call-healing.d.ts +71 -0
  182. package/dist/types/utils/tool-choice-capability.d.ts +41 -0
  183. package/dist/types/utils/tool-choice.d.ts +50 -0
  184. package/dist/types/utils/validation.d.ts +17 -0
  185. package/dist/types/utils.d.ts +89 -0
  186. package/package.json +25 -24
@@ -0,0 +1,104 @@
1
+ import type OpenAI from "openai";
2
+ import type { ResponseInput, ResponseInputContent, ResponseOutputItem } from "openai/resources/responses/responses";
3
+ import { type Api, type AssistantMessage, type ImageContent, type Model, type ServiceTier, type StopReason, type StreamOptions, type TextContent, type TextSignatureV1, type ToolCall, type ToolResultMessage } from "../types";
4
+ import type { AssistantMessageEventStream } from "../utils/event-stream";
5
+ export declare function encodeTextSignatureV1(id: string, phase?: TextSignatureV1["phase"]): string;
6
+ export declare function parseTextSignature(signature: string | undefined): {
7
+ id: string;
8
+ phase?: TextSignatureV1["phase"];
9
+ } | undefined;
10
+ export declare function encodeResponsesToolCallId(callId: string, itemId: string | null | undefined): string;
11
+ export declare function normalizeResponsesToolCallIdForTransform(id: string, model?: Model<Api>, source?: AssistantMessage): string;
12
+ export declare function collectKnownCallIds(messages: ResponseInput): Set<string>;
13
+ /** Scan replay items for call_ids that were originally custom tool calls. */
14
+ export declare function collectCustomCallIds(messages: ResponseInput): Set<string>;
15
+ /**
16
+ * Convert orphan `function_call_output` / `custom_tool_call_output` items —
17
+ * those whose `call_id` has no matching preceding `function_call` /
18
+ * `custom_tool_call` in the same input — into assistant text notes.
19
+ *
20
+ * The Responses API rejects unpaired outputs with
21
+ * `400 No tool call found for function call output with call_id …`. Orphans
22
+ * sneak in through two paths today:
23
+ *
24
+ * - A previous turn's `providerPayload` snapshot replaces the input array via
25
+ * the `dt: false` splice (see {@link convertConversationMessages}), wiping
26
+ * the matching `function_call` while leaving the matching
27
+ * `function_call_output` queued in a later `toolResult`.
28
+ * - A locally-rejected tool call (argument-validation failure, hook reject,
29
+ * aborted turn before the call streamed) produces a tool result without a
30
+ * `function_call` ever landing in any persisted provider payload.
31
+ *
32
+ * Dropping the result loses information the model needs to recover; sending
33
+ * it as-is 400s the request. Folding it into an assistant `message` preserves
34
+ * the payload (call_id + truncated output) while staying within the Responses
35
+ * input grammar. Matches the behavior of {@link transformRequestBody} in the
36
+ * OpenAI code backend provider — issue #1351 / regression of #472.
37
+ */
38
+ export declare function repairOrphanResponsesToolOutputs(input: ResponseInput): ResponseInput;
39
+ export declare function convertResponsesInputContent(content: string | Array<TextContent | ImageContent>, supportsImages: boolean): ResponseInputContent[] | undefined;
40
+ export declare function convertResponsesAssistantMessage<TApi extends Api>(assistantMsg: AssistantMessage, model: Model<TApi>, msgIndex: number, knownCallIds: Set<string>, includeThinkingSignatures?: boolean, customCallIds?: Set<string>): ResponseInput;
41
+ export declare function appendResponsesToolResultMessages<TApi extends Api>(messages: ResponseInput, toolResult: ToolResultMessage, model: Model<TApi>, strictResponsesPairing: boolean, knownCallIds: ReadonlySet<string>, customCallIds?: ReadonlySet<string>): void;
42
+ export interface ProcessResponsesStreamOptions {
43
+ onFirstToken?: () => void;
44
+ onOutputItemDone?: (item: ResponseOutputItem) => void;
45
+ }
46
+ export declare function processResponsesStream<TApi extends Api>(openaiStream: AsyncIterable<OpenAI.Responses.ResponseStreamEvent>, output: AssistantMessage, stream: AssistantMessageEventStream, model: Model<TApi>, options?: ProcessResponsesStreamOptions): Promise<void>;
47
+ /**
48
+ * Mark tool-call blocks left incomplete by a length-truncated response so the
49
+ * agent loop rejects them instead of executing a best-effort partial parse.
50
+ *
51
+ * The universal signal is finalization: a call that never received its terminal
52
+ * `output_item.done` (passed in via `isFinalized`) was cut off mid-arguments.
53
+ * This covers both JSON function calls and raw-input custom tools without
54
+ * mis-flagging a *completed* custom tool whose raw input is not valid JSON. As a
55
+ * defensive secondary, a finalized JSON function call whose buffered arguments
56
+ * still don't parse (e.g. a misbehaving relay) is flagged too. No-op unless the
57
+ * turn stopped for length.
58
+ *
59
+ * Shared by both Responses providers (`openai-responses`, `openai-codex-responses`).
60
+ */
61
+ export declare function flagTruncatedToolCalls(output: AssistantMessage, stopReason: StopReason, isFinalized: (block: ToolCall) => boolean): void;
62
+ export declare function mapOpenAIResponsesStopReason(status: OpenAI.Responses.ResponseStatus | undefined): StopReason;
63
+ /** Initial empty `AssistantMessage` that streaming providers accumulate into. */
64
+ export declare function createInitialResponsesAssistantMessage(api: Api, provider: string, modelId: string): AssistantMessage;
65
+ /** Extension fields we add on top of `ResponseCreateParamsStreaming` across the Responses-family providers. */
66
+ export type ResponsesSamplingParamsExtras = {
67
+ top_p?: number;
68
+ top_k?: number;
69
+ min_p?: number;
70
+ presence_penalty?: number;
71
+ repetition_penalty?: number;
72
+ };
73
+ type CommonResponsesParams = OpenAI.Responses.ResponseCreateParamsStreaming & ResponsesSamplingParamsExtras;
74
+ type CommonSamplingOptions = Pick<StreamOptions, "temperature" | "topP" | "topK" | "minP" | "presencePenalty" | "repetitionPenalty" | "maxTokens"> & {
75
+ serviceTier?: ServiceTier;
76
+ };
77
+ /**
78
+ * Apply the common `StreamOptions` → Responses sampling-parameter mapping (max output tokens,
79
+ * temperature, top-p/k, min-p, presence/repetition penalties, service tier). Mutates `params`.
80
+ */
81
+ export declare function applyCommonResponsesSamplingParams<P extends CommonResponsesParams>(params: P, options: CommonSamplingOptions | undefined, provider: string): void;
82
+ type ReasoningOptions = {
83
+ reasoning?: string;
84
+ reasoningSummary?: "auto" | "detailed" | "concise" | null;
85
+ };
86
+ /**
87
+ * Apply reasoning-related Responses parameters: enable encrypted reasoning content for replay,
88
+ * set effort/summary when requested, and otherwise inject the GPT-5 "Juice: 0" no-reasoning hack.
89
+ * Mutates `params` and may push a developer message into `messages`.
90
+ */
91
+ export declare function applyResponsesReasoningParams<P extends OpenAI.Responses.ResponseCreateParamsStreaming>(params: P, model: Model<Api>, options: ReasoningOptions | undefined, messages: ResponseInput, mapEffort?: (effort: string) => string): void;
92
+ /** Populate `output.usage` from a Responses-API `response.usage` payload. Does not invoke `calculateCost`. */
93
+ export declare function populateResponsesUsageFromResponse(output: AssistantMessage, usage: {
94
+ input_tokens?: number | null;
95
+ output_tokens?: number | null;
96
+ total_tokens?: number | null;
97
+ input_tokens_details?: {
98
+ cached_tokens?: number | null;
99
+ } | null;
100
+ output_tokens_details?: {
101
+ reasoning_tokens?: number | null;
102
+ } | null;
103
+ } | null | undefined): void;
104
+ export {};
@@ -0,0 +1,32 @@
1
+ import type { Tool as OpenAITool } from "openai/resources/responses/responses";
2
+ import type { Model, ServiceTier, StreamFunction, StreamOptions, Tool, ToolChoice } from "../types";
3
+ import { type OpenAIResponsesToolChoice } from "../utils/tool-choice";
4
+ export declare function normalizeOpenAIResponsesPromptCacheKey(sessionId: string | undefined): string | undefined;
5
+ export interface OpenAIResponsesOptions extends StreamOptions {
6
+ reasoning?: "minimal" | "low" | "medium" | "high" | "xhigh" | "max";
7
+ reasoningSummary?: "auto" | "detailed" | "concise" | null;
8
+ serviceTier?: ServiceTier;
9
+ toolChoice?: ToolChoice;
10
+ /**
11
+ * Enforce strict tool call/result pairing when building Responses API inputs.
12
+ * Azure OpenAI and GitHub Copilot Responses paths require tool results to match prior tool calls.
13
+ */
14
+ strictResponsesPairing?: boolean;
15
+ }
16
+ /**
17
+ * Generate function for OpenAI Responses API
18
+ */
19
+ export declare const streamOpenAIResponses: StreamFunction<"openai-responses">;
20
+ export declare function supportsDeveloperRole(modelOrBaseUrl: Pick<Model, "provider" | "baseUrl"> | string): boolean;
21
+ /**
22
+ * Whether this model should get the OpenAI custom-tool grammar variant
23
+ * for `apply_patch`. The generated model catalog sets
24
+ * `model.applyPatchToolType` for first-party GPT-5 Responses models; this
25
+ * runtime path only consumes that metadata.
26
+ * @internal Exported for tests.
27
+ */
28
+ export declare function supportsFreeformApplyPatch(model: Model<"openai-responses">): boolean;
29
+ /** @internal Exported for tests. */
30
+ export declare function mapOpenAIResponsesToolChoiceForTools(choice: ToolChoice | undefined, tools: Tool[], model: Model<"openai-responses">): OpenAIResponsesToolChoice;
31
+ /** @internal Exported for tests. */
32
+ export declare function convertTools(tools: Tool[], strictMode: boolean, model: Model<"openai-responses">): OpenAITool[];
@@ -0,0 +1,13 @@
1
+ import type { Api, AssistantMessageEventStream as AssistantMessageEventStreamType, Context, Model, SimpleStreamOptions } from "../types";
2
+ /**
3
+ * Stream a turn through an `skc auth-gateway` over the pi-native protocol.
4
+ *
5
+ * The returned {@link AssistantMessageEventStream} receives each parsed
6
+ * `AssistantMessageEvent` verbatim from the gateway; the terminal `done` /
7
+ * `error` event resolves `.result()` automatically via the base class's
8
+ * completion check. Non-streaming consumers just call `.result()` and pay
9
+ * for SSE framing they don't use — that overhead is dominated by provider
10
+ * latency, so we always stream rather than maintaining a parallel
11
+ * non-streaming path.
12
+ */
13
+ export declare function streamPiNative<TApi extends Api>(model: Model<TApi>, context: Context, options?: SimpleStreamOptions): AssistantMessageEventStreamType;
@@ -0,0 +1,60 @@
1
+ /**
2
+ * Pi-native wire format for the auth-gateway.
3
+ *
4
+ * Where the OpenAI / Anthropic / Responses route modules translate foreign
5
+ * wire shapes through pi-ai's canonical {@link Context}, this module accepts
6
+ * the canonical shape *directly* — for clients that already speak pi-ai
7
+ * (containerized SKC deployments and sidecar auth gateways).
8
+ * Skipping the wire-format → Context → wire-format round-trip cuts
9
+ * per-request CPU but, more importantly, avoids the quantization that those
10
+ * translations impose on first-class pi-ai fields (service tier, cache
11
+ * markers, thinking budgets, tool-choice variants, …).
12
+ *
13
+ * The streaming wire is {@link AssistantMessageEvent} serialized as SSE. Public
14
+ * projections omit private raw reasoning and serialized Responses reasoning
15
+ * signatures while preserving provider-displayable summaries and genuine opaque
16
+ * signatures. Including `partial: AssistantMessage` on every delta is O(N²) in
17
+ * turn length on the wire — acceptable for the loopback / sidecar topology this
18
+ * transport is designed for; provider latency dominates the actual cost.
19
+ *
20
+ * Endpoint contract:
21
+ * POST /v1/pi/stream
22
+ * body: { modelId, context, options?, stream? } // `stream` defaults to true
23
+ * 200 SSE: stream of `AssistantMessageEvent` (terminated by `data: [DONE]`)
24
+ * 200 JSON (stream=false): { message: AssistantMessage }
25
+ * 4xx/5xx: { error: { type, message } }
26
+ */
27
+ import type { AssistantMessageEventStream, Context, SimpleStreamOptions } from "../types";
28
+ export interface PiNativeParsedRequest {
29
+ modelId: string;
30
+ context: Context;
31
+ options: SimpleStreamOptions;
32
+ stream: boolean;
33
+ }
34
+ /**
35
+ * Parse a pi-native request body. Validation is intentionally minimal — only
36
+ * the shape the gateway itself reads is checked (`modelId`, `context.messages`
37
+ * array, options is an object). Everything downstream is the canonical pi-ai
38
+ * type surface; mis-shaped values surface as a `502 upstream_error` from
39
+ * `streamSimple` rather than being re-validated here.
40
+ *
41
+ * Accepts both `{ modelId: string }` and `{ model: { id: string } }` so the
42
+ * existing `streamProxy` client (which sends the full Model object) can target
43
+ * the gateway with only a URL swap.
44
+ */
45
+ export declare function parseRequest(body: unknown, _headers?: Headers): PiNativeParsedRequest;
46
+ /**
47
+ * Ship only public-safe {@link AssistantMessageEvent} projections. Unknown
48
+ * thinking blocks remain buffered until their terminal partial establishes that
49
+ * the provider-native block is safe; raw and mixed blocks never reach SSE.
50
+ */
51
+ export declare function encodeStream(events: AssistantMessageEventStream): ReadableStream<Uint8Array>;
52
+ /**
53
+ * Pi-native error envelope:
54
+ * `{ error: { type, message } }`
55
+ *
56
+ * Mirrors OpenAI's outer shape (which clients/SDKs already parse) without the
57
+ * provider-specific status taxonomy — pi-native callers consume `type`
58
+ * directly.
59
+ */
60
+ export declare function formatError(status: number, type: string, message: string): Response;
@@ -0,0 +1,31 @@
1
+ /**
2
+ * Lazy provider module loading.
3
+ *
4
+ * Each provider module is loaded only when its stream function is first called.
5
+ * This avoids eagerly importing heavy SDK dependencies (e.g., @anthropic-ai/sdk,
6
+ * openai) at startup. The loaded module promise is cached so subsequent calls
7
+ * reuse the same import.
8
+ *
9
+ * stream.ts imports its provider stream functions from this module (see the
10
+ * lazy wrappers below), so this file IS the main streaming path's provider
11
+ * loader: heavy SDKs stay out of the CLI startup parse graph.
12
+ */
13
+ import type { AssistantMessageEventStream, Context, Model, OptionsForApi } from "../types";
14
+ import { AssistantMessageEventStream as EventStreamImpl } from "../utils/event-stream";
15
+ import type { BedrockOptions } from "./amazon-bedrock";
16
+ interface BedrockProviderModule {
17
+ streamBedrock: (model: Model<"bedrock-converse-stream">, context: Context, options: BedrockOptions) => AssistantMessageEventStream;
18
+ }
19
+ export declare function setBedrockProviderModule(module: BedrockProviderModule): void;
20
+ export declare const streamAnthropic: (model: Model<"anthropic-messages">, context: Context, options: OptionsForApi<"anthropic-messages">) => EventStreamImpl;
21
+ export declare const streamAzureOpenAIResponses: (model: Model<"azure-openai-responses">, context: Context, options: OptionsForApi<"azure-openai-responses">) => EventStreamImpl;
22
+ export declare const streamGoogle: (model: Model<"google-generative-ai">, context: Context, options: OptionsForApi<"google-generative-ai">) => EventStreamImpl;
23
+ export declare const streamGoogleGeminiCli: (model: Model<"google-gemini-cli">, context: Context, options: OptionsForApi<"google-gemini-cli">) => EventStreamImpl;
24
+ export declare const streamGoogleVertex: (model: Model<"google-vertex">, context: Context, options: OptionsForApi<"google-vertex">) => EventStreamImpl;
25
+ export declare const streamOpenAICodexResponses: (model: Model<"openai-codex-responses">, context: Context, options: OptionsForApi<"openai-codex-responses">) => EventStreamImpl;
26
+ export declare const streamOpenAICompletions: (model: Model<"openai-completions">, context: Context, options: OptionsForApi<"openai-completions">) => EventStreamImpl;
27
+ export declare const streamOpenAIResponses: (model: Model<"openai-responses">, context: Context, options: OptionsForApi<"openai-responses">) => EventStreamImpl;
28
+ export declare const streamCursor: (model: Model<"cursor-agent">, context: Context, options: OptionsForApi<"cursor-agent">) => EventStreamImpl;
29
+ export declare const streamOllama: (model: Model<"ollama-chat">, context: Context, options: OptionsForApi<"ollama-chat">) => EventStreamImpl;
30
+ export declare const streamBedrock: (model: Model<"bedrock-converse-stream">, context: Context, options: OptionsForApi<"bedrock-converse-stream">) => EventStreamImpl;
31
+ export {};
@@ -0,0 +1,26 @@
1
+ /**
2
+ * Synthetic provider - wraps OpenAI or Anthropic API based on format setting.
3
+ *
4
+ * Synthetic offers both OpenAI-compatible and Anthropic-compatible APIs:
5
+ * - OpenAI: https://api.synthetic.new/openai/v1/chat/completions
6
+ * - Anthropic: https://api.synthetic.new/anthropic/v1/messages
7
+ *
8
+ * @see https://dev.synthetic.new/docs/api/overview
9
+ */
10
+ import type { Api, Context, Model } from "../types";
11
+ import type { AssistantMessageEventStream } from "../utils/event-stream";
12
+ import { type OpenAIAnthropicApiFormat, type OpenAIAnthropicShimOptions } from "./openai-anthropic-shim";
13
+ export type SyntheticApiFormat = OpenAIAnthropicApiFormat;
14
+ export interface SyntheticOptions extends OpenAIAnthropicShimOptions {
15
+ /** API format: "openai" or "anthropic". Default: "openai" */
16
+ format?: SyntheticApiFormat;
17
+ }
18
+ /**
19
+ * Stream from Synthetic, routing to either OpenAI or Anthropic API based on format.
20
+ * Returns synchronously like other providers - async processing happens internally.
21
+ */
22
+ export declare function streamSynthetic(model: Model<"openai-completions">, context: Context, options?: SyntheticOptions): AssistantMessageEventStream;
23
+ /**
24
+ * Check if a model is a Synthetic model.
25
+ */
26
+ export declare function isSyntheticModel(model: Model<Api>): boolean;
@@ -0,0 +1,14 @@
1
+ import type { Api, AssistantMessage, Message, Model } from "../types";
2
+ /**
3
+ * Normalize tool call ID for cross-provider compatibility.
4
+ * OpenAI Responses API generates IDs that are 450+ chars with special characters like `|`.
5
+ * Anthropic APIs require IDs matching ^[a-zA-Z0-9_-]+$ (max 64 chars).
6
+ *
7
+ * For aborted/errored turns, this function:
8
+ * - Preserves tool call structure (unlike converting to text summaries)
9
+ * - Injects synthetic "aborted" tool results
10
+ * - Adds a <turn-aborted> guidance marker for the model
11
+ */
12
+ export declare function transformMessages<TApi extends Api>(messages: Message[], model: Model<TApi>, normalizeToolCallId?: (id: string, model: Model<TApi>, source: AssistantMessage) => string, options?: {
13
+ repairLatestAssistantThinking?: boolean;
14
+ }): Message[];
@@ -0,0 +1,8 @@
1
+ import type { ImageContent, TextContent } from "../types";
2
+ export declare const NON_VISION_IMAGE_PLACEHOLDER = "[image omitted: model does not support vision]";
3
+ export declare function partitionVisionContent(content: ReadonlyArray<TextContent | ImageContent>, supportsImages: boolean): {
4
+ textBlocks: TextContent[];
5
+ imageBlocks: ImageContent[];
6
+ omittedImages: boolean;
7
+ };
8
+ export declare function joinTextWithImagePlaceholder(text: string, omittedImages: boolean): string;
@@ -0,0 +1,19 @@
1
+ /**
2
+ * Rate limit reason classification and backoff calculation utilities.
3
+ * Ported from opencode-antigravity-auth plugin for consistency.
4
+ */
5
+ export type RateLimitReason = "QUOTA_EXHAUSTED" | "RATE_LIMIT_EXCEEDED" | "MODEL_CAPACITY_EXHAUSTED" | "SERVER_ERROR" | "UNKNOWN";
6
+ /**
7
+ * Classify a rate-limit error message into a reason category.
8
+ * Priority order: MODEL_CAPACITY > RATE_LIMIT > QUOTA > SERVER_ERROR > UNKNOWN.
9
+ *
10
+ * "resource exhausted" maps to MODEL_CAPACITY (transient, short wait)
11
+ * "quota exceeded" maps to QUOTA_EXHAUSTED (long wait, switch account)
12
+ */
13
+ export declare function parseRateLimitReason(errorMessage: string): RateLimitReason;
14
+ /**
15
+ * Calculate backoff delay in ms for a given rate limit reason.
16
+ * MODEL_CAPACITY gets jitter to prevent thundering herd.
17
+ */
18
+ export declare function calculateRateLimitBackoffMs(reason: RateLimitReason): number;
19
+ export declare function isUsageLimitError(errorMessage: string): boolean;
@@ -0,0 +1,43 @@
1
+ import type { Effort } from "./model-thinking";
2
+ import type { AnthropicOptions } from "./providers/anthropic";
3
+ import type { Api, AssistantMessage, Context, Model, OptionsForApi, SimpleStreamOptions, ToolChoice } from "./types";
4
+ import { AssistantMessageEventStream } from "./utils/event-stream";
5
+ /**
6
+ * Get API key for provider from known environment variables, e.g. OPENAI_API_KEY.
7
+ *
8
+ * Provider authentication intentionally excludes cwd/.env values. Project dotenv files are
9
+ * loaded into $env for app/tool execution, but must not silently fund SKC model requests.
10
+ */
11
+ export declare function getEnvApiKey(provider: string): string | undefined;
12
+ /**
13
+ * Enumerate every provider that has an env-var fallback for `getEnvApiKey`.
14
+ * Used by `skc auth-broker migrate --include-env` to discover env-sourced keys
15
+ * that should be uploaded to the broker.
16
+ */
17
+ export declare function listProvidersWithEnvKey(): string[];
18
+ /**
19
+ * Provider-specific credential guidance appended to "no credential" errors.
20
+ *
21
+ * Headless SKC has no interactive `/login` TUI, so a bare "No API key" /
22
+ * "No credentials" error left users — OpenCode Go subscribers especially
23
+ * (#755) — unsure what signal SKC actually reads. OpenCode subscriptions are
24
+ * themselves API keys, so this names the env var SKC reads for the provider,
25
+ * warns that a project `.env` is intentionally ignored for provider
26
+ * credentials, and points OpenCode users at one-time interactive CLI credential capture.
27
+ *
28
+ * Returns an empty string when the provider has no env-var key and no special
29
+ * handling, so callers can append it unconditionally.
30
+ */
31
+ export declare function formatProviderCredentialHint(provider: string): string;
32
+ /**
33
+ * Build an actionable "missing API key" error for a provider, used by the
34
+ * low-level `stream`/`complete` entry points (#755).
35
+ */
36
+ export declare function formatMissingApiKeyError(provider: string): string;
37
+ export declare function stream<TApi extends Api>(model: Model<TApi>, context: Context, options?: OptionsForApi<TApi>): AssistantMessageEventStream;
38
+ export declare function complete<TApi extends Api>(model: Model<TApi>, context: Context, options?: OptionsForApi<TApi>): Promise<AssistantMessage>;
39
+ export declare function streamSimple<TApi extends Api>(model: Model<TApi>, context: Context, options?: SimpleStreamOptions): AssistantMessageEventStream;
40
+ export declare function completeSimple<TApi extends Api>(model: Model<TApi>, context: Context, options?: SimpleStreamOptions): Promise<AssistantMessage>;
41
+ export declare const OUTPUT_FALLBACK_BUFFER = 4000;
42
+ export declare const ANTHROPIC_THINKING: Record<Effort, number>;
43
+ export declare function mapAnthropicToolChoice(choice?: ToolChoice): AnthropicOptions["toolChoice"];