@selesai/code 0.9.10 → 0.9.11

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (302) hide show
  1. package/CHANGELOG.md +260 -255
  2. package/dist/cli/args.js +13 -2
  3. package/dist/cli/args.test.js +8 -0
  4. package/dist/cli/credential-print.js +4 -4
  5. package/dist/cli.js +0 -0
  6. package/dist/core/agent-session-auto-handoff.test.js +24 -1
  7. package/dist/core/agent-session.d.ts +11 -4
  8. package/dist/core/agent-session.js +96 -35
  9. package/dist/core/compaction/branch-summarization.js +30 -30
  10. package/dist/core/compaction/compaction.js +81 -81
  11. package/dist/core/compaction/utils.js +2 -2
  12. package/dist/core/defaults.d.ts +1 -0
  13. package/dist/core/defaults.js +9 -0
  14. package/dist/core/export-html/template.css +1066 -1066
  15. package/dist/core/export-html/template.html +55 -55
  16. package/dist/core/export-html/template.js +1864 -1864
  17. package/dist/core/export-html/vendor/highlight.min.js +1212 -1212
  18. package/dist/core/export-html/vendor/marked.min.js +78 -78
  19. package/dist/core/extensions/loader.js +108 -35
  20. package/dist/core/extensions/loader.test.d.ts +1 -0
  21. package/dist/core/extensions/loader.test.js +21 -0
  22. package/dist/core/extensions/types.d.ts +12 -2
  23. package/dist/core/keybindings.d.ts +2 -2
  24. package/dist/core/messages.js +7 -7
  25. package/dist/core/model-resolver.d.ts +1 -0
  26. package/dist/core/model-resolver.js +10 -4
  27. package/dist/core/package-manager.js +4 -4
  28. package/dist/core/sdk.d.ts +5 -5
  29. package/dist/core/sdk.js +16 -4
  30. package/dist/core/settings-manager.d.ts +6 -0
  31. package/dist/core/settings-manager.js +22 -0
  32. package/dist/core/tools/bash.d.ts +18 -1
  33. package/dist/core/tools/bash.js +41 -23
  34. package/dist/core/tools/index.d.ts +4 -1
  35. package/dist/core/tools/index.js +9 -1
  36. package/dist/core/tools/powershell.d.ts +15 -0
  37. package/dist/core/tools/powershell.js +38 -0
  38. package/dist/core/tools/powershell.test.d.ts +1 -0
  39. package/dist/core/tools/powershell.test.js +20 -0
  40. package/dist/extensions/context-compaction-reminder.test.ts +82 -82
  41. package/dist/extensions/context-compaction-reminder.ts +28 -28
  42. package/dist/extensions/pi-intercom/LICENSE +21 -21
  43. package/dist/extensions/pi-intercom/broker/client.test.ts +83 -83
  44. package/dist/extensions/pi-intercom/broker/extension.test.ts +387 -387
  45. package/dist/extensions/pi-intercom/broker/framing.test.ts +114 -114
  46. package/dist/extensions/pi-intercom/broker/paths.test.ts +153 -153
  47. package/dist/extensions/pi-intercom/broker/paths.ts +134 -134
  48. package/dist/extensions/pi-intercom/broker/runtime-claim.test.ts +34 -34
  49. package/dist/extensions/pi-intercom/broker/runtime-claim.ts +21 -21
  50. package/dist/extensions/pi-intercom/cwd.test.ts +40 -40
  51. package/dist/extensions/pi-intercom/cwd.ts +31 -31
  52. package/dist/extensions/pi-intercom/extension-api.ts +44 -44
  53. package/dist/extensions/pi-intercom/format-context.test.ts +31 -31
  54. package/dist/extensions/pi-intercom/format-context.ts +32 -32
  55. package/dist/extensions/pi-intercom/test/overlay-width.test.ts +66 -66
  56. package/dist/extensions/pi-intercom/ui/compose.ts +143 -143
  57. package/dist/extensions/pi-intercom/ui/session-list.ts +166 -166
  58. package/dist/extensions/pi-powerline-footer/bash-mode/shell-session.ts +286 -286
  59. package/dist/extensions/pi-powerline-footer/bash-mode/transcript.ts +108 -108
  60. package/dist/extensions/pi-powerline-footer/separators.ts +57 -57
  61. package/dist/extensions/pi-powerline-footer/tests/session-usage.test.ts +47 -47
  62. package/dist/extensions/pi-powerline-footer/tests/tps.test.ts +39 -39
  63. package/dist/extensions/pi-powerline-footer/theme.example.json +24 -24
  64. package/dist/extensions/pi-powerline-footer/theme.json +12 -12
  65. package/dist/extensions/pi-powerline-footer/tps.ts +345 -345
  66. package/dist/extensions/pi-rewind-hook/README.md +245 -245
  67. package/dist/extensions/pi-rewind-hook/index.ts +1445 -1445
  68. package/dist/extensions/pi-rewind-hook/package.json +29 -29
  69. package/dist/extensions/pi-subagents/install.mjs +0 -0
  70. package/dist/extensions/pi-web-agent/package.json +31 -31
  71. package/dist/extensions/pi-web-agent/src/backends/config.ts +205 -205
  72. package/dist/extensions/pi-web-agent/src/backends/doctor.ts +136 -136
  73. package/dist/extensions/pi-web-agent/src/backends/factory.ts +152 -152
  74. package/dist/extensions/pi-web-agent/src/backends/settings-reader.ts +25 -25
  75. package/dist/extensions/pi-web-agent/src/cache/ttl-cache.ts +28 -28
  76. package/dist/extensions/pi-web-agent/src/changelog-notice.ts +136 -136
  77. package/dist/extensions/pi-web-agent/src/commands/web-agent-config.ts +946 -946
  78. package/dist/extensions/pi-web-agent/src/extension.ts +126 -126
  79. package/dist/extensions/pi-web-agent/src/extract/readability.ts +118 -118
  80. package/dist/extensions/pi-web-agent/src/fetch/browser-resolution.ts +199 -199
  81. package/dist/extensions/pi-web-agent/src/fetch/firecrawl-fetch.ts +100 -100
  82. package/dist/extensions/pi-web-agent/src/fetch/headless-fetch.ts +117 -117
  83. package/dist/extensions/pi-web-agent/src/fetch/http-fetch.ts +67 -67
  84. package/dist/extensions/pi-web-agent/src/orchestration/answer-synthesizer.ts +60 -60
  85. package/dist/extensions/pi-web-agent/src/orchestration/candidate-selector.ts +50 -50
  86. package/dist/extensions/pi-web-agent/src/orchestration/direct-url.ts +52 -52
  87. package/dist/extensions/pi-web-agent/src/orchestration/evidence-quality.ts +105 -105
  88. package/dist/extensions/pi-web-agent/src/orchestration/evidence-ranker.ts +45 -45
  89. package/dist/extensions/pi-web-agent/src/orchestration/index.ts +28 -28
  90. package/dist/extensions/pi-web-agent/src/orchestration/query-planner.ts +47 -47
  91. package/dist/extensions/pi-web-agent/src/orchestration/research-orchestrator.ts +376 -376
  92. package/dist/extensions/pi-web-agent/src/orchestration/research-types.ts +64 -64
  93. package/dist/extensions/pi-web-agent/src/orchestration/research-worker.ts +181 -181
  94. package/dist/extensions/pi-web-agent/src/orchestration/source-profile.ts +101 -101
  95. package/dist/extensions/pi-web-agent/src/orchestration/stop-decider.ts +81 -81
  96. package/dist/extensions/pi-web-agent/src/presentation/config-store.ts +210 -210
  97. package/dist/extensions/pi-web-agent/src/presentation/config.ts +75 -75
  98. package/dist/extensions/pi-web-agent/src/presentation/explore-presentation.ts +61 -61
  99. package/dist/extensions/pi-web-agent/src/presentation/fetch-presentation.ts +54 -54
  100. package/dist/extensions/pi-web-agent/src/presentation/search-presentation.ts +41 -41
  101. package/dist/extensions/pi-web-agent/src/presentation/select-view.ts +20 -20
  102. package/dist/extensions/pi-web-agent/src/presentation/types.ts +63 -63
  103. package/dist/extensions/pi-web-agent/src/search/brave.ts +114 -114
  104. package/dist/extensions/pi-web-agent/src/search/duckduckgo.ts +72 -72
  105. package/dist/extensions/pi-web-agent/src/search/searxng.ts +96 -96
  106. package/dist/extensions/pi-web-agent/src/tools/web-explore.ts +71 -71
  107. package/dist/extensions/pi-web-agent/src/tools/web-fetch-headless.ts +31 -31
  108. package/dist/extensions/pi-web-agent/src/tools/web-fetch.ts +31 -31
  109. package/dist/extensions/pi-web-agent/src/tools/web-search.ts +157 -157
  110. package/dist/extensions/pi-web-agent/src/types.ts +88 -88
  111. package/dist/extensions/ponytail/package.json +8 -8
  112. package/dist/extensions/question/batch.ts +103 -103
  113. package/dist/extensions/question/constants.ts +30 -30
  114. package/dist/extensions/question/helpers.ts +58 -58
  115. package/dist/extensions/question/navigation.ts +14 -14
  116. package/dist/extensions/question/package.json +19 -19
  117. package/dist/extensions/question/schemas.ts +43 -43
  118. package/dist/extensions/question/selection-mode.ts +46 -46
  119. package/dist/extensions/question/shortcuts.ts +44 -44
  120. package/dist/extensions/question/types.ts +132 -132
  121. package/dist/extensions/test-resolve-hook-impl.mjs +6 -6
  122. package/dist/extensions/test-resolve-hook.mjs +6 -6
  123. package/dist/extensions/web-agent-onboarding.ts +222 -222
  124. package/dist/extensions/workflow/package.json +17 -17
  125. package/dist/package-manager-cli.js +71 -71
  126. package/dist/rpc-entry.js +0 -0
  127. package/dist/skills/agent-browser/SKILL.md +52 -52
  128. package/dist/skills/batch-grill-me/SKILL.md +19 -19
  129. package/dist/skills/grill-me/SKILL.md +10 -10
  130. package/dist/skills/handoff/SKILL.md +16 -16
  131. package/dist/skills/handoff-text/SKILL.md +14 -14
  132. package/dist/skills/implanger/SKILL.md +69 -69
  133. package/dist/skills/improve-codebase/REFERENCE.md +78 -78
  134. package/dist/skills/improve-codebase/SKILL.md +178 -178
  135. package/dist/skills/planger/SKILL.md +166 -166
  136. package/dist/skills/ponytail/SKILL.md +116 -116
  137. package/dist/skills/ponytail-audit/SKILL.md +42 -42
  138. package/dist/skills/ponytail-debt/SKILL.md +45 -45
  139. package/dist/skills/ponytail-gain/SKILL.md +51 -51
  140. package/dist/skills/ponytail-help/SKILL.md +70 -70
  141. package/dist/skills/ponytail-review/SKILL.md +58 -58
  142. package/dist/skills/selesai-handoff/SKILL.md +20 -20
  143. package/dist/skills/workflow-creation/SKILL.md +73 -73
  144. package/dist/themes/powerline-footer/theme.json +33 -33
  145. package/dist/utils/shell.d.ts +3 -0
  146. package/dist/utils/shell.js +17 -5
  147. package/docs/compaction.md +396 -396
  148. package/docs/containerization.md +111 -111
  149. package/docs/development.md +71 -71
  150. package/docs/docs.json +164 -164
  151. package/docs/environment-variables.md +86 -86
  152. package/docs/index.md +83 -83
  153. package/docs/json.md +82 -82
  154. package/docs/models.md +502 -502
  155. package/docs/packages.md +227 -227
  156. package/docs/plans/subagent-delegation/phase-0-correctness.md +264 -264
  157. package/docs/plans/subagent-delegation/phase-1-behavioral-contract.md +485 -485
  158. package/docs/plans/subagent-delegation/phase-2-context-controls.md +281 -281
  159. package/docs/plans/subagent-delegation/phase-3-advisory-routing.md +361 -361
  160. package/docs/plans/subagent-delegation/phase-4-optional-enforcement.md +380 -380
  161. package/docs/prompt-templates.md +95 -95
  162. package/docs/providers.md +293 -293
  163. package/docs/sdk.md +1144 -1143
  164. package/docs/security.md +59 -59
  165. package/docs/session-format.md +414 -414
  166. package/docs/sessions.md +145 -145
  167. package/docs/shared-host-extensions.md +109 -109
  168. package/docs/shell-aliases.md +13 -13
  169. package/docs/skills.md +231 -231
  170. package/docs/terminal-setup.md +142 -142
  171. package/docs/termux.md +127 -127
  172. package/docs/themes.md +295 -295
  173. package/docs/tmux.md +63 -63
  174. package/docs/tui.md +927 -927
  175. package/docs/windows.md +17 -17
  176. package/examples/README.md +25 -25
  177. package/examples/extensions/README.md +211 -211
  178. package/examples/extensions/auto-commit-on-exit.ts +49 -49
  179. package/examples/extensions/bash-spawn-hook.ts +30 -30
  180. package/examples/extensions/bookmark.ts +50 -50
  181. package/examples/extensions/border-status-editor.ts +150 -150
  182. package/examples/extensions/built-in-tool-renderer.ts +249 -249
  183. package/examples/extensions/claude-rules.ts +86 -86
  184. package/examples/extensions/commands.ts +72 -72
  185. package/examples/extensions/confirm-destructive.ts +59 -59
  186. package/examples/extensions/custom-compaction.ts +130 -130
  187. package/examples/extensions/custom-footer.ts +64 -64
  188. package/examples/extensions/custom-header.ts +73 -73
  189. package/examples/extensions/custom-provider-anthropic/index.ts +610 -610
  190. package/examples/extensions/custom-provider-anthropic/package-lock.json +24 -24
  191. package/examples/extensions/custom-provider-anthropic/package.json +19 -19
  192. package/examples/extensions/custom-provider-gitlab-duo/index.ts +404 -404
  193. package/examples/extensions/custom-provider-gitlab-duo/package.json +16 -16
  194. package/examples/extensions/custom-provider-gitlab-duo/test.ts +82 -82
  195. package/examples/extensions/dirty-repo-guard.ts +56 -56
  196. package/examples/extensions/doom-overlay/README.md +46 -46
  197. package/examples/extensions/doom-overlay/doom/build/doom.js +21 -21
  198. package/examples/extensions/doom-overlay/doom/build/doom.wasm +0 -0
  199. package/examples/extensions/doom-overlay/doom/build.sh +152 -152
  200. package/examples/extensions/doom-overlay/doom/doomgeneric_pi.c +72 -72
  201. package/examples/extensions/doom-overlay/doom-component.ts +132 -132
  202. package/examples/extensions/doom-overlay/doom-engine.ts +173 -173
  203. package/examples/extensions/doom-overlay/doom-keys.ts +104 -104
  204. package/examples/extensions/doom-overlay/index.ts +74 -74
  205. package/examples/extensions/doom-overlay/wad-finder.ts +51 -51
  206. package/examples/extensions/dynamic-resources/SKILL.md +8 -8
  207. package/examples/extensions/dynamic-resources/dynamic.json +79 -79
  208. package/examples/extensions/dynamic-resources/dynamic.md +5 -5
  209. package/examples/extensions/dynamic-resources/index.ts +15 -15
  210. package/examples/extensions/dynamic-tools.ts +74 -74
  211. package/examples/extensions/event-bus.ts +43 -43
  212. package/examples/extensions/file-trigger.ts +41 -41
  213. package/examples/extensions/git-checkpoint.ts +53 -53
  214. package/examples/extensions/git-merge-and-resolve.ts +115 -115
  215. package/examples/extensions/github-issue-autocomplete.ts +185 -185
  216. package/examples/extensions/gondolin/index.ts +531 -531
  217. package/examples/extensions/gondolin/package-lock.json +185 -185
  218. package/examples/extensions/gondolin/package.json +19 -19
  219. package/examples/extensions/handoff.ts +199 -199
  220. package/examples/extensions/hello.ts +26 -26
  221. package/examples/extensions/hidden-thinking-label.ts +53 -53
  222. package/examples/extensions/inline-bash.ts +94 -94
  223. package/examples/extensions/input-transform-streaming.ts +39 -39
  224. package/examples/extensions/input-transform.ts +43 -43
  225. package/examples/extensions/interactive-shell.ts +196 -196
  226. package/examples/extensions/mac-system-theme.ts +47 -47
  227. package/examples/extensions/message-renderer.ts +59 -59
  228. package/examples/extensions/minimal-mode.ts +426 -426
  229. package/examples/extensions/modal-editor.ts +85 -85
  230. package/examples/extensions/model-status.ts +31 -31
  231. package/examples/extensions/notify.ts +55 -55
  232. package/examples/extensions/overlay-qa-tests.ts +1450 -1450
  233. package/examples/extensions/overlay-test.ts +153 -153
  234. package/examples/extensions/permission-gate.ts +34 -34
  235. package/examples/extensions/pirate.ts +47 -47
  236. package/examples/extensions/plan-mode/README.md +66 -66
  237. package/examples/extensions/plan-mode/index.ts +390 -390
  238. package/examples/extensions/plan-mode/utils.ts +168 -168
  239. package/examples/extensions/preset.ts +436 -436
  240. package/examples/extensions/project-trust.ts +64 -64
  241. package/examples/extensions/prompt-customizer.ts +97 -97
  242. package/examples/extensions/protected-paths.ts +30 -30
  243. package/examples/extensions/provider-payload.ts +18 -18
  244. package/examples/extensions/qna.ts +122 -122
  245. package/examples/extensions/question.ts +285 -285
  246. package/examples/extensions/questionnaire.ts +448 -448
  247. package/examples/extensions/rainbow-editor.ts +88 -88
  248. package/examples/extensions/reload-runtime.ts +37 -37
  249. package/examples/extensions/rpc-demo.ts +118 -118
  250. package/examples/extensions/sandbox/index.ts +321 -321
  251. package/examples/extensions/sandbox/package-lock.json +92 -92
  252. package/examples/extensions/sandbox/package.json +19 -19
  253. package/examples/extensions/send-user-message.ts +97 -97
  254. package/examples/extensions/session-name.ts +27 -27
  255. package/examples/extensions/shutdown-command.ts +63 -63
  256. package/examples/extensions/snake.ts +343 -343
  257. package/examples/extensions/space-invaders.ts +560 -560
  258. package/examples/extensions/ssh.ts +220 -220
  259. package/examples/extensions/status-line.ts +32 -32
  260. package/examples/extensions/structured-output.ts +65 -65
  261. package/examples/extensions/subagent/README.md +175 -175
  262. package/examples/extensions/subagent/agents/planner.md +37 -37
  263. package/examples/extensions/subagent/agents/reviewer.md +35 -35
  264. package/examples/extensions/subagent/agents/scout.md +50 -50
  265. package/examples/extensions/subagent/agents/worker.md +24 -24
  266. package/examples/extensions/subagent/agents.ts +126 -126
  267. package/examples/extensions/subagent/index.ts +1015 -1015
  268. package/examples/extensions/subagent/prompts/implement-and-review.md +10 -10
  269. package/examples/extensions/subagent/prompts/implement.md +10 -10
  270. package/examples/extensions/subagent/prompts/scout-and-plan.md +9 -9
  271. package/examples/extensions/summarize.ts +209 -209
  272. package/examples/extensions/system-prompt-header.ts +17 -17
  273. package/examples/extensions/tic-tac-toe.ts +1008 -1008
  274. package/examples/extensions/timed-confirm.ts +70 -70
  275. package/examples/extensions/titlebar-spinner.ts +58 -58
  276. package/examples/extensions/todo.ts +297 -297
  277. package/examples/extensions/tool-override.ts +144 -144
  278. package/examples/extensions/tools.ts +146 -146
  279. package/examples/extensions/trigger-compact.ts +50 -50
  280. package/examples/extensions/truncated-tool.ts +195 -195
  281. package/examples/extensions/widget-placement.ts +9 -9
  282. package/examples/extensions/with-deps/index.ts +32 -32
  283. package/examples/extensions/with-deps/package-lock.json +31 -31
  284. package/examples/extensions/with-deps/package.json +22 -22
  285. package/examples/extensions/working-indicator.ts +123 -123
  286. package/examples/extensions/working-message-test.ts +25 -25
  287. package/examples/rpc-extension-ui.ts +632 -632
  288. package/examples/sdk/01-minimal.ts +26 -26
  289. package/examples/sdk/02-custom-model.ts +53 -53
  290. package/examples/sdk/03-custom-prompt.ts +75 -75
  291. package/examples/sdk/04-skills.ts +55 -55
  292. package/examples/sdk/05-tools.ts +48 -48
  293. package/examples/sdk/06-extensions.ts +99 -99
  294. package/examples/sdk/07-context-files.ts +47 -47
  295. package/examples/sdk/08-prompt-templates.ts +51 -51
  296. package/examples/sdk/09-api-keys-and-oauth.ts +52 -52
  297. package/examples/sdk/10-settings.ts +53 -53
  298. package/examples/sdk/11-sessions.ts +52 -52
  299. package/examples/sdk/12-full-control.ts +79 -79
  300. package/examples/sdk/13-session-runtime.ts +67 -67
  301. package/examples/sdk/README.md +144 -144
  302. package/package.json +4 -4
package/docs/models.md CHANGED
@@ -1,502 +1,502 @@
1
- # Custom Models
2
-
3
- Add custom providers and models (Ollama, vLLM, LM Studio, proxies) via `~/.pi/agent/models.json`.
4
-
5
- ## Table of Contents
6
-
7
- - [Minimal Example](#minimal-example)
8
- - [Full Example](#full-example)
9
- - [Supported APIs](#supported-apis)
10
- - [Provider Configuration](#provider-configuration)
11
- - [Model Configuration](#model-configuration)
12
- - [Overriding Built-in Providers](#overriding-built-in-providers)
13
- - [Per-model Overrides](#per-model-overrides)
14
- - [Anthropic Messages Compatibility](#anthropic-messages-compatibility)
15
- - [OpenAI Compatibility](#openai-compatibility)
16
-
17
- ## Minimal Example
18
-
19
- For local models (Ollama, LM Studio, vLLM), only `id` is required per model:
20
-
21
- ```json
22
- {
23
- "providers": {
24
- "ollama": {
25
- "baseUrl": "http://localhost:11434/v1",
26
- "api": "openai-completions",
27
- "apiKey": "ollama",
28
- "models": [
29
- { "id": "llama3.1:8b" },
30
- { "id": "qwen2.5-coder:7b" }
31
- ]
32
- }
33
- }
34
- }
35
- ```
36
-
37
- The `apiKey` value is a placeholder because Ollama ignores it. pi still treats models as requiring auth before they appear in `/model`, so keyless local servers should keep a dummy value, save a key for that provider with `/login`, or pass `--api-key` when selecting the model.
38
-
39
- Some OpenAI-compatible servers do not understand the `developer` role used for reasoning-capable models. For those providers, set `compat.supportsDeveloperRole` to `false` so pi sends the system prompt as a `system` message instead. If the server also does not support `reasoning_effort`, set `compat.supportsReasoningEffort` to `false` too.
40
-
41
- You can set `compat` at the provider level to apply to all models, or at the model level to override a specific model. This commonly applies to Ollama, vLLM, SGLang, and similar OpenAI-compatible servers.
42
-
43
- ```json
44
- {
45
- "providers": {
46
- "ollama": {
47
- "baseUrl": "http://localhost:11434/v1",
48
- "api": "openai-completions",
49
- "apiKey": "ollama",
50
- "compat": {
51
- "supportsDeveloperRole": false,
52
- "supportsReasoningEffort": false
53
- },
54
- "models": [
55
- {
56
- "id": "gpt-oss:20b",
57
- "reasoning": true
58
- }
59
- ]
60
- }
61
- }
62
- }
63
- ```
64
-
65
- ## Full Example
66
-
67
- Override defaults when you need specific values:
68
-
69
- ```json
70
- {
71
- "providers": {
72
- "ollama": {
73
- "baseUrl": "http://localhost:11434/v1",
74
- "api": "openai-completions",
75
- "apiKey": "ollama",
76
- "models": [
77
- {
78
- "id": "llama3.1:8b",
79
- "name": "Llama 3.1 8B (Local)",
80
- "reasoning": false,
81
- "input": ["text"],
82
- "contextWindow": 128000,
83
- "maxTokens": 32000,
84
- "cost": { "input": 0, "output": 0, "cacheRead": 0, "cacheWrite": 0 }
85
- }
86
- ]
87
- }
88
- }
89
- }
90
- ```
91
-
92
- The file reloads each time you open `/model`. Edit during session; no restart needed.
93
-
94
- ## Google AI Studio Example
95
-
96
- Use `google-generative-ai` with a `baseUrl` to add models from Google AI Studio, including custom Gemma 4 entries:
97
-
98
- ```json
99
- {
100
- "providers": {
101
- "my-google": {
102
- "baseUrl": "https://generativelanguage.googleapis.com/v1beta",
103
- "api": "google-generative-ai",
104
- "apiKey": "$GEMINI_API_KEY",
105
- "models": [
106
- {
107
- "id": "gemma-4-31b-it",
108
- "name": "Gemma 4 31B",
109
- "input": ["text", "image"],
110
- "contextWindow": 262144,
111
- "reasoning": true
112
- }
113
- ]
114
- }
115
- }
116
- }
117
- ```
118
-
119
- The `baseUrl` is required when adding custom models to the `google-generative-ai` API type.
120
-
121
- ## Supported APIs
122
-
123
- | API | Description |
124
- |-----|-------------|
125
- | `openai-completions` | OpenAI Chat Completions (most compatible) |
126
- | `openai-responses` | OpenAI Responses API |
127
- | `anthropic-messages` | Anthropic Messages API |
128
- | `google-generative-ai` | Google Generative AI |
129
-
130
- Set `api` at provider level (default for all models) or model level (override per model).
131
-
132
- ## Provider Configuration
133
-
134
- | Field | Description |
135
- |-------|-------------|
136
- | `baseUrl` | API endpoint URL |
137
- | `api` | API type (see above) |
138
- | `apiKey` | Optional API key config (see value resolution below). Omit it when auth is provided by `/login`/`auth.json` or CLI `--api-key`. |
139
- | `headers` | Custom headers (see value resolution below) |
140
- | `authHeader` | Set `true` to add `Authorization: Bearer <apiKey>` automatically |
141
- | `models` | Array of model configurations |
142
- | `modelOverrides` | Per-model overrides for built-in models on this provider |
143
-
144
- For providers with `models`, non-built-in provider configs need `baseUrl` and an `api` value at either provider or model level. `apiKey` is not required to load the file: models become available when auth is configured through `/login`/`auth.json`, CLI `--api-key`, or provider `apiKey`. If no auth is configured, the models load but stay unavailable in `/model` and `--list-models`.
145
-
146
- ### Value Resolution
147
-
148
- The `apiKey` and `headers` fields support command execution, environment interpolation, and literals:
149
-
150
- - **Shell command:** `"!command"` at the start executes the whole value as a command and uses stdout
151
- ```json
152
- "apiKey": "!security find-generic-password -ws 'anthropic'"
153
- "apiKey": "!op read 'op://vault/item/credential'"
154
- ```
155
- - **Environment interpolation:** `"$ENV_VAR"` or `"${ENV_VAR}"` uses the value of the named variable. Interpolation works inside larger literals.
156
- ```json
157
- "apiKey": "$MY_API_KEY"
158
- "apiKey": "${KEY_PREFIX}_${KEY_SUFFIX}"
159
- ```
160
- `$FOO_BAR` is the variable `FOO_BAR`; use `${FOO}_BAR` when `BAR` is literal text. Missing environment variables make the value unresolved.
161
- - **Escapes:** `"$$"` emits a literal `"$"`; `"$!"` emits a literal `"!"` without triggering command execution.
162
- ```json
163
- "apiKey": "$$literal-dollar-prefix"
164
- "apiKey": "$!literal-bang-prefix"
165
- ```
166
- - **Literal value:** Used directly. Plain uppercase strings such as `MY_API_KEY` are literals; use `$MY_API_KEY` for environment variables.
167
- ```json
168
- "apiKey": "sk-..."
169
- ```
170
-
171
- For `models.json`, shell commands are resolved at request time. pi intentionally does not apply built-in TTL, stale reuse, or recovery logic for arbitrary commands. Different commands need different caching and failure strategies, and pi cannot infer the right one.
172
-
173
- If your command is slow, expensive, rate-limited, or should keep using a previous value on transient failures, wrap it in your own script or command that implements the caching or TTL behavior you want.
174
-
175
- `/model` availability checks use configured auth presence and do not execute shell commands.
176
-
177
- ### Custom Headers
178
-
179
- ```json
180
- {
181
- "providers": {
182
- "custom-proxy": {
183
- "baseUrl": "https://proxy.example.com/v1",
184
- "apiKey": "$MY_API_KEY",
185
- "api": "anthropic-messages",
186
- "headers": {
187
- "x-portkey-api-key": "$PORTKEY_API_KEY",
188
- "x-secret": "!op read 'op://vault/item/secret'"
189
- },
190
- "models": [...]
191
- }
192
- }
193
- }
194
- ```
195
-
196
- ## Model Configuration
197
-
198
- | Field | Required | Default | Description |
199
- |-------|----------|---------|-------------|
200
- | `id` | Yes | — | Model identifier (passed to the API) |
201
- | `name` | No | `id` | Human-readable model label. Used for matching (`--model` patterns) and shown as secondary model detail text. |
202
- | `api` | No | provider's `api` | Override provider's API for this model |
203
- | `reasoning` | No | `false` | Supports extended thinking |
204
- | `thinkingLevelMap` | No | omitted | Maps pi thinking levels to provider values and marks unsupported levels (see below) |
205
- | `input` | No | `["text"]` | Input types: `["text"]` or `["text", "image"]` |
206
- | `contextWindow` | No | `128000` | Context window size in tokens |
207
- | `maxTokens` | No | `16384` | Maximum output tokens |
208
- | `cost` | No | all zeros | `{"input": 0, "output": 0, "cacheRead": 0, "cacheWrite": 0}` (per million tokens) |
209
- | `compat` | No | provider `compat` | Provider compatibility overrides. Merged with provider-level `compat` when both are set. |
210
-
211
- Current behavior:
212
- - `/model`, `--list-models`, and the interactive footer display entries by model `id`.
213
- - The configured `name` is used for model matching and secondary model detail text. It does not replace the footer/status-bar model id.
214
-
215
- ### Thinking Level Map
216
-
217
- Use `thinkingLevelMap` on a model to describe model-specific thinking controls. Keys are pi thinking levels: `off`, `minimal`, `low`, `medium`, `high`, `xhigh`.
218
-
219
- Values are tristate:
220
-
221
- | Value | Meaning |
222
- |-------|---------|
223
- | omitted | Level is supported and uses the provider's default mapping |
224
- | string | Level is supported and this value is sent to the provider |
225
- | `null` | Level is unsupported and hidden/skipped/clamped away |
226
-
227
- Example for a model that only supports off, high, and max reasoning:
228
-
229
- ```json
230
- {
231
- "id": "deepseek-v4-pro",
232
- "reasoning": true,
233
- "thinkingLevelMap": {
234
- "minimal": null,
235
- "low": null,
236
- "medium": null,
237
- "high": "high",
238
- "xhigh": "max"
239
- }
240
- }
241
- ```
242
-
243
- Example for a model where thinking cannot be disabled:
244
-
245
- ```json
246
- {
247
- "id": "always-thinking-model",
248
- "reasoning": true,
249
- "thinkingLevelMap": {
250
- "off": null
251
- }
252
- }
253
- ```
254
-
255
- Migration: older configs that used `compat.reasoningEffortMap` should move that mapping to model-level `thinkingLevelMap`. Use `null` for levels that should not appear in the UI.
256
-
257
- ## Overriding Built-in Providers
258
-
259
- Route a built-in provider through a proxy without redefining models:
260
-
261
- ```json
262
- {
263
- "providers": {
264
- "anthropic": {
265
- "baseUrl": "https://my-proxy.example.com/v1"
266
- }
267
- }
268
- }
269
- ```
270
-
271
- All built-in Anthropic models remain available. Existing OAuth or API key auth continues to work.
272
-
273
- To merge custom models into a built-in provider, include the `models` array:
274
-
275
- ```json
276
- {
277
- "providers": {
278
- "anthropic": {
279
- "baseUrl": "https://my-proxy.example.com/v1",
280
- "apiKey": "$ANTHROPIC_API_KEY",
281
- "api": "anthropic-messages",
282
- "models": [...]
283
- }
284
- }
285
- }
286
- ```
287
-
288
- Merge semantics:
289
- - Built-in models are kept.
290
- - Custom models are upserted by `id` within the provider.
291
- - If a custom model `id` matches a built-in model `id`, the custom model replaces that built-in model.
292
- - If a custom model `id` is new, it is added alongside built-in models.
293
-
294
- ## Per-model Overrides
295
-
296
- Use `modelOverrides` to customize specific built-in models without replacing the provider's full model list.
297
-
298
- ```json
299
- {
300
- "providers": {
301
- "openrouter": {
302
- "modelOverrides": {
303
- "anthropic/claude-sonnet-4": {
304
- "name": "Claude Sonnet 4 (Bedrock Route)",
305
- "compat": {
306
- "openRouterRouting": {
307
- "only": ["amazon-bedrock"]
308
- }
309
- }
310
- }
311
- }
312
- }
313
- }
314
- }
315
- ```
316
-
317
- `modelOverrides` supports these fields per model: `name`, `reasoning`, `input`, `cost` (partial), `contextWindow`, `maxTokens`, `headers`, `compat`.
318
-
319
- Behavior notes:
320
- - `modelOverrides` are applied to built-in provider models.
321
- - Unknown model IDs are ignored.
322
- - You can combine provider-level `baseUrl`/`headers` with `modelOverrides`.
323
- - Overriding `name` changes model matching and secondary detail text only; the footer and primary model lists continue to show the model `id`.
324
- - If `models` is also defined for a provider, custom models are merged after built-in overrides. A custom model with the same `id` replaces the overridden built-in model entry.
325
-
326
- ## Anthropic Messages Compatibility
327
-
328
- For providers or proxies using `api: "anthropic-messages"`, use `compat` to control Anthropic-specific request compatibility.
329
-
330
- By default pi sends per-tool `eager_input_streaming: true`. If a proxy or Anthropic-compatible backend rejects that field, set `supportsEagerToolInputStreaming` to `false`. Pi will omit `tools[].eager_input_streaming` and send the legacy `fine-grained-tool-streaming-2025-05-14` beta header for tool-enabled requests instead.
331
-
332
- Some Anthropic models require adaptive thinking (`thinking.type: "adaptive"` plus `output_config.effort`) instead of the legacy budget-based thinking payload. Built-in models set this automatically. For custom providers or aliases that route to those models, set `forceAdaptiveThinking` to `true`.
333
-
334
- Some Anthropic-compatible providers emit thinking blocks with empty signatures and still expect them on replay. Set `allowEmptySignature` to `true` only for those providers; real Anthropic rejects empty thinking signatures.
335
-
336
- Built-in Anthropic models enable `supportsStrictTools` in their model metadata. Custom Anthropic-compatible models must set it to `true` when their endpoint accepts strict JSON-schema tool definitions.
337
-
338
- ```json
339
- {
340
- "providers": {
341
- "anthropic-proxy": {
342
- "baseUrl": "https://proxy.example.com",
343
- "api": "anthropic-messages",
344
- "apiKey": "$ANTHROPIC_PROXY_KEY",
345
- "compat": {
346
- "supportsEagerToolInputStreaming": false,
347
- "supportsLongCacheRetention": true,
348
- "forceAdaptiveThinking": true,
349
- "allowEmptySignature": true
350
- },
351
- "models": [
352
- {
353
- "id": "claude-opus-4-7",
354
- "reasoning": true,
355
- "input": ["text", "image"]
356
- }
357
- ]
358
- }
359
- }
360
- }
361
- ```
362
-
363
- | Field | Description |
364
- |-------|-------------|
365
- | `supportsEagerToolInputStreaming` | Whether the provider accepts per-tool `eager_input_streaming`. Default: `true`. Set to `false` to omit that field and use the legacy fine-grained tool streaming beta header on tool-enabled requests. |
366
- | `supportsLongCacheRetention` | Whether the provider accepts Anthropic long cache retention (`cache_control.ttl: "1h"`) when cache retention is `long`. Default: `true`. |
367
- | `sendSessionAffinityHeaders` | Whether to send `x-session-affinity` from the session id when caching is enabled. Default: auto-detected for known providers. |
368
- | `supportsCacheControlOnTools` | Whether the provider accepts Anthropic-style `cache_control` markers on tool definitions. Default: `true`. |
369
- | `forceAdaptiveThinking` | Whether to send adaptive thinking (`thinking.type: "adaptive"` plus `output_config.effort`) for this model. Built-in adaptive models set this automatically. Default: `false`. |
370
- | `allowEmptySignature` | Whether to replay empty thinking signatures as `signature: ""` instead of converting thinking to text. Default: `false`. |
371
- | `supportsStrictTools` | Whether the provider accepts strict JSON-schema tool definitions. Default: `false`; built-in Anthropic models enable it in generated metadata. |
372
-
373
- ## OpenAI Compatibility
374
-
375
- For providers with partial OpenAI compatibility, use the `compat` field.
376
-
377
- - Provider-level `compat` applies defaults to all models under that provider.
378
- - Model-level `compat` overrides provider-level values for that model.
379
-
380
- ```json
381
- {
382
- "providers": {
383
- "local-llm": {
384
- "baseUrl": "http://localhost:8080/v1",
385
- "api": "openai-completions",
386
- "compat": {
387
- "supportsUsageInStreaming": false,
388
- "maxTokensField": "max_tokens"
389
- },
390
- "models": [...]
391
- }
392
- }
393
- }
394
- ```
395
-
396
- | Field | Description |
397
- |-------|-------------|
398
- | `supportsStore` | Provider supports `store` field |
399
- | `supportsDeveloperRole` | Use `developer` vs `system` role |
400
- | `supportsReasoningEffort` | Support for `reasoning_effort` parameter |
401
- | `supportsUsageInStreaming` | Supports `stream_options: { include_usage: true }` (default: `true`) |
402
- | `maxTokensField` | Use `max_completion_tokens` or `max_tokens` |
403
- | `requiresToolResultName` | Include `name` on tool result messages |
404
- | `requiresAssistantAfterToolResult` | Insert an assistant message before a user message after tool results |
405
- | `requiresThinkingAsText` | Convert thinking blocks to plain text |
406
- | `requiresReasoningContentOnAssistantMessages` | Include empty `reasoning_content` on all replayed assistant messages when reasoning is enabled |
407
- | `thinkingFormat` | Use `reasoning_effort`, `openrouter`, `deepseek`, `together`, `zai`, `qwen`, `chat-template`, or `qwen-chat-template` thinking parameters |
408
- | `chatTemplateKwargs` | `chat_template_kwargs` values for `thinkingFormat: "chat-template"`; use `{ "$var": "thinking.enabled" }` or `{ "$var": "thinking.effort" }` for pi-controlled thinking values |
409
- | `cacheControlFormat` | Use Anthropic-style `cache_control` markers on the system prompt, last tool definition, and last user, assistant, or tool-result text content. Currently only `anthropic` is supported. |
410
- | `sendSessionAffinityHeaders` | For `openai-completions`, send session-affinity headers from the session id when caching is enabled. Default: `false`. |
411
- | `sessionAffinityFormat` | For `openai-completions` and `openai-responses`, the session-affinity header format: `openai` sends `session_id`/`x-client-request-id` (completions also `x-session-affinity`), `openai-nosession` omits the underscore-containing `session_id` header, `openrouter` sends `x-session-id`. Does not affect the `prompt_cache_key` body param. Default: auto-detected. |
412
- | `supportsStrictMode` | Whether the provider accepts strict JSON-schema function tool definitions. Defaults depend on the API; built-in OpenAI models carry explicit capability metadata. |
413
- | `supportsOpenAIGrammarTools` | Whether OpenAI-compatible APIs emit custom Lark/regex grammar tools. When `false`, grammar-constrained tools fall back to normal function tools. Default: `false`; the built-in model catalog enables it for GPT-5+ models on OpenAI, OpenAI Codex, Azure OpenAI, GitHub Copilot, opencode, and Cloudflare AI Gateway. |
414
- | `deferredToolsMode` | Use provider-specific deferred tool serialization. Currently only `"kimi"` is supported for Kimi's OpenAI-compatible Chat Completions format. |
415
- | `supportsLongCacheRetention` | Whether the provider accepts long cache retention when cache retention is `long`: `prompt_cache_retention: "24h"` for OpenAI prompt caching, or `cache_control.ttl: "1h"` when `cacheControlFormat` is `anthropic`. Default: `true`. |
416
- | `openRouterRouting` | OpenRouter provider routing preferences. This object is sent as-is in the `provider` field of the [OpenRouter API request](https://openrouter.ai/docs/guides/routing/provider-selection). |
417
- | `vercelGatewayRouting` | Vercel AI Gateway routing config for provider selection (`only`, `order`) |
418
-
419
- `openrouter` uses `reasoning: { effort }`. `together` uses `reasoning: { enabled }` and also `reasoning_effort` when `supportsReasoningEffort` is enabled. `qwen` uses top-level `enable_thinking`. Use `qwen-chat-template` for local Qwen-compatible servers that require `chat_template_kwargs.enable_thinking` and `preserve_thinking`. Use `chat-template` for vLLM/Hugging Face chat templates that need configurable `chat_template_kwargs`, such as `chatTemplateKwargs: { "thinking": { "$var": "thinking.enabled" } }` for DeepSeek V3.x templates.
420
-
421
- `cacheControlFormat: "anthropic"` is for OpenAI-compatible providers that expose Anthropic-style prompt caching through `cache_control` markers on text content and tool definitions.
422
-
423
- Example:
424
-
425
- ```json
426
- {
427
- "providers": {
428
- "openrouter": {
429
- "baseUrl": "https://openrouter.ai/api/v1",
430
- "apiKey": "$OPENROUTER_API_KEY",
431
- "api": "openai-completions",
432
- "models": [
433
- {
434
- "id": "openrouter/anthropic/claude-3.5-sonnet",
435
- "name": "OpenRouter Claude 3.5 Sonnet",
436
- "compat": {
437
- "openRouterRouting": {
438
- "allow_fallbacks": true,
439
- "require_parameters": false,
440
- "data_collection": "deny",
441
- "zdr": true,
442
- "enforce_distillable_text": false,
443
- "order": ["anthropic", "amazon-bedrock", "google-vertex"],
444
- "only": ["anthropic", "amazon-bedrock"],
445
- "ignore": ["gmicloud", "friendli"],
446
- "quantizations": ["fp16", "bf16"],
447
- "sort": {
448
- "by": "price",
449
- "partition": "model"
450
- },
451
- "max_price": {
452
- "prompt": 10,
453
- "completion": 20
454
- },
455
- "preferred_min_throughput": {
456
- "p50": 100,
457
- "p90": 50
458
- },
459
- "preferred_max_latency": {
460
- "p50": 1,
461
- "p90": 3,
462
- "p99": 5
463
- }
464
- }
465
- }
466
- }
467
- ]
468
- }
469
- }
470
- }
471
- ```
472
-
473
- Vercel AI Gateway example:
474
-
475
- ```json
476
- {
477
- "providers": {
478
- "vercel-ai-gateway": {
479
- "baseUrl": "https://ai-gateway.vercel.sh/v1",
480
- "apiKey": "$AI_GATEWAY_API_KEY",
481
- "api": "openai-completions",
482
- "models": [
483
- {
484
- "id": "moonshotai/kimi-k2.5",
485
- "name": "Kimi K2.5 (Fireworks via Vercel)",
486
- "reasoning": true,
487
- "input": ["text", "image"],
488
- "cost": { "input": 0.6, "output": 3, "cacheRead": 0, "cacheWrite": 0 },
489
- "contextWindow": 262144,
490
- "maxTokens": 262144,
491
- "compat": {
492
- "vercelGatewayRouting": {
493
- "only": ["fireworks", "novita"],
494
- "order": ["fireworks", "novita"]
495
- }
496
- }
497
- }
498
- ]
499
- }
500
- }
501
- }
502
- ```
1
+ # Custom Models
2
+
3
+ Add custom providers and models (Ollama, vLLM, LM Studio, proxies) via `~/.pi/agent/models.json`.
4
+
5
+ ## Table of Contents
6
+
7
+ - [Minimal Example](#minimal-example)
8
+ - [Full Example](#full-example)
9
+ - [Supported APIs](#supported-apis)
10
+ - [Provider Configuration](#provider-configuration)
11
+ - [Model Configuration](#model-configuration)
12
+ - [Overriding Built-in Providers](#overriding-built-in-providers)
13
+ - [Per-model Overrides](#per-model-overrides)
14
+ - [Anthropic Messages Compatibility](#anthropic-messages-compatibility)
15
+ - [OpenAI Compatibility](#openai-compatibility)
16
+
17
+ ## Minimal Example
18
+
19
+ For local models (Ollama, LM Studio, vLLM), only `id` is required per model:
20
+
21
+ ```json
22
+ {
23
+ "providers": {
24
+ "ollama": {
25
+ "baseUrl": "http://localhost:11434/v1",
26
+ "api": "openai-completions",
27
+ "apiKey": "ollama",
28
+ "models": [
29
+ { "id": "llama3.1:8b" },
30
+ { "id": "qwen2.5-coder:7b" }
31
+ ]
32
+ }
33
+ }
34
+ }
35
+ ```
36
+
37
+ The `apiKey` value is a placeholder because Ollama ignores it. pi still treats models as requiring auth before they appear in `/model`, so keyless local servers should keep a dummy value, save a key for that provider with `/login`, or pass `--api-key` when selecting the model.
38
+
39
+ Some OpenAI-compatible servers do not understand the `developer` role used for reasoning-capable models. For those providers, set `compat.supportsDeveloperRole` to `false` so pi sends the system prompt as a `system` message instead. If the server also does not support `reasoning_effort`, set `compat.supportsReasoningEffort` to `false` too.
40
+
41
+ You can set `compat` at the provider level to apply to all models, or at the model level to override a specific model. This commonly applies to Ollama, vLLM, SGLang, and similar OpenAI-compatible servers.
42
+
43
+ ```json
44
+ {
45
+ "providers": {
46
+ "ollama": {
47
+ "baseUrl": "http://localhost:11434/v1",
48
+ "api": "openai-completions",
49
+ "apiKey": "ollama",
50
+ "compat": {
51
+ "supportsDeveloperRole": false,
52
+ "supportsReasoningEffort": false
53
+ },
54
+ "models": [
55
+ {
56
+ "id": "gpt-oss:20b",
57
+ "reasoning": true
58
+ }
59
+ ]
60
+ }
61
+ }
62
+ }
63
+ ```
64
+
65
+ ## Full Example
66
+
67
+ Override defaults when you need specific values:
68
+
69
+ ```json
70
+ {
71
+ "providers": {
72
+ "ollama": {
73
+ "baseUrl": "http://localhost:11434/v1",
74
+ "api": "openai-completions",
75
+ "apiKey": "ollama",
76
+ "models": [
77
+ {
78
+ "id": "llama3.1:8b",
79
+ "name": "Llama 3.1 8B (Local)",
80
+ "reasoning": false,
81
+ "input": ["text"],
82
+ "contextWindow": 128000,
83
+ "maxTokens": 32000,
84
+ "cost": { "input": 0, "output": 0, "cacheRead": 0, "cacheWrite": 0 }
85
+ }
86
+ ]
87
+ }
88
+ }
89
+ }
90
+ ```
91
+
92
+ The file reloads each time you open `/model`. Edit during session; no restart needed.
93
+
94
+ ## Google AI Studio Example
95
+
96
+ Use `google-generative-ai` with a `baseUrl` to add models from Google AI Studio, including custom Gemma 4 entries:
97
+
98
+ ```json
99
+ {
100
+ "providers": {
101
+ "my-google": {
102
+ "baseUrl": "https://generativelanguage.googleapis.com/v1beta",
103
+ "api": "google-generative-ai",
104
+ "apiKey": "$GEMINI_API_KEY",
105
+ "models": [
106
+ {
107
+ "id": "gemma-4-31b-it",
108
+ "name": "Gemma 4 31B",
109
+ "input": ["text", "image"],
110
+ "contextWindow": 262144,
111
+ "reasoning": true
112
+ }
113
+ ]
114
+ }
115
+ }
116
+ }
117
+ ```
118
+
119
+ The `baseUrl` is required when adding custom models to the `google-generative-ai` API type.
120
+
121
+ ## Supported APIs
122
+
123
+ | API | Description |
124
+ |-----|-------------|
125
+ | `openai-completions` | OpenAI Chat Completions (most compatible) |
126
+ | `openai-responses` | OpenAI Responses API |
127
+ | `anthropic-messages` | Anthropic Messages API |
128
+ | `google-generative-ai` | Google Generative AI |
129
+
130
+ Set `api` at provider level (default for all models) or model level (override per model).
131
+
132
+ ## Provider Configuration
133
+
134
+ | Field | Description |
135
+ |-------|-------------|
136
+ | `baseUrl` | API endpoint URL |
137
+ | `api` | API type (see above) |
138
+ | `apiKey` | Optional API key config (see value resolution below). Omit it when auth is provided by `/login`/`auth.json` or CLI `--api-key`. |
139
+ | `headers` | Custom headers (see value resolution below) |
140
+ | `authHeader` | Set `true` to add `Authorization: Bearer <apiKey>` automatically |
141
+ | `models` | Array of model configurations |
142
+ | `modelOverrides` | Per-model overrides for built-in models on this provider |
143
+
144
+ For providers with `models`, non-built-in provider configs need `baseUrl` and an `api` value at either provider or model level. `apiKey` is not required to load the file: models become available when auth is configured through `/login`/`auth.json`, CLI `--api-key`, or provider `apiKey`. If no auth is configured, the models load but stay unavailable in `/model` and `--list-models`.
145
+
146
+ ### Value Resolution
147
+
148
+ The `apiKey` and `headers` fields support command execution, environment interpolation, and literals:
149
+
150
+ - **Shell command:** `"!command"` at the start executes the whole value as a command and uses stdout
151
+ ```json
152
+ "apiKey": "!security find-generic-password -ws 'anthropic'"
153
+ "apiKey": "!op read 'op://vault/item/credential'"
154
+ ```
155
+ - **Environment interpolation:** `"$ENV_VAR"` or `"${ENV_VAR}"` uses the value of the named variable. Interpolation works inside larger literals.
156
+ ```json
157
+ "apiKey": "$MY_API_KEY"
158
+ "apiKey": "${KEY_PREFIX}_${KEY_SUFFIX}"
159
+ ```
160
+ `$FOO_BAR` is the variable `FOO_BAR`; use `${FOO}_BAR` when `BAR` is literal text. Missing environment variables make the value unresolved.
161
+ - **Escapes:** `"$$"` emits a literal `"$"`; `"$!"` emits a literal `"!"` without triggering command execution.
162
+ ```json
163
+ "apiKey": "$$literal-dollar-prefix"
164
+ "apiKey": "$!literal-bang-prefix"
165
+ ```
166
+ - **Literal value:** Used directly. Plain uppercase strings such as `MY_API_KEY` are literals; use `$MY_API_KEY` for environment variables.
167
+ ```json
168
+ "apiKey": "sk-..."
169
+ ```
170
+
171
+ For `models.json`, shell commands are resolved at request time. pi intentionally does not apply built-in TTL, stale reuse, or recovery logic for arbitrary commands. Different commands need different caching and failure strategies, and pi cannot infer the right one.
172
+
173
+ If your command is slow, expensive, rate-limited, or should keep using a previous value on transient failures, wrap it in your own script or command that implements the caching or TTL behavior you want.
174
+
175
+ `/model` availability checks use configured auth presence and do not execute shell commands.
176
+
177
+ ### Custom Headers
178
+
179
+ ```json
180
+ {
181
+ "providers": {
182
+ "custom-proxy": {
183
+ "baseUrl": "https://proxy.example.com/v1",
184
+ "apiKey": "$MY_API_KEY",
185
+ "api": "anthropic-messages",
186
+ "headers": {
187
+ "x-portkey-api-key": "$PORTKEY_API_KEY",
188
+ "x-secret": "!op read 'op://vault/item/secret'"
189
+ },
190
+ "models": [...]
191
+ }
192
+ }
193
+ }
194
+ ```
195
+
196
+ ## Model Configuration
197
+
198
+ | Field | Required | Default | Description |
199
+ |-------|----------|---------|-------------|
200
+ | `id` | Yes | — | Model identifier (passed to the API) |
201
+ | `name` | No | `id` | Human-readable model label. Used for matching (`--model` patterns) and shown as secondary model detail text. |
202
+ | `api` | No | provider's `api` | Override provider's API for this model |
203
+ | `reasoning` | No | `false` | Supports extended thinking |
204
+ | `thinkingLevelMap` | No | omitted | Maps pi thinking levels to provider values and marks unsupported levels (see below) |
205
+ | `input` | No | `["text"]` | Input types: `["text"]` or `["text", "image"]` |
206
+ | `contextWindow` | No | `128000` | Context window size in tokens |
207
+ | `maxTokens` | No | `16384` | Maximum output tokens |
208
+ | `cost` | No | all zeros | `{"input": 0, "output": 0, "cacheRead": 0, "cacheWrite": 0}` (per million tokens) |
209
+ | `compat` | No | provider `compat` | Provider compatibility overrides. Merged with provider-level `compat` when both are set. |
210
+
211
+ Current behavior:
212
+ - `/model`, `--list-models`, and the interactive footer display entries by model `id`.
213
+ - The configured `name` is used for model matching and secondary model detail text. It does not replace the footer/status-bar model id.
214
+
215
+ ### Thinking Level Map
216
+
217
+ Use `thinkingLevelMap` on a model to describe model-specific thinking controls. Keys are pi thinking levels: `off`, `minimal`, `low`, `medium`, `high`, `xhigh`.
218
+
219
+ Values are tristate:
220
+
221
+ | Value | Meaning |
222
+ |-------|---------|
223
+ | omitted | Level is supported and uses the provider's default mapping |
224
+ | string | Level is supported and this value is sent to the provider |
225
+ | `null` | Level is unsupported and hidden/skipped/clamped away |
226
+
227
+ Example for a model that only supports off, high, and max reasoning:
228
+
229
+ ```json
230
+ {
231
+ "id": "deepseek-v4-pro",
232
+ "reasoning": true,
233
+ "thinkingLevelMap": {
234
+ "minimal": null,
235
+ "low": null,
236
+ "medium": null,
237
+ "high": "high",
238
+ "xhigh": "max"
239
+ }
240
+ }
241
+ ```
242
+
243
+ Example for a model where thinking cannot be disabled:
244
+
245
+ ```json
246
+ {
247
+ "id": "always-thinking-model",
248
+ "reasoning": true,
249
+ "thinkingLevelMap": {
250
+ "off": null
251
+ }
252
+ }
253
+ ```
254
+
255
+ Migration: older configs that used `compat.reasoningEffortMap` should move that mapping to model-level `thinkingLevelMap`. Use `null` for levels that should not appear in the UI.
256
+
257
+ ## Overriding Built-in Providers
258
+
259
+ Route a built-in provider through a proxy without redefining models:
260
+
261
+ ```json
262
+ {
263
+ "providers": {
264
+ "anthropic": {
265
+ "baseUrl": "https://my-proxy.example.com/v1"
266
+ }
267
+ }
268
+ }
269
+ ```
270
+
271
+ All built-in Anthropic models remain available. Existing OAuth or API key auth continues to work.
272
+
273
+ To merge custom models into a built-in provider, include the `models` array:
274
+
275
+ ```json
276
+ {
277
+ "providers": {
278
+ "anthropic": {
279
+ "baseUrl": "https://my-proxy.example.com/v1",
280
+ "apiKey": "$ANTHROPIC_API_KEY",
281
+ "api": "anthropic-messages",
282
+ "models": [...]
283
+ }
284
+ }
285
+ }
286
+ ```
287
+
288
+ Merge semantics:
289
+ - Built-in models are kept.
290
+ - Custom models are upserted by `id` within the provider.
291
+ - If a custom model `id` matches a built-in model `id`, the custom model replaces that built-in model.
292
+ - If a custom model `id` is new, it is added alongside built-in models.
293
+
294
+ ## Per-model Overrides
295
+
296
+ Use `modelOverrides` to customize specific built-in models without replacing the provider's full model list.
297
+
298
+ ```json
299
+ {
300
+ "providers": {
301
+ "openrouter": {
302
+ "modelOverrides": {
303
+ "anthropic/claude-sonnet-4": {
304
+ "name": "Claude Sonnet 4 (Bedrock Route)",
305
+ "compat": {
306
+ "openRouterRouting": {
307
+ "only": ["amazon-bedrock"]
308
+ }
309
+ }
310
+ }
311
+ }
312
+ }
313
+ }
314
+ }
315
+ ```
316
+
317
+ `modelOverrides` supports these fields per model: `name`, `reasoning`, `input`, `cost` (partial), `contextWindow`, `maxTokens`, `headers`, `compat`.
318
+
319
+ Behavior notes:
320
+ - `modelOverrides` are applied to built-in provider models.
321
+ - Unknown model IDs are ignored.
322
+ - You can combine provider-level `baseUrl`/`headers` with `modelOverrides`.
323
+ - Overriding `name` changes model matching and secondary detail text only; the footer and primary model lists continue to show the model `id`.
324
+ - If `models` is also defined for a provider, custom models are merged after built-in overrides. A custom model with the same `id` replaces the overridden built-in model entry.
325
+
326
+ ## Anthropic Messages Compatibility
327
+
328
+ For providers or proxies using `api: "anthropic-messages"`, use `compat` to control Anthropic-specific request compatibility.
329
+
330
+ By default pi sends per-tool `eager_input_streaming: true`. If a proxy or Anthropic-compatible backend rejects that field, set `supportsEagerToolInputStreaming` to `false`. Pi will omit `tools[].eager_input_streaming` and send the legacy `fine-grained-tool-streaming-2025-05-14` beta header for tool-enabled requests instead.
331
+
332
+ Some Anthropic models require adaptive thinking (`thinking.type: "adaptive"` plus `output_config.effort`) instead of the legacy budget-based thinking payload. Built-in models set this automatically. For custom providers or aliases that route to those models, set `forceAdaptiveThinking` to `true`.
333
+
334
+ Some Anthropic-compatible providers emit thinking blocks with empty signatures and still expect them on replay. Set `allowEmptySignature` to `true` only for those providers; real Anthropic rejects empty thinking signatures.
335
+
336
+ Built-in Anthropic models enable `supportsStrictTools` in their model metadata. Custom Anthropic-compatible models must set it to `true` when their endpoint accepts strict JSON-schema tool definitions.
337
+
338
+ ```json
339
+ {
340
+ "providers": {
341
+ "anthropic-proxy": {
342
+ "baseUrl": "https://proxy.example.com",
343
+ "api": "anthropic-messages",
344
+ "apiKey": "$ANTHROPIC_PROXY_KEY",
345
+ "compat": {
346
+ "supportsEagerToolInputStreaming": false,
347
+ "supportsLongCacheRetention": true,
348
+ "forceAdaptiveThinking": true,
349
+ "allowEmptySignature": true
350
+ },
351
+ "models": [
352
+ {
353
+ "id": "claude-opus-4-7",
354
+ "reasoning": true,
355
+ "input": ["text", "image"]
356
+ }
357
+ ]
358
+ }
359
+ }
360
+ }
361
+ ```
362
+
363
+ | Field | Description |
364
+ |-------|-------------|
365
+ | `supportsEagerToolInputStreaming` | Whether the provider accepts per-tool `eager_input_streaming`. Default: `true`. Set to `false` to omit that field and use the legacy fine-grained tool streaming beta header on tool-enabled requests. |
366
+ | `supportsLongCacheRetention` | Whether the provider accepts Anthropic long cache retention (`cache_control.ttl: "1h"`) when cache retention is `long`. Default: `true`. |
367
+ | `sendSessionAffinityHeaders` | Whether to send `x-session-affinity` from the session id when caching is enabled. Default: auto-detected for known providers. |
368
+ | `supportsCacheControlOnTools` | Whether the provider accepts Anthropic-style `cache_control` markers on tool definitions. Default: `true`. |
369
+ | `forceAdaptiveThinking` | Whether to send adaptive thinking (`thinking.type: "adaptive"` plus `output_config.effort`) for this model. Built-in adaptive models set this automatically. Default: `false`. |
370
+ | `allowEmptySignature` | Whether to replay empty thinking signatures as `signature: ""` instead of converting thinking to text. Default: `false`. |
371
+ | `supportsStrictTools` | Whether the provider accepts strict JSON-schema tool definitions. Default: `false`; built-in Anthropic models enable it in generated metadata. |
372
+
373
+ ## OpenAI Compatibility
374
+
375
+ For providers with partial OpenAI compatibility, use the `compat` field.
376
+
377
+ - Provider-level `compat` applies defaults to all models under that provider.
378
+ - Model-level `compat` overrides provider-level values for that model.
379
+
380
+ ```json
381
+ {
382
+ "providers": {
383
+ "local-llm": {
384
+ "baseUrl": "http://localhost:8080/v1",
385
+ "api": "openai-completions",
386
+ "compat": {
387
+ "supportsUsageInStreaming": false,
388
+ "maxTokensField": "max_tokens"
389
+ },
390
+ "models": [...]
391
+ }
392
+ }
393
+ }
394
+ ```
395
+
396
+ | Field | Description |
397
+ |-------|-------------|
398
+ | `supportsStore` | Provider supports `store` field |
399
+ | `supportsDeveloperRole` | Use `developer` vs `system` role |
400
+ | `supportsReasoningEffort` | Support for `reasoning_effort` parameter |
401
+ | `supportsUsageInStreaming` | Supports `stream_options: { include_usage: true }` (default: `true`) |
402
+ | `maxTokensField` | Use `max_completion_tokens` or `max_tokens` |
403
+ | `requiresToolResultName` | Include `name` on tool result messages |
404
+ | `requiresAssistantAfterToolResult` | Insert an assistant message before a user message after tool results |
405
+ | `requiresThinkingAsText` | Convert thinking blocks to plain text |
406
+ | `requiresReasoningContentOnAssistantMessages` | Include empty `reasoning_content` on all replayed assistant messages when reasoning is enabled |
407
+ | `thinkingFormat` | Use `reasoning_effort`, `openrouter`, `deepseek`, `together`, `zai`, `qwen`, `chat-template`, or `qwen-chat-template` thinking parameters |
408
+ | `chatTemplateKwargs` | `chat_template_kwargs` values for `thinkingFormat: "chat-template"`; use `{ "$var": "thinking.enabled" }` or `{ "$var": "thinking.effort" }` for pi-controlled thinking values |
409
+ | `cacheControlFormat` | Use Anthropic-style `cache_control` markers on the system prompt, last tool definition, and last user, assistant, or tool-result text content. Currently only `anthropic` is supported. |
410
+ | `sendSessionAffinityHeaders` | For `openai-completions`, send session-affinity headers from the session id when caching is enabled. Default: `false`. |
411
+ | `sessionAffinityFormat` | For `openai-completions` and `openai-responses`, the session-affinity header format: `openai` sends `session_id`/`x-client-request-id` (completions also `x-session-affinity`), `openai-nosession` omits the underscore-containing `session_id` header, `openrouter` sends `x-session-id`. Does not affect the `prompt_cache_key` body param. Default: auto-detected. |
412
+ | `supportsStrictMode` | Whether the provider accepts strict JSON-schema function tool definitions. Defaults depend on the API; built-in OpenAI models carry explicit capability metadata. |
413
+ | `supportsOpenAIGrammarTools` | Whether OpenAI-compatible APIs emit custom Lark/regex grammar tools. When `false`, grammar-constrained tools fall back to normal function tools. Default: `false`; the built-in model catalog enables it for GPT-5+ models on OpenAI, OpenAI Codex, Azure OpenAI, GitHub Copilot, opencode, and Cloudflare AI Gateway. |
414
+ | `deferredToolsMode` | Use provider-specific deferred tool serialization. Currently only `"kimi"` is supported for Kimi's OpenAI-compatible Chat Completions format. |
415
+ | `supportsLongCacheRetention` | Whether the provider accepts long cache retention when cache retention is `long`: `prompt_cache_retention: "24h"` for OpenAI prompt caching, or `cache_control.ttl: "1h"` when `cacheControlFormat` is `anthropic`. Default: `true`. |
416
+ | `openRouterRouting` | OpenRouter provider routing preferences. This object is sent as-is in the `provider` field of the [OpenRouter API request](https://openrouter.ai/docs/guides/routing/provider-selection). |
417
+ | `vercelGatewayRouting` | Vercel AI Gateway routing config for provider selection (`only`, `order`) |
418
+
419
+ `openrouter` uses `reasoning: { effort }`. `together` uses `reasoning: { enabled }` and also `reasoning_effort` when `supportsReasoningEffort` is enabled. `qwen` uses top-level `enable_thinking`. Use `qwen-chat-template` for local Qwen-compatible servers that require `chat_template_kwargs.enable_thinking` and `preserve_thinking`. Use `chat-template` for vLLM/Hugging Face chat templates that need configurable `chat_template_kwargs`, such as `chatTemplateKwargs: { "thinking": { "$var": "thinking.enabled" } }` for DeepSeek V3.x templates.
420
+
421
+ `cacheControlFormat: "anthropic"` is for OpenAI-compatible providers that expose Anthropic-style prompt caching through `cache_control` markers on text content and tool definitions.
422
+
423
+ Example:
424
+
425
+ ```json
426
+ {
427
+ "providers": {
428
+ "openrouter": {
429
+ "baseUrl": "https://openrouter.ai/api/v1",
430
+ "apiKey": "$OPENROUTER_API_KEY",
431
+ "api": "openai-completions",
432
+ "models": [
433
+ {
434
+ "id": "openrouter/anthropic/claude-3.5-sonnet",
435
+ "name": "OpenRouter Claude 3.5 Sonnet",
436
+ "compat": {
437
+ "openRouterRouting": {
438
+ "allow_fallbacks": true,
439
+ "require_parameters": false,
440
+ "data_collection": "deny",
441
+ "zdr": true,
442
+ "enforce_distillable_text": false,
443
+ "order": ["anthropic", "amazon-bedrock", "google-vertex"],
444
+ "only": ["anthropic", "amazon-bedrock"],
445
+ "ignore": ["gmicloud", "friendli"],
446
+ "quantizations": ["fp16", "bf16"],
447
+ "sort": {
448
+ "by": "price",
449
+ "partition": "model"
450
+ },
451
+ "max_price": {
452
+ "prompt": 10,
453
+ "completion": 20
454
+ },
455
+ "preferred_min_throughput": {
456
+ "p50": 100,
457
+ "p90": 50
458
+ },
459
+ "preferred_max_latency": {
460
+ "p50": 1,
461
+ "p90": 3,
462
+ "p99": 5
463
+ }
464
+ }
465
+ }
466
+ }
467
+ ]
468
+ }
469
+ }
470
+ }
471
+ ```
472
+
473
+ Vercel AI Gateway example:
474
+
475
+ ```json
476
+ {
477
+ "providers": {
478
+ "vercel-ai-gateway": {
479
+ "baseUrl": "https://ai-gateway.vercel.sh/v1",
480
+ "apiKey": "$AI_GATEWAY_API_KEY",
481
+ "api": "openai-completions",
482
+ "models": [
483
+ {
484
+ "id": "moonshotai/kimi-k2.5",
485
+ "name": "Kimi K2.5 (Fireworks via Vercel)",
486
+ "reasoning": true,
487
+ "input": ["text", "image"],
488
+ "cost": { "input": 0.6, "output": 3, "cacheRead": 0, "cacheWrite": 0 },
489
+ "contextWindow": 262144,
490
+ "maxTokens": 262144,
491
+ "compat": {
492
+ "vercelGatewayRouting": {
493
+ "only": ["fireworks", "novita"],
494
+ "order": ["fireworks", "novita"]
495
+ }
496
+ }
497
+ }
498
+ ]
499
+ }
500
+ }
501
+ }
502
+ ```