@vellumai/assistant 0.8.12 → 0.9.0-staging.1

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (414) hide show
  1. package/AGENTS.md +0 -14
  2. package/ARCHITECTURE.md +45 -45
  3. package/README.md +1 -1
  4. package/bun.lock +200 -154
  5. package/docs/architecture/integrations.md +3 -3
  6. package/docs/architecture/memory.md +2 -2
  7. package/docs/architecture/security.md +10 -10
  8. package/docs/runbook-trusted-contacts.md +12 -12
  9. package/docs/skills.md +6 -6
  10. package/docs/workflows-testing.md +221 -0
  11. package/docs/workflows.md +510 -0
  12. package/examples/plugins/echo/README.md +5 -5
  13. package/knip.json +2 -0
  14. package/openapi.yaml +6935 -5545
  15. package/package.json +8 -4
  16. package/scripts/generate-openapi.ts +66 -114
  17. package/src/__tests__/access-request-seed-content-blocks.test.ts +213 -0
  18. package/src/__tests__/adaptive-thinking-repair.test.ts +32 -3
  19. package/src/__tests__/agent-loop-output-hooks.test.ts +183 -0
  20. package/src/__tests__/agent-loop-regrowth-guard.test.ts +506 -0
  21. package/src/__tests__/agent-wake-disk-pressure-callsite.test.ts +2 -0
  22. package/src/__tests__/agent-wake-override-profile.test.ts +77 -0
  23. package/src/__tests__/app-dir-path-guard.test.ts +27 -3
  24. package/src/__tests__/approval-cascade.test.ts +0 -5
  25. package/src/__tests__/approval-routes-http.test.ts +91 -0
  26. package/src/__tests__/assistant-stream-state.test.ts +107 -0
  27. package/src/__tests__/browser-fill-credential.test.ts +3 -3
  28. package/src/__tests__/bundled-skill-retrieval-guard.test.ts +1 -1
  29. package/src/__tests__/compaction-events.test.ts +63 -7
  30. package/src/__tests__/compaction-trail-store.test.ts +74 -1
  31. package/src/__tests__/compaction.benchmark.test.ts +63 -41
  32. package/src/__tests__/compactor-low-watermark-cut.test.ts +349 -0
  33. package/src/__tests__/context-window-manager-compact-retry.test.ts +64 -0
  34. package/src/__tests__/conversation-abort-tool-results.test.ts +0 -5
  35. package/src/__tests__/conversation-confirmation-signals.test.ts +0 -5
  36. package/src/__tests__/conversation-history-web-search.test.ts +7 -0
  37. package/src/__tests__/conversation-process-callsite.test.ts +0 -5
  38. package/src/__tests__/conversation-provider-retry-repair.test.ts +0 -5
  39. package/src/__tests__/conversation-queue.test.ts +0 -5
  40. package/src/__tests__/conversation-slash-queue.test.ts +0 -5
  41. package/src/__tests__/conversation-slash-unknown.test.ts +0 -5
  42. package/src/__tests__/conversation-speed-override.test.ts +0 -5
  43. package/src/__tests__/conversation-surfaces-task-progress.test.ts +67 -0
  44. package/src/__tests__/conversation-usage.test.ts +2 -0
  45. package/src/__tests__/conversation-workspace-injection.test.ts +0 -5
  46. package/src/__tests__/conversation-workspace-tool-tracking.test.ts +0 -5
  47. package/src/__tests__/credential-broker-browser-fill.test.ts +2 -2
  48. package/src/__tests__/credential-broker-server-use.test.ts +2 -2
  49. package/src/__tests__/credential-broker.test.ts +1 -1
  50. package/src/__tests__/credential-prompt-route.test.ts +417 -0
  51. package/src/__tests__/credential-security-invariants.test.ts +1 -0
  52. package/src/__tests__/credential-vault.test.ts +37 -0
  53. package/src/__tests__/db-schedule-syntax-migration.test.ts +24 -0
  54. package/src/__tests__/dynamic-page-surface.test.ts +125 -0
  55. package/src/__tests__/empty-state-greeting-cache.test.ts +94 -0
  56. package/src/__tests__/gateway-flag-listener.test.ts +24 -7
  57. package/src/__tests__/guardian-action-sweep.test.ts +56 -219
  58. package/src/__tests__/guardian-routing-invariants.test.ts +138 -0
  59. package/src/__tests__/helpers/channel-test-adapter.ts +0 -2
  60. package/src/__tests__/list-messages-hidden-metadata.test.ts +99 -0
  61. package/src/__tests__/llm-request-log-source-clickhouse.test.ts +87 -1
  62. package/src/__tests__/llm-resolver.test.ts +115 -0
  63. package/src/__tests__/managed-profile-guard.test.ts +6 -5
  64. package/src/__tests__/max-tokens-continue-hook.test.ts +184 -0
  65. package/src/__tests__/media-generate-image.test.ts +20 -9
  66. package/src/__tests__/model-intents.test.ts +1 -1
  67. package/src/__tests__/normalize-onboarding.test.ts +26 -0
  68. package/src/__tests__/notification-decision-strategy.test.ts +4 -2
  69. package/src/__tests__/notification-telegram-adapter.test.ts +21 -3
  70. package/src/__tests__/pending-interactions-resolved-event.test.ts +62 -0
  71. package/src/__tests__/post-turn-tool-result-truncation.test.ts +72 -18
  72. package/src/__tests__/require-fresh-approval.test.ts +425 -1
  73. package/src/__tests__/runtime-events-sse-parity.test.ts +2 -0
  74. package/src/__tests__/schedule-routes-workflow-validation.test.ts +408 -0
  75. package/src/__tests__/schedule-routes.test.ts +257 -4
  76. package/src/__tests__/schedule-store.test.ts +60 -0
  77. package/src/__tests__/schedule-tools.test.ts +247 -2
  78. package/src/__tests__/skill-execute-input.test.ts +85 -0
  79. package/src/__tests__/skill-secret-handling-guard.test.ts +21 -20
  80. package/src/__tests__/skills.test.ts +3 -3
  81. package/src/__tests__/slack-app-setup-skill-regression.test.ts +1 -1
  82. package/src/__tests__/subagent-tool-filtering.test.ts +50 -0
  83. package/src/__tests__/subagent-tool-gate-mode.test.ts +547 -0
  84. package/src/__tests__/system-prompt.test.ts +1 -1
  85. package/src/__tests__/task-progress-nudge-hook.test.ts +372 -0
  86. package/src/__tests__/task-scheduler.test.ts +299 -0
  87. package/src/__tests__/tool-approval-seed-content-blocks.test.ts +209 -0
  88. package/src/__tests__/tool-result-spool.test.ts +3 -1
  89. package/src/__tests__/workspace-migration-102-preserve-heartbeat-enabled-for-existing-workspaces.test.ts +181 -0
  90. package/src/__tests__/workspace-migration-103-upgrade-quality-profile-to-opus-4-8.test.ts +174 -0
  91. package/src/agent/compaction-circuit.ts +11 -0
  92. package/src/agent/loop.ts +181 -12
  93. package/src/api/constants/call-sites.ts +12 -0
  94. package/src/api/events/assistant-thinking-delta.ts +10 -0
  95. package/src/api/events/usage-update.ts +7 -0
  96. package/src/api/index.ts +4 -1
  97. package/src/api/responses/memory-v3-selection-log.ts +18 -11
  98. package/src/approvals/approval-primitive.ts +2 -2
  99. package/src/background-wake/background-wake-routes.test.ts +5 -2
  100. package/src/bundler/compiler-tools.ts +1 -1
  101. package/src/calls/call-domain.ts +1 -1
  102. package/src/calls/guardian-action-sweep.ts +16 -93
  103. package/src/cli/AGENTS.md +4 -0
  104. package/src/cli/commands/__tests__/schedules.test.ts +430 -1
  105. package/src/cli/commands/credentials.ts +28 -24
  106. package/src/cli/commands/image-generation.ts +23 -9
  107. package/src/cli/commands/notifications.ts +1 -1
  108. package/src/cli/commands/plugins.ts +89 -46
  109. package/src/cli/commands/schedules.ts +384 -11
  110. package/src/cli/lib/__tests__/inspect-plugin.test.ts +69 -5
  111. package/src/cli/lib/__tests__/install-from-github.test.ts +15 -0
  112. package/src/cli/lib/__tests__/upgrade-plugin.test.ts +81 -4
  113. package/src/cli/lib/inspect-plugin.ts +62 -1
  114. package/src/cli/lib/install-from-github.ts +52 -5
  115. package/src/cli/lib/upgrade-plugin.ts +18 -0
  116. package/src/config/__tests__/workflows-schema.test.ts +60 -0
  117. package/src/config/bundled-skills/acp/SKILL.md +2 -2
  118. package/src/config/bundled-skills/image-studio/SKILL.md +66 -19
  119. package/src/config/bundled-skills/image-studio/TOOLS.json +1 -6
  120. package/src/config/bundled-skills/image-studio/tools/media-generate-image.ts +22 -3
  121. package/src/config/bundled-skills/personal-page/SKILL.md +57 -0
  122. package/src/config/bundled-skills/personal-page/TOOLS.json +27 -0
  123. package/src/config/bundled-skills/personal-page/tools/app-refresh.ts +17 -0
  124. package/src/config/bundled-skills/schedule/SKILL.md +7 -2
  125. package/src/config/bundled-skills/schedule/TOOLS.json +48 -4
  126. package/src/config/bundled-skills/workflows/SKILL.md +214 -0
  127. package/src/config/bundled-skills/workflows/TOOLS.json +84 -0
  128. package/src/config/bundled-skills/workflows/tools/manage-workflows.ts +12 -0
  129. package/src/config/bundled-skills/workflows/tools/run-workflow.ts +12 -0
  130. package/src/config/bundled-tool-registry.ts +14 -2
  131. package/src/config/call-site-defaults.ts +5 -0
  132. package/src/config/feature-flag-registry.json +12 -4
  133. package/src/config/llm-context-resolution.ts +8 -0
  134. package/src/config/llm-resolver.ts +30 -0
  135. package/src/config/preloaded-apps/personal-page/src/components/About.tsx +22 -0
  136. package/src/config/preloaded-apps/personal-page/src/components/App.tsx +16 -0
  137. package/src/config/preloaded-apps/personal-page/src/components/Features.tsx +77 -0
  138. package/src/config/preloaded-apps/personal-page/src/components/Hero.tsx +57 -0
  139. package/src/config/preloaded-apps/personal-page/src/components/Pending.tsx +28 -0
  140. package/src/config/preloaded-apps/personal-page/src/components/animations.tsx +234 -0
  141. package/src/config/preloaded-apps/personal-page/src/components/icons.tsx +48 -0
  142. package/src/config/preloaded-apps/personal-page/src/components/media.ts +16 -0
  143. package/src/config/preloaded-apps/personal-page/src/index.html +20 -0
  144. package/src/config/preloaded-apps/personal-page/src/main.tsx +7 -0
  145. package/src/config/preloaded-apps/personal-page/src/profile-data.ts +82 -0
  146. package/src/config/preloaded-apps/personal-page/src/styles.css +759 -0
  147. package/src/config/schema.ts +2 -0
  148. package/src/config/schemas/call-site-catalog.ts +7 -0
  149. package/src/config/schemas/heartbeat.ts +4 -1
  150. package/src/config/schemas/llm.ts +33 -26
  151. package/src/config/schemas/memory-retrospective.ts +19 -0
  152. package/src/config/schemas/platform.ts +8 -0
  153. package/src/config/schemas/services.ts +5 -2
  154. package/src/config/schemas/workflows.ts +42 -0
  155. package/src/config/skills.ts +3 -3
  156. package/src/context/compactor.ts +273 -39
  157. package/src/context/post-turn-tool-result-truncation.ts +23 -6
  158. package/src/context/tool-result-spool.ts +12 -17
  159. package/src/credential-execution/executable-discovery.ts +1 -1
  160. package/src/credential-execution/process-manager.ts +37 -3
  161. package/src/credential-execution/prompted-credential.ts +205 -0
  162. package/src/daemon/conversation-agent-loop-handlers.ts +14 -0
  163. package/src/daemon/conversation-process.ts +11 -2
  164. package/src/daemon/conversation-surfaces.ts +75 -0
  165. package/src/daemon/conversation-tool-setup.ts +103 -26
  166. package/src/daemon/conversation-usage.ts +2 -0
  167. package/src/daemon/conversation.ts +115 -11
  168. package/src/daemon/handlers/shared.ts +26 -14
  169. package/src/daemon/host-cu-proxy.ts +15 -12
  170. package/src/daemon/host-file-proxy.ts +15 -12
  171. package/src/daemon/host-transfer-proxy.ts +30 -24
  172. package/src/daemon/lifecycle.ts +40 -3
  173. package/src/daemon/message-protocol.ts +3 -0
  174. package/src/daemon/message-types/messages.ts +2 -10
  175. package/src/daemon/message-types/workflows.ts +49 -0
  176. package/src/daemon/parse-actual-tokens-from-error.test.ts +62 -1
  177. package/src/daemon/parse-actual-tokens-from-error.ts +43 -4
  178. package/src/daemon/process-message.ts +6 -0
  179. package/src/daemon/tool-setup-types.ts +57 -0
  180. package/src/daemon/wake-conversation-ops.ts +18 -0
  181. package/src/heartbeat/heartbeat-run-store.ts +8 -2
  182. package/src/home/feed-types.ts +1 -1
  183. package/src/ipc/gateway-flag-listener.ts +28 -6
  184. package/src/mcp/mcp-auth-state.ts +8 -20
  185. package/src/media/__tests__/image-models.test.ts +57 -0
  186. package/src/media/image-models.ts +66 -0
  187. package/src/memory/__tests__/auto-analysis-enqueue.test.ts +38 -0
  188. package/src/memory/__tests__/find-most-recent-retrospective-for.test.ts +12 -2
  189. package/src/memory/__tests__/memory-retrospective-job.test.ts +911 -34
  190. package/src/memory/__tests__/memory-retrospective-startup-cleanup.test.ts +227 -5
  191. package/src/memory/__tests__/memory-retrospective-state.test.ts +195 -0
  192. package/src/memory/__tests__/preloaded-apps.test.ts +85 -0
  193. package/src/memory/auto-analysis-enqueue.ts +14 -1
  194. package/src/memory/compaction-log-store-clickhouse.ts +6 -4
  195. package/src/memory/conversation-crud.ts +9 -2
  196. package/src/memory/conversation-disk-view.ts +1 -1
  197. package/src/memory/conversation-queries.ts +22 -7
  198. package/src/memory/db-init.ts +20 -0
  199. package/src/memory/db-maintenance.ts +16 -0
  200. package/src/memory/embedding-runtime-manager.ts +1 -1
  201. package/src/memory/llm-request-log-source-clickhouse.ts +112 -14
  202. package/src/memory/llm-request-log-source-local.ts +19 -1
  203. package/src/memory/llm-request-log-source.ts +35 -6
  204. package/src/memory/llm-request-log-store.ts +90 -2
  205. package/src/memory/memory-retrospective-constants.ts +9 -0
  206. package/src/memory/memory-retrospective-enqueue.ts +3 -6
  207. package/src/memory/memory-retrospective-fork-boundary.ts +94 -0
  208. package/src/memory/memory-retrospective-job.ts +500 -208
  209. package/src/memory/memory-retrospective-startup-cleanup.ts +97 -19
  210. package/src/memory/memory-retrospective-state.ts +85 -2
  211. package/src/memory/migrations/281-memory-retrospective-remembered-log.ts +40 -0
  212. package/src/memory/migrations/282-schedule-inference-profile.test.ts +77 -0
  213. package/src/memory/migrations/282-schedule-inference-profile.ts +26 -0
  214. package/src/memory/migrations/283-memory-v3-selections-message-id-and-sections.test.ts +102 -0
  215. package/src/memory/migrations/283-memory-v3-selections-message-id-and-sections.ts +53 -0
  216. package/src/memory/migrations/284-workflow-runs.ts +51 -0
  217. package/src/memory/migrations/285-schedule-workflow-mode.ts +26 -0
  218. package/src/memory/migrations/286-workflow-run-trust.ts +27 -0
  219. package/src/memory/migrations/287-conversation-origin-channel-index.ts +15 -0
  220. package/src/memory/migrations/288-backfill-origin-channel-from-bindings.ts +43 -0
  221. package/src/memory/migrations/289-contact-channels-unique-ext-user.ts +115 -0
  222. package/src/memory/migrations/290-schedule-capabilities.test.ts +77 -0
  223. package/src/memory/migrations/290-schedule-capabilities.ts +25 -0
  224. package/src/memory/migrations/__tests__/281-memory-retrospective-remembered-log.test.ts +96 -0
  225. package/src/memory/migrations/__tests__/289-contact-channels-unique-ext-user.test.ts +571 -0
  226. package/src/memory/migrations/index.ts +10 -0
  227. package/src/memory/preloaded-apps.ts +116 -0
  228. package/src/memory/schema/infrastructure.ts +4 -0
  229. package/src/memory/schema/memory-core.ts +4 -0
  230. package/src/memory/v2/__tests__/concept-page-frontmatter-schema.test.ts +45 -0
  231. package/src/memory/v2/__tests__/frontmatter-sweep.test.ts +11 -7
  232. package/src/memory/v2/__tests__/page-store.test.ts +13 -2
  233. package/src/memory/v2/__tests__/qdrant.test.ts +24 -0
  234. package/src/memory/v2/frontmatter-sweep.ts +7 -6
  235. package/src/memory/v2/page-store.ts +4 -3
  236. package/src/memory/v2/qdrant.ts +42 -3
  237. package/src/memory/v2/types.ts +16 -10
  238. package/src/messaging/draft-store.ts +1 -1
  239. package/src/notifications/access-request-copy.ts +200 -113
  240. package/src/notifications/adapters/slack.ts +250 -111
  241. package/src/notifications/adapters/telegram.ts +7 -44
  242. package/src/notifications/approval-card-builder.ts +93 -0
  243. package/src/notifications/broadcaster.ts +74 -0
  244. package/src/notifications/conversation-pairing.ts +8 -6
  245. package/src/notifications/copy-composer.ts +32 -26
  246. package/src/notifications/decision-engine.ts +59 -7
  247. package/src/notifications/guardian-question-mode.ts +145 -155
  248. package/src/notifications/home-feed-side-effect.ts +1 -1
  249. package/src/notifications/notification-utils.ts +66 -0
  250. package/src/notifications/signal.ts +6 -0
  251. package/src/notifications/tool-approval-copy.ts +142 -0
  252. package/src/notifications/types.ts +19 -0
  253. package/src/permissions/threshold.ts +11 -0
  254. package/src/plugin-api/types.ts +16 -4
  255. package/src/plugins/defaults/compaction/window-manager.ts +44 -0
  256. package/src/plugins/defaults/index.ts +46 -0
  257. package/src/plugins/defaults/max-tokens-continue/continue-state-store.ts +53 -0
  258. package/src/plugins/defaults/max-tokens-continue/hooks/post-model-call.ts +80 -0
  259. package/src/plugins/defaults/max-tokens-continue/hooks/stop.ts +20 -0
  260. package/src/plugins/defaults/max-tokens-continue/package.json +14 -0
  261. package/src/plugins/defaults/memory-retrieval/hooks/__tests__/user-prompt-submit.test.ts +37 -0
  262. package/src/plugins/defaults/memory-retrieval/hooks/user-prompt-submit.ts +29 -1
  263. package/src/plugins/defaults/memory-v3-shadow/__tests__/carry-integration.test.ts +8 -3
  264. package/src/plugins/defaults/memory-v3-shadow/__tests__/injection.test.ts +4 -2
  265. package/src/plugins/defaults/memory-v3-shadow/__tests__/pool-select.test.ts +28 -18
  266. package/src/plugins/defaults/memory-v3-shadow/__tests__/section-dense-store.test.ts +67 -0
  267. package/src/plugins/defaults/memory-v3-shadow/__tests__/selection-log-store.test.ts +122 -22
  268. package/src/plugins/defaults/memory-v3-shadow/__tests__/shadow-integration.test.ts +2 -0
  269. package/src/plugins/defaults/memory-v3-shadow/__tests__/shadow-plugin.test.ts +63 -1
  270. package/src/plugins/defaults/memory-v3-shadow/injector.ts +61 -18
  271. package/src/plugins/defaults/memory-v3-shadow/orchestrate.ts +1 -1
  272. package/src/plugins/defaults/memory-v3-shadow/pool-select.ts +39 -10
  273. package/src/plugins/defaults/memory-v3-shadow/section-dense-store.ts +34 -1
  274. package/src/plugins/defaults/memory-v3-shadow/selection-log-store.ts +112 -47
  275. package/src/plugins/defaults/memory-v3-shadow/shadow-plugin.ts +78 -15
  276. package/src/plugins/defaults/task-progress-nudge/hooks/post-tool-use.ts +206 -0
  277. package/src/plugins/defaults/task-progress-nudge/package.json +15 -0
  278. package/src/prompts/__tests__/system-prompt.test.ts +100 -1
  279. package/src/prompts/__tests__/task-progress-hint-section.test.ts +5 -7
  280. package/src/prompts/normalize-onboarding.ts +2 -0
  281. package/src/prompts/persona-resolver.ts +3 -0
  282. package/src/prompts/system-prompt.ts +51 -2
  283. package/src/prompts/templates/BOOTSTRAP-ACTIVATION-RAIL.md +3 -1
  284. package/src/prompts/templates/system-sections.ts +8 -3
  285. package/src/providers/call-site-routing.ts +6 -3
  286. package/src/providers/fireworks/client.ts +3 -0
  287. package/src/providers/inference/auth.ts +52 -46
  288. package/src/providers/model-intents.ts +2 -2
  289. package/src/providers/openai/__tests__/coerce-object-args.test.ts +105 -0
  290. package/src/providers/openai/chat-completions-provider.ts +47 -9
  291. package/src/providers/openai/coerce-object-args.ts +104 -0
  292. package/src/providers/retry.ts +8 -5
  293. package/src/providers/types.ts +10 -0
  294. package/src/runtime/__tests__/agent-wake.test.ts +629 -7
  295. package/src/runtime/access-request-helper.ts +16 -9
  296. package/src/runtime/agent-wake.ts +302 -51
  297. package/src/runtime/assistant-stream-state.ts +141 -8
  298. package/src/runtime/background-job-runner.ts +9 -0
  299. package/src/runtime/channel-approval-types.ts +1 -0
  300. package/src/runtime/guardian-action-message-composer.ts +0 -54
  301. package/src/runtime/http-server.ts +6 -14
  302. package/src/runtime/http-types.ts +0 -1
  303. package/src/runtime/message-composer-types.ts +0 -9
  304. package/src/runtime/middleware/__tests__/rate-limiter.test.ts +63 -0
  305. package/src/runtime/middleware/auth.ts +27 -3
  306. package/src/runtime/middleware/rate-limiter.ts +28 -1
  307. package/src/runtime/migrations/vbundle-builder.ts +6 -5
  308. package/src/runtime/pending-interactions.ts +20 -1
  309. package/src/runtime/routes/__tests__/conversation-compaction-routes.test.ts +232 -173
  310. package/src/runtime/routes/__tests__/plugins-routes.test.ts +18 -0
  311. package/src/runtime/routes/__tests__/retrospective-routes.test.ts +436 -0
  312. package/src/runtime/routes/__tests__/surface-action-routes.test.ts +11 -0
  313. package/src/runtime/routes/approval-routes.ts +35 -8
  314. package/src/runtime/routes/approval-strategies/guardian-callback-strategy.ts +2 -1
  315. package/src/runtime/routes/btw-routes.ts +37 -1
  316. package/src/runtime/routes/channel-route-shared.ts +3 -1
  317. package/src/runtime/routes/consolidation-routes.ts +17 -13
  318. package/src/runtime/routes/conversation-compaction-routes.ts +159 -119
  319. package/src/runtime/routes/conversation-list-routes.ts +41 -4
  320. package/src/runtime/routes/conversation-query-routes.ts +196 -11
  321. package/src/runtime/routes/conversation-routes.ts +15 -1
  322. package/src/runtime/routes/credential-prompt-routes.ts +39 -17
  323. package/src/runtime/routes/empty-state-greeting-cache.ts +65 -0
  324. package/src/runtime/routes/heartbeat-routes.ts +18 -13
  325. package/src/runtime/routes/image-generation-routes.ts +20 -2
  326. package/src/runtime/routes/inbound-stages/acl-enforcement.ts +26 -1
  327. package/src/runtime/routes/index.ts +6 -0
  328. package/src/runtime/routes/log-export/AGENTS.md +1 -1
  329. package/src/runtime/routes/log-export/workspace-allowlist.ts +1 -1
  330. package/src/runtime/routes/migration-routes.ts +5 -9
  331. package/src/runtime/routes/plugins-routes.ts +26 -0
  332. package/src/runtime/routes/ps-routes.ts +10 -8
  333. package/src/runtime/routes/retrospective-routes.ts +235 -0
  334. package/src/runtime/routes/runs-pagination.ts +75 -0
  335. package/src/runtime/routes/schedule-routes.ts +367 -59
  336. package/src/runtime/routes/sounds-config-routes.ts +239 -0
  337. package/src/runtime/routes/surface-action-routes.ts +84 -4
  338. package/src/runtime/routes/workflow-routes.test.ts +372 -0
  339. package/src/runtime/routes/workflow-routes.ts +363 -0
  340. package/src/runtime/routes/workspace-routes.test.ts +61 -1
  341. package/src/runtime/routes/workspace-routes.ts +26 -1
  342. package/src/runtime/services/__tests__/analyze-conversation.test.ts +38 -0
  343. package/src/runtime/services/analyze-conversation.ts +26 -13
  344. package/src/schedule/inference-profile.ts +28 -0
  345. package/src/schedule/schedule-store.ts +122 -4
  346. package/src/schedule/scheduler-types.ts +6 -0
  347. package/src/schedule/scheduler.ts +96 -0
  348. package/src/security/secret-allowlist.ts +1 -1
  349. package/src/skills/path-classifier.ts +1 -1
  350. package/src/tools/browser/browser-execution.ts +9 -11
  351. package/src/tools/credentials/broker.ts +4 -4
  352. package/src/tools/credentials/vault.ts +26 -137
  353. package/src/tools/executor.ts +69 -0
  354. package/src/tools/flag-gated-tools.test.ts +76 -0
  355. package/src/tools/permission-checker.ts +8 -1
  356. package/src/tools/registry.ts +51 -0
  357. package/src/tools/schedule/create.ts +77 -1
  358. package/src/tools/schedule/list.ts +1 -0
  359. package/src/tools/schedule/update.ts +73 -1
  360. package/src/tools/skills/execute.ts +56 -0
  361. package/src/tools/terminal/shell.ts +1 -1
  362. package/src/tools/ui-surface/definitions.ts +48 -2
  363. package/src/tools/workflows/manage-workflows.ts +183 -0
  364. package/src/tools/workflows/run-workflow.test.ts +442 -0
  365. package/src/tools/workflows/run-workflow.ts +88 -0
  366. package/src/types/onboarding-context.ts +2 -0
  367. package/src/usage/attribution.ts +24 -0
  368. package/src/util/canonicalize-identity.ts +12 -3
  369. package/src/util/platform.ts +17 -17
  370. package/src/watcher/__tests__/engine.test.ts +24 -0
  371. package/src/watcher/__tests__/telemetry.test.ts +135 -0
  372. package/src/watcher/engine.ts +7 -0
  373. package/src/watcher/telemetry.ts +74 -0
  374. package/src/workflows/capabilities.test.ts +365 -0
  375. package/src/workflows/capabilities.ts +359 -0
  376. package/src/workflows/deterministic-stringify.ts +27 -0
  377. package/src/workflows/engine-integration.test.ts +656 -0
  378. package/src/workflows/engine.test.ts +1144 -0
  379. package/src/workflows/engine.ts +1078 -0
  380. package/src/workflows/fanout-load.test.ts +168 -0
  381. package/src/workflows/journal-store.test.ts +369 -0
  382. package/src/workflows/journal-store.ts +470 -0
  383. package/src/workflows/leaf-runner.test.ts +704 -0
  384. package/src/workflows/leaf-runner.ts +589 -0
  385. package/src/workflows/library.test.ts +134 -0
  386. package/src/workflows/library.ts +124 -0
  387. package/src/workflows/run-manager.test.ts +711 -0
  388. package/src/workflows/run-manager.ts +593 -0
  389. package/src/workflows/sandbox-escape.test.ts +339 -0
  390. package/src/workflows/sandbox.test.ts +251 -0
  391. package/src/workflows/sandbox.ts +447 -0
  392. package/src/workspace/adaptive-thinking-repair.ts +33 -11
  393. package/src/workspace/migrations/021-move-signals-to-workspace.ts +1 -1
  394. package/src/workspace/migrations/022-move-hooks-to-workspace.ts +1 -1
  395. package/src/workspace/migrations/026-backfill-install-meta.ts +1 -1
  396. package/src/workspace/migrations/030-seed-pkb-autoinject.ts +2 -1
  397. package/src/workspace/migrations/031-drop-user-md.ts +1 -4
  398. package/src/workspace/migrations/048-remove-workspace-hooks.ts +1 -1
  399. package/src/workspace/migrations/056-release-notes-inference-profile-reordering.ts +5 -2
  400. package/src/workspace/migrations/061-move-backup-key-to-workspace.ts +1 -1
  401. package/src/workspace/migrations/082-backfill-managed-profile-labels.ts +8 -2
  402. package/src/workspace/migrations/097-enable-adaptive-thinking-managed-profiles.ts +41 -14
  403. package/src/workspace/migrations/102-preserve-heartbeat-enabled-for-existing-workspaces.ts +69 -0
  404. package/src/workspace/migrations/103-upgrade-quality-profile-to-opus-4-8.ts +83 -0
  405. package/src/workspace/migrations/104-recheck-adaptive-thinking-model-implied-anthropic.ts +133 -0
  406. package/src/workspace/migrations/registry.ts +6 -0
  407. package/src/workspace/migrations/runner.ts +1 -1
  408. package/tsconfig.json +1 -1
  409. package/src/__tests__/guardian-action-copy-generator.test.ts +0 -200
  410. package/src/__tests__/guardian-action-grant-mint-consume.test.ts +0 -579
  411. package/src/__tests__/guardian-action-store.test.ts +0 -106
  412. package/src/daemon/guardian-action-generators.ts +0 -71
  413. package/src/memory/guardian-action-store.ts +0 -484
  414. package/src/runtime/guardian-action-grant-minter.ts +0 -150
@@ -55,7 +55,7 @@ The `claude-agent-acp` adapter requires a Claude OAuth token. Store it once in t
55
55
  assistant credentials set --service acp --field claude_oauth_token <token>
56
56
  ```
57
57
 
58
- When the token is missing, do NOT ask the user to paste it into chat. Collect it via the secret-request flow instead: `credential_store` with action `prompt`, service `acp`, field `claude_oauth_token`. That prompts the user through a secure UI so the token never enters the conversation or the workspace config. Users generate the token by running `claude setup-token` on a machine where they are logged in to Claude.
58
+ When the token is missing, do NOT ask the user to paste it into chat. Collect it via the secure prompt instead: `assistant credentials prompt --service acp --field claude_oauth_token --label "Claude OAuth Token"`. That prompts the user through a secure UI so the token never enters the conversation or the workspace config. Users generate the token by running `claude setup-token` on a machine where they are logged in to Claude.
59
59
 
60
60
  ## Codex setup
61
61
 
@@ -78,7 +78,7 @@ Gemini CLI speaks ACP natively (`gemini --acp`) - there is no separate adapter b
78
78
  assistant credentials set --service acp --field gemini_api_key <key>
79
79
  ```
80
80
 
81
- Or collect the key via the secret-request flow: `credential_store` with action `prompt`, service `acp`, field `gemini_api_key`. Either way the key never appears in chat or workspace config, and every spawn injects it as `GEMINI_API_KEY` automatically. The key is optional - a spawn proceeds without it when the vault has no entry.
81
+ Or collect the key via the secure prompt: `assistant credentials prompt --service acp --field gemini_api_key --label "Gemini API Key"`. Either way the key never appears in chat or workspace config, and every spawn injects it as `GEMINI_API_KEY` automatically. The key is optional - a spawn proceeds without it when the vault has no entry.
82
82
 
83
83
  The alternative is browser OAuth: run `gemini` once interactively and complete the sign-in flow. This is impractical on hosted assistants (no browser), so prefer the credential store there.
84
84
 
@@ -1,6 +1,6 @@
1
1
  ---
2
2
  name: image-studio
3
- description: Generate and edit images using AI
3
+ description: Create images from a text description, or edit photos and graphics the user provides (remove backgrounds or watermarks, retouch, restyle, in-paint). Can produce multiple variants when the user wants options to choose from.
4
4
  compatibility: "Designed for Vellum personal assistants"
5
5
  metadata:
6
6
  emoji: "🎨"
@@ -9,36 +9,83 @@ metadata:
9
9
  category: "content"
10
10
  activation-hints:
11
11
  - "User asks to generate, draw, or create an image from a text prompt"
12
- - "User wants to edit an existing image — background removal, in-painting, style change, retouching"
12
+ - "User wants to edit an existing image: background removal, watermark removal, in-painting, style change, retouching"
13
13
  - "User wants multiple variations of a visual (logo concepts, mood boards, illustration options)"
14
14
  ---
15
15
 
16
- You are an image generation assistant. When the user asks you to create or edit images, use the `media_generate_image` tool.
17
-
18
- ## Usage
19
-
20
- - **Text-to-image**: "Generate an image of a sunset over the ocean"
21
- - **Image editing**: "Remove the background from this image" (requires providing source image file paths)
22
- - **Multiple variants**: "Generate 3 variations of a logo for a coffee shop"
16
+ Use the `media_generate_image` tool via `skill_execute` to create or edit images.
23
17
 
24
18
  ## Modes
25
19
 
26
20
  - **generate** (default): Create a new image from a text prompt.
27
- - **edit**: Modify an existing image based on a text prompt. Requires one or more source images via `source_paths` (file paths on disk).
21
+ - **edit**: Modify an existing image. Requires one or more source images via `source_paths`.
28
22
 
29
23
  ## Models
30
24
 
31
- - `gemini-3.1-flash-image-preview` (default) - Nano Banana 2, fast, good quality
32
- - `gemini-3-pro-image-preview` - Nano Banana Pro, higher quality, slower
33
- - `gpt-image-2` - OpenAI GPT Image 2, high fidelity, slower
25
+ Do not pass the `model` parameter unless you need a specific tier. Omitting it uses the configured default, which is correct for most requests.
26
+
27
+ When you do need to choose, use an alias, not a concrete model ID. Aliases always resolve to the current model for that tier:
28
+
29
+ - `fast`: quickest, good quality (default tier)
30
+ - `quality`: higher fidelity, slower
31
+ - `openai`: OpenAI's model; most permissive on photo edits
32
+
33
+ Pass a concrete model ID only if the user names one explicitly. If the tool rejects an unknown model ID, the error lists the currently available models and aliases.
34
+
35
+ ## Example calls
36
+
37
+ Generate (no model parameter, default is correct):
38
+
39
+ ```json
40
+ { "tool": "media_generate_image", "input": { "prompt": "A sunset over the ocean, golden hour, soft haze, 35mm photo style", "variants": 2 } }
41
+ ```
42
+
43
+ Edit:
44
+
45
+ ```json
46
+ { "tool": "media_generate_image", "input": { "prompt": "Remove the watermark text from the background. Keep the subject, framing, lighting, and colors exactly identical. Change nothing else.", "mode": "edit", "source_paths": ["conversations/<conv-id>/attachments/photo.jpeg"], "model": "openai" } }
47
+ ```
48
+
49
+ `source_paths` is a flat array of file path strings. Do NOT pass objects:
50
+
51
+ - Wrong: `"source_paths": [{ "path": "img.jpeg" }]` → schema validation error
52
+ - Right: `"source_paths": ["img.jpeg"]`
34
53
 
35
- ## Tips
54
+ ## Source images for edit mode
36
55
 
37
- - Be descriptive in your prompts for better results. Include details about style, composition, lighting, and mood.
38
- - When editing images, clearly describe what changes you want made to the source image.
39
- - Use the `variants` parameter (1-4) to generate multiple options and pick the best one.
40
- - If no API key is configured for the selected model's provider (Gemini or OpenAI), the tool will return an error - ask the user to set one up.
56
+ - Paths resolve inside the workspace. Conversation attachments live under `conversations/<conversation-id>/attachments/`; prefer that path for images the user attached.
57
+ - Host paths (e.g. `~/Desktop/photo.jpg`) only work if the file arrived as an attachment; the tool falls back to the stored workspace copy. If the user references a host file that was never attached, pull it into the workspace first, then pass the workspace path.
58
+
59
+ ## Prompting
60
+
61
+ - Generate: describe style, composition, lighting, and mood, not just the subject.
62
+ - Edit: name the change AND what must stay the same. Models re-render the whole image, so without preservation language ("keep subject, framing, and lighting identical; only change X") they drift on crop and color.
63
+ - Aspect ratio and size have no parameter today. State them in the prompt ("16:9 widescreen banner") and verify the output.
64
+ - Use `variants` (1 to 4) when the user wants options. In **edit mode always use `variants: 1`**: edits run 60-90 seconds per variant, and two variants can exceed the tool execution timeout (`timeouts.toolExecutionTimeoutSec`, default 120s). If the user wants multiple edit options, make separate sequential calls.
65
+
66
+ ## Timing
67
+
68
+ Edits on large photos are slow (1 to 2 minutes). If the tool reports a timeout ("timed out after Ns"), the result is lost; do not wait for it to appear. Retry with `variants: 1`, or if it already was 1, fall back to the CLI which writes files to disk: `assistant image-generation generate --prompt "..." --mode edit --source <path> --model openai --output-dir <dir>`.
69
+
70
+ ## Output handling
71
+
72
+ Images return as inline content blocks in the tool result; they are not written to disk automatically.
73
+
74
+ - If the user just wants to see the image, the inline result is enough.
75
+ - If the user wants a file or you need to iterate on the result, save it to disk and deliver it through the conversation's attachment mechanism.
41
76
 
42
77
  ## Error handling
43
78
 
44
- When image generation fails, report the error to the user as-is. **Do not** attempt to fix the error by changing service configuration (e.g. switching between "managed" and "your-own" mode, or changing the provider/model). Service configuration changes should only be made at the user's explicit request via Settings.
79
+ Two kinds of failure. Treat them differently:
80
+
81
+ 1. **Configuration errors** (missing API key, provider not set up): report the error to the user as-is. Do NOT change service configuration (managed vs your-own mode, default provider, or default model in Settings). Configuration changes happen only at the user's explicit request.
82
+ 2. **Generation failures** (any other error: "invalid", content policy, safety rejection, provider error). Do not diagnose the cause; switch providers. The error message names the model that failed:
83
+ - If the error names a `gemini-*` model (or no model) → retry ONCE with `model: "openai"`.
84
+ - If the error names `gpt-image-2` or another `gpt-*` model → retry ONCE with `model: "quality"`.
85
+ - If the retry also fails → stop and report both errors to the user.
86
+
87
+ Do NOT rephrase the prompt and retry on the same model, even if the error suggests checking the prompt. One provider switch, then stop.
88
+
89
+ ## Complete when
90
+
91
+ The tool has returned at least one image and the user can see it: either the inline result in chat or an attached saved file. An error report counts as complete only after the retry path in Error handling has been exhausted.
@@ -27,12 +27,7 @@
27
27
  },
28
28
  "model": {
29
29
  "type": "string",
30
- "enum": [
31
- "gemini-3.1-flash-image-preview",
32
- "gemini-3-pro-image-preview",
33
- "gpt-image-2"
34
- ],
35
- "description": "Which model to use for generation. If omitted, uses the user's configured preference."
30
+ "description": "Model tier alias (fast | quality | openai) or a concrete model ID. If omitted, uses the user's configured preference. Unknown values return an error listing currently available models."
36
31
  },
37
32
  "variants": {
38
33
  "type": "number",
@@ -1,5 +1,9 @@
1
1
  import { getConfig } from "../../../../config/loader.js";
2
2
  import { resolveImageGenCredentials } from "../../../../media/image-credentials.js";
3
+ import {
4
+ describeImageModels,
5
+ resolveImageModel,
6
+ } from "../../../../media/image-models.js";
3
7
  import {
4
8
  generateImage,
5
9
  mapImageGenError,
@@ -19,7 +23,20 @@ export async function run(
19
23
  ): Promise<ToolExecutionResult> {
20
24
  const config = getConfig();
21
25
  const svc = config.services["image-generation"];
22
- const modelOverride = input.model;
26
+ let modelOverride = input.model;
27
+ // Resolve tier aliases (fast, quality, openai) to concrete model IDs via
28
+ // the registry. Unknown values get an error listing the current catalog so
29
+ // callers can self-correct without a stale schema enum.
30
+ if (typeof modelOverride === "string" && modelOverride) {
31
+ const entry = resolveImageModel(modelOverride);
32
+ if (!entry) {
33
+ return {
34
+ content: `Unknown model "${modelOverride}". Available models and aliases:\n${describeImageModels()}\n\nRetry with one of the aliases above, or omit the model parameter to use the configured default.`,
35
+ isError: true,
36
+ };
37
+ }
38
+ modelOverride = entry.id;
39
+ }
23
40
  // Derive provider from the explicit model when supplied so that requesting
24
41
  // e.g. `gpt-image-2` while config.provider === "gemini" routes to OpenAI
25
42
  // instead of silently falling back to the Gemini default model.
@@ -30,7 +47,7 @@ export async function run(
30
47
  });
31
48
  if (!credentials) {
32
49
  return {
33
- content: `${errorHint ?? "Image generation is not configured."}\n\nReport this error to the user. Do not change service configuration (mode, provider, or model) to try to fix it.`,
50
+ content: `${errorHint ?? "Image generation is not configured."}\n\nReport this error to the user as-is. Do not change service configuration (managed/your-own mode or default provider/model settings) to try to fix it.`,
34
51
  isError: true,
35
52
  };
36
53
  }
@@ -130,8 +147,10 @@ export async function run(
130
147
  contentBlocks,
131
148
  };
132
149
  } catch (error) {
150
+ // Echo the model that failed so callers (including the skill's retry
151
+ // branch) can key off the error text instead of remembering their input.
133
152
  return {
134
- content: `${mapImageGenError(provider, error)}\n\nReport this error to the user. Do not change service configuration (mode, provider, or model) to try to fix it.`,
153
+ content: `${mapImageGenError(provider, error)}\n\nFailed model: ${model}\n\nDo not change service configuration (managed/your-own mode or default provider/model settings) to try to fix it. Retrying this call once with a different model parameter is allowed; follow the skill's error handling instructions.`,
135
154
  isError: true,
136
155
  };
137
156
  }
@@ -0,0 +1,57 @@
1
+ ---
2
+ name: personal-page
3
+ description: Build or update the user's personal page — a cinematic landing page about THE USER themselves (their bio, work, life areas). Use whenever the user asks for a personal page, a page about them, or to update what's on their page. NOT for general landing pages or other apps (that's app-builder). The page already exists in the workspace; this skill is how to fill it.
4
+ metadata:
5
+ emoji: "✨"
6
+ vellum:
7
+ display-name: "Personal Page"
8
+ category: "content"
9
+ activation-hints:
10
+ - "User asks for a personal page, profile page, or landing page about themselves"
11
+ - "User asks to change, refresh, or add something to their personal page"
12
+ - "Onboarding asks you to turn research about the user into a page"
13
+ avoid-when:
14
+ - "User wants a landing page for a product, company, or someone else — use app-builder"
15
+ ---
16
+
17
+ You fill in the user's personal page — a finished, cinematic dark landing page
18
+ that already lives in the workspace as the app `personal-page`. Your job is
19
+ **content, not construction**: the design is done.
20
+
21
+ If `/workspace/data/apps/personal-page/` does not exist, this skill does not
22
+ apply (the preseeded page is rolled out gradually) — build whatever the user
23
+ asked for with app-builder instead.
24
+
25
+ **Present it as yours.** To the user this is "the page you filled in for
26
+ them" — fine to acknowledge the page shell exists, but never narrate the
27
+ mechanics below (file paths, refresh calls, contracts) or your progress in
28
+ chat. Do the work, then one short line.
29
+
30
+ ## How to populate it
31
+
32
+ 1. Read `/workspace/data/apps/personal-page/src/profile-data.ts`. The contract
33
+ comment at the top explains every field and shows a worked example.
34
+ 2. Overwrite the `profile` export with what you know or learned about the
35
+ user. Set `status: "ready"`. Touch ONLY this file — never the layout,
36
+ styles, animations, or media. Do NOT use `app_create`; the app exists.
37
+ 3. Call `app_refresh` (app_id: `personal-page`). If it reports compile errors,
38
+ fix your profile-data.ts edit and refresh again. **`app_refresh` is the
39
+ only build step — never compile manually (no esbuild, bundlers, or build
40
+ commands in the terminal); the platform's compiler has the correct JSX
41
+ setup and a manual build will produce a broken page.**
42
+ 4. As soon as the refresh is clean, call `app_open` (app_id: `personal-page`)
43
+ so the finished page is up and waiting — don't wait to be asked.
44
+
45
+ ## Writing the content
46
+
47
+ - Facts beat adjectives: names, numbers, places. One fact per bullet.
48
+ - Only what you actually verified; write around gaps rather than inventing.
49
+ - Third person, warm but grounded. The page should feel like a profile in a
50
+ design magazine, not a résumé.
51
+
52
+ ## Later edits
53
+
54
+ - Content tweaks ("add my marathon PR", "fix my title") → same flow: edit
55
+ `profile-data.ts`, `app_refresh`.
56
+ - Only if the user explicitly asks to change the **design or layout** may you
57
+ edit the other source files, following app-builder conventions.
@@ -0,0 +1,27 @@
1
+ {
2
+ "version": 1,
3
+ "tools": [
4
+ {
5
+ "name": "app_refresh",
6
+ "description": "Trigger recompilation and surface refresh for an app after making file changes. Call this once after you finish editing app files with generic file tools.",
7
+ "category": "apps",
8
+ "risk": "low",
9
+ "input_schema": {
10
+ "type": "object",
11
+ "properties": {
12
+ "app_id": {
13
+ "type": "string",
14
+ "description": "The ID of the app to refresh"
15
+ },
16
+ "change_summary": {
17
+ "type": "string",
18
+ "description": "Short summary of what changed, using git conventional commit format (e.g. 'feat: add dark mode toggle', 'fix: form validation'). Used as the version history entry."
19
+ }
20
+ },
21
+ "required": ["app_id"]
22
+ },
23
+ "executor": "tools/app-refresh.ts",
24
+ "execution_target": "host"
25
+ }
26
+ ]
27
+ }
@@ -0,0 +1,17 @@
1
+ import { setAppCommitMessage } from "../../../../memory/app-git-service.js";
2
+ import * as appStore from "../../../../memory/app-store.js";
3
+ import { executeAppRefresh } from "../../../../tools/apps/executors.js";
4
+ import type {
5
+ ToolContext,
6
+ ToolExecutionResult,
7
+ } from "../../../../tools/types.js";
8
+
9
+ export async function run(
10
+ input: Record<string, unknown>,
11
+ context: ToolContext,
12
+ ): Promise<ToolExecutionResult> {
13
+ if (typeof input.change_summary === "string" && input.change_summary.trim()) {
14
+ setAppCommitMessage(context.conversationId, input.change_summary.trim());
15
+ }
16
+ return executeAppRefresh({ app_id: input.app_id as string }, appStore);
17
+ }
@@ -17,7 +17,7 @@ metadata:
17
17
  - "User wants to act immediately or run a quick command that completes within the conversation — schedule is only for deferred or recurring execution"
18
18
  ---
19
19
 
20
- Manage scheduled automations. Schedules can be **recurring** (cron or RRULE expression) or **one-shot** (a single `fire_at` timestamp). Schedules support three modes: **execute** (run a message through the assistant), **notify** (send a notification to the user), and **script** (run a shell command directly without LLM involvement).
20
+ Manage scheduled automations. Schedules can be **recurring** (cron or RRULE expression) or **one-shot** (a single `fire_at` timestamp). Schedules support four modes: **execute** (run a message through the assistant), **notify** (send a notification to the user), **script** (run a shell command directly without LLM involvement), and **workflow** (run a saved multi-agent workflow by name).
21
21
 
22
22
  ## Schedule Syntax
23
23
 
@@ -96,8 +96,13 @@ The `mode` parameter controls what happens when a schedule fires:
96
96
  - **execute** (default) - sends the schedule's message to a background assistant conversation for autonomous handling. The assistant processes the message as if the user sent it.
97
97
  - **notify** - sends a notification to the user via the notification pipeline. No assistant processing occurs.
98
98
  - **script** - runs the `script` field as a shell command directly. No LLM invoked, no conversation created. stdout/stderr are captured in the schedule run record. Exit code 0 = success, non-zero = error. Commands run in the workspace directory with a 60-second timeout by default. Override the timeout per schedule with `timeout_ms` (range 1000–1800000 ms) when a script needs more or less time; pass `timeout_ms: null` on update to revert to the default. The guardian can also adjust this from the /assistant/settings/schedules page.
99
+ - **workflow** - runs a saved workflow (by `workflow_name`) at trigger time, optionally with `workflow_args`. Requires the `workflows` feature flag; `workflow_name` is required. Use this to run a previously saved multi-agent workflow on a schedule (e.g. "run my inbox-triage workflow every morning at 8am"). Optionally pass `capabilities` (the run's single consent point) to grant the scheduled run's leaves side-effecting tools or host functions beyond the read-only baseline; declaring any prompts the guardian for approval once at creation.
99
100
 
100
- Use `notify` for simple reminders ("remind me to take medicine at 9am"), `execute` for tasks that need assistant action ("check my calendar at 8am and send me a digest"), and `script` for lightweight shell automations that don't need LLM involvement ("refresh a cache", "poll an API", "rotate logs").
101
+ Use `notify` for simple reminders ("remind me to take medicine at 9am"), `execute` for tasks that need assistant action ("check my calendar at 8am and send me a digest"), `script` for lightweight shell automations that don't need LLM involvement ("refresh a cache", "poll an API", "rotate logs"), and `workflow` to run a saved workflow on a schedule.
102
+
103
+ ## Inference Profile
104
+
105
+ Execute-mode runs use the default `mainAgent` model selection unless the schedule pins an `inference_profile` (a key from `llm.profiles`). Pin a profile when a recurring task should run on a specific model — e.g. a cost-optimized profile for a high-frequency digest. Pass `inference_profile: null` on update to revert to the default. The pinned profile is shown on the schedule's details page in settings.
101
106
 
102
107
  ## Conversation Reuse
103
108
 
@@ -48,8 +48,36 @@
48
48
  },
49
49
  "mode": {
50
50
  "type": "string",
51
- "enum": ["notify", "execute", "script"],
52
- "description": "Whether to notify the user, execute autonomously via LLM, or run a shell command directly. Defaults to \"execute\"."
51
+ "enum": ["notify", "execute", "script", "workflow"],
52
+ "description": "Whether to notify the user, execute autonomously via LLM, run a shell command directly (script), or run a saved workflow (workflow). Defaults to \"execute\"."
53
+ },
54
+ "workflow_name": {
55
+ "type": "string",
56
+ "description": "Name of a saved workflow to run when the schedule triggers. Required for workflow mode."
57
+ },
58
+ "workflow_args": {
59
+ "type": "object",
60
+ "description": "Arguments object passed to the saved workflow as `args`. Optional; only used in workflow mode."
61
+ },
62
+ "capabilities": {
63
+ "type": "object",
64
+ "description": "Workflow mode only. Per-run capability grant (the single consent point) persisted with the schedule and applied to each triggered run. Leaves get a read-only baseline by default. Declaring any side-effecting tool or host function requires a fresh approval at creation.",
65
+ "properties": {
66
+ "tools": {
67
+ "type": "array",
68
+ "items": { "type": "string" },
69
+ "description": "Side-effecting tool names granted to leaves on top of the read-only baseline."
70
+ },
71
+ "hostFunctions": {
72
+ "type": "array",
73
+ "items": { "type": "string" },
74
+ "description": "Host-function names the run may invoke."
75
+ },
76
+ "persona": {
77
+ "type": "boolean",
78
+ "description": "Grant leaves access to persona (identity + memory) context."
79
+ }
80
+ }
53
81
  },
54
82
  "routing_intent": {
55
83
  "type": "string",
@@ -79,6 +107,10 @@
79
107
  "timeout_ms": {
80
108
  "type": "integer",
81
109
  "description": "For script mode: maximum time in milliseconds the shell command may run before it is killed. Defaults to 60000 (60s). Allowed range: 1000 to 1800000."
110
+ },
111
+ "inference_profile": {
112
+ "type": "string",
113
+ "description": "Inference profile (a key from llm.profiles) the schedule's LLM-executed runs should use, e.g. to run a heavy task on a cheaper model. When omitted, runs use the default mainAgent model selection. Must name a configured profile."
82
114
  }
83
115
  },
84
116
  "required": ["name", "description"]
@@ -155,8 +187,16 @@
155
187
  },
156
188
  "mode": {
157
189
  "type": "string",
158
- "enum": ["notify", "execute", "script"],
159
- "description": "Whether to notify the user, execute autonomously via LLM, or run a shell command directly"
190
+ "enum": ["notify", "execute", "script", "workflow"],
191
+ "description": "Whether to notify the user, execute autonomously via LLM, run a shell command directly (script), or run a saved workflow (workflow)"
192
+ },
193
+ "workflow_name": {
194
+ "type": "string",
195
+ "description": "Name of a saved workflow to run when the schedule triggers. Required when the schedule's resulting mode is workflow."
196
+ },
197
+ "workflow_args": {
198
+ "type": "object",
199
+ "description": "Arguments object passed to the saved workflow as `args`. Only used in workflow mode."
160
200
  },
161
201
  "routing_intent": {
162
202
  "type": "string",
@@ -186,6 +226,10 @@
186
226
  "timeout_ms": {
187
227
  "type": ["integer", "null"],
188
228
  "description": "For script mode: maximum time in milliseconds the shell command may run before it is killed (default 60000, range 1000 to 1800000). Pass null to clear a custom timeout and revert to the default."
229
+ },
230
+ "inference_profile": {
231
+ "type": ["string", "null"],
232
+ "description": "Inference profile (a key from llm.profiles) the schedule's LLM-executed runs should use. Pass null to clear the override and revert to the default mainAgent model selection. Must name a configured profile."
189
233
  }
190
234
  },
191
235
  "required": ["job_id"]
@@ -0,0 +1,214 @@
1
+ ---
2
+ name: workflows
3
+ description: Author and run autonomous multi-agent workflows that fan work across parallel leaf agents
4
+ compatibility: "Designed for Vellum personal assistants"
5
+ metadata:
6
+ emoji: "⚙️"
7
+ vellum:
8
+ display-name: "Workflows"
9
+ category: "system"
10
+ feature-flag: "workflows"
11
+ activation-hints:
12
+ - "A task decomposes into many similar sub-tasks that can run concurrently (score every item, extract a field from each of many documents, draft-then-verify a batch)"
13
+ - "You want fan-out orchestrated deterministically and the result reported back when the whole run finishes"
14
+ avoid-when:
15
+ - "The task is a single tool call or a quick lookup — do it inline"
16
+ - "The work needs interactive, conversational back-and-forth rather than unattended fan-out"
17
+ ---
18
+
19
+ A workflow is a short JS/TS script you author that runs in a sandbox and fans work
20
+ out across many short-lived **leaf agents**, orchestrated deterministically. Launch
21
+ one with `run_workflow` (inline `script` OR saved `name`, exactly one). It returns a
22
+ `runId` immediately; the run is asynchronous and you are notified in this
23
+ conversation when it completes — **do NOT poll**.
24
+
25
+ Reach for one when a task decomposes into many similar small sub-tasks that can run
26
+ concurrently. For a single task or a quick lookup, do it inline.
27
+
28
+ ## The script model
29
+
30
+ These are the load-bearing invariants. Get them wrong and the run misbehaves silently.
31
+
32
+ ### Scripts are SYNCHRONOUS — never `await`
33
+
34
+ Host functions block and return their result directly. Write straight-line code.
35
+
36
+ ```js
37
+ const r = agent("Summarize this thread."); // r is the result, right here
38
+ ```
39
+
40
+ Do **not** write `await agent(...)`, and do **not** make the script `async`. An
41
+ `async` script deadlocks on its second host call — the sandbox can suspend the main
42
+ evaluation stack but not a promise continuation.
43
+
44
+ ### Every script begins with a literal `meta`
45
+
46
+ The first statement must be a pure-literal export — no computed values, template
47
+ strings, or concatenation:
48
+
49
+ ```js
50
+ export const meta = {
51
+ name: "triage-inbox",
52
+ description: "Triage and label inbox messages",
53
+ };
54
+ ```
55
+
56
+ `meta` is extracted **statically**, without executing the script, so it must be a
57
+ plain object literal with string `name` and `description`. The `name` is how a saved
58
+ workflow is referenced by `workflow(name)` and the scheduler.
59
+
60
+ ### You must `return` the result
61
+
62
+ The script body runs as a function. Its result is whatever it `return`s at the top
63
+ level — a bare trailing expression (e.g. `result;`) is **discarded** and the run
64
+ finishes with no result. Always `return` the value you want surfaced.
65
+
66
+ ```js
67
+ const result = agent(`Write the final summary: ${JSON.stringify(parts)}`);
68
+ return result;
69
+ ```
70
+
71
+ ### Determinism (this is what makes runs resumable)
72
+
73
+ Every leaf call is journaled by sequence number and input hash, so a resumed run can
74
+ replay the unchanged prefix instead of re-spawning agents. That only holds if the
75
+ script is deterministic, so `Date.now()`, `Math.random()`, and argless `new Date()`
76
+ **throw**. Pass any timestamps or random seeds in through `args`.
77
+
78
+ ## Host API
79
+
80
+ All functions are synchronous from the script's perspective.
81
+
82
+ | Function | Returns | Notes |
83
+ | ---------------------------- | ---------------------------------------------- | ----------------------------------------------------------------------------------------------------------- |
84
+ | `agent(prompt, opts?)` | the leaf's result | Runs ONE leaf. **Throws** on leaf failure (fails the whole run). |
85
+ | `leaf(prompt, opts?)` | a leaf descriptor | Runs nothing on its own; used inside `parallel`/`map`/`pipeline`. |
86
+ | `parallel(specs)` | `results[]` | Runs an array of `leaf(...)` descriptors concurrently, results in input order. A failed leaf becomes `null` (never throws). |
87
+ | `map(items, build)` | `results[]` | `build(item, i)` returns a `leaf(...)` descriptor per item; runs them like `parallel`. |
88
+ | `pipeline(items, ...stages)` | `results[]` | Each `stage(prev, i)` returns a `leaf(...)` descriptor (run an agent) OR a plain value (pass through unchanged — filter/transform locally, no agent spent). Per-stage barrier: stage N+1 starts only after all of stage N finishes. |
89
+ | `phase(title)` | — | Marks a named phase for progress reporting. |
90
+ | `log(msg)` | — | Emits a progress log line. |
91
+ | `usage()` | `{ agentsSpawned, inputTokens, outputTokens }` | Live snapshot so a script can self-moderate. |
92
+ | `workflow(name, args?)` | the child's result | Runs a SAVED workflow inline, depth 1 only (a child may not call `workflow()`). |
93
+ | `args` | the run input | The `args` object passed to `run_workflow`. |
94
+
95
+ Use `agent` for a single sequential leaf (throws on failure). Use `parallel`/`map`/
96
+ `pipeline` for fan-out (a failed leaf is `null`, so a batch survives a few bad items).
97
+
98
+ ## Leaf options (`opts` for `agent` / `leaf`)
99
+
100
+ | Option | Type | Effect |
101
+ | --------- | -------------------------- | --------------------------------------------------------------------------------------- |
102
+ | `schema` | JSON Schema object literal | Forces structured output via a tool. A schema leaf runs with **no tools** (pure judge/extractor). Use a plain JSON Schema literal, not Zod. |
103
+ | `label` | string | Short display/diagnostic label for the leaf. |
104
+ | `profile` | string | Overrides the model profile. Must exist in `llm.profiles` or the leaf throws. See [Listing profiles](#listing-available-profiles). |
105
+ | `persona` | boolean | `true` makes the leaf speak AS the assistant (identity + memory) — use for output meant to be in the assistant's voice. Default is anonymous — use for impartial judging/extraction. |
106
+
107
+ `persona: true` is the costly path (it runs the full memory-injection pipeline). Use
108
+ it only for the few leaves whose output must be in the assistant's voice; keep bulk
109
+ judging/extraction anonymous.
110
+
111
+ ## Capabilities — the single consent point
112
+
113
+ The `capabilities` argument to `run_workflow` declares **once**, up front, what the
114
+ run's leaves may do. There are **no per-call permission prompts inside a running
115
+ workflow**.
116
+
117
+ ```jsonc
118
+ {
119
+ "tools": ["file_write", "gmail_send"], // side-effecting tools granted to leaves
120
+ "hostFunctions": [], // host-function names the run may invoke
121
+ "persona": true // grant leaves persona (identity + memory)
122
+ }
123
+ ```
124
+
125
+ - **Read-only baseline** (always available, no declaration, no launch prompt):
126
+ `file_read`, `file_list`, `recall`, `web_search`.
127
+ - **`web_fetch` is NOT in the baseline** — an outbound fetch is side-effecting (its
128
+ URL can exfiltrate read data), so a leaf that must fetch a URL has to declare
129
+ `"web_fetch"` in `capabilities.tools`.
130
+ - **Declaring ANY side-effecting tool** (writes, sends, shell, `web_fetch`, …) or
131
+ host function makes the LAUNCH prompt the user for approval **once** — that single
132
+ approval covers the whole run. A read-only run (no declared tools) launches with no
133
+ prompt. Declare the minimum you need.
134
+
135
+ Runs are autonomous but BOUNDED by a per-run agent cap — spend is structurally
136
+ capped and you cannot exceed it.
137
+
138
+ ## Listing available profiles
139
+
140
+ Before choosing a `profile` for a leaf, look up the valid values rather than guessing
141
+ — an unknown profile throws.
142
+
143
+ - **Preferred (model-accessible):** call `manage_workflows` with action
144
+ `list_profiles`. It returns the profile names defined in `llm.profiles` plus the
145
+ workspace-wide active profile.
146
+ - The same data is served by the daemon route `GET config/llm/profiles` (operationId
147
+ `llm_profiles_list`), which clients use to populate profile dropdowns.
148
+
149
+ Omit `profile` to use the default: a persona leaf mirrors the main agent (the active
150
+ profile floats above the call-site default); an anonymous leaf uses the cost-optimized
151
+ `workflowLeaf` default.
152
+
153
+ ## Run management
154
+
155
+ Use `manage_workflows` to inspect and control runs:
156
+
157
+ | Action | Requires | Purpose |
158
+ | --------------- | ---------- | ---------------------------------------------------------------- |
159
+ | `status` | `run_id` | Status + agent/token counts for one run (NOT the result). |
160
+ | `get_result` | `run_id` | The full result of a finished run. |
161
+ | `list_runs` | — | Recent runs, newest first. |
162
+ | `abort` | `run_id` | Signal an in-flight run to abort. |
163
+ | `resume` | `run_id` | Resume an `interrupted` run (see below). |
164
+ | `list_profiles` | — | List defined profiles + the active profile (for leaf `profile`).|
165
+
166
+ The completion notification injected when a run finishes carries a **truncated**
167
+ preview of the result (large results are cut off). To read the complete result,
168
+ call `manage_workflows` with action `get_result` and the `run_id` — `status`
169
+ deliberately omits the result to stay lightweight.
170
+
171
+ ### Crash recovery / resume
172
+
173
+ Resume is **not automatic**. If the assistant restarts mid-run, the run is reconciled
174
+ to status `interrupted` (the agent/token accounting is preserved so the agent cap
175
+ still carries across the restart). It sits there until you explicitly resume it by
176
+ `run_id` via `manage_workflows` action `resume`. Resuming re-invokes the engine with
177
+ the same `runId`: the journal **replays the completed prefix** without re-spawning (or
178
+ re-paying for) finished leaves, then continues from the first unfinished leaf under the
179
+ run's originally-declared capabilities. Only `interrupted` runs are resumable; a
180
+ `completed` / `failed` / `aborted` run is terminal.
181
+
182
+ ## Worked example
183
+
184
+ Score each inbox item in parallel (anonymous schema leaves), then write one summary in
185
+ the assistant's voice (a single persona leaf). The item list comes in via `args` —
186
+ never fetched inside the script.
187
+
188
+ ```js
189
+ export const meta = {
190
+ name: "triage-inbox",
191
+ description: "Score and summarize inbox items",
192
+ };
193
+
194
+ phase("score");
195
+ const scored = map(args.items, (item) =>
196
+ leaf(`Rate this message's urgency 0-10 with a one-line reason:\n${item.subject}\n${item.body}`, {
197
+ label: `score:${item.id}`,
198
+ schema: {
199
+ type: "object",
200
+ properties: { urgency: { type: "number" }, reason: { type: "string" } },
201
+ required: ["urgency", "reason"],
202
+ },
203
+ }),
204
+ );
205
+
206
+ phase("summarize");
207
+ const summary = agent(
208
+ `Here are scored inbox items. Write a short triage summary, highlighting anything urgent:\n${JSON.stringify(scored)}`,
209
+ { persona: true },
210
+ );
211
+ return summary;
212
+ ```
213
+
214
+ A failed scoring leaf shows up as `null` in `scored`; the run continues.