@gobing-ai/spur 0.3.43 → 0.3.44

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (68) hide show
  1. package/.claude-plugin/marketplace.json +1 -1
  2. package/config/templates/AGENTS.md +5 -0
  3. package/config/workflows/idea-pipeline.yaml +205 -33
  4. package/config/workflows/task-pipeline.yaml +1 -1
  5. package/package.json +1 -1
  6. package/plugins/sp/commands/dev-runall.md +3 -1
  7. package/plugins/sp/plugin.json +1 -1
  8. package/plugins/sp/skills/spur-cli/SKILL.md +3 -0
  9. package/plugins/sp/skills/spur-cli/references/features/hierarchy-mece.md +8 -3
  10. package/plugins/sp/skills/spur-cli/references/tasks.md +22 -0
  11. package/plugins/sp/skills/spur-dev/SKILL.md +12 -4
  12. package/plugins/sp/skills/spur-dev/references/dev-operations.md +1 -1
  13. package/plugins/sp/skills/spur-dev/references/execution-batch.md +12 -3
  14. package/plugins/sp/skills/spur-dev/references/flag-glossary.md +1 -1
  15. package/plugins/sp/skills/spur-dev/references/planning-workflow.md +53 -0
  16. package/spur.js +147 -67
  17. package/web/_astro/{BoardApp.8sc6Rb1t.js → BoardApp.BnhkWEOG.js} +99 -99
  18. package/web/_astro/BoardApp.CxOS4HSo.js +1 -0
  19. package/web/_astro/{TaskDetail.DOKJP-Kp.js → TaskDetail.B6yiT-U7.js} +1 -1
  20. package/web/_astro/{arc.CTqmqspU.js → arc.PbmgYm3_.js} +1 -1
  21. package/web/_astro/{architectureDiagram-3BPJPVTR.B0iBdehs.js → architectureDiagram-3BPJPVTR.BkOMSdmD.js} +1 -1
  22. package/web/_astro/{blockDiagram-GPEHLZMM.BHE5UFj5.js → blockDiagram-GPEHLZMM.BtomoUdy.js} +1 -1
  23. package/web/_astro/{c4Diagram-AAUBKEIU.DQApkgGM.js → c4Diagram-AAUBKEIU.UzrYwJnF.js} +1 -1
  24. package/web/_astro/channel.DRRsElzO.js +1 -0
  25. package/web/_astro/{chunk-2J33WTMH.fPrKLM8v.js → chunk-2J33WTMH.nNChLHkw.js} +1 -1
  26. package/web/_astro/{chunk-4BX2VUAB.DExvwBhn.js → chunk-4BX2VUAB.DLYAecPo.js} +1 -1
  27. package/web/_astro/{chunk-55IACEB6.EjNkjuD6.js → chunk-55IACEB6.t5bzuj1J.js} +1 -1
  28. package/web/_astro/{chunk-727SXJPM.Dl8xy5rT.js → chunk-727SXJPM.DBV62bIy.js} +1 -1
  29. package/web/_astro/{chunk-AQP2D5EJ.BSzMBgBs.js → chunk-AQP2D5EJ.BdaPhTPs.js} +1 -1
  30. package/web/_astro/{chunk-FMBD7UC4.CAeOHpbb.js → chunk-FMBD7UC4.eoZ88KMf.js} +1 -1
  31. package/web/_astro/{chunk-ND2GUHAM.Bbeu_ZOH.js → chunk-ND2GUHAM.QdNDyeSh.js} +1 -1
  32. package/web/_astro/{chunk-QZHKN3VN.Da_GaIAZ.js → chunk-QZHKN3VN.DPaToLfx.js} +1 -1
  33. package/web/_astro/{classDiagram-4FO5ZUOK.sMfFYSQp.js → classDiagram-4FO5ZUOK.By9BMe7b.js} +1 -1
  34. package/web/_astro/{classDiagram-v2-Q7XG4LA2.sMfFYSQp.js → classDiagram-v2-Q7XG4LA2.By9BMe7b.js} +1 -1
  35. package/web/_astro/{cose-bilkent-S5V4N54A.PS3IwOH4.js → cose-bilkent-S5V4N54A.DMypl6q_.js} +1 -1
  36. package/web/_astro/{dagre-BM42HDAG.EQCx8a6D.js → dagre-BM42HDAG.p26m9grg.js} +1 -1
  37. package/web/_astro/{diagram-2AECGRRQ.aEFQl8v8.js → diagram-2AECGRRQ.C6m3oTIu.js} +1 -1
  38. package/web/_astro/{diagram-5GNKFQAL.jyN4tqFX.js → diagram-5GNKFQAL.JCxivn2z.js} +1 -1
  39. package/web/_astro/{diagram-KO2AKTUF.DaTVHnme.js → diagram-KO2AKTUF.DifbEd3P.js} +1 -1
  40. package/web/_astro/{diagram-LMA3HP47.DWsGRLAN.js → diagram-LMA3HP47.Bffga-cf.js} +1 -1
  41. package/web/_astro/{diagram-OG6HWLK6.DYxmOViu.js → diagram-OG6HWLK6.CUKx5_Cx.js} +1 -1
  42. package/web/_astro/{erDiagram-TEJ5UH35.BKM75coa.js → erDiagram-TEJ5UH35.DaJ9KZR0.js} +1 -1
  43. package/web/_astro/{flowDiagram-I6XJVG4X.B7qpqDW5.js → flowDiagram-I6XJVG4X.Do6l5pzg.js} +1 -1
  44. package/web/_astro/{ganttDiagram-6RSMTGT7.CYBL-hwy.js → ganttDiagram-6RSMTGT7.DqBKgB-3.js} +1 -1
  45. package/web/_astro/{gitGraphDiagram-PVQCEYII.Cz4ly3Kq.js → gitGraphDiagram-PVQCEYII.B3VLSQ22.js} +1 -1
  46. package/web/_astro/{index.yAse9IaO.css → index.QfZ9SC3X.css} +1 -1
  47. package/web/_astro/{infoDiagram-5YYISTIA.CIS7cXrJ.js → infoDiagram-5YYISTIA.hK8ZC1oa.js} +1 -1
  48. package/web/_astro/{ishikawaDiagram-YF4QCWOH.DQ-pmaOG.js → ishikawaDiagram-YF4QCWOH.C8DXwh9Z.js} +1 -1
  49. package/web/_astro/{journeyDiagram-JHISSGLW.DVS6F7UM.js → journeyDiagram-JHISSGLW.mj5aHQDb.js} +1 -1
  50. package/web/_astro/{kanban-definition-UN3LZRKU.D8t-G-U6.js → kanban-definition-UN3LZRKU.Dwlhi49r.js} +1 -1
  51. package/web/_astro/{linear.IO8rTd_n.js → linear.7zBHdlwl.js} +1 -1
  52. package/web/_astro/{mermaid.core.CjVnvFNf.js → mermaid.core.AYdA0EJr.js} +4 -4
  53. package/web/_astro/{mindmap-definition-RKZ34NQL.DZoEMTnU.js → mindmap-definition-RKZ34NQL.BSnniHEq.js} +1 -1
  54. package/web/_astro/{pieDiagram-4H26LBE5.86r1QWnX.js → pieDiagram-4H26LBE5.B7EVpA3J.js} +1 -1
  55. package/web/_astro/{quadrantDiagram-W4KKPZXB.DzFA9uPk.js → quadrantDiagram-W4KKPZXB.rg4i8YOY.js} +1 -1
  56. package/web/_astro/{requirementDiagram-4Y6WPE33.dQvdvfMr.js → requirementDiagram-4Y6WPE33.C4pr792x.js} +1 -1
  57. package/web/_astro/{sankeyDiagram-5OEKKPKP.BvC-VNI8.js → sankeyDiagram-5OEKKPKP.DTXU996O.js} +1 -1
  58. package/web/_astro/{sequenceDiagram-3UESZ5HK.DDQNwHs8.js → sequenceDiagram-3UESZ5HK.D73DOjzi.js} +1 -1
  59. package/web/_astro/{stateDiagram-AJRCARHV.BKwYy-tQ.js → stateDiagram-AJRCARHV.CVkeSaEH.js} +1 -1
  60. package/web/_astro/{stateDiagram-v2-BHNVJYJU.KxCiRYYf.js → stateDiagram-v2-BHNVJYJU.DkBGwoN5.js} +1 -1
  61. package/web/_astro/{timeline-definition-PNZ67QCA.Bh9NDjOx.js → timeline-definition-PNZ67QCA.DnJ9Yh7G.js} +1 -1
  62. package/web/_astro/{vennDiagram-CIIHVFJN.DIM-n6us.js → vennDiagram-CIIHVFJN.DN1tvR4Y.js} +1 -1
  63. package/web/_astro/{wardley-L42UT6IY.DW15qOFQ.js → wardley-L42UT6IY.D54jJ-7-.js} +1 -1
  64. package/web/_astro/{wardleyDiagram-YWT4CUSO.UTIvjsG9.js → wardleyDiagram-YWT4CUSO.C2grUBbE.js} +1 -1
  65. package/web/_astro/{xychartDiagram-2RQKCTM6.CQYhkVDS.js → xychartDiagram-2RQKCTM6.D05okBnV.js} +1 -1
  66. package/web/index.html +2 -2
  67. package/web/_astro/BoardApp.ChKnY1JD.js +0 -1
  68. package/web/_astro/channel.B0q9WGrH.js +0 -1
@@ -6,7 +6,7 @@
6
6
  },
7
7
  "plugins": [
8
8
  {
9
- "version": "0.3.43",
9
+ "version": "0.3.44",
10
10
  "name": "sp",
11
11
  "source": "./plugins/sp",
12
12
  "description": "Spur - my harness toolkits"
@@ -69,6 +69,11 @@ All product development work goes through the harness by default.
69
69
  auto`, a name, parallel/headless execution, direct `spur agent run`, and engine-driven workflow
70
70
  `agent.run` remain subprocess surfaces.
71
71
 
72
+ **Task lookup fast path:** Given a WBS, do not search `docs/tasks*` or guess `--folder`. Use
73
+ `spur task show <wbs> --json` for task metadata and content; its response also includes `filePath`.
74
+ Use `spur task path <wbs> --json` only when a filesystem consumer needs the absolute path. Both
75
+ commands resolve across configured task folders. Reuse the first `show` response within the run.
76
+
72
77
  **Platform fallback:** Platforms without slash commands and/or subagents still use the harness.
73
78
  Install the plugin through Superskill for the target platform, then use skills `sp:spur-dev`,
74
79
  `sp:spur-cli`, `sp:code-verification` (and related) plus the `spur` CLI. Do not invent a parallel
@@ -10,7 +10,7 @@
10
10
  # Shape: start -> discovery -> idea-eval -> feature-create -> ac-generate -> feature-check
11
11
  # -> system-design (conditional: needs_design signal)
12
12
  # -> design-approval (taste HITL gate; not auto-clicked by --auto)
13
- # -> decompose -> batch-create -> handoff
13
+ # -> decompose -> batch-create -> handoff-finalize -> handoff
14
14
  # (idea-eval rejection -> cancelled)
15
15
  # (feature-check failure routes back to ac-generate; batch-create failure to decompose).
16
16
  #
@@ -128,20 +128,36 @@ states:
128
128
  - id: feature-create
129
129
  description: >
130
130
  Create a feature via spur feature create, or select an existing feature id.
131
- The agent writes the feature id to .spur/run/${vars.__runId}-idea-feature-id.txt for subsequent
132
- states to read. Prefer the enhanced idea from .spur/run/${vars.__runId}-idea-eval-report.md as
131
+ The agent writes the feature id to .spur/run/${vars.__runId}-idea-feature-id.txt and body-only
132
+ intent artifacts .spur/run/${vars.__runId}-idea-goal.md / -idea-scope.md. Goal carries
133
+ concise intent only; Scope carries explicit in/out boundaries — task breakdowns and
134
+ checklists never enter Goal. Shell actions then persist both sections through
135
+ `spur feature update --section Goal|Scope --from-file ...`; a missing/empty artifact
136
+ stops the state. Prefer the enhanced idea from .spur/run/${vars.__runId}-idea-eval-report.md as
133
137
  context; do not overwrite vars.idea.
134
138
  onEnter:
135
139
  - kind: agent.run
136
140
  options:
137
141
  agent: ${vars.agent}
138
- input: "Create a feature for the idea: ${vars.idea}. Read .spur/run/${vars.__runId}-idea-eval-report.md if present for the enhanced idea and scores. Use 'spur feature create \"<name>\" --json' to create it. Write the feature id to .spur/run/${vars.__runId}-idea-feature-id.txt. If an existing feature is appropriate, use its id instead."
142
+ input: "Create a feature for the idea: ${vars.idea}. Read .spur/run/${vars.__runId}-idea-eval-report.md if present for the enhanced idea and scores. Use 'spur feature create \"<name>\" --json' to create it. Write the feature id to .spur/run/${vars.__runId}-idea-feature-id.txt. If an existing feature is appropriate, use its id instead. Also write two body-only intent artifacts: .spur/run/${vars.__runId}-idea-goal.md with concise Goal intent only (a short statement of what the feature achieves; never task breakdowns, checklists, or how-to steps), and .spur/run/${vars.__runId}-idea-scope.md with explicit in-scope and out-of-scope boundary bullets."
139
143
  expectFile: .spur/run/${vars.__runId}-idea-feature-id.txt
140
144
  timeoutMs: ${vars.stepTimeoutMs}
141
145
  - kind: file.read.into-var
142
146
  options:
143
147
  path: .spur/run/${vars.__runId}-idea-feature-id.txt
144
148
  var: featureId
149
+ - kind: shell
150
+ options:
151
+ command: >-
152
+ mkdir -p .spur/run &&
153
+ test -s .spur/run/$__runId-idea-goal.md &&
154
+ $spurBin feature update "$featureId" --section Goal --from-file .spur/run/$__runId-idea-goal.md
155
+ - kind: shell
156
+ options:
157
+ command: >-
158
+ mkdir -p .spur/run &&
159
+ test -s .spur/run/$__runId-idea-scope.md &&
160
+ $spurBin feature update "$featureId" --section Scope --from-file .spur/run/$__runId-idea-scope.md
145
161
 
146
162
  - id: ac-generate
147
163
  description: >
@@ -177,7 +193,7 @@ states:
177
193
  options:
178
194
  command: >-
179
195
  if test -s .spur/run/$__runId-idea-ac-content.md; then
180
- $spurBin feature update $featureId --section "Acceptance Criteria" --from-file .spur/run/$__runId-idea-ac-content.md && date -u +%Y-%m-%dT%H:%M:%SZ > .spur/run/$__runId-idea-ac-done.txt;
196
+ $spurBin feature update "$featureId" --section "Acceptance Criteria" --from-file .spur/run/$__runId-idea-ac-content.md && date -u +%Y-%m-%dT%H:%M:%SZ > .spur/run/$__runId-idea-ac-done.txt;
181
197
  else
182
198
  exit 0;
183
199
  fi
@@ -206,13 +222,33 @@ states:
206
222
  Dispatch sp:sys-architecture to produce the system design. This state runs
207
223
  only when needs_design=true (design=auto signal path). The agent creates ADR
208
224
  entries, architecture updates, and design satellites through constitution rules.
209
- The design-approval state follows (taste gate).
225
+ A run-scoped design-review artifact (.spur/run/${vars.__runId}-idea-design-review.md)
226
+ with fixed headings `## Proposed design`, `## Operator feedback`, and `## Reconciliation`
227
+ carries operator rejection feedback back into this state: on the first pass the agent
228
+ fills `## Proposed design`; on retry (operator feedback present) it revises the design,
229
+ records the reconciliation, and updates invalidated Acceptance Criteria through
230
+ `spur feature update` before design can exit. The design-approval state follows (taste gate).
210
231
  onEnter:
232
+ - kind: shell
233
+ options:
234
+ command: >-
235
+ mkdir -p .spur/run &&
236
+ REVIEW=".spur/run/$__runId-idea-design-review.md" &&
237
+ test -f "$REVIEW" || printf '## Proposed design\n\n## Operator feedback\n\n## Reconciliation\n' > "$REVIEW"
211
238
  - kind: agent.run
212
239
  options:
213
240
  agent: ${vars.agent}
214
- input: "Run sp:sys-architecture for feature ${vars.featureId}. Read the brainstorm artifact and feature AC. Produce ADR entries, architecture updates, and design satellites (docs/design/<slug>.md) following the constitution edit rules. Do not write task or feature corpus files directly."
241
+ input: "Run sp:sys-architecture for feature ${vars.featureId}. Read the brainstorm artifact, feature AC, and .spur/run/${vars.__runId}-idea-design-review.md. Produce ADR entries, architecture updates, and design satellites (docs/design/<slug>.md) following the constitution edit rules. Do not write task or feature corpus files directly. Design-review contract (.spur/run/${vars.__runId}-idea-design-review.md, fixed headings `## Proposed design`, `## Operator feedback`, `## Reconciliation`): on the first pass write the proposed design summary under `## Proposed design` and leave `## Operator feedback` empty; on retry with operator feedback present, revise the design/ADR artifacts, document the changes under `## Reconciliation`, and when the feedback invalidates an Acceptance Criteria scenario write the revised AC section body to a file and persist it via `$spurBin feature update \"$featureId\" --section \"Acceptance Criteria\" --from-file <file>` — never edit feature corpus files directly."
242
+ expectFile: .spur/run/${vars.__runId}-idea-design-review.md
215
243
  timeoutMs: ${vars.stepTimeoutMs}
244
+ # expectFile proves existence only, and the onEnter skeleton pre-creates the file — so an
245
+ # agent no-op would still pass it (0515 P3-2). Fail closed: the `## Proposed design` section
246
+ # must carry non-whitespace content before design can exit to approval/decompose.
247
+ - kind: shell
248
+ options:
249
+ command: >-
250
+ REVIEW=".spur/run/$__runId-idea-design-review.md" &&
251
+ awk '/^## Proposed design/{f=1; next} /^## Operator feedback/{f=0} f' "$REVIEW" | grep -q '[^[:space:]]'
216
252
 
217
253
  - id: design-approval
218
254
  description: >
@@ -226,6 +262,11 @@ states:
226
262
  agent passes. The onEnter shell increments a reject counter; the `no`
227
263
  edge to system-design is capped at 1 revise, after which the run
228
264
  terminates as `failed` naming this gate.
265
+
266
+ Rejection must carry operator feedback: before answering `no`, the operator
267
+ records the concrete issue(s) in the run-scoped design-review artifact
268
+ (.spur/run/${vars.__runId}-idea-design-review.md) under `## Operator feedback` —
269
+ the revision pass reads it to reconcile invalidated AC before re-approval.
229
270
  pause: true
230
271
  onEnter:
231
272
  - kind: shell
@@ -233,23 +274,45 @@ states:
233
274
  command: 'mkdir -p .spur/run && count=$(cat .spur/run/$__runId-idea-design-reject-count 2>/dev/null || echo 0); echo $((count + 1)) > .spur/run/$__runId-idea-design-reject-count'
234
275
  - kind: hitl.confirm
235
276
  options:
236
- prompt: "Review the system design for feature ${vars.featureId}. Approve to proceed to decomposition?"
277
+ prompt: "Review the system design for feature ${vars.featureId}. Approve to proceed to decomposition? To reject, first record your feedback in .spur/run/${vars.__runId}-idea-design-review.md under the `## Operator feedback` heading, then answer no."
237
278
 
238
279
  - id: decompose
239
280
  description: >
240
281
  Dispatch sp:spec-decomposition with the brainstorm artifact, feature AC, and
241
282
  design doc as input. The agent produces a task-batch JSON file at
242
- .spur/run/${vars.__runId}-idea-task-batch.json, validated against task-batch.schema.json.
283
+ .spur/run/${vars.__runId}-idea-task-batch.json, validated against task-batch.schema.json,
284
+ plus the private task-order sidecar .spur/run/${vars.__runId}-idea-task-order.json
285
+ (a JSON array of { name, depends_on_names[] }; `[]` when no ordering exists).
286
+ A post-agent shell action validates the sidecar shape before decomposition
287
+ can exit — missing or ambiguous title-to-WBS resolution fails before handoff.
243
288
  onEnter:
244
289
  - kind: shell
245
290
  options:
246
- command: 'mkdir -p .spur/run && count=$(cat .spur/run/$__runId-idea-decompose-retry-count 2>/dev/null || echo 0); echo $((count + 1)) > .spur/run/$__runId-idea-decompose-retry-count; rm -f .spur/run/$__runId-idea-task-batch.json .spur/run/$__runId-idea-batch-create.done .spur/run/$__runId-idea-batch-create.failed'
291
+ command: 'mkdir -p .spur/run && count=$(cat .spur/run/$__runId-idea-decompose-retry-count 2>/dev/null || echo 0); echo $((count + 1)) > .spur/run/$__runId-idea-decompose-retry-count; rm -f .spur/run/$__runId-idea-task-batch.json .spur/run/$__runId-idea-task-order.json .spur/run/$__runId-idea-batch-create.done .spur/run/$__runId-idea-batch-create.failed .spur/run/$__runId-idea-batch-create-result.json .spur/run/$__runId-idea-handoff.md'
247
292
  - kind: agent.run
248
293
  options:
249
294
  agent: ${vars.agent}
250
- input: "Run sp:spec-decomposition for feature ${vars.featureId}. Read the brainstorm artifact, feature AC, and design doc. SIZING FIRST, before any JSON: apply the skill's `Default to NOT decomposing` rubric to the whole unit of work — if it scores 0-2 the correct output is a ONE-entry batch, not many. Scenario count is not task count: merge scenarios that one task delivers (same file surface, same subsystem, or unreadable apart in review), and list every scenario a task covers in its background. Merging never costs AC coverage — one task may carry several scenarios. Do not emit one entry per scenario or per requirement by reflex. Then produce a task-batch JSON array at .spur/run/${vars.__runId}-idea-task-batch.json, validated against apps/cli/schemas/task-batch.schema.json. Schema-permitted fields per entry: `name`, `background`, `requirements`, `feature_id`, `parent_wbs`, `priority`, `tags`, `template` — schema validation rejects anything else. Acceptance Criteria, Design, and Plan sections are filled in by the per-task refine step after batch-create, NOT at decompose time. Validate locally against the schema before emitting."
295
+ input: "Run sp:spec-decomposition for feature ${vars.featureId}. Read the brainstorm artifact, feature AC, and design doc. SIZING FIRST, before any JSON: apply the skill's `Default to NOT decomposing` rubric to the whole unit of work — if it scores 0-2 the correct output is a ONE-entry batch, not many. Scenario count is not task count: merge scenarios that one task delivers (same file surface, same subsystem, or unreadable apart in review), and list every scenario a task covers in its background. Merging never costs AC coverage — one task may carry several scenarios. Do not emit one entry per scenario or per requirement by reflex. Then produce a task-batch JSON array at .spur/run/${vars.__runId}-idea-task-batch.json, validated against apps/cli/schemas/task-batch.schema.json. Schema-permitted fields per entry: `name`, `background`, `requirements`, `feature_id`, `parent_wbs`, `priority`, `tags`, `template` — schema validation rejects anything else. Acceptance Criteria, Design, and Plan sections are filled in by the per-task refine step after batch-create, NOT at decompose time. Validate locally against the schema before emitting. Also emit the private task-order sidecar at .spur/run/${vars.__runId}-idea-task-order.json: a JSON array (one entry per batch item) of `{ name: <exact batch item name>, depends_on_names: [<batch item names>] }` declaring ordering/dependencies between the batch items; use `[]` when no ordering exists. Every `name` and every dependency must match exactly one batch item `name` — it is private workflow data, not part of task-batch.schema.json."
251
296
  expectFile: .spur/run/${vars.__runId}-idea-task-batch.json
252
297
  timeoutMs: ${vars.stepTimeoutMs}
298
+ # R1 (0518): the task-order sidecar is the ordering contract for handoff-finalize.
299
+ # Validate it fails closed: sidecar must be an array, batch names unique, sidecar names
300
+ # unique, every sidecar name/dependency must refer to exactly one batch name, and — the
301
+ # converse (F2, 0518 verify) — every batch name must appear in the sidecar, so a partial
302
+ # sidecar can never silently skip `task deps` for an unlisted item (`[]` is valid).
303
+ - kind: shell
304
+ options:
305
+ command: >-
306
+ BATCH=".spur/run/$__runId-idea-task-batch.json" &&
307
+ ORDER=".spur/run/$__runId-idea-task-order.json" &&
308
+ jq -e --slurpfile b "$BATCH" '
309
+ (type == "array") and
310
+ (($b[0] | map(.name) | length) == ($b[0] | map(.name) | unique | length)) and
311
+ ((map(.name) | length) == (map(.name) | unique | length)) and
312
+ ((map(.name) - ($b[0] | map(.name))) | length == 0) and
313
+ ((($b[0] | map(.name)) - map(.name)) | length == 0) and
314
+ (((map(.depends_on_names // []) | flatten | unique) - ($b[0] | map(.name))) | length == 0)
315
+ ' "$ORDER" >/dev/null
253
316
 
254
317
  - id: batch-create
255
318
  description: >
@@ -269,21 +332,116 @@ states:
269
332
  description: >
270
333
  Executes task batch creation exactly once per decomposition attempt. The
271
334
  side effect lives in onEnter and records a sentinel; transition guards only
272
- inspect sentinel/retry state.
335
+ inspect sentinel/retry state. The `--json` result is captured atomically in
336
+ .spur/run/${vars.__runId}-idea-batch-create-result.json (temp file + mv); the done
337
+ sentinel is written only after the JSON parses and `created == (.wbs | length)`,
338
+ so handoff-finalize can zip batch names to the returned WBS list. Any failure
339
+ (CLI error or malformed result) writes the failed sentinel, preserving the
340
+ existing retry behavior.
273
341
  onEnter:
274
342
  - kind: shell
275
343
  options:
276
- command: 'if test -f .spur/run/$__runId-idea-batch-create.done; then exit 0; fi; rm -f .spur/run/$__runId-idea-batch-create.failed; if $spurBin task batch-create --file .spur/run/$__runId-idea-task-batch.json; then date -u +%Y-%m-%dT%H:%M:%SZ > .spur/run/$__runId-idea-batch-create.done; else date -u +%Y-%m-%dT%H:%M:%SZ > .spur/run/$__runId-idea-batch-create.failed; fi'
344
+ command: >-
345
+ if test -f .spur/run/$__runId-idea-batch-create.done; then exit 0; fi;
346
+ rm -f .spur/run/$__runId-idea-batch-create.failed .spur/run/$__runId-idea-batch-create-result.json .spur/run/$__runId-idea-batch-create-result.json.tmp;
347
+ if $spurBin task batch-create --file .spur/run/$__runId-idea-task-batch.json --json > .spur/run/$__runId-idea-batch-create-result.json.tmp && jq -e ".created == (.wbs | length)" .spur/run/$__runId-idea-batch-create-result.json.tmp >/dev/null 2>&1; then
348
+ mv .spur/run/$__runId-idea-batch-create-result.json.tmp .spur/run/$__runId-idea-batch-create-result.json &&
349
+ date -u +%Y-%m-%dT%H:%M:%SZ > .spur/run/$__runId-idea-batch-create.done;
350
+ else
351
+ rm -f .spur/run/$__runId-idea-batch-create-result.json.tmp;
352
+ date -u +%Y-%m-%dT%H:%M:%SZ > .spur/run/$__runId-idea-batch-create.failed;
353
+ fi
354
+
355
+ - id: handoff-finalize
356
+ description: >
357
+ Post-create finalization (0518): zip batch item names to the created WBS
358
+ values from the captured batch-create result (equal-length/unique-name
359
+ checks), apply every non-empty depends_on_names list through
360
+ `spur task deps`, refresh the feature roster, check every created task, and
361
+ write the run-scoped handoff report
362
+ (.spur/run/${vars.__runId}-idea-handoff.md) with exactly one next command.
363
+ A mapping or CLI error fails the run before terminal handoff; an unready
364
+ task is a successful planning outcome recorded as a refineall recommendation.
365
+ onEnter:
366
+ # F1 (0518 verify): the per-task check loop fails closed. A stderr-only (non-JSON)
367
+ # `task check --json` exception would otherwise drop that task from the JSONL results
368
+ # without failing the run (the loop's exit status is its last body command), silently
369
+ # flipping the recommendation to runall — violating R3's any-fail=>refineall invariant.
370
+ # Each `jq -c … >> "$CHECKS"` appends under `|| exit 1`, and a row-count assertion
371
+ # (CHECKS lines == WBS count) gates the recommendation computation.
372
+ - kind: shell
373
+ options:
374
+ command: >-
375
+ BATCH=".spur/run/$__runId-idea-task-batch.json" &&
376
+ RESULT=".spur/run/$__runId-idea-batch-create-result.json" &&
377
+ ORDER=".spur/run/$__runId-idea-task-order.json" &&
378
+ REPORT=".spur/run/$__runId-idea-handoff.md" &&
379
+ DEPMAP=".spur/run/$__runId-idea-dep-map.tsv" &&
380
+ CHECKS=".spur/run/$__runId-idea-check-results.jsonl" &&
381
+ rm -f "$DEPMAP" "$CHECKS" &&
382
+ jq -e --slurpfile b "$BATCH" '(.wbs | length) == ($b[0] | length) and (($b[0] | map(.name) | unique | length) == ($b[0] | length))' "$RESULT" >/dev/null &&
383
+ jq -r --slurpfile b "$BATCH" --slurpfile r "$RESULT" --slurpfile o "$ORDER" '
384
+ ($b[0] | map(.name)) as $names | ($r[0].wbs) as $wbs |
385
+ $o[0] | .[] |
386
+ (.name as $n | $names | index($n)) as $i |
387
+ (if $i == null or $wbs[$i] == null then "MISSING" else $wbs[$i] end) as $own |
388
+ ((.depends_on_names // []) | map(. as $dep | $names | index($dep) as $d | if $d == null then "MISSING" else $wbs[$d] end) | join(" ")) as $deps |
389
+ [.name, $own, $deps] | @tsv
390
+ ' "$ORDER" > "$DEPMAP" &&
391
+ while IFS="$(printf '\t')" read -r name own deps; do
392
+ test "$own" != MISSING || exit 1;
393
+ if test -n "$deps"; then
394
+ for d in $deps; do test "$d" != MISSING || exit 1; done;
395
+ $spurBin task deps "$own" set $deps --json >/dev/null || exit 1;
396
+ fi;
397
+ done < "$DEPMAP" &&
398
+ $spurBin feature refresh --feature "$featureId" --json >/dev/null &&
399
+ WBS_LIST=$(jq -r '.wbs | join(" ")' "$RESULT") &&
400
+ for wbs in $WBS_LIST; do
401
+ TMP=".spur/run/$__runId-idea-check-$wbs.tmp";
402
+ if $spurBin task check "$wbs" --json > "$TMP" 2>&1; then P=true; else P=false; fi;
403
+ jq -c --arg w "$wbs" --argjson p "$P" 'if (type == "array" and length > 0) then {wbs: $w, pass: $p, status: (.[0].status // "unknown")} else {wbs: $w, pass: $p, status: "missing"} end' "$TMP" >> "$CHECKS" || exit 1;
404
+ rm -f "$TMP";
405
+ done &&
406
+ CHECK_ROWS=$(wc -l < "$CHECKS" | tr -d ' ') &&
407
+ WBS_COUNT=$(printf '%s\n' $WBS_LIST | wc -l | tr -d ' ') &&
408
+ test "$CHECK_ROWS" = "$WBS_COUNT" &&
409
+ NEXT=$(jq -r --arg feature "$featureId" --slurpfile c "$CHECKS" 'if ($c | any(.[]; .pass == false)) then "/sp:dev-refineall --feature \($feature) --auto --depth ready" else "/sp:dev-runall --feature \($feature) --auto" end' "$RESULT") &&
410
+ {
411
+ echo "# Idea pipeline handoff report";
412
+ echo;
413
+ echo "Feature: $featureId";
414
+ echo "Run ID: $__runId";
415
+ echo;
416
+ echo "## Created tasks";
417
+ printf '%s\n' $WBS_LIST | sed 's/^/ - /';
418
+ echo;
419
+ echo "## Per-task readiness (spur task check)";
420
+ echo;
421
+ echo "| WBS | Outcome |";
422
+ echo "|-----|---------|";
423
+ while IFS= read -r line; do
424
+ w=$(printf '%s' "$line" | jq -r '.wbs');
425
+ o=$(printf '%s' "$line" | jq -r 'if .pass then "PASS" else "FAIL" end');
426
+ printf '| %s | %s |\n' "$w" "$o";
427
+ done < "$CHECKS";
428
+ echo;
429
+ echo "## Next command";
430
+ echo;
431
+ echo "$NEXT";
432
+ } > "$REPORT"
277
433
 
278
434
  - id: handoff
279
435
  description: >
280
- Terminal — idea pipeline complete. Output: feature id, task WBS list, and
281
- the next command to run (/sp:dev-run <first-wbs> or /sp:dev-runall --feature <id> or --tasks feature:<id>).
282
- No task execution the pipeline stops here.
436
+ Terminal — idea pipeline complete. Ordering applied, roster refreshed, and
437
+ the handoff report written at .spur/run/${vars.__runId}-idea-handoff.md:
438
+ feature id, task WBS list, per-task readiness, and exactly one recommended
439
+ next command (ready-depth refineall when any task is unready, else auto
440
+ runall). No task execution — the pipeline stops here.
283
441
  onEnter:
284
442
  - kind: note
285
443
  options:
286
- message: "Idea pipeline handoff. Feature: ${vars.featureId}. Tasks created. Next: /sp:dev-runall --feature ${vars.featureId}"
444
+ message: "Idea pipeline handoff. Feature: ${vars.featureId}. See .spur/run/${vars.__runId}-idea-handoff.md for the created task list and recommended next command."
287
445
  # Checkpoint write: record session state for resume (Phase 4, task 0171 R3)
288
446
  - kind: shell
289
447
  options:
@@ -374,7 +532,7 @@ transitions:
374
532
  guard:
375
533
  kind: shell
376
534
  options:
377
- command: 'test "$profile" = auto && $spurBin feature check $featureId && test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" != false'
535
+ command: 'test "$profile" = auto && $spurBin feature check "$featureId" && test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" != false'
378
536
  # Auto-skip 2: pass + skip-design route → decompose
379
537
  - from: ac-generate
380
538
  to: decompose
@@ -382,7 +540,7 @@ transitions:
382
540
  guard:
383
541
  kind: shell
384
542
  options:
385
- command: 'test "$profile" = auto && $spurBin feature check $featureId && (test "$design" = skip || (test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" = false))'
543
+ command: 'test "$profile" = auto && $spurBin feature check "$featureId" && (test "$design" = skip || (test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" = false))'
386
544
  # Auto-skip 3: check failed, retry < 3 → loop back to ac-generate (self-loop)
387
545
  - from: ac-generate
388
546
  to: ac-generate
@@ -390,7 +548,7 @@ transitions:
390
548
  guard:
391
549
  kind: shell
392
550
  options:
393
- command: 'test "$profile" = auto && ! $spurBin feature check $featureId && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -lt 3'
551
+ command: 'test "$profile" = auto && ! $spurBin feature check "$featureId" && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -lt 3'
394
552
  # Auto-skip 4: check failed, retry cap reached → escalate to failed
395
553
  - from: ac-generate
396
554
  to: failed
@@ -398,7 +556,7 @@ transitions:
398
556
  guard:
399
557
  kind: shell
400
558
  options:
401
- command: 'test "$profile" = auto && ! $spurBin feature check $featureId && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -ge 3'
559
+ command: 'test "$profile" = auto && ! $spurBin feature check "$featureId" && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -ge 3'
402
560
  # Interactive fallback: enter feature-check HITL gate
403
561
  - from: ac-generate
404
562
  to: feature-check
@@ -414,28 +572,28 @@ transitions:
414
572
  guard:
415
573
  kind: shell
416
574
  options:
417
- command: 'test "$__hitlAnswer" = yes && $spurBin feature check $featureId && test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" != false'
575
+ command: 'test "$__hitlAnswer" = yes && $spurBin feature check "$featureId" && test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" != false'
418
576
  - from: feature-check
419
577
  to: decompose
420
578
  description: Feature check passed, skip-design route — go directly to decompose.
421
579
  guard:
422
580
  kind: shell
423
581
  options:
424
- command: 'test "$__hitlAnswer" = yes && $spurBin feature check $featureId && (test "$design" = skip || (test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" = false))'
582
+ command: 'test "$__hitlAnswer" = yes && $spurBin feature check "$featureId" && (test "$design" = skip || (test "$design" = auto && test "$(jq -r .needs_design .spur/run/$__runId-idea-needs-design.json 2>/dev/null)" = false))'
425
583
  - from: feature-check
426
584
  to: ac-generate
427
585
  description: "Feature check failed — revise AC (retry cap: 3)."
428
586
  guard:
429
587
  kind: shell
430
588
  options:
431
- command: '(test "$__hitlAnswer" = no || ! $spurBin feature check $featureId) && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -lt 3'
589
+ command: '(test "$__hitlAnswer" = no || ! $spurBin feature check "$featureId") && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -lt 3'
432
590
  - from: feature-check
433
591
  to: failed
434
592
  description: Feature check failed after 3 retries — escalate to failed.
435
593
  guard:
436
594
  kind: shell
437
595
  options:
438
- command: '(test "$__hitlAnswer" = no || ! $spurBin feature check $featureId) && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -ge 3'
596
+ command: '(test "$__hitlAnswer" = no || ! $spurBin feature check "$featureId") && test "$(cat .spur/run/$__runId-idea-ac-retry-count 2>/dev/null || echo 0)" -ge 3'
439
597
  - from: feature-check
440
598
  to: cancelled
441
599
  description: Operator cancelled the feature-check gate.
@@ -448,25 +606,32 @@ transitions:
448
606
  # Declaration order: auto-skip guard tried FIRST (like task-pipeline's review->verify pattern).
449
607
  - from: system-design
450
608
  to: decompose
451
- description: profile=auto AND design_approved=true — skip design approval gate.
609
+ description: profile=auto AND design_approved=true AND feature check passes — skip design approval gate.
452
610
  guard:
453
611
  kind: shell
454
612
  options:
455
- command: 'test "$profile" = auto && test "$design_approved" = true'
613
+ command: 'test "$profile" = auto && test "$design_approved" = true && $spurBin feature check "$featureId"'
456
614
  - from: system-design
457
615
  to: design-approval
458
616
  description: System design done — gate on design approval.
459
617
  guard:
460
618
  kind: always
461
619
 
462
- # ── design-approval -> decompose (after operator approval) ──
620
+ # ── design-approval -> decompose (after operator approval), AC-failure -> feature-check ──
463
621
  - from: design-approval
464
622
  to: decompose
465
- description: Design approved — proceed to decomposition.
623
+ description: Design approved and feature check passes — proceed to decomposition.
466
624
  guard:
467
625
  kind: shell
468
626
  options:
469
- command: 'test "$__hitlAnswer" = yes'
627
+ command: 'test "$__hitlAnswer" = yes && $spurBin feature check "$featureId"'
628
+ - from: design-approval
629
+ to: feature-check
630
+ description: "Design approved but feature check fails — route back through the AC gate (revise AC, retry cap: 3)."
631
+ guard:
632
+ kind: shell
633
+ options:
634
+ command: 'test "$__hitlAnswer" = yes && ! $spurBin feature check "$featureId"'
470
635
  - from: design-approval
471
636
  to: system-design
472
637
  description: "Design rejected - revise system design (cap: 1 revise)."
@@ -526,10 +691,10 @@ transitions:
526
691
  options:
527
692
  command: 'test "$__hitlAnswer" = cancel'
528
693
 
529
- # ── batch-create-run: success -> handoff, failure -> decompose (retry cap: 3), failure+cap -> failed ──
694
+ # ── batch-create-run: success -> handoff-finalize, failure -> decompose (retry cap: 3), failure+cap -> failed ──
530
695
  - from: batch-create-run
531
- to: handoff
532
- description: Batch created — handoff.
696
+ to: handoff-finalize
697
+ description: Batch created and result captured apply ordering, refresh roster, write handoff report.
533
698
  guard:
534
699
  kind: shell
535
700
  options:
@@ -548,3 +713,10 @@ transitions:
548
713
  kind: shell
549
714
  options:
550
715
  command: 'test -f .spur/run/$__runId-idea-batch-create.failed && test "$(cat .spur/run/$__runId-idea-decompose-retry-count 2>/dev/null || echo 0)" -ge 3'
716
+
717
+ # ── handoff-finalize: -> handoff (terminal) ──
718
+ - from: handoff-finalize
719
+ to: handoff
720
+ description: Ordering applied, roster refreshed, handoff report written — terminal handoff.
721
+ guard:
722
+ kind: always
@@ -451,7 +451,7 @@ states:
451
451
  - kind: agent.run
452
452
  options:
453
453
  agent: ${vars.agent}
454
- input: /sp:dev-verify ${vars.wbs} --auto --fix all
454
+ input: /sp:dev-verify ${vars.wbs} --auto --fix all --focus all
455
455
  timeoutMs: ${vars.stepTimeoutMs}
456
456
  answerFile: .spur/run/${vars.wbs}-verify-answer.txt
457
457
  - kind: shell
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@gobing-ai/spur",
3
- "version": "0.3.43",
3
+ "version": "0.3.44",
4
4
  "description": "Spur CLI — local-first harness for mainstream coding agents: constraint checking, workflow orchestration, agent health, and history analytics. Bun-native; exposes the `spur` command.",
5
5
  "keywords": [
6
6
  "spur",
@@ -36,7 +36,9 @@ Flags: `--tasks <selector>` (required — explicit WBS list, status pseudo-list,
36
36
  or `ready`), `--feature <id>` (sugar for `feature:<id>`; when the effective selector is
37
37
  feature-derived, `dev-runall` runs `spur feature check <id> --strict --json` **once** before
38
38
  resolving tasks — a non-zero strict check aborts the batch with verdict `aborted` and the
39
- structured findings, before any task pipeline action, task 0510 R2), `--mode`
39
+ structured findings, before any task pipeline action, task 0510 R2; scoped: `L4.scenario-unverified`
40
+ (the expected pre-run state of any not-yet-run feature) is reported verbatim but does not abort —
41
+ any other strict error aborts), `--mode`
40
42
  `<sequential|parallel>` (default `sequential`; `parallel` fans out a proven-independent subset —
41
43
  see `execution-batch.md` § Parallel Execution), `--keep-going`
42
44
  (batch failure policy — skip a failed task's in-batch dependents, continue independents; default
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "sp",
3
- "version": "0.3.43",
3
+ "version": "0.3.44",
4
4
  "description": "Spur — a local-first harness engineering toolkit that wraps mainstream coding agents with constraint checking, workflow orchestration, and history analytics.",
5
5
  "extensions": {
6
6
  "pi": ["./hooks/pi/guard-extension.ts"]
@@ -101,6 +101,9 @@ the whole point of this facade is that the CLI surface has a single, scalable ho
101
101
  prose checks here.
102
102
  - **Not a competency.** Design, implementation, testing, and review are competency skills, not CLI
103
103
  verbs — they are not documented here.
104
+ - **Not the orchestration owner.** The ADR-054 boundary: this facade owns CLI noun/verb/flag
105
+ semantics — including task and feature status-transition verbs — while multi-step lifecycle
106
+ orchestration belongs to `sp:spur-dev`.
104
107
 
105
108
  ## See also
106
109
 
@@ -117,15 +117,15 @@ If any check fails → **child of best parent** or **task under existing feature
117
117
  | --- | --- | --- |
118
118
  | Board module + UX slice as child | `F8` Features board → `F81` detail action group | **Good** — same product surface, finer grain |
119
119
  | Restructure tooling under Feature CLI | `F3` Feature management CLI → `F31` restructure kit | **Good** — not a new root letter (was root `S`, moved) |
120
- | Runner vs plugin integration | `B` Agent execution (`spur agent`) vs `H` Agent integration (`plugins/sp`) | **Keep separate** — Goals at `B_agent-execution.md:14` vs `H_agent-integration.md:14`; **reject merge on “agent” alone** |
120
+ | Runtime vs plugin harness | `B` Agent execution (`spur agent`) vs `I` sp plugin (`plugins/sp`) | **Keep separate** — B owns runner/process/session/executor behavior; I owns skills, commands, subagents, hooks, and orchestration guidance. H is frozen history. |
121
121
  | Observability UX as sibling roots | `J` board + `K` System Events table redesign + `L` payload enrichment | **Reparent K,L under J** — not body-merge; not peer roots (audit 0356) |
122
- | Plugin epics as new letters | `N` dev-next UX, `O` token architecture as roots beside `H` | **Reparent under H** — same plugin plane, finer grain |
122
+ | Plugin epics as new letters | Historical `N` dev-next UX and `O` token architecture | **Historical move under H is retained; new plugin work goes under I** — do not extend frozen H |
123
123
  | Workflow observability as board root | `P` workflow run observability as peer of `D` Workflows | **Reparent under D** — object is `spur workflow run`, not board J |
124
124
  | Planning validation as root | `Q` AC-verifiable gates as peer of `F` Planning | **Reparent under F** |
125
125
  | Status feedback as root | `R` feature status loop as peer of `F` | **Reparent under F** (corpus status is planning) |
126
126
  | CLI backbone vs board product | `G` Collaboration (message/team CLI) vs `M` Teams board | **Keep both** — different surfaces; do not merge |
127
127
  | One-off polish as root | New letter for “fix button loading” | **Bad** — task under existing feature |
128
- | Done historical epic as root | `I` sp plugin hands-off ready (`done`) | **Keep** as historical root OK; optional later neatness under H |
128
+ | Done plugin epic as root | `I1` sp plugin hands-off ready (`done`) | **Good after reparent** completed delivery under durable plugin root I |
129
129
 
130
130
  ### Audit 0356 snapshot (A–R dispositions)
131
131
 
@@ -144,6 +144,11 @@ Authoritative seed for restructure mapping lives in task 0356 Solution and
144
144
 
145
145
  **Rejected merges:** B∪H (name-only overlap); J∪K body-merge (use reparent for K/L).
146
146
 
147
+ **Ownership amendment (2026-08-11):** audit 0356 is a historical snapshot. The current active
148
+ boundary is B = runtime agent execution, I = `sp` plugin harness, H = frozen mixed history. The
149
+ applied map is `docs/plans/2026-08-11-sp-plugin-feature-tree-restructure-map.md`; do not place new
150
+ work under H.
151
+
147
152
  ---
148
153
 
149
154
  ## Checklist: before `spur feature create`
@@ -18,6 +18,21 @@ pipeline run) lives in **`sp:spur-dev`** — do not reimplement that loop here.
18
18
  *drive* a task through its lifecycle, reach for `sp:spur-dev`; when you need to know *which verb
19
19
  does what*, this skill.
20
20
 
21
+ ## WBS lookup fast path
22
+
23
+ Start from the WBS, not the corpus layout:
24
+
25
+ ```bash
26
+ spur task show <wbs> --json # metadata + full content + filePath
27
+ spur task path <wbs> --json # absolute path only
28
+ ```
29
+
30
+ Do not search `docs/tasks*` or guess `--folder` to locate a known WBS. `show` is the default when an
31
+ agent needs to read a task; use `path` only when another filesystem command needs the absolute path.
32
+ Both commands resolve across configured task folders. Add `--folder` only when deliberately limiting
33
+ the lookup to one non-default corpus. Capture `show` once per run and reuse its response instead of
34
+ re-reading or re-tokenizing the task.
35
+
21
36
  ## Verb map
22
37
 
23
38
  | Verb | Purpose | Key flags |
@@ -219,6 +234,13 @@ spur task check --strict --json # elevate ALL warnings to failures
219
234
  spur task check 0040 --strict-core # the testing→done gate variant
220
235
  ```
221
236
 
237
+ **Folder resolution (task 0522):** a WBS-targeted check (`<wbs>` present, no `--folder`) resolves
238
+ the task across **all configured task folders** — the same resolution as `task show` / `task path` /
239
+ `task update` — so a task in an inactive configured folder is checked, not reported missing. An
240
+ explicit `--folder <path>` is normalized to an absolute path and restricts lookup to that single
241
+ directory (relative and absolute spellings are equivalent). Unscoped checks (no WBS) and
242
+ `task list` remain active-folder-only.
243
+
222
244
  `--json` emits the structured matrix — per-task findings (missing sections, broken feature edges,
223
245
  AC-coverage orphans via L4 traceability) keyed by WBS, plus a per-task `pass` verdict. **Query this,
224
246
  do not re-derive it**: parse the JSON to answer "which tasks are ready?", "what's blocking 0040?",
@@ -45,6 +45,10 @@ the lifecycle*; the competency skills know *how to do each job*; the CLI knows *
45
45
  The skill was decomposed **by function** (ADR-028): design, decomposition, implementation, testing,
46
46
  and verification each became a standalone competency skill, leaving this spine to orchestrate them.
47
47
 
48
+ **Ownership (ADR-054).** This spine owns multi-step lifecycle orchestration — intake, gates,
49
+ decomposition, pipeline runs, HITL pauses. CLI noun/verb/flag semantics, including
50
+ status-transition verbs, are the facade's (`sp:spur-cli`), never this skill's.
51
+
48
52
  **The competencies the spine dispatches:**
49
53
 
50
54
  | Unit of work | Competency skill |
@@ -154,13 +158,17 @@ CLI does.
154
158
  you ship corrupted corpus.
155
159
  2. **The pipeline, not you, writes results.** `## Testing` and `## Review` sections are
156
160
  filled by the pipeline's `record` step. Do not edit them directly during execution.
157
- 3. **Check before every write.** Run `spur task check <wbs> --json` to know what sections
161
+ 3. **Resolve task IDs through the CLI.** Read a known WBS with `spur task show <wbs> --json`; it
162
+ returns metadata, full content, and `filePath` across configured task folders. Use `spur task
163
+ path <wbs> --json` only when another tool needs the absolute path. Never search `docs/tasks*` or
164
+ guess `--folder`; reuse the first `show` response throughout the run.
165
+ 4. **Check before every write.** Run `spur task check <wbs> --json` to know what sections
158
166
  the task needs at its current status. Guessing produces matrix violations.
159
- 4. **AC titles are identity keys.** Renaming a scenario after tasks are created breaks
167
+ 5. **AC titles are identity keys.** Renaming a scenario after tasks are created breaks
160
168
  traceability edges. If you must rename, update the task's scenario references too.
161
- 5. **Batch-create is atomic.** A single schema violation rejects the entire batch. Validate
169
+ 6. **Batch-create is atomic.** A single schema violation rejects the entire batch. Validate
162
170
  locally against `task-batch.schema.json` before invoking the CLI.
163
- 6. **Two-halves seam.** The planning and execution halves share this skill today but are
171
+ 7. **Two-halves seam.** The planning and execution halves share this skill today but are
164
172
  designed to split cleanly. Keep new logic in one half or the other — never straddle the
165
173
  seam with cross-half dependencies.
166
174
 
@@ -288,7 +288,7 @@ must not be changed without updating the backing skill.
288
288
  ### 13. runall
289
289
 
290
290
  - **Purpose:** Run a batch of tasks through their pipelines in dependency-correct order — resolve a set, topo-sort, run each via `task-pipeline.yaml`, inspect verdicts, apply the failure policy, emit a batch report.
291
- - **Inputs:** `--tasks <selector>` (required — explicit WBS list, status pseudo-list, `feature:<id>`, or `ready`). `--mode <sequential|parallel>` (default `sequential`; `parallel` fans out a proven-independent subset per `execution-batch.md` § Parallel Execution). `--keep-going` skips a failed task's in-batch dependents and continues independents (default halts on first failure). `--auto` sets `profile=auto` on each per-task run (skips the HITL approve gate). `--feature <id>` is sugar for `--tasks feature:<id>`; when the effective selector is feature-derived, the batch runs `spur feature check <id> --strict --json` **once** before task resolution — a non-zero strict check aborts with verdict `aborted`, zero attempted tasks, and the structured findings (task 0510 R2). Explicit WBS/status/`ready` selectors add no feature check. The orchestrator loop itself continues in this session; interactive sequential omit/`inline` uses the host driver — host-controlled with eligible `agent.run` stages dispatching once to a native subagent and host fallback (task 0508) — while `--agent auto`/a name, parallel mode, and headless invocation keep the isolated per-task workflow subprocess boundary. `--agent <inline|auto|name>` pins the executor for the per-task stages (see [SSOT](cross-cutting.md#inline-default-execution-surface)); the value crosses into per-task `vars.agent`, not the orchestrator. `--json` emits the report as JSON. `--wrap` triggers `wrapup-pipeline.yaml` after the batch completes, `--next` chains each task to terminal status then runs the wrap hop **once for the batch**, `--continue` resumes from checkpoint.
291
+ - **Inputs:** `--tasks <selector>` (required — explicit WBS list, status pseudo-list, `feature:<id>`, or `ready`). `--mode <sequential|parallel>` (default `sequential`; `parallel` fans out a proven-independent subset per `execution-batch.md` § Parallel Execution). `--keep-going` skips a failed task's in-batch dependents and continues independents (default halts on first failure). `--auto` sets `profile=auto` on each per-task run (skips the HITL approve gate). `--feature <id>` is sugar for `--tasks feature:<id>`; when the effective selector is feature-derived, the batch runs `spur feature check <id> --strict --json` **once** before task resolution — a non-zero strict check aborts with verdict `aborted`, zero attempted tasks, and the structured findings (task 0510 R2); scoped: `L4.scenario-unverified` (expected pre-run state of any not-yet-run feature) is reported verbatim but does not abort — any other strict error aborts. Explicit WBS/status/`ready` selectors add no feature check. The orchestrator loop itself continues in this session; interactive sequential omit/`inline` uses the host driver — host-controlled with eligible `agent.run` stages dispatching once to a native subagent and host fallback (task 0508) — while `--agent auto`/a name, parallel mode, and headless invocation keep the isolated per-task workflow subprocess boundary. `--agent <inline|auto|name>` pins the executor for the per-task stages (see [SSOT](cross-cutting.md#inline-default-execution-surface)); the value crosses into per-task `vars.agent`, not the orchestrator. `--json` emits the report as JSON. `--wrap` triggers `wrapup-pipeline.yaml` after the batch completes, `--next` chains each task to terminal status then runs the wrap hop **once for the batch**, `--continue` resumes from checkpoint.
292
292
  - **Three orthogonal axes (do not confuse):** `--keep-going` = batch failure policy (halt vs skip dependents); `--continue` = resume from checkpoint (pick up an interrupted batch); `--next` = per-task lifecycle chaining (advance status on a verdict — `dev-verify`/`dev-verifyall` only). `routing-table.md` offers `--continue` and `--next` as competing options for the same situation only when the batch was interrupted mid-run; otherwise they address different problems.
293
293
  - **Backing:** `sp:spur-dev` skill, `runall` operation → delegates the driver loop to the **`sp:super-planner`** agent (the batch orchestrator).
294
294
  - **Behavior:** The orchestrator reads [execution-batch.md](execution-batch.md) and drives: resolve selector → freeze set → topo-sort by `dependencies[]` (Kahn, WBS-ascending tie-break; cycle aborts) → resolve out-of-set deps by status (done → allow, else → block subtree) → run each task via `spur workflow run task-pipeline.yaml --async` + `spur workflow trace` polling → inspect terminal state + `.spur/run/<wbs>-verdict.json` → stop-the-batch default or `--keep-going` subtree skip → emit batch report. Per-task pipeline is invoked **verbatim** — no new FSM, no step edits. `--auto`/`--agent` are the only flags that cross the orchestrator→pipeline boundary (both into per-task `--vars`).