novahiz 0.2.2 → 0.2.5

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (565) hide show
  1. package/NOTICE.md +25 -15
  2. package/README.md +353 -355
  3. package/adapters/README.md +1 -1
  4. package/adapters/opencode/agent/novahiz.md +1 -2
  5. package/adapters/opencode/commands/novahiz-init.md +10 -0
  6. package/adapters/opencode/instructions.md +1 -1
  7. package/adapters/opencode/novahiz.ts +182 -4
  8. package/bin/novahiz.cmd +2 -2
  9. package/catalog/categories.json +1085 -1093
  10. package/catalog/overrides.json +12 -13
  11. package/catalog/providers.json +18 -34
  12. package/catalog/rules.json +15 -22
  13. package/docs/ARCHITECTURE.md +2 -2
  14. package/docs/CATALOG.md +7 -6
  15. package/docs/CONFIGURATION.md +5 -5
  16. package/docs/CONSTITUTION.md +4 -4
  17. package/docs/GATE.md +7 -11
  18. package/docs/HARNESSES.md +11 -0
  19. package/docs/INSTALL.md +7 -5
  20. package/docs/PLUGIN.md +24 -6
  21. package/docs/PROVIDERS.md +12 -10
  22. package/docs/ROADMAPS.md +4 -3
  23. package/docs/RULES.md +11 -10
  24. package/install/bootstrap.mjs +9 -19
  25. package/install/install.mjs +23 -41
  26. package/mcp/novahiz-tools/index.mjs +35 -6
  27. package/novahiz.config.example.json +1 -5
  28. package/package.json +6 -7
  29. package/skills/apple-ui-audit/SKILL.md +95 -0
  30. package/skills/apple-ui-audit/references/screen-walkthrough.md +29 -0
  31. package/skills/browser-session/SKILL.md +61 -0
  32. package/skills/browser-session/references/session-run-sheet.md +30 -0
  33. package/skills/browser-session/scripts/browser_session_preflight.mjs +74 -0
  34. package/skills/browser-session/scripts/sample-run.json +5 -0
  35. package/skills/code-standards/SKILL.md +85 -0
  36. package/skills/code-standards/references/pr-review-worksheet.md +35 -0
  37. package/skills/design-token-pipeline/SKILL.md +85 -0
  38. package/skills/design-token-pipeline/references/dtcg-sample.json +37 -0
  39. package/skills/llm-threat-review/SKILL.md +99 -0
  40. package/skills/llm-threat-review/references/review-worksheet.md +34 -0
  41. package/skills/llm-threat-review/scripts/llm_risk_lint.mjs +105 -0
  42. package/skills/novahiz-audit/SKILL.md +1 -1
  43. package/skills/novahiz-gate/SKILL.md +7 -6
  44. package/skills/novahiz-humanizer/SKILL.md +1 -1
  45. package/skills/novahiz-implement/SKILL.md +1 -1
  46. package/skills/novahiz-init/SKILL.md +112 -0
  47. package/skills/novahiz-planner/SKILL.md +1 -2
  48. package/skills/openapi-mcp-server/SKILL.md +76 -0
  49. package/skills/openapi-mcp-server/references/mcp-ship-checklist.md +30 -0
  50. package/skills/openapi-mcp-server/references/tool-worksheet.md +30 -0
  51. package/skills/openapi-mcp-server/scripts/openapi_mcp_lint.mjs +133 -0
  52. package/skills/package-risk-audit/SKILL.md +83 -0
  53. package/skills/package-risk-audit/references/intake-worksheet.md +19 -0
  54. package/skills/package-risk-audit/scripts/dep_risk_lint.mjs +131 -0
  55. package/skills/secrets-hygiene/SKILL.md +62 -0
  56. package/skills/secrets-hygiene/references/rotation-playbook.md +27 -0
  57. package/skills/secrets-hygiene/scripts/secret_scan.mjs +120 -0
  58. package/skills/skill-authoring/SKILL.md +93 -0
  59. package/skills/skill-authoring/references/skill-template.md +59 -0
  60. package/skills/skill-authoring/scripts/skill_frontmatter_lint.mjs +65 -0
  61. package/skills/skill-eval-loop/SKILL.md +103 -0
  62. package/skills/skill-eval-loop/references/negative-patterns.md +12 -0
  63. package/skills/skill-eval-loop/references/run-sheet.md +40 -0
  64. package/skills/skill-eval-loop/scripts/skill_structure_check.mjs +100 -0
  65. package/skills/ui-craft-rules/SKILL.md +96 -0
  66. package/skills/ui-slop-remover/SKILL.md +110 -0
  67. package/skills/ui-slop-remover/references/tell-catalogue.md +58 -0
  68. package/skills/ui-slop-remover/scripts/slop_lint.mjs +119 -0
  69. package/src/autodocs.ts +192 -0
  70. package/src/cli.ts +8 -1
  71. package/src/commands/autodocs.ts +250 -0
  72. package/src/commands/doctor.ts +226 -217
  73. package/src/commands/gate.ts +16 -5
  74. package/src/commands/init.ts +356 -0
  75. package/src/commands/task.ts +20 -7
  76. package/src/gate.ts +401 -401
  77. package/src/ledger.ts +39 -0
  78. package/src/targets.ts +43 -7
  79. package/bundled-skills/ab-testing/SKILL.md +0 -353
  80. package/bundled-skills/ab-testing/evals/evals.json +0 -105
  81. package/bundled-skills/ab-testing/references/sample-size-guide.md +0 -263
  82. package/bundled-skills/ab-testing/references/test-templates.md +0 -277
  83. package/bundled-skills/ad-creative/SKILL.md +0 -425
  84. package/bundled-skills/ad-creative/assets/creative-review-template.html +0 -400
  85. package/bundled-skills/ad-creative/evals/evals.json +0 -217
  86. package/bundled-skills/ad-creative/references/creative-review-page.md +0 -106
  87. package/bundled-skills/ad-creative/references/creative-roadmap.md +0 -118
  88. package/bundled-skills/ad-creative/references/generative-tools.md +0 -637
  89. package/bundled-skills/ad-creative/references/hook-system.md +0 -115
  90. package/bundled-skills/ad-creative/references/imessage-video-ads.md +0 -201
  91. package/bundled-skills/ad-creative/references/meta-creative-formats.md +0 -114
  92. package/bundled-skills/ad-creative/references/motion-video-ads.md +0 -126
  93. package/bundled-skills/ad-creative/references/platform-specs.md +0 -213
  94. package/bundled-skills/ad-creative/references/short-form-video-specs.md +0 -213
  95. package/bundled-skills/ad-creative/references/static-ad-templates.md +0 -308
  96. package/bundled-skills/ads/SKILL.md +0 -499
  97. package/bundled-skills/ads/evals/evals.json +0 -157
  98. package/bundled-skills/ads/references/abm-playbook.md +0 -93
  99. package/bundled-skills/ads/references/ad-copy-templates.md +0 -207
  100. package/bundled-skills/ads/references/audience-targeting.md +0 -243
  101. package/bundled-skills/ads/references/audit-guardrails.md +0 -83
  102. package/bundled-skills/ads/references/b2b-paid-playbook.md +0 -115
  103. package/bundled-skills/ads/references/conversion-tracking.md +0 -361
  104. package/bundled-skills/ads/references/creative-research-automation.md +0 -103
  105. package/bundled-skills/ads/references/google-ads-audit-checklist.md +0 -84
  106. package/bundled-skills/ads/references/google-search-playbook.md +0 -120
  107. package/bundled-skills/ads/references/linkedin-b2b-playbook.md +0 -107
  108. package/bundled-skills/ads/references/meta-decision-system.md +0 -176
  109. package/bundled-skills/ads/references/payback-period.md +0 -69
  110. package/bundled-skills/ads/references/platform-setup-checklists.md +0 -277
  111. package/bundled-skills/ads/references/rsa-output-spec.md +0 -88
  112. package/bundled-skills/agent-device/SKILL.md +0 -20
  113. package/bundled-skills/ai-seo/SKILL.md +0 -488
  114. package/bundled-skills/ai-seo/evals/evals.json +0 -135
  115. package/bundled-skills/ai-seo/references/agent-readiness.md +0 -57
  116. package/bundled-skills/ai-seo/references/citations-vs-recommendations.md +0 -85
  117. package/bundled-skills/ai-seo/references/content-patterns.md +0 -287
  118. package/bundled-skills/ai-seo/references/content-types.md +0 -71
  119. package/bundled-skills/ai-seo/references/format-volatility.md +0 -98
  120. package/bundled-skills/ai-seo/references/okf.md +0 -104
  121. package/bundled-skills/ai-seo/references/platform-ranking-factors.md +0 -154
  122. package/bundled-skills/ai-seo/references/youtube-ai-citations.md +0 -60
  123. package/bundled-skills/analytics/SKILL.md +0 -310
  124. package/bundled-skills/analytics/evals/evals.json +0 -90
  125. package/bundled-skills/analytics/references/event-library.md +0 -260
  126. package/bundled-skills/analytics/references/ga4-implementation.md +0 -300
  127. package/bundled-skills/analytics/references/gtm-implementation.md +0 -390
  128. package/bundled-skills/android-emulator/SKILL.md +0 -32
  129. package/bundled-skills/android-reverse-engineering/LICENSE +0 -190
  130. package/bundled-skills/android-reverse-engineering/SKILL.md +0 -313
  131. package/bundled-skills/android-reverse-engineering/references/api-extraction-patterns.md +0 -197
  132. package/bundled-skills/android-reverse-engineering/references/call-flow-analysis.md +0 -210
  133. package/bundled-skills/android-reverse-engineering/references/fernflower-usage.md +0 -115
  134. package/bundled-skills/android-reverse-engineering/references/jadx-usage.md +0 -116
  135. package/bundled-skills/android-reverse-engineering/references/kotlin-name-recovery.md +0 -108
  136. package/bundled-skills/android-reverse-engineering/references/setup-guide.md +0 -221
  137. package/bundled-skills/android-reverse-engineering/references/third_party_hosts.txt +0 -122
  138. package/bundled-skills/android-reverse-engineering/scripts/check-deps.ps1 +0 -165
  139. package/bundled-skills/android-reverse-engineering/scripts/check-deps.sh +0 -129
  140. package/bundled-skills/android-reverse-engineering/scripts/decompile.ps1 +0 -400
  141. package/bundled-skills/android-reverse-engineering/scripts/decompile.sh +0 -578
  142. package/bundled-skills/android-reverse-engineering/scripts/find-api-calls.ps1 +0 -121
  143. package/bundled-skills/android-reverse-engineering/scripts/find-api-calls.sh +0 -339
  144. package/bundled-skills/android-reverse-engineering/scripts/fingerprint.sh +0 -241
  145. package/bundled-skills/android-reverse-engineering/scripts/install-dep.ps1 +0 -345
  146. package/bundled-skills/android-reverse-engineering/scripts/install-dep.sh +0 -448
  147. package/bundled-skills/android-reverse-engineering/scripts/lookup-name.sh +0 -85
  148. package/bundled-skills/android-reverse-engineering/scripts/recover-kotlin-names.sh +0 -140
  149. package/bundled-skills/build-graph/SKILL.md +0 -38
  150. package/bundled-skills/claude-history-ingest/SKILL.md +0 -463
  151. package/bundled-skills/claude-history-ingest/references/claude-data-format.md +0 -118
  152. package/bundled-skills/codex-history-ingest/SKILL.md +0 -248
  153. package/bundled-skills/codex-history-ingest/references/codex-data-format.md +0 -82
  154. package/bundled-skills/cold-email/SKILL.md +0 -159
  155. package/bundled-skills/cold-email/evals/evals.json +0 -94
  156. package/bundled-skills/cold-email/references/benchmarks.md +0 -83
  157. package/bundled-skills/cold-email/references/follow-up-sequences.md +0 -81
  158. package/bundled-skills/cold-email/references/frameworks.md +0 -90
  159. package/bundled-skills/cold-email/references/personalization.md +0 -79
  160. package/bundled-skills/cold-email/references/subject-lines.md +0 -53
  161. package/bundled-skills/competitor-profiling/SKILL.md +0 -415
  162. package/bundled-skills/competitor-profiling/evals/evals.json +0 -85
  163. package/bundled-skills/competitor-profiling/references/templates.md +0 -167
  164. package/bundled-skills/competitor-profiling/references/tool-reference.md +0 -179
  165. package/bundled-skills/computer-use/SKILL.md +0 -41
  166. package/bundled-skills/content-strategy/SKILL.md +0 -439
  167. package/bundled-skills/content-strategy/evals/evals.json +0 -116
  168. package/bundled-skills/content-strategy/references/content-distribution.md +0 -83
  169. package/bundled-skills/content-strategy/references/headless-cms.md +0 -194
  170. package/bundled-skills/copilot-history-ingest/SKILL.md +0 -376
  171. package/bundled-skills/copilot-history-ingest/references/copilot-data-format.md +0 -321
  172. package/bundled-skills/copy-editing/SKILL.md +0 -457
  173. package/bundled-skills/copy-editing/evals/evals.json +0 -89
  174. package/bundled-skills/copy-editing/references/checklist.md +0 -66
  175. package/bundled-skills/copy-editing/references/content-refresh.md +0 -38
  176. package/bundled-skills/copy-editing/references/plain-english-alternatives.md +0 -394
  177. package/bundled-skills/copywriting/SKILL.md +0 -256
  178. package/bundled-skills/copywriting/evals/evals.json +0 -126
  179. package/bundled-skills/copywriting/references/copy-frameworks.md +0 -433
  180. package/bundled-skills/copywriting/references/natural-transitions.md +0 -272
  181. package/bundled-skills/cro/SKILL.md +0 -187
  182. package/bundled-skills/cro/evals/evals.json +0 -111
  183. package/bundled-skills/cro/references/experiments.md +0 -248
  184. package/bundled-skills/cro/references/form.md +0 -422
  185. package/bundled-skills/cross-linker/SKILL.md +0 -309
  186. package/bundled-skills/customer-research/SKILL.md +0 -305
  187. package/bundled-skills/customer-research/evals/evals.json +0 -176
  188. package/bundled-skills/customer-research/references/interviews-and-surveys.md +0 -155
  189. package/bundled-skills/customer-research/references/source-guides.md +0 -401
  190. package/bundled-skills/daily-update/SKILL.md +0 -223
  191. package/bundled-skills/dart-add-unit-test/SKILL.md +0 -122
  192. package/bundled-skills/dart-build-cli-app/SKILL.md +0 -185
  193. package/bundled-skills/dart-collect-coverage/SKILL.md +0 -141
  194. package/bundled-skills/dart-fix-runtime-errors/SKILL.md +0 -166
  195. package/bundled-skills/dart-generate-test-mocks/SKILL.md +0 -155
  196. package/bundled-skills/dart-migrate-to-checks-package/SKILL.md +0 -528
  197. package/bundled-skills/dart-resolve-package-conflicts/SKILL.md +0 -116
  198. package/bundled-skills/dart-run-static-analysis/SKILL.md +0 -104
  199. package/bundled-skills/dart-setup-ffi-assets/SKILL.md +0 -419
  200. package/bundled-skills/dart-use-doc-examples/SKILL.md +0 -96
  201. package/bundled-skills/dart-use-ffigen/SKILL.md +0 -218
  202. package/bundled-skills/dart-use-pattern-matching/SKILL.md +0 -146
  203. package/bundled-skills/dart-use-primary-constructors/SKILL.md +0 -262
  204. package/bundled-skills/dart-write-documentation/SKILL.md +0 -136
  205. package/bundled-skills/debug-issue/SKILL.md +0 -28
  206. package/bundled-skills/dogfood/SKILL.md +0 -27
  207. package/bundled-skills/eas-app-stores/SKILL.md +0 -176
  208. package/bundled-skills/eas-app-stores/agents/openai.yaml +0 -4
  209. package/bundled-skills/eas-app-stores/references/app-store-metadata.md +0 -497
  210. package/bundled-skills/eas-app-stores/references/ios-app-store.md +0 -376
  211. package/bundled-skills/eas-app-stores/references/native-ios.md +0 -167
  212. package/bundled-skills/eas-app-stores/references/play-store.md +0 -244
  213. package/bundled-skills/eas-app-stores/references/testflight.md +0 -62
  214. package/bundled-skills/eas-app-stores/references/workflows.md +0 -120
  215. package/bundled-skills/eas-hosting/SKILL.md +0 -431
  216. package/bundled-skills/eas-hosting/agents/openai.yaml +0 -4
  217. package/bundled-skills/eas-observe/SKILL.md +0 -54
  218. package/bundled-skills/eas-observe/agents/openai.yaml +0 -4
  219. package/bundled-skills/eas-observe/references/metrics.md +0 -98
  220. package/bundled-skills/eas-observe/references/queries.md +0 -403
  221. package/bundled-skills/eas-observe/references/setup.md +0 -476
  222. package/bundled-skills/eas-observe/references/third-party.md +0 -136
  223. package/bundled-skills/eas-simulator/SKILL.md +0 -208
  224. package/bundled-skills/eas-simulator/agents/openai.yaml +0 -4
  225. package/bundled-skills/eas-simulator/references/controllers.md +0 -106
  226. package/bundled-skills/eas-simulator/references/run-your-app.md +0 -224
  227. package/bundled-skills/eas-simulator/references/troubleshooting.md +0 -44
  228. package/bundled-skills/eas-update/SKILL.md +0 -146
  229. package/bundled-skills/eas-update/agents/openai.yaml +0 -4
  230. package/bundled-skills/eas-update-insights/SKILL.md +0 -238
  231. package/bundled-skills/eas-update-insights/agents/openai.yaml +0 -4
  232. package/bundled-skills/eas-update-insights/references/channel-insights-schema.md +0 -47
  233. package/bundled-skills/eas-update-insights/references/update-insights-schema.md +0 -69
  234. package/bundled-skills/eas-workflows/SKILL.md +0 -99
  235. package/bundled-skills/eas-workflows/agents/openai.yaml +0 -4
  236. package/bundled-skills/eas-workflows/scripts/fetch.js +0 -109
  237. package/bundled-skills/eas-workflows/scripts/package.json +0 -6
  238. package/bundled-skills/emails/SKILL.md +0 -311
  239. package/bundled-skills/emails/evals/evals.json +0 -93
  240. package/bundled-skills/emails/references/copy-guidelines.md +0 -113
  241. package/bundled-skills/emails/references/email-types.md +0 -515
  242. package/bundled-skills/emails/references/sequence-templates.md +0 -168
  243. package/bundled-skills/explore-codebase/SKILL.md +0 -29
  244. package/bundled-skills/expo-animation/LICENSE +0 -21
  245. package/bundled-skills/expo-animation/RECIPES.md +0 -385
  246. package/bundled-skills/expo-animation/SKILL.md +0 -267
  247. package/bundled-skills/expo-animation/agents/openai.yaml +0 -4
  248. package/bundled-skills/expo-app-clip/SKILL.md +0 -290
  249. package/bundled-skills/expo-app-clip/agents/openai.yaml +0 -4
  250. package/bundled-skills/expo-app-clip/references/native-module.md +0 -96
  251. package/bundled-skills/expo-brownfield/SKILL.md +0 -62
  252. package/bundled-skills/expo-brownfield/agents/openai.yaml +0 -4
  253. package/bundled-skills/expo-brownfield/references/brownfield-integrated.md +0 -526
  254. package/bundled-skills/expo-brownfield/references/brownfield-isolated.md +0 -451
  255. package/bundled-skills/expo-brownfield/references/comparison.md +0 -63
  256. package/bundled-skills/expo-brownfield/references/troubleshooting.md +0 -88
  257. package/bundled-skills/expo-data-fetching/SKILL.md +0 -478
  258. package/bundled-skills/expo-data-fetching/agents/openai.yaml +0 -4
  259. package/bundled-skills/expo-data-fetching/references/expo-router-loaders.md +0 -341
  260. package/bundled-skills/expo-data-fetching/references/offline-and-cancellation.md +0 -68
  261. package/bundled-skills/expo-design-system/SKILL.md +0 -376
  262. package/bundled-skills/expo-design-system/agents/openai.yaml +0 -4
  263. package/bundled-skills/expo-design-system/references/audit.md +0 -190
  264. package/bundled-skills/expo-design-system/references/native-slop.md +0 -74
  265. package/bundled-skills/expo-dev-client/SKILL.md +0 -182
  266. package/bundled-skills/expo-dev-client/agents/openai.yaml +0 -4
  267. package/bundled-skills/expo-dom/SKILL.md +0 -425
  268. package/bundled-skills/expo-dom/agents/openai.yaml +0 -4
  269. package/bundled-skills/expo-examples/SKILL.md +0 -106
  270. package/bundled-skills/expo-examples/agents/openai.yaml +0 -4
  271. package/bundled-skills/expo-examples/references/catalog.md +0 -105
  272. package/bundled-skills/expo-migrate-module/SKILL.md +0 -113
  273. package/bundled-skills/expo-migrate-module/agents/openai.yaml +0 -4
  274. package/bundled-skills/expo-migrate-module/references/compatibility.md +0 -73
  275. package/bundled-skills/expo-migrate-module/references/example.md +0 -212
  276. package/bundled-skills/expo-migrate-module/references/migration-map.md +0 -306
  277. package/bundled-skills/expo-module/SKILL.md +0 -151
  278. package/bundled-skills/expo-module/agents/openai.yaml +0 -4
  279. package/bundled-skills/expo-module/references/config-plugin.md +0 -90
  280. package/bundled-skills/expo-module/references/create-expo-module.md +0 -206
  281. package/bundled-skills/expo-module/references/lifecycle.md +0 -127
  282. package/bundled-skills/expo-module/references/module-config.md +0 -48
  283. package/bundled-skills/expo-module/references/native-module.md +0 -286
  284. package/bundled-skills/expo-module/references/native-view.md +0 -171
  285. package/bundled-skills/expo-native-ui/SKILL.md +0 -201
  286. package/bundled-skills/expo-native-ui/agents/openai.yaml +0 -4
  287. package/bundled-skills/expo-native-ui/references/controls.md +0 -231
  288. package/bundled-skills/expo-native-ui/references/gradients.md +0 -106
  289. package/bundled-skills/expo-native-ui/references/icons.md +0 -232
  290. package/bundled-skills/expo-native-ui/references/media.md +0 -193
  291. package/bundled-skills/expo-native-ui/references/storage.md +0 -121
  292. package/bundled-skills/expo-native-ui/references/visual-effects.md +0 -198
  293. package/bundled-skills/expo-native-ui/references/webgpu-three.md +0 -605
  294. package/bundled-skills/expo-overview/SKILL.md +0 -115
  295. package/bundled-skills/expo-overview/agents/openai.yaml +0 -4
  296. package/bundled-skills/expo-project-structure/SKILL.md +0 -114
  297. package/bundled-skills/expo-project-structure/agents/openai.yaml +0 -4
  298. package/bundled-skills/expo-router/SKILL.md +0 -240
  299. package/bundled-skills/expo-router/agents/openai.yaml +0 -4
  300. package/bundled-skills/expo-router/references/form-sheet.md +0 -253
  301. package/bundled-skills/expo-router/references/route-structure.md +0 -229
  302. package/bundled-skills/expo-router/references/search.md +0 -248
  303. package/bundled-skills/expo-router/references/tabs.md +0 -433
  304. package/bundled-skills/expo-router/references/toolbar-and-headers.md +0 -284
  305. package/bundled-skills/expo-router/references/zoom-transitions.md +0 -158
  306. package/bundled-skills/expo-skill-eval/SKILL.md +0 -314
  307. package/bundled-skills/expo-skill-eval/agents/visual-grader.md +0 -52
  308. package/bundled-skills/expo-skill-eval/references/design-rubric.md +0 -45
  309. package/bundled-skills/expo-skill-eval/references/runtime-matrix.md +0 -34
  310. package/bundled-skills/expo-skill-eval/scripts/check-static.sh +0 -64
  311. package/bundled-skills/expo-skill-eval/scripts/clean-fixture.sh +0 -56
  312. package/bundled-skills/expo-skill-eval/scripts/generate_viewer.py +0 -450
  313. package/bundled-skills/expo-skill-eval/scripts/latest-sdk.sh +0 -46
  314. package/bundled-skills/expo-skill-eval/scripts/make-fixture.sh +0 -95
  315. package/bundled-skills/expo-skill-eval/scripts/make-workspace.sh +0 -31
  316. package/bundled-skills/expo-skill-eval/scripts/snapshot-android.sh +0 -284
  317. package/bundled-skills/expo-skill-eval/scripts/snapshot-ios.sh +0 -132
  318. package/bundled-skills/expo-skill-eval/scripts/snapshot-web.sh +0 -61
  319. package/bundled-skills/expo-skill-feedback/SKILL.md +0 -87
  320. package/bundled-skills/expo-skill-feedback/agents/openai.yaml +0 -4
  321. package/bundled-skills/expo-skill-feedback/scripts/skill-event.cjs +0 -180
  322. package/bundled-skills/expo-skill-feedback/scripts/telemetry.cjs +0 -65
  323. package/bundled-skills/expo-skill-feedback/scripts/telemetry_common.cjs +0 -171
  324. package/bundled-skills/expo-ui/SKILL.md +0 -101
  325. package/bundled-skills/expo-ui/agents/openai.yaml +0 -4
  326. package/bundled-skills/expo-ui/references/drop-in-replacements.md +0 -27
  327. package/bundled-skills/expo-ui/references/jetpack-compose.md +0 -73
  328. package/bundled-skills/expo-ui/references/swift-ui.md +0 -73
  329. package/bundled-skills/expo-ui/references/universal.md +0 -104
  330. package/bundled-skills/expo-ui/scripts/list-components.js +0 -193
  331. package/bundled-skills/expo-upgrade/SKILL.md +0 -150
  332. package/bundled-skills/expo-upgrade/agents/openai.yaml +0 -4
  333. package/bundled-skills/expo-upgrade/references/expo-av-to-audio.md +0 -132
  334. package/bundled-skills/expo-upgrade/references/expo-av-to-video.md +0 -160
  335. package/bundled-skills/expo-upgrade/references/native-tabs.md +0 -124
  336. package/bundled-skills/expo-upgrade/references/new-architecture.md +0 -79
  337. package/bundled-skills/expo-upgrade/references/react-19.md +0 -79
  338. package/bundled-skills/expo-upgrade/references/react-compiler.md +0 -59
  339. package/bundled-skills/expo-upgrade/references/react-navigation-to-expo-router.md +0 -61
  340. package/bundled-skills/expo-web-to-native/SKILL.md +0 -91
  341. package/bundled-skills/expo-web-to-native/agents/openai.yaml +0 -4
  342. package/bundled-skills/expo-web-to-native/references/false-friends.md +0 -119
  343. package/bundled-skills/expo-web-to-native/references/native-patterns.md +0 -39
  344. package/bundled-skills/expo-web-to-native/references/run-as-goal.md +0 -48
  345. package/bundled-skills/expo-web-to-native/references/verify-on-device.md +0 -45
  346. package/bundled-skills/find-skills/SKILL.md +0 -141
  347. package/bundled-skills/flutter-add-integration-test/SKILL.md +0 -163
  348. package/bundled-skills/flutter-add-widget-preview/SKILL.md +0 -145
  349. package/bundled-skills/flutter-add-widget-test/SKILL.md +0 -154
  350. package/bundled-skills/flutter-apply-architecture-best-practices/SKILL.md +0 -162
  351. package/bundled-skills/flutter-build-responsive-layout/SKILL.md +0 -139
  352. package/bundled-skills/flutter-fix-layout-issues/SKILL.md +0 -130
  353. package/bundled-skills/flutter-implement-json-serialization/SKILL.md +0 -153
  354. package/bundled-skills/flutter-setup-declarative-routing/SKILL.md +0 -255
  355. package/bundled-skills/flutter-setup-localization/SKILL.md +0 -210
  356. package/bundled-skills/flutter-use-http-package/SKILL.md +0 -174
  357. package/bundled-skills/graph-colorize/SKILL.md +0 -180
  358. package/bundled-skills/hermes-history-ingest/SKILL.md +0 -239
  359. package/bundled-skills/hermes-history-ingest/references/hermes-data-format.md +0 -131
  360. package/bundled-skills/impl-validator/SKILL.md +0 -118
  361. package/bundled-skills/ios-simulator/SKILL.md +0 -32
  362. package/bundled-skills/json-canvas/SKILL.md +0 -244
  363. package/bundled-skills/json-canvas/references/EXAMPLES.md +0 -329
  364. package/bundled-skills/launch/SKILL.md +0 -381
  365. package/bundled-skills/launch/evals/evals.json +0 -106
  366. package/bundled-skills/llm-wiki/SKILL.md +0 -661
  367. package/bundled-skills/llm-wiki/references/WRITING.md +0 -28
  368. package/bundled-skills/llm-wiki/references/karpathy-pattern.md +0 -45
  369. package/bundled-skills/marketing-plan/SKILL.md +0 -300
  370. package/bundled-skills/marketing-plan/evals/evals.json +0 -112
  371. package/bundled-skills/marketing-plan/references/aarrr-framework.md +0 -180
  372. package/bundled-skills/marketing-plan/references/budget-planning.md +0 -168
  373. package/bundled-skills/marketing-plan/references/client-types.md +0 -373
  374. package/bundled-skills/marketing-plan/references/current-state-rubric.md +0 -255
  375. package/bundled-skills/marketing-plan/references/example-quietude.md +0 -972
  376. package/bundled-skills/marketing-plan/references/funding-stage-unlocks.md +0 -230
  377. package/bundled-skills/marketing-plan/references/growth-patterns.md +0 -190
  378. package/bundled-skills/marketing-plan/references/idea-cross-reference.md +0 -265
  379. package/bundled-skills/marketing-plan/references/measurement-framework.md +0 -213
  380. package/bundled-skills/marketing-plan/references/methodology.md +0 -363
  381. package/bundled-skills/marketing-plan/references/ops-stack-mapping.md +0 -197
  382. package/bundled-skills/marketing-plan/references/plan-template.md +0 -494
  383. package/bundled-skills/marketing-plan/references/team-and-agency-model.md +0 -278
  384. package/bundled-skills/marketing-psychology/SKILL.md +0 -455
  385. package/bundled-skills/marketing-psychology/evals/evals.json +0 -88
  386. package/bundled-skills/memory-bridge/SKILL.md +0 -163
  387. package/bundled-skills/memory-save/SKILL.md +0 -88
  388. package/bundled-skills/obsidian-bases/SKILL.md +0 -499
  389. package/bundled-skills/obsidian-bases/references/FUNCTIONS_REFERENCE.md +0 -173
  390. package/bundled-skills/obsidian-cli/SKILL.md +0 -106
  391. package/bundled-skills/obsidian-layout-adjustment/SKILL.md +0 -226
  392. package/bundled-skills/obsidian-layout-adjustment/evals/evals.json +0 -53
  393. package/bundled-skills/obsidian-layout-adjustment/references/workflow-reference.md +0 -212
  394. package/bundled-skills/obsidian-markdown/SKILL.md +0 -196
  395. package/bundled-skills/obsidian-markdown/references/CALLOUTS.md +0 -58
  396. package/bundled-skills/obsidian-markdown/references/EMBEDS.md +0 -70
  397. package/bundled-skills/obsidian-markdown/references/PROPERTIES.md +0 -61
  398. package/bundled-skills/offers/SKILL.md +0 -155
  399. package/bundled-skills/offers/evals/evals.json +0 -21
  400. package/bundled-skills/offers/references/bonus-stacking.md +0 -150
  401. package/bundled-skills/offers/references/examples.md +0 -215
  402. package/bundled-skills/offers/references/guarantee-design.md +0 -172
  403. package/bundled-skills/offers/references/offer-anatomy.md +0 -203
  404. package/bundled-skills/offers/references/offer-formats.md +0 -240
  405. package/bundled-skills/offers/references/saas-offers.md +0 -116
  406. package/bundled-skills/offers/references/scarcity-urgency.md +0 -175
  407. package/bundled-skills/offers/references/value-equation.md +0 -134
  408. package/bundled-skills/onboarding/SKILL.md +0 -261
  409. package/bundled-skills/onboarding/evals/evals.json +0 -108
  410. package/bundled-skills/onboarding/references/activation-models.md +0 -57
  411. package/bundled-skills/onboarding/references/experiments.md +0 -258
  412. package/bundled-skills/onboarding/references/minimum-path-to-value.md +0 -49
  413. package/bundled-skills/openclaw-history-ingest/SKILL.md +0 -257
  414. package/bundled-skills/openclaw-history-ingest/references/openclaw-data-format.md +0 -154
  415. package/bundled-skills/orchestration/SKILL.md +0 -69
  416. package/bundled-skills/pi-history-ingest/SKILL.md +0 -312
  417. package/bundled-skills/pricing/SKILL.md +0 -295
  418. package/bundled-skills/pricing/evals/evals.json +0 -119
  419. package/bundled-skills/pricing/references/pricing-models.md +0 -66
  420. package/bundled-skills/pricing/references/pricing-page-teardown.md +0 -85
  421. package/bundled-skills/pricing/references/research-methods.md +0 -152
  422. package/bundled-skills/pricing/references/tier-structure.md +0 -232
  423. package/bundled-skills/product-marketing/SKILL.md +0 -255
  424. package/bundled-skills/product-marketing/evals/evals.json +0 -98
  425. package/bundled-skills/prospecting/SKILL.md +0 -262
  426. package/bundled-skills/prospecting/evals/evals.json +0 -122
  427. package/bundled-skills/prospecting/references/b2b-prospecting.md +0 -106
  428. package/bundled-skills/prospecting/references/compliance.md +0 -123
  429. package/bundled-skills/prospecting/references/data-sources.md +0 -287
  430. package/bundled-skills/prospecting/references/demand-signals.md +0 -129
  431. package/bundled-skills/prospecting/references/local-prospecting.md +0 -165
  432. package/bundled-skills/prospecting/references/saas-prospecting.md +0 -123
  433. package/bundled-skills/refactor-safely/SKILL.md +0 -29
  434. package/bundled-skills/review-delta/SKILL.md +0 -46
  435. package/bundled-skills/review-pr/SKILL.md +0 -66
  436. package/bundled-skills/sales-enablement/SKILL.md +0 -359
  437. package/bundled-skills/sales-enablement/evals/evals.json +0 -91
  438. package/bundled-skills/sales-enablement/references/deck-frameworks.md +0 -263
  439. package/bundled-skills/sales-enablement/references/demo-scripts.md +0 -355
  440. package/bundled-skills/sales-enablement/references/objection-library.md +0 -270
  441. package/bundled-skills/sales-enablement/references/one-pager-templates.md +0 -208
  442. package/bundled-skills/seo-audit/SKILL.md +0 -499
  443. package/bundled-skills/seo-audit/evals/evals.json +0 -136
  444. package/bundled-skills/seo-audit/references/ai-writing-detection.md +0 -200
  445. package/bundled-skills/seo-audit/references/international-seo.md +0 -230
  446. package/bundled-skills/session-brain/SKILL.md +0 -101
  447. package/bundled-skills/session-search/SKILL.md +0 -97
  448. package/bundled-skills/social/SKILL.md +0 -413
  449. package/bundled-skills/social/evals/evals.json +0 -107
  450. package/bundled-skills/social/references/carousel-frameworks.md +0 -141
  451. package/bundled-skills/social/references/listening-sources-template.md +0 -123
  452. package/bundled-skills/social/references/listening.md +0 -283
  453. package/bundled-skills/social/references/platform-limits.md +0 -110
  454. package/bundled-skills/social/references/platforms.md +0 -170
  455. package/bundled-skills/social/references/post-templates.md +0 -179
  456. package/bundled-skills/social/references/reverse-engineering.md +0 -195
  457. package/bundled-skills/social/references/short-form-video.md +0 -237
  458. package/bundled-skills/tag-taxonomy/SKILL.md +0 -218
  459. package/bundled-skills/vault-skill-factory/SKILL.md +0 -143
  460. package/bundled-skills/video/SKILL.md +0 -346
  461. package/bundled-skills/video/evals/evals.json +0 -103
  462. package/bundled-skills/video/references/ai-video-prompting.md +0 -175
  463. package/bundled-skills/video/references/edit-anatomy.md +0 -84
  464. package/bundled-skills/wiki-agent/SKILL.md +0 -322
  465. package/bundled-skills/wiki-capture/SKILL.md +0 -337
  466. package/bundled-skills/wiki-capture/references/RAW-FORMAT.md +0 -115
  467. package/bundled-skills/wiki-context-pack/SKILL.md +0 -89
  468. package/bundled-skills/wiki-dashboard/SKILL.md +0 -472
  469. package/bundled-skills/wiki-dedup/SKILL.md +0 -320
  470. package/bundled-skills/wiki-digest/SKILL.md +0 -243
  471. package/bundled-skills/wiki-export/SKILL.md +0 -445
  472. package/bundled-skills/wiki-history-ingest/SKILL.md +0 -61
  473. package/bundled-skills/wiki-import/SKILL.md +0 -276
  474. package/bundled-skills/wiki-ingest/SKILL.md +0 -563
  475. package/bundled-skills/wiki-ingest/references/ingest-prompts.md +0 -54
  476. package/bundled-skills/wiki-ingest/references/pageindex.md +0 -72
  477. package/bundled-skills/wiki-ingest/references/url-sources.md +0 -295
  478. package/bundled-skills/wiki-lint/SKILL.md +0 -643
  479. package/bundled-skills/wiki-narrate/SKILL.md +0 -120
  480. package/bundled-skills/wiki-narrate/references/voices.md +0 -29
  481. package/bundled-skills/wiki-query/SKILL.md +0 -307
  482. package/bundled-skills/wiki-rebuild/SKILL.md +0 -212
  483. package/bundled-skills/wiki-research/SKILL.md +0 -319
  484. package/bundled-skills/wiki-setup/SKILL.md +0 -359
  485. package/bundled-skills/wiki-stage-commit/SKILL.md +0 -164
  486. package/bundled-skills/wiki-status/SKILL.md +0 -556
  487. package/bundled-skills/wiki-switch/SKILL.md +0 -115
  488. package/bundled-skills/wiki-synthesize/SKILL.md +0 -212
  489. package/bundled-skills/wiki-update/SKILL.md +0 -271
  490. package/skills/ai-security/SKILL.md +0 -132
  491. package/skills/ai-security/references/atlas-coverage.md +0 -60
  492. package/skills/ai-security/scripts/ai_threat_scanner.py +0 -373
  493. package/skills/anti-AI-design/SKILL.md +0 -72
  494. package/skills/apple-hig-audit/SKILL.md +0 -64
  495. package/skills/apple-hig-audit/references/accessibility.md +0 -78
  496. package/skills/apple-hig-audit/references/platform-specifics.md +0 -57
  497. package/skills/apple-hig-audit/references/visual-design.md +0 -68
  498. package/skills/apple-hig-audit/scripts/hig_checker.py +0 -115
  499. package/skills/apple-hig-audit/templates/hig-audit-template.md +0 -114
  500. package/skills/dependency-auditor/README.md +0 -98
  501. package/skills/dependency-auditor/SKILL.md +0 -114
  502. package/skills/dependency-auditor/assets/sample_go.mod +0 -53
  503. package/skills/dependency-auditor/assets/sample_package.json +0 -72
  504. package/skills/dependency-auditor/assets/sample_requirements.txt +0 -71
  505. package/skills/dependency-auditor/expected_outputs/sample_license_report.txt +0 -37
  506. package/skills/dependency-auditor/expected_outputs/sample_upgrade_plan.txt +0 -59
  507. package/skills/dependency-auditor/expected_outputs/sample_vulnerability_report.json +0 -71
  508. package/skills/dependency-auditor/references/dependency_management_best_practices.md +0 -46
  509. package/skills/dependency-auditor/references/license_compatibility_matrix.md +0 -57
  510. package/skills/dependency-auditor/references/vulnerability_assessment_guide.md +0 -62
  511. package/skills/dependency-auditor/scripts/dep_scanner.py +0 -595
  512. package/skills/dependency-auditor/scripts/license_checker.py +0 -442
  513. package/skills/dependency-auditor/scripts/upgrade_planner.py +0 -490
  514. package/skills/dependency-auditor/test-inventory.json +0 -421
  515. package/skills/dependency-auditor/test-project/package.json +0 -72
  516. package/skills/design-system-tokens/SKILL.md +0 -90
  517. package/skills/design-system-tokens/assets/design_system_doc_template.md +0 -129
  518. package/skills/design-system-tokens/references/component-architecture.md +0 -88
  519. package/skills/design-system-tokens/references/developer-handoff.md +0 -61
  520. package/skills/design-system-tokens/references/responsive-calculations.md +0 -72
  521. package/skills/design-system-tokens/references/token-generation.md +0 -62
  522. package/skills/design-system-tokens/scripts/design_token_generator.py +0 -315
  523. package/skills/engineering-code-standards/SKILL.md +0 -174
  524. package/skills/env-secrets-manager/SKILL.md +0 -102
  525. package/skills/env-secrets-manager/references/secret-patterns.md +0 -56
  526. package/skills/env-secrets-manager/references/validation-detection-rotation.md +0 -76
  527. package/skills/env-secrets-manager/scripts/env_auditor.py +0 -340
  528. package/skills/frontend-design-taste/SKILL.md +0 -62
  529. package/skills/mcp-server-builder/README.md +0 -42
  530. package/skills/mcp-server-builder/SKILL.md +0 -144
  531. package/skills/mcp-server-builder/references/openapi-extraction-guide.md +0 -34
  532. package/skills/mcp-server-builder/references/production-hardening-guide.md +0 -80
  533. package/skills/mcp-server-builder/references/python-server-template.md +0 -22
  534. package/skills/mcp-server-builder/references/spec-compatibility.md +0 -129
  535. package/skills/mcp-server-builder/references/typescript-server-template.md +0 -19
  536. package/skills/mcp-server-builder/references/validation-checklist.md +0 -30
  537. package/skills/mcp-server-builder/scripts/mcp_validator.py +0 -180
  538. package/skills/mcp-server-builder/scripts/openapi_to_mcp.py +0 -279
  539. package/skills/playwright-agent/SKILL.md +0 -147
  540. package/skills/skill-creator/LICENSE.txt +0 -14
  541. package/skills/skill-creator/SKILL.md +0 -471
  542. package/skills/skill-creator/agents/analyzer.md +0 -274
  543. package/skills/skill-creator/agents/comparator.md +0 -202
  544. package/skills/skill-creator/agents/grader.md +0 -223
  545. package/skills/skill-creator/assets/eval_review.html +0 -143
  546. package/skills/skill-creator/eval-viewer/generate_review.py +0 -471
  547. package/skills/skill-creator/eval-viewer/viewer.html +0 -1321
  548. package/skills/skill-creator/references/schemas.md +0 -430
  549. package/skills/skill-creator/scripts/__init__.py +0 -0
  550. package/skills/skill-creator/scripts/aggregate_benchmark.py +0 -401
  551. package/skills/skill-creator/scripts/generate_report.py +0 -322
  552. package/skills/skill-creator/scripts/improve_description.py +0 -247
  553. package/skills/skill-creator/scripts/package_skill.py +0 -136
  554. package/skills/skill-creator/scripts/quick_validate.py +0 -103
  555. package/skills/skill-creator/scripts/run_eval.py +0 -310
  556. package/skills/skill-creator/scripts/run_loop.py +0 -328
  557. package/skills/skill-creator/scripts/utils.py +0 -47
  558. package/skills/write-a-skill/SKILL.md +0 -111
  559. package/skills/write-a-skill/references/companion_tooling.md +0 -61
  560. package/skills/write-a-skill/references/description_design_patterns.md +0 -136
  561. package/skills/write-a-skill/references/progressive_disclosure_principles.md +0 -87
  562. package/skills/write-a-skill/references/quality_gates_for_skills.md +0 -158
  563. package/skills/write-a-skill/scripts/skill_description_validator.py +0 -241
  564. package/skills/write-a-skill/scripts/skill_review_checklist_runner.py +0 -257
  565. package/skills/write-a-skill/scripts/skill_structure_validator.py +0 -265
@@ -1,328 +0,0 @@
1
- #!/usr/bin/env python3
2
- """Run the eval + improve loop until all pass or max iterations reached.
3
-
4
- Combines run_eval.py and improve_description.py in a loop, tracking history
5
- and returning the best description found. Supports train/test split to prevent
6
- overfitting.
7
- """
8
-
9
- import argparse
10
- import json
11
- import random
12
- import sys
13
- import tempfile
14
- import time
15
- import webbrowser
16
- from pathlib import Path
17
-
18
- from scripts.generate_report import generate_html
19
- from scripts.improve_description import improve_description
20
- from scripts.run_eval import find_project_root, run_eval
21
- from scripts.utils import parse_skill_md
22
-
23
-
24
- def split_eval_set(eval_set: list[dict], holdout: float, seed: int = 42) -> tuple[list[dict], list[dict]]:
25
- """Split eval set into train and test sets, stratified by should_trigger."""
26
- random.seed(seed)
27
-
28
- # Separate by should_trigger
29
- trigger = [e for e in eval_set if e["should_trigger"]]
30
- no_trigger = [e for e in eval_set if not e["should_trigger"]]
31
-
32
- # Shuffle each group
33
- random.shuffle(trigger)
34
- random.shuffle(no_trigger)
35
-
36
- # Calculate split points
37
- n_trigger_test = max(1, int(len(trigger) * holdout))
38
- n_no_trigger_test = max(1, int(len(no_trigger) * holdout))
39
-
40
- # Split
41
- test_set = trigger[:n_trigger_test] + no_trigger[:n_no_trigger_test]
42
- train_set = trigger[n_trigger_test:] + no_trigger[n_no_trigger_test:]
43
-
44
- return train_set, test_set
45
-
46
-
47
- def run_loop(
48
- eval_set: list[dict],
49
- skill_path: Path,
50
- description_override: str | None,
51
- num_workers: int,
52
- timeout: int,
53
- max_iterations: int,
54
- runs_per_query: int,
55
- trigger_threshold: float,
56
- holdout: float,
57
- model: str,
58
- verbose: bool,
59
- live_report_path: Path | None = None,
60
- log_dir: Path | None = None,
61
- ) -> dict:
62
- """Run the eval + improvement loop."""
63
- project_root = find_project_root()
64
- name, original_description, content = parse_skill_md(skill_path)
65
- current_description = description_override or original_description
66
-
67
- # Split into train/test if holdout > 0
68
- if holdout > 0:
69
- train_set, test_set = split_eval_set(eval_set, holdout)
70
- if verbose:
71
- print(f"Split: {len(train_set)} train, {len(test_set)} test (holdout={holdout})", file=sys.stderr)
72
- else:
73
- train_set = eval_set
74
- test_set = []
75
-
76
- history = []
77
- exit_reason = "unknown"
78
-
79
- for iteration in range(1, max_iterations + 1):
80
- if verbose:
81
- print(f"\n{'='*60}", file=sys.stderr)
82
- print(f"Iteration {iteration}/{max_iterations}", file=sys.stderr)
83
- print(f"Description: {current_description}", file=sys.stderr)
84
- print(f"{'='*60}", file=sys.stderr)
85
-
86
- # Evaluate train + test together in one batch for parallelism
87
- all_queries = train_set + test_set
88
- t0 = time.time()
89
- all_results = run_eval(
90
- eval_set=all_queries,
91
- skill_name=name,
92
- description=current_description,
93
- num_workers=num_workers,
94
- timeout=timeout,
95
- project_root=project_root,
96
- runs_per_query=runs_per_query,
97
- trigger_threshold=trigger_threshold,
98
- model=model,
99
- )
100
- eval_elapsed = time.time() - t0
101
-
102
- # Split results back into train/test by matching queries
103
- train_queries_set = {q["query"] for q in train_set}
104
- train_result_list = [r for r in all_results["results"] if r["query"] in train_queries_set]
105
- test_result_list = [r for r in all_results["results"] if r["query"] not in train_queries_set]
106
-
107
- train_passed = sum(1 for r in train_result_list if r["pass"])
108
- train_total = len(train_result_list)
109
- train_summary = {"passed": train_passed, "failed": train_total - train_passed, "total": train_total}
110
- train_results = {"results": train_result_list, "summary": train_summary}
111
-
112
- if test_set:
113
- test_passed = sum(1 for r in test_result_list if r["pass"])
114
- test_total = len(test_result_list)
115
- test_summary = {"passed": test_passed, "failed": test_total - test_passed, "total": test_total}
116
- test_results = {"results": test_result_list, "summary": test_summary}
117
- else:
118
- test_results = None
119
- test_summary = None
120
-
121
- history.append({
122
- "iteration": iteration,
123
- "description": current_description,
124
- "train_passed": train_summary["passed"],
125
- "train_failed": train_summary["failed"],
126
- "train_total": train_summary["total"],
127
- "train_results": train_results["results"],
128
- "test_passed": test_summary["passed"] if test_summary else None,
129
- "test_failed": test_summary["failed"] if test_summary else None,
130
- "test_total": test_summary["total"] if test_summary else None,
131
- "test_results": test_results["results"] if test_results else None,
132
- # For backward compat with report generator
133
- "passed": train_summary["passed"],
134
- "failed": train_summary["failed"],
135
- "total": train_summary["total"],
136
- "results": train_results["results"],
137
- })
138
-
139
- # Write live report if path provided
140
- if live_report_path:
141
- partial_output = {
142
- "original_description": original_description,
143
- "best_description": current_description,
144
- "best_score": "in progress",
145
- "iterations_run": len(history),
146
- "holdout": holdout,
147
- "train_size": len(train_set),
148
- "test_size": len(test_set),
149
- "history": history,
150
- }
151
- live_report_path.write_text(generate_html(partial_output, auto_refresh=True, skill_name=name))
152
-
153
- if verbose:
154
- def print_eval_stats(label, results, elapsed):
155
- pos = [r for r in results if r["should_trigger"]]
156
- neg = [r for r in results if not r["should_trigger"]]
157
- tp = sum(r["triggers"] for r in pos)
158
- pos_runs = sum(r["runs"] for r in pos)
159
- fn = pos_runs - tp
160
- fp = sum(r["triggers"] for r in neg)
161
- neg_runs = sum(r["runs"] for r in neg)
162
- tn = neg_runs - fp
163
- total = tp + tn + fp + fn
164
- precision = tp / (tp + fp) if (tp + fp) > 0 else 1.0
165
- recall = tp / (tp + fn) if (tp + fn) > 0 else 1.0
166
- accuracy = (tp + tn) / total if total > 0 else 0.0
167
- print(f"{label}: {tp+tn}/{total} correct, precision={precision:.0%} recall={recall:.0%} accuracy={accuracy:.0%} ({elapsed:.1f}s)", file=sys.stderr)
168
- for r in results:
169
- status = "PASS" if r["pass"] else "FAIL"
170
- rate_str = f"{r['triggers']}/{r['runs']}"
171
- print(f" [{status}] rate={rate_str} expected={r['should_trigger']}: {r['query'][:60]}", file=sys.stderr)
172
-
173
- print_eval_stats("Train", train_results["results"], eval_elapsed)
174
- if test_summary:
175
- print_eval_stats("Test ", test_results["results"], 0)
176
-
177
- if train_summary["failed"] == 0:
178
- exit_reason = f"all_passed (iteration {iteration})"
179
- if verbose:
180
- print(f"\nAll train queries passed on iteration {iteration}!", file=sys.stderr)
181
- break
182
-
183
- if iteration == max_iterations:
184
- exit_reason = f"max_iterations ({max_iterations})"
185
- if verbose:
186
- print(f"\nMax iterations reached ({max_iterations}).", file=sys.stderr)
187
- break
188
-
189
- # Improve the description based on train results
190
- if verbose:
191
- print(f"\nImproving description...", file=sys.stderr)
192
-
193
- t0 = time.time()
194
- # Strip test scores from history so improvement model can't see them
195
- blinded_history = [
196
- {k: v for k, v in h.items() if not k.startswith("test_")}
197
- for h in history
198
- ]
199
- new_description = improve_description(
200
- skill_name=name,
201
- skill_content=content,
202
- current_description=current_description,
203
- eval_results=train_results,
204
- history=blinded_history,
205
- model=model,
206
- log_dir=log_dir,
207
- iteration=iteration,
208
- )
209
- improve_elapsed = time.time() - t0
210
-
211
- if verbose:
212
- print(f"Proposed ({improve_elapsed:.1f}s): {new_description}", file=sys.stderr)
213
-
214
- current_description = new_description
215
-
216
- # Find the best iteration by TEST score (or train if no test set)
217
- if test_set:
218
- best = max(history, key=lambda h: h["test_passed"] or 0)
219
- best_score = f"{best['test_passed']}/{best['test_total']}"
220
- else:
221
- best = max(history, key=lambda h: h["train_passed"])
222
- best_score = f"{best['train_passed']}/{best['train_total']}"
223
-
224
- if verbose:
225
- print(f"\nExit reason: {exit_reason}", file=sys.stderr)
226
- print(f"Best score: {best_score} (iteration {best['iteration']})", file=sys.stderr)
227
-
228
- return {
229
- "exit_reason": exit_reason,
230
- "original_description": original_description,
231
- "best_description": best["description"],
232
- "best_score": best_score,
233
- "best_train_score": f"{best['train_passed']}/{best['train_total']}",
234
- "best_test_score": f"{best['test_passed']}/{best['test_total']}" if test_set else None,
235
- "final_description": current_description,
236
- "iterations_run": len(history),
237
- "holdout": holdout,
238
- "train_size": len(train_set),
239
- "test_size": len(test_set),
240
- "history": history,
241
- }
242
-
243
-
244
- def main():
245
- parser = argparse.ArgumentParser(description="Run eval + improve loop")
246
- parser.add_argument("--eval-set", required=True, help="Path to eval set JSON file")
247
- parser.add_argument("--skill-path", required=True, help="Path to skill directory")
248
- parser.add_argument("--description", default=None, help="Override starting description")
249
- parser.add_argument("--num-workers", type=int, default=10, help="Number of parallel workers")
250
- parser.add_argument("--timeout", type=int, default=30, help="Timeout per query in seconds")
251
- parser.add_argument("--max-iterations", type=int, default=5, help="Max improvement iterations")
252
- parser.add_argument("--runs-per-query", type=int, default=3, help="Number of runs per query")
253
- parser.add_argument("--trigger-threshold", type=float, default=0.5, help="Trigger rate threshold")
254
- parser.add_argument("--holdout", type=float, default=0.4, help="Fraction of eval set to hold out for testing (0 to disable)")
255
- parser.add_argument("--model", required=True, help="Model for improvement")
256
- parser.add_argument("--verbose", action="store_true", help="Print progress to stderr")
257
- parser.add_argument("--report", default="auto", help="Generate HTML report at this path (default: 'auto' for temp file, 'none' to disable)")
258
- parser.add_argument("--results-dir", default=None, help="Save all outputs (results.json, report.html, log.txt) to a timestamped subdirectory here")
259
- args = parser.parse_args()
260
-
261
- eval_set = json.loads(Path(args.eval_set).read_text())
262
- skill_path = Path(args.skill_path)
263
-
264
- if not (skill_path / "SKILL.md").exists():
265
- print(f"Error: No SKILL.md found at {skill_path}", file=sys.stderr)
266
- sys.exit(1)
267
-
268
- name, _, _ = parse_skill_md(skill_path)
269
-
270
- # Set up live report path
271
- if args.report != "none":
272
- if args.report == "auto":
273
- timestamp = time.strftime("%Y%m%d_%H%M%S")
274
- live_report_path = Path(tempfile.gettempdir()) / f"skill_description_report_{skill_path.name}_{timestamp}.html"
275
- else:
276
- live_report_path = Path(args.report)
277
- # Open the report immediately so the user can watch
278
- live_report_path.write_text("<html><body><h1>Starting optimization loop...</h1><meta http-equiv='refresh' content='5'></body></html>")
279
- webbrowser.open(str(live_report_path))
280
- else:
281
- live_report_path = None
282
-
283
- # Determine output directory (create before run_loop so logs can be written)
284
- if args.results_dir:
285
- timestamp = time.strftime("%Y-%m-%d_%H%M%S")
286
- results_dir = Path(args.results_dir) / timestamp
287
- results_dir.mkdir(parents=True, exist_ok=True)
288
- else:
289
- results_dir = None
290
-
291
- log_dir = results_dir / "logs" if results_dir else None
292
-
293
- output = run_loop(
294
- eval_set=eval_set,
295
- skill_path=skill_path,
296
- description_override=args.description,
297
- num_workers=args.num_workers,
298
- timeout=args.timeout,
299
- max_iterations=args.max_iterations,
300
- runs_per_query=args.runs_per_query,
301
- trigger_threshold=args.trigger_threshold,
302
- holdout=args.holdout,
303
- model=args.model,
304
- verbose=args.verbose,
305
- live_report_path=live_report_path,
306
- log_dir=log_dir,
307
- )
308
-
309
- # Save JSON output
310
- json_output = json.dumps(output, indent=2)
311
- print(json_output)
312
- if results_dir:
313
- (results_dir / "results.json").write_text(json_output)
314
-
315
- # Write final HTML report (without auto-refresh)
316
- if live_report_path:
317
- live_report_path.write_text(generate_html(output, auto_refresh=False, skill_name=name))
318
- print(f"\nReport: {live_report_path}", file=sys.stderr)
319
-
320
- if results_dir and live_report_path:
321
- (results_dir / "report.html").write_text(generate_html(output, auto_refresh=False, skill_name=name))
322
-
323
- if results_dir:
324
- print(f"Results saved to: {results_dir}", file=sys.stderr)
325
-
326
-
327
- if __name__ == "__main__":
328
- main()
@@ -1,47 +0,0 @@
1
- """Shared utilities for skill-creator scripts."""
2
-
3
- from pathlib import Path
4
-
5
-
6
-
7
- def parse_skill_md(skill_path: Path) -> tuple[str, str, str]:
8
- """Parse a SKILL.md file, returning (name, description, full_content)."""
9
- content = (skill_path / "SKILL.md").read_text()
10
- lines = content.split("\n")
11
-
12
- if lines[0].strip() != "---":
13
- raise ValueError("SKILL.md missing frontmatter (no opening ---)")
14
-
15
- end_idx = None
16
- for i, line in enumerate(lines[1:], start=1):
17
- if line.strip() == "---":
18
- end_idx = i
19
- break
20
-
21
- if end_idx is None:
22
- raise ValueError("SKILL.md missing frontmatter (no closing ---)")
23
-
24
- name = ""
25
- description = ""
26
- frontmatter_lines = lines[1:end_idx]
27
- i = 0
28
- while i < len(frontmatter_lines):
29
- line = frontmatter_lines[i]
30
- if line.startswith("name:"):
31
- name = line[len("name:"):].strip().strip('"').strip("'")
32
- elif line.startswith("description:"):
33
- value = line[len("description:"):].strip()
34
- # Handle YAML multiline indicators (>, |, >-, |-)
35
- if value in (">", "|", ">-", "|-"):
36
- continuation_lines: list[str] = []
37
- i += 1
38
- while i < len(frontmatter_lines) and (frontmatter_lines[i].startswith(" ") or frontmatter_lines[i].startswith("\t")):
39
- continuation_lines.append(frontmatter_lines[i].strip())
40
- i += 1
41
- description = " ".join(continuation_lines)
42
- continue
43
- else:
44
- description = value.strip('"').strip("'")
45
- i += 1
46
-
47
- return name, description, content
@@ -1,111 +0,0 @@
1
- ---
2
- name: write-a-skill
3
- description: Create new agent skills with proper structure, progressive disclosure, and bundled resources. Use when user wants to create, write, build, or author a new skill.
4
- license: Apache-2.0
5
- metadata:
6
- author: Novahiz
7
- organization: Novahiz
8
- version: "2.0.0"
9
- date: September 2026
10
- ---
11
-
12
- # Writing skills
13
-
14
- ## Process
15
-
16
- 1. **Clarify the need.** Ask about the domain, the concrete cases the skill must handle, whether scripts are required, and what reference material should ship with it.
17
- 2. **Draft the files.** Write SKILL.md first. Push anything past 100 lines into a reference file. Add scripts only for deterministic work the model should not re-derive each time.
18
- 3. **Review with the author.** Present the draft. Ask what is missing, what is unclear, and which sections are too heavy or too thin.
19
-
20
- ## Layout
21
-
22
- ```
23
- skill-name/
24
- ├── SKILL.md # required entry point
25
- ├── REFERENCE.md # optional long-form detail
26
- ├── EXAMPLES.md # optional worked cases
27
- └── scripts/ # optional helpers
28
- └── helper.js
29
- ```
30
-
31
- ## SKILL.md skeleton
32
-
33
- ```md
34
- ---
35
- name: skill-name
36
- description: What the skill does. Use when [specific triggers].
37
- ---
38
-
39
- # Skill name
40
-
41
- ## Quick start
42
-
43
- [Minimal working example]
44
-
45
- ## Workflows
46
-
47
- [Step-by-step processes for complex tasks]
48
-
49
- ## Advanced features
50
-
51
- [See REFERENCE.md](REFERENCE.md)
52
- ```
53
-
54
- ## Description rules
55
-
56
- The description is the only text the agent sees when it decides whether to load the skill. It sits in the system prompt next to every other installed skill, and the agent picks from it alone.
57
-
58
- Two facts must land:
59
-
60
- 1. What the skill does.
61
- 2. When to fire: keywords, contexts, file types.
62
-
63
- Constraints:
64
-
65
- - Hard cap at 1024 characters.
66
- - Third person throughout.
67
- - Sentence one: the action.
68
- - Sentence two: `Use when [specific triggers]`.
69
-
70
- Good:
71
-
72
- ```
73
- Extract text and tables from PDF files, fill forms, merge documents. Use when working with PDF files or when user mentions PDFs, forms, or document extraction.
74
- ```
75
-
76
- Bad:
77
-
78
- ```
79
- Helps with documents.
80
- ```
81
-
82
- The weak version gives the agent nothing to match against. Every document skill in the list looks the same from it.
83
-
84
- ## When scripts belong in the skill
85
-
86
- Ship a script when the work is deterministic (validation, formatting), when the same code would otherwise be generated over and over, or when failure needs explicit handling. Scripts cut token use and remove the chance of two runs inventing two different implementations.
87
-
88
- ## When to split the file
89
-
90
- Split when SKILL.md crosses 100 lines, when the content covers distinct domains (finance schemas next to sales schemas), or when advanced paths are rarely needed. The agent should read SKILL.md in full and stop there unless a pointer sends it deeper.
91
-
92
- ## Review checklist
93
-
94
- - [ ] Description carries triggers (`Use when ...`)
95
- - [ ] SKILL.md stays under 100 lines
96
- - [ ] No dates, versions, or `as of YYYY` claims
97
- - [ ] One word per concept throughout
98
- - [ ] At least one concrete example
99
- - [ ] References sit one level deep
100
-
101
- ## Tooling
102
-
103
- Three stdlib validators sit next to this skill:
104
-
105
- ```
106
- python scripts/skill_description_validator.py path/to/SKILL.md
107
- python scripts/skill_structure_validator.py path/to/skill-folder
108
- python scripts/skill_review_checklist_runner.py path/to/skill-folder
109
- ```
110
-
111
- Catalogue and companions: [references/companion_tooling.md](references/companion_tooling.md).
@@ -1,61 +0,0 @@
1
- # Companion Tooling
2
-
3
- Validation tools layered on top of the core skill-authoring process. Use these when authoring a new skill in this repo.
4
-
5
- ## Validation Tools (stdlib Python)
6
-
7
- | Tool | Purpose | Run before |
8
- |---|---|---|
9
- | `scripts/skill_description_validator.py` | Validates description: ≤1024 chars, third person, "Use when" trigger, action verb in first sentence | First draft of SKILL.md |
10
- | `scripts/skill_structure_validator.py` | Validates folder structure: SKILL.md present, ≤100 lines, references one level deep, no circular refs | Pre-commit |
11
- | `scripts/skill_review_checklist_runner.py` | Runs all 6 review-checklist items against a skill folder | Final check before PR |
12
-
13
- All three tools:
14
- - Stdlib-only (no external dependencies)
15
- - Run with embedded sample if no path provided
16
- - Output text or JSON (`--output json`)
17
- - Exit code: 0 if PASS, 1 if FAIL/WARN
18
-
19
- ## cs-skill-author Persona Agent
20
-
21
- Lives at `../agents/cs-skill-author.md`. Voice: forcing-question interrogator. Surfaces the skill-authoring workflow as an interrogation before any new skill commit.
22
-
23
- **Opening question:** "What capability does this skill provide, and what's the trigger phrase that distinguishes it from existing skills?"
24
-
25
- **Six forcing questions** (matches the review checklist):
26
- 1. What's the description? Is it ≤1024 chars + third person + has "Use when ..."?
27
- 2. Is SKILL.md under 100 lines? If not, where will the split land (REFERENCE.md / EXAMPLES.md / references/)?
28
- 3. Are there time-sensitive claims (dates, "as of YYYY")?
29
- 4. Is terminology consistent: same word for the same concept throughout?
30
- 5. Concrete examples: at least 1 code block, ideally good/bad contrast?
31
- 6. References one level deep, no circular refs?
32
-
33
- ## `/cs:write-a-skill` Slash Command
34
-
35
- Lives at `../commands/cs-write-a-skill.md`. Three-step flow:
36
-
37
- 1. Run `cs-skill-author` interrogation (6 questions)
38
- 2. Draft skill files per the structure pattern
39
- 3. Run all 3 validation tools; show verdict; fix until PASS
40
-
41
- Use when: starting a new skill in this repo from scratch.
42
-
43
- ## Why Wrap the Core Process
44
-
45
- The core write-a-skill process is a tight, principled skill: perfect as-is for individual authoring sessions. The wrapper layers add three things this repo benefits from at scale:
46
-
47
- 1. **Programmatic enforcement** of the review checklist (the validation tools): prevents human review-checklist drift across 100+ skills.
48
- 2. **Forcing-question interrogation** (the cs-skill-author persona): adapts the "review with user" phase to the cs-* persona pattern used elsewhere in this repo.
49
- 3. **Citation-backed references**: the core process links to principles; the wrapper adds 5+ authoritative external sources per reference for newcomers learning the pattern.
50
-
51
- ---
52
-
53
- **Source authorities (non-exhaustive):**
54
-
55
- - **Novahiz skill-authoring standard** (this repo): the core process
56
- - **Skills documentation**: official guidance on skill structure
57
- - **Engineering blog: Skills patterns** (continuously updated): patterns for skill authoring
58
- - **Karpathy, A.: "Software 3.0" + LLM coding pitfalls** (X.com posts 2024-2025): discipline reference applied throughout this repo's karpathy-coder skill
59
- - **Pareto principle applied to documentation**: concise = trustworthy; 80% of value in 20% of words
60
- - **Hyrum's Law** as applied to skill descriptions: once a description shape is observed, downstream agents depend on it
61
- - **Conway's Law as applied to skill libraries**: skill organization mirrors team responsibilities; progressive disclosure mirrors information needs across team boundaries
@@ -1,136 +0,0 @@
1
- # Description Design Patterns for Skills
2
-
3
- This reference answers exactly one decision: how do we write a skill description that an agent actually picks correctly when faced with a long skill list?
4
-
5
- Pair with `scripts/skill_description_validator.py` for automated enforcement.
6
-
7
- ## The Core Rule
8
-
9
- The description is the only thing your agent sees when deciding which skill to load.
10
-
11
- Implication: the description is not marketing copy. It is a routing signal for the agent. Every word competes with every other skill's description for activation attention.
12
-
13
- ## The Four Format Rules
14
-
15
- 1. **Max 1024 chars**: beyond this, agents lose the early sentences when condensing context
16
- 2. **Third person**: first-person ("I help with...") confuses agent self-identification; second-person ("You can...") confuses pronoun reference
17
- 3. **First sentence: what it does**: front-load the verb + object
18
- 4. **Second sentence: "Use when [specific triggers]"**: agent's most reliable activation cue
19
-
20
- ## Good vs Bad Examples
21
-
22
- **Good:**
23
-
24
- ```
25
- Extract text and tables from PDF files, fill forms, merge documents. Use when working with PDF files or when user mentions PDFs, forms, or document extraction.
26
- ```
27
-
28
- **Why good:**
29
- - Front-loaded verbs: Extract, fill, merge
30
- - Concrete objects: text, tables, PDF files, forms
31
- - Explicit trigger: "Use when working with PDF files"
32
- - Specific keywords for matching: "PDFs", "forms", "document extraction"
33
-
34
- **Bad:**
35
-
36
- ```
37
- Helps with documents.
38
- ```
39
-
40
- **Why bad:**
41
- - "Helps" is content-free
42
- - "Documents" is generic: every doc skill has this
43
- - No trigger
44
- - No keyword variety
45
-
46
- **Bad in a different way (over-specified):**
47
-
48
- ```
49
- This skill performs comprehensive PDF document processing including but not limited to extraction, manipulation, format conversion, content analysis, metadata management, and security operations on PDF files, with support for various PDF versions and embedded media types.
50
- ```
51
-
52
- **Why bad:** verbose, no triggers, agent can't extract the key keywords from the wall of text.
53
-
54
- ## The Trigger Sentence Pattern
55
-
56
- The "Use when" sentence is the highest-impact part of the description. Patterns that work:
57
-
58
- **Keyword triggers** (when user types specific words):
59
- ```
60
- Use when user mentions PDFs, forms, or document extraction.
61
- ```
62
-
63
- **File-type triggers** (when agent sees specific files):
64
- ```
65
- Use when working with `.tsx` files or React component tests.
66
- ```
67
-
68
- **Context triggers** (when agent is in a specific state):
69
- ```
70
- Use when the user requests a code review of a pull request.
71
- ```
72
-
73
- **Workflow triggers** (when agent is mid-workflow):
74
- ```
75
- Use after running tests and before committing changes.
76
- ```
77
-
78
- ## Vocabulary Selection
79
-
80
- The description's words must overlap with words users + agents naturally use for the task.
81
-
82
- | Bad keyword | Better keyword | Why |
83
- |---|---|---|
84
- | "documents" | "PDF files" / "Word docs" | More specific = less collision |
85
- | "improve" | "refactor" / "fix" / "optimize" | Specific verb = clearer routing |
86
- | "various" | (delete; just list them) | Hedge language = no info |
87
- | "modern" | (cite the actual tool/version) | Trend words age badly |
88
- | "comprehensive" | (delete; just list capabilities) | Adjective inflation |
89
-
90
- ## Length Optimization
91
-
92
- Below 1024 chars, shorter is usually better. Target: 100-300 chars for most skills.
93
-
94
- Where complexity demands more chars, prioritize:
95
- 1. The verb-object pair (what it does): never compress
96
- 2. The trigger phrase: never compress
97
- 3. Keyword variety (different ways users describe it): expand here if space allows
98
- 4. Anti-keyword (what it does NOT do): only if there's a frequently-confused sibling skill
99
-
100
- ## Anti-Patterns to Avoid
101
-
102
- 1. **First-person voice**: "I extract PDFs": confuses agent self-reference
103
- 2. **Marketing language**: "fast, powerful, intuitive": agent doesn't care, ignores adjectives
104
- 3. **Trigger-less descriptions**: every skill needs "Use when X"
105
- 4. **Multi-purpose dumping**: if your skill does 10 unrelated things, it's probably 10 skills
106
- 5. **Pronouns and hedges**: "you can also use this if you want to": drop entirely
107
- 6. **Recursive descriptions**: "Use this skill when you need this skill": adds nothing
108
- 7. **Implementation details**: "Built on Python + stdlib": agent doesn't care; matters for README, not description
109
-
110
- ## Pre-Commit Discipline
111
-
112
- Run before every skill PR:
113
-
114
- ```bash
115
- python scripts/skill_description_validator.py path/to/SKILL.md
116
- ```
117
-
118
- If validator returns FAIL, fix before merging. If WARN, justify and document the trade-off.
119
-
120
- ## When This Reference Doesn't Help
121
-
122
- - **Naming the skill itself**: different concern; see naming-conventions guidance per-repo
123
- - **Skill discovery in marketplaces**: different audience (humans browsing), different rules
124
- - **System-prompt design for the agent that loads skills**: separate concern
125
-
126
- ---
127
-
128
- **Source authorities (non-exhaustive):**
129
-
130
- - **Novahiz authoring standard** (this repo): the 4 format rules + good/bad example pattern
131
- - **Skill registry documentation**: how a skill-loader uses descriptions for routing
132
- - **Effective system prompt writing guides**: same principles applied to system-prompt design
133
- - **Karpathy, A.: public commentary on LLM prompt design**: emphasis on specificity + lack of ambiguity
134
- - **Garrett, J.J.: "The Elements of User Experience"** (2002) + information architecture principles: labels must match user mental models
135
- - **Nielsen Norman Group: Microcontent guidelines**: applies to skill descriptions: front-load value, hard-cap length, scannable structure
136
- - **Search-engine + SEO patterns adapted for agent routing**: keyword density, intent matching, semantic field coverage