dsh-aris-panel 0.1.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (442) hide show
  1. package/LICENSE +21 -0
  2. package/README.md +98 -0
  3. package/README_CN.md +87 -0
  4. package/dsh/checkout.patch.yml +38 -0
  5. package/dsh/client.js +634 -0
  6. package/dsh/cordis.patch.yml +44 -0
  7. package/dsh/index.mjs +76 -0
  8. package/dsh/run-status.mjs +182 -0
  9. package/dsh/scope-limits.mjs +50 -0
  10. package/dsh/workbench.mjs +291 -0
  11. package/mcp-servers/claude-review/README.md +93 -0
  12. package/mcp-servers/claude-review/run_with_claude_aws.sh +49 -0
  13. package/mcp-servers/claude-review/server.py +718 -0
  14. package/mcp-servers/codex-image2/README.md +65 -0
  15. package/mcp-servers/codex-image2/server.py +893 -0
  16. package/mcp-servers/feishu-bridge/requirements.txt +1 -0
  17. package/mcp-servers/feishu-bridge/server.py +240 -0
  18. package/mcp-servers/gemini-review/README.md +171 -0
  19. package/mcp-servers/gemini-review/server.py +1856 -0
  20. package/mcp-servers/llm-chat/requirements.txt +1 -0
  21. package/mcp-servers/llm-chat/server.py +664 -0
  22. package/mcp-servers/manual-review/README.md +133 -0
  23. package/mcp-servers/manual-review/server.py +910 -0
  24. package/mcp-servers/manual-review/ui.html +279 -0
  25. package/mcp-servers/minimax-chat/requirements.txt +1 -0
  26. package/mcp-servers/minimax-chat/server.py +381 -0
  27. package/package.json +51 -0
  28. package/skills/ablation-planner/SKILL.md +123 -0
  29. package/skills/alphaxiv/SKILL.md +196 -0
  30. package/skills/analyze-results/SKILL.md +46 -0
  31. package/skills/arxiv/SKILL.md +248 -0
  32. package/skills/auto-paper-improvement-loop/SKILL.md +651 -0
  33. package/skills/auto-review-loop/SKILL.md +1137 -0
  34. package/skills/auto-review-loop-llm/SKILL.md +259 -0
  35. package/skills/auto-review-loop-minimax/SKILL.md +302 -0
  36. package/skills/citation-audit/SKILL.md +502 -0
  37. package/skills/claims-drafting/SKILL.md +227 -0
  38. package/skills/comm-lit-review/SKILL.md +297 -0
  39. package/skills/deepxiv/SKILL.md +263 -0
  40. package/skills/dse-loop/SKILL.md +296 -0
  41. package/skills/embodiment-description/SKILL.md +129 -0
  42. package/skills/exa-search/SKILL.md +205 -0
  43. package/skills/experiment-audit/SKILL.md +311 -0
  44. package/skills/experiment-bridge/SKILL.md +376 -0
  45. package/skills/experiment-plan/SKILL.md +249 -0
  46. package/skills/experiment-queue/SKILL.md +431 -0
  47. package/skills/experiment-queue/scripts/build_manifest.py +142 -0
  48. package/skills/experiment-queue/scripts/queue_manager.py +433 -0
  49. package/skills/feishu-notify/SKILL.md +156 -0
  50. package/skills/figure-description/SKILL.md +138 -0
  51. package/skills/figure-spec/SKILL.md +262 -0
  52. package/skills/figure-spec/scripts/figure_renderer.py +799 -0
  53. package/skills/formula-derivation/SKILL.md +280 -0
  54. package/skills/gemini-search/SKILL.md +231 -0
  55. package/skills/grant-proposal/SKILL.md +698 -0
  56. package/skills/idea-creator/SKILL.md +542 -0
  57. package/skills/idea-discovery/SKILL.md +521 -0
  58. package/skills/idea-discovery-robot/SKILL.md +363 -0
  59. package/skills/integrity-forensics/SKILL.md +284 -0
  60. package/skills/interview-cheatsheet/SKILL.md +245 -0
  61. package/skills/invention-structuring/SKILL.md +188 -0
  62. package/skills/jurisdiction-format/SKILL.md +192 -0
  63. package/skills/kill-argument/SKILL.md +437 -0
  64. package/skills/mermaid-diagram/SKILL.md +419 -0
  65. package/skills/meta-apply/SKILL.md +141 -0
  66. package/skills/meta-optimize/SKILL.md +437 -0
  67. package/skills/monitor-experiment/SKILL.md +140 -0
  68. package/skills/novelty-check/SKILL.md +101 -0
  69. package/skills/openalex/SKILL.md +237 -0
  70. package/skills/overleaf-sync/SKILL.md +220 -0
  71. package/skills/paper-claim-audit/SKILL.md +348 -0
  72. package/skills/paper-compile/SKILL.md +266 -0
  73. package/skills/paper-figure/SKILL.md +312 -0
  74. package/skills/paper-illustration/SKILL.md +736 -0
  75. package/skills/paper-illustration-image2/SKILL.md +391 -0
  76. package/skills/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  77. package/skills/paper-plan/SKILL.md +386 -0
  78. package/skills/paper-poster/SKILL.md +19 -0
  79. package/skills/paper-poster-html/DESIGN_FINAL.md +176 -0
  80. package/skills/paper-poster-html/IMPLEMENTATION_CONVENTIONS.md +161 -0
  81. package/skills/paper-poster-html/LICENSES/posterly-MIT.txt +21 -0
  82. package/skills/paper-poster-html/NOTICE.md +57 -0
  83. package/skills/paper-poster-html/SKILL.md +323 -0
  84. package/skills/paper-poster-html/scripts/_posterly/__init__.py +0 -0
  85. package/skills/paper-poster-html/scripts/_posterly/canvas.py +200 -0
  86. package/skills/paper-poster-html/scripts/_posterly/measure.py +588 -0
  87. package/skills/paper-poster-html/scripts/_posterly/polish.py +498 -0
  88. package/skills/paper-poster-html/scripts/_posterly/preflight.py +489 -0
  89. package/skills/paper-poster-html/scripts/_posterly/render.py +215 -0
  90. package/skills/paper-poster-html/scripts/_posterly/textutil.py +16 -0
  91. package/skills/paper-poster-html/scripts/_posterly/verify_final.py +171 -0
  92. package/skills/paper-poster-html/scripts/asset_check.py +897 -0
  93. package/skills/paper-poster-html/scripts/extract_pdf_figures.py +666 -0
  94. package/skills/paper-poster-html/scripts/poster_check.py +251 -0
  95. package/skills/paper-poster-html/scripts/preprocess_figures.py +238 -0
  96. package/skills/paper-poster-html/scripts/render_preview.py +217 -0
  97. package/skills/paper-poster-html/scripts/run_gates.py +556 -0
  98. package/skills/paper-poster-html/scripts/style_check.py +1324 -0
  99. package/skills/paper-poster-html/templates/COMPONENTS.md +462 -0
  100. package/skills/paper-poster-html/templates/README.md +170 -0
  101. package/skills/paper-poster-html/templates/landscape_4col.html +1032 -0
  102. package/skills/paper-poster-html/templates/landscape_hero.html +1046 -0
  103. package/skills/paper-poster-html/templates/portrait_2col.html +947 -0
  104. package/skills/paper-poster-html/templates/tokens/acl.json +9 -0
  105. package/skills/paper-poster-html/templates/tokens/cvpr.json +9 -0
  106. package/skills/paper-poster-html/templates/tokens/generic.json +9 -0
  107. package/skills/paper-poster-html/templates/tokens/iclr.json +9 -0
  108. package/skills/paper-poster-html/templates/tokens/icml.json +9 -0
  109. package/skills/paper-poster-html/templates/tokens/neurips.json +9 -0
  110. package/skills/paper-slides/SKILL.md +635 -0
  111. package/skills/paper-talk/SKILL.md +381 -0
  112. package/skills/paper-write/SKILL.md +604 -0
  113. package/skills/paper-write/templates/IEEEtran.bst +2409 -0
  114. package/skills/paper-write/templates/IEEEtran.cls +6347 -0
  115. package/skills/paper-write/templates/iclr2026.tex +84 -0
  116. package/skills/paper-write/templates/icml2025.tex +87 -0
  117. package/skills/paper-write/templates/ieee_conference.tex +89 -0
  118. package/skills/paper-write/templates/ieee_journal.tex +93 -0
  119. package/skills/paper-write/templates/math_commands.tex +48 -0
  120. package/skills/paper-write/templates/neurips2025.tex +80 -0
  121. package/skills/paper-writing/SKILL.md +916 -0
  122. package/skills/patent-novelty-check/SKILL.md +153 -0
  123. package/skills/patent-pipeline/SKILL.md +344 -0
  124. package/skills/patent-review/SKILL.md +203 -0
  125. package/skills/pixel-art/SKILL.md +137 -0
  126. package/skills/prior-art-search/SKILL.md +146 -0
  127. package/skills/proof-checker/SKILL.md +866 -0
  128. package/skills/proof-orchestrator/NOTICE.md +24 -0
  129. package/skills/proof-orchestrator/SKILL.md +254 -0
  130. package/skills/proof-orchestrator/references/audit-output-contract.md +126 -0
  131. package/skills/proof-orchestrator/references/deepseek-routing.md +74 -0
  132. package/skills/proof-orchestrator/references/dispatch-prompts.md +227 -0
  133. package/skills/proof-orchestrator/references/notation-audit.md +135 -0
  134. package/skills/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  135. package/skills/proof-orchestrator/references/stress-tests.md +38 -0
  136. package/skills/proof-writer/SKILL.md +223 -0
  137. package/skills/qzcli/SKILL.md +324 -0
  138. package/skills/rebuttal/SKILL.md +376 -0
  139. package/skills/render-html/SKILL.md +316 -0
  140. package/skills/render-html/scripts/render_html.py +1006 -0
  141. package/skills/render-html/scripts/templates/academic.html +703 -0
  142. package/skills/render-html/scripts/templates/dashboard.html +333 -0
  143. package/skills/research-lit/SKILL.md +756 -0
  144. package/skills/research-pipeline/SKILL.md +384 -0
  145. package/skills/research-refine/SKILL.md +770 -0
  146. package/skills/research-refine-pipeline/SKILL.md +186 -0
  147. package/skills/research-review/SKILL.md +198 -0
  148. package/skills/research-wiki/SKILL.md +461 -0
  149. package/skills/resubmit-pipeline/SKILL.md +447 -0
  150. package/skills/result-to-claim/SKILL.md +311 -0
  151. package/skills/run-experiment/SKILL.md +313 -0
  152. package/skills/semantic-scholar/SKILL.md +236 -0
  153. package/skills/serverless-modal/SKILL.md +335 -0
  154. package/skills/shared-references/acceptance-gate.md +324 -0
  155. package/skills/shared-references/assurance-contract.md +248 -0
  156. package/skills/shared-references/capture-antipatterns.md +78 -0
  157. package/skills/shared-references/citation-discipline.md +583 -0
  158. package/skills/shared-references/compute-env-contract.md +163 -0
  159. package/skills/shared-references/effort-contract.md +183 -0
  160. package/skills/shared-references/evidence-precheck.md +65 -0
  161. package/skills/shared-references/experiment-integrity.md +49 -0
  162. package/skills/shared-references/external-cadence.md +326 -0
  163. package/skills/shared-references/fan-out-pattern.md +366 -0
  164. package/skills/shared-references/injection-hygiene.md +127 -0
  165. package/skills/shared-references/integration-contract.md +461 -0
  166. package/skills/shared-references/output-composition.md +93 -0
  167. package/skills/shared-references/output-language.md +45 -0
  168. package/skills/shared-references/output-manifest.md +49 -0
  169. package/skills/shared-references/output-versioning.md +111 -0
  170. package/skills/shared-references/patent-format-cn.md +199 -0
  171. package/skills/shared-references/patent-format-ep.md +173 -0
  172. package/skills/shared-references/patent-format-us.md +161 -0
  173. package/skills/shared-references/patent-writing-principles.md +197 -0
  174. package/skills/shared-references/prior-art-databases.md +141 -0
  175. package/skills/shared-references/resumable-runs.md +109 -0
  176. package/skills/shared-references/review-scope-limits.md +81 -0
  177. package/skills/shared-references/review-tracing.md +391 -0
  178. package/skills/shared-references/reviewer-independence.md +79 -0
  179. package/skills/shared-references/reviewer-routing.md +852 -0
  180. package/skills/shared-references/skill-governance.md +104 -0
  181. package/skills/shared-references/taste-calibration.md +85 -0
  182. package/skills/shared-references/venue-checklists.md +114 -0
  183. package/skills/shared-references/wiki-helper-resolution.md +134 -0
  184. package/skills/shared-references/writing-principles.md +525 -0
  185. package/skills/skills-codex/README.md +102 -0
  186. package/skills/skills-codex/README_CN.md +100 -0
  187. package/skills/skills-codex/ablation-planner/SKILL.md +126 -0
  188. package/skills/skills-codex/alphaxiv/SKILL.md +186 -0
  189. package/skills/skills-codex/analyze-results/SKILL.md +45 -0
  190. package/skills/skills-codex/arxiv/SKILL.md +210 -0
  191. package/skills/skills-codex/auto-paper-improvement-loop/SKILL.md +574 -0
  192. package/skills/skills-codex/auto-review-loop/SKILL.md +500 -0
  193. package/skills/skills-codex/auto-review-loop-llm/SKILL.md +247 -0
  194. package/skills/skills-codex/auto-review-loop-minimax/SKILL.md +290 -0
  195. package/skills/skills-codex/citation-audit/SKILL.md +504 -0
  196. package/skills/skills-codex/claims-drafting/SKILL.md +239 -0
  197. package/skills/skills-codex/comm-lit-review/SKILL.md +299 -0
  198. package/skills/skills-codex/comm-lit-review/references/domain-taxonomy.md +57 -0
  199. package/skills/skills-codex/comm-lit-review/references/output-template.md +37 -0
  200. package/skills/skills-codex/comm-lit-review/references/source-policy.md +99 -0
  201. package/skills/skills-codex/comm-lit-review/references/venue-tiering.md +112 -0
  202. package/skills/skills-codex/deepxiv/SKILL.md +142 -0
  203. package/skills/skills-codex/dse-loop/SKILL.md +285 -0
  204. package/skills/skills-codex/embodiment-description/SKILL.md +129 -0
  205. package/skills/skills-codex/exa-search/SKILL.md +192 -0
  206. package/skills/skills-codex/experiment-audit/SKILL.md +286 -0
  207. package/skills/skills-codex/experiment-bridge/SKILL.md +356 -0
  208. package/skills/skills-codex/experiment-plan/SKILL.md +249 -0
  209. package/skills/skills-codex/experiment-queue/SKILL.md +401 -0
  210. package/skills/skills-codex/feishu-notify/SKILL.md +155 -0
  211. package/skills/skills-codex/figure-description/SKILL.md +138 -0
  212. package/skills/skills-codex/figure-spec/SKILL.md +252 -0
  213. package/skills/skills-codex/formula-derivation/SKILL.md +280 -0
  214. package/skills/skills-codex/gemini-search/SKILL.md +205 -0
  215. package/skills/skills-codex/grant-proposal/SKILL.md +626 -0
  216. package/skills/skills-codex/idea-creator/SKILL.md +405 -0
  217. package/skills/skills-codex/idea-discovery/SKILL.md +475 -0
  218. package/skills/skills-codex/idea-discovery-robot/SKILL.md +362 -0
  219. package/skills/skills-codex/integrity-forensics/SKILL.md +106 -0
  220. package/skills/skills-codex/interview-cheatsheet/SKILL.md +245 -0
  221. package/skills/skills-codex/invention-structuring/SKILL.md +188 -0
  222. package/skills/skills-codex/jurisdiction-format/SKILL.md +192 -0
  223. package/skills/skills-codex/kill-argument/SKILL.md +403 -0
  224. package/skills/skills-codex/mermaid-diagram/SKILL.md +379 -0
  225. package/skills/skills-codex/meta-apply/SKILL.md +154 -0
  226. package/skills/skills-codex/meta-optimize/SKILL.md +348 -0
  227. package/skills/skills-codex/monitor-experiment/SKILL.md +98 -0
  228. package/skills/skills-codex/novelty-check/SKILL.md +89 -0
  229. package/skills/skills-codex/openalex/SKILL.md +228 -0
  230. package/skills/skills-codex/overleaf-sync/SKILL.md +220 -0
  231. package/skills/skills-codex/paper-claim-audit/SKILL.md +350 -0
  232. package/skills/skills-codex/paper-compile/SKILL.md +253 -0
  233. package/skills/skills-codex/paper-figure/SKILL.md +311 -0
  234. package/skills/skills-codex/paper-illustration/SKILL.md +690 -0
  235. package/skills/skills-codex/paper-illustration-image2/SKILL.md +383 -0
  236. package/skills/skills-codex/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  237. package/skills/skills-codex/paper-plan/SKILL.md +278 -0
  238. package/skills/skills-codex/paper-poster/SKILL.md +19 -0
  239. package/skills/skills-codex/paper-poster-html/SKILL.md +377 -0
  240. package/skills/skills-codex/paper-slides/SKILL.md +571 -0
  241. package/skills/skills-codex/paper-talk/SKILL.md +381 -0
  242. package/skills/skills-codex/paper-write/SKILL.md +411 -0
  243. package/skills/skills-codex/paper-write/templates/IEEEtran.bst +2409 -0
  244. package/skills/skills-codex/paper-write/templates/IEEEtran.cls +6347 -0
  245. package/skills/skills-codex/paper-write/templates/aaai2026.bst +1493 -0
  246. package/skills/skills-codex/paper-write/templates/aaai2026.sty +315 -0
  247. package/skills/skills-codex/paper-write/templates/aaai2026.tex +952 -0
  248. package/skills/skills-codex/paper-write/templates/acl.sty +312 -0
  249. package/skills/skills-codex/paper-write/templates/acl2026.tex +377 -0
  250. package/skills/skills-codex/paper-write/templates/acl_natbib.bst +1940 -0
  251. package/skills/skills-codex/paper-write/templates/acm.bst +3081 -0
  252. package/skills/skills-codex/paper-write/templates/acm_mm2026.tex +204 -0
  253. package/skills/skills-codex/paper-write/templates/acmart.cls +3520 -0
  254. package/skills/skills-codex/paper-write/templates/cvpr.bst +1448 -0
  255. package/skills/skills-codex/paper-write/templates/cvpr.sty +508 -0
  256. package/skills/skills-codex/paper-write/templates/cvpr2026.tex +63 -0
  257. package/skills/skills-codex/paper-write/templates/iclr2026.tex +84 -0
  258. package/skills/skills-codex/paper-write/templates/iclr2026_conference.bst +1440 -0
  259. package/skills/skills-codex/paper-write/templates/iclr2026_conference.sty +246 -0
  260. package/skills/skills-codex/paper-write/templates/icml2026.sty +767 -0
  261. package/skills/skills-codex/paper-write/templates/icml2026.tex +662 -0
  262. package/skills/skills-codex/paper-write/templates/ieee_conference.tex +89 -0
  263. package/skills/skills-codex/paper-write/templates/ieee_journal.tex +93 -0
  264. package/skills/skills-codex/paper-write/templates/math_commands.tex +48 -0
  265. package/skills/skills-codex/paper-write/templates/neurips2026.tex +493 -0
  266. package/skills/skills-codex/paper-write/templates/neurips_2026.sty +437 -0
  267. package/skills/skills-codex/paper-writing/SKILL.md +731 -0
  268. package/skills/skills-codex/patent-novelty-check/SKILL.md +153 -0
  269. package/skills/skills-codex/patent-pipeline/SKILL.md +344 -0
  270. package/skills/skills-codex/patent-review/SKILL.md +202 -0
  271. package/skills/skills-codex/pixel-art/SKILL.md +139 -0
  272. package/skills/skills-codex/prior-art-search/SKILL.md +146 -0
  273. package/skills/skills-codex/proof-checker/SKILL.md +554 -0
  274. package/skills/skills-codex/proof-orchestrator/SKILL.md +260 -0
  275. package/skills/skills-codex/proof-orchestrator/references/audit-output-contract.md +126 -0
  276. package/skills/skills-codex/proof-orchestrator/references/deepseek-routing.md +76 -0
  277. package/skills/skills-codex/proof-orchestrator/references/dispatch-prompts.md +227 -0
  278. package/skills/skills-codex/proof-orchestrator/references/notation-audit.md +135 -0
  279. package/skills/skills-codex/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  280. package/skills/skills-codex/proof-orchestrator/references/stress-tests.md +38 -0
  281. package/skills/skills-codex/proof-writer/SKILL.md +222 -0
  282. package/skills/skills-codex/qzcli/SKILL.md +324 -0
  283. package/skills/skills-codex/rebuttal/SKILL.md +305 -0
  284. package/skills/skills-codex/render-html/SKILL.md +305 -0
  285. package/skills/skills-codex/render-html/scripts/__pycache__/render_html.cpython-314.pyc +0 -0
  286. package/skills/skills-codex/render-html/scripts/render_html.py +909 -0
  287. package/skills/skills-codex/render-html/scripts/templates/academic.html +342 -0
  288. package/skills/skills-codex/render-html/scripts/templates/dashboard.html +333 -0
  289. package/skills/skills-codex/research-lit/SKILL.md +464 -0
  290. package/skills/skills-codex/research-pipeline/SKILL.md +340 -0
  291. package/skills/skills-codex/research-refine/SKILL.md +721 -0
  292. package/skills/skills-codex/research-refine-pipeline/SKILL.md +186 -0
  293. package/skills/skills-codex/research-review/SKILL.md +135 -0
  294. package/skills/skills-codex/research-wiki/SKILL.md +421 -0
  295. package/skills/skills-codex/resubmit-pipeline/SKILL.md +444 -0
  296. package/skills/skills-codex/result-to-claim/SKILL.md +246 -0
  297. package/skills/skills-codex/run-experiment/SKILL.md +236 -0
  298. package/skills/skills-codex/semantic-scholar/SKILL.md +219 -0
  299. package/skills/skills-codex/serverless-modal/SKILL.md +335 -0
  300. package/skills/skills-codex/shared-references/acceptance-gate.md +336 -0
  301. package/skills/skills-codex/shared-references/assurance-contract.md +139 -0
  302. package/skills/skills-codex/shared-references/capture-antipatterns.md +84 -0
  303. package/skills/skills-codex/shared-references/citation-discipline.md +452 -0
  304. package/skills/skills-codex/shared-references/compute-env-contract.md +163 -0
  305. package/skills/skills-codex/shared-references/effort-contract.md +143 -0
  306. package/skills/skills-codex/shared-references/evidence-precheck.md +73 -0
  307. package/skills/skills-codex/shared-references/experiment-integrity.md +49 -0
  308. package/skills/skills-codex/shared-references/external-cadence.md +334 -0
  309. package/skills/skills-codex/shared-references/fan-out-pattern.md +375 -0
  310. package/skills/skills-codex/shared-references/injection-hygiene.md +132 -0
  311. package/skills/skills-codex/shared-references/integration-contract.md +372 -0
  312. package/skills/skills-codex/shared-references/output-composition.md +98 -0
  313. package/skills/skills-codex/shared-references/output-language.md +45 -0
  314. package/skills/skills-codex/shared-references/output-manifest.md +40 -0
  315. package/skills/skills-codex/shared-references/output-versioning.md +111 -0
  316. package/skills/skills-codex/shared-references/patent-format-cn.md +199 -0
  317. package/skills/skills-codex/shared-references/patent-format-ep.md +173 -0
  318. package/skills/skills-codex/shared-references/patent-format-us.md +161 -0
  319. package/skills/skills-codex/shared-references/patent-writing-principles.md +197 -0
  320. package/skills/skills-codex/shared-references/prior-art-databases.md +141 -0
  321. package/skills/skills-codex/shared-references/resumable-runs.md +125 -0
  322. package/skills/skills-codex/shared-references/review-scope-limits.md +81 -0
  323. package/skills/skills-codex/shared-references/review-tracing.md +144 -0
  324. package/skills/skills-codex/shared-references/reviewer-independence.md +66 -0
  325. package/skills/skills-codex/shared-references/reviewer-routing.md +128 -0
  326. package/skills/skills-codex/shared-references/skill-governance.md +119 -0
  327. package/skills/skills-codex/shared-references/taste-calibration.md +90 -0
  328. package/skills/skills-codex/shared-references/venue-checklists.md +73 -0
  329. package/skills/skills-codex/shared-references/wiki-helper-resolution.md +69 -0
  330. package/skills/skills-codex/shared-references/writing-principles.md +525 -0
  331. package/skills/skills-codex/slides-polish/SKILL.md +563 -0
  332. package/skills/skills-codex/specification-writing/SKILL.md +211 -0
  333. package/skills/skills-codex/system-profile/SKILL.md +103 -0
  334. package/skills/skills-codex/training-check/SKILL.md +83 -0
  335. package/skills/skills-codex/vast-gpu/SKILL.md +394 -0
  336. package/skills/skills-codex/web-debug-search/SKILL.md +334 -0
  337. package/skills/skills-codex/wiki-enrich/SKILL.md +255 -0
  338. package/skills/skills-codex/writing-systems-papers/SKILL.md +184 -0
  339. package/skills/skills-codex-claude-review/README.md +79 -0
  340. package/skills/skills-codex-claude-review/README_CN.md +78 -0
  341. package/skills/skills-codex-claude-review/auto-paper-improvement-loop/SKILL.md +581 -0
  342. package/skills/skills-codex-claude-review/auto-review-loop/SKILL.md +510 -0
  343. package/skills/skills-codex-claude-review/novelty-check/SKILL.md +102 -0
  344. package/skills/skills-codex-claude-review/paper-figure/SKILL.md +319 -0
  345. package/skills/skills-codex-claude-review/paper-plan/SKILL.md +287 -0
  346. package/skills/skills-codex-claude-review/paper-write/SKILL.md +420 -0
  347. package/skills/skills-codex-claude-review/research-refine/SKILL.md +732 -0
  348. package/skills/skills-codex-claude-review/research-review/SKILL.md +149 -0
  349. package/skills/skills-codex-gemini-review/README.md +176 -0
  350. package/skills/skills-codex-gemini-review/README_CN.md +175 -0
  351. package/skills/skills-codex-gemini-review/auto-paper-improvement-loop/SKILL.md +331 -0
  352. package/skills/skills-codex-gemini-review/auto-review-loop/SKILL.md +304 -0
  353. package/skills/skills-codex-gemini-review/grant-proposal/SKILL.md +630 -0
  354. package/skills/skills-codex-gemini-review/idea-creator/SKILL.md +263 -0
  355. package/skills/skills-codex-gemini-review/idea-discovery/SKILL.md +275 -0
  356. package/skills/skills-codex-gemini-review/idea-discovery-robot/SKILL.md +365 -0
  357. package/skills/skills-codex-gemini-review/novelty-check/SKILL.md +92 -0
  358. package/skills/skills-codex-gemini-review/paper-figure/SKILL.md +289 -0
  359. package/skills/skills-codex-gemini-review/paper-plan/SKILL.md +265 -0
  360. package/skills/skills-codex-gemini-review/paper-poster-html/SKILL.md +102 -0
  361. package/skills/skills-codex-gemini-review/paper-slides/SKILL.md +582 -0
  362. package/skills/skills-codex-gemini-review/paper-write/SKILL.md +346 -0
  363. package/skills/skills-codex-gemini-review/paper-writing/SKILL.md +312 -0
  364. package/skills/skills-codex-gemini-review/research-refine/SKILL.md +674 -0
  365. package/skills/skills-codex-gemini-review/research-review/SKILL.md +112 -0
  366. package/skills/slides-polish/SKILL.md +565 -0
  367. package/skills/specification-writing/SKILL.md +211 -0
  368. package/skills/system-profile/SKILL.md +103 -0
  369. package/skills/training-check/SKILL.md +132 -0
  370. package/skills/vast-gpu/SKILL.md +394 -0
  371. package/skills/web-debug-search/SKILL.md +334 -0
  372. package/skills/wiki-enrich/SKILL.md +257 -0
  373. package/skills/writing-systems-papers/SKILL.md +184 -0
  374. package/templates/CLAUDE_MD_TEMPLATE.md +29 -0
  375. package/templates/EXPERIMENT_LOG_TEMPLATE.md +47 -0
  376. package/templates/EXPERIMENT_PLAN_TEMPLATE.md +51 -0
  377. package/templates/EXPERIMENT_PLAN_TEMPLATE_CN.md +53 -0
  378. package/templates/FINDINGS_TEMPLATE.md +52 -0
  379. package/templates/IDEA_CANDIDATES_TEMPLATE.md +47 -0
  380. package/templates/IDEA_CANDIDATES_TEMPLATE_CN.md +47 -0
  381. package/templates/INVENTION_BRIEF_TEMPLATE.md +87 -0
  382. package/templates/MANIFEST_TEMPLATE.md +7 -0
  383. package/templates/NARRATIVE_REPORT_TEMPLATE.md +49 -0
  384. package/templates/PAPER_PLAN_TEMPLATE.md +47 -0
  385. package/templates/PATENT_CLAIMS_TEMPLATE.md +78 -0
  386. package/templates/PATENT_SPECIFICATION_TEMPLATE.md +68 -0
  387. package/templates/README.md +57 -0
  388. package/templates/RESEARCH_BRIEF_TEMPLATE.md +35 -0
  389. package/templates/RESEARCH_BRIEF_TEMPLATE_CN.md +41 -0
  390. package/templates/RESEARCH_CONTRACT_TEMPLATE.md +60 -0
  391. package/templates/claude-hooks/corpus_write_guard.json +16 -0
  392. package/templates/claude-hooks/corpus_write_guard.py +85 -0
  393. package/templates/claude-hooks/meta_logging.json +74 -0
  394. package/templates/gitignore-trace.txt +3 -0
  395. package/tools/__pycache__/check_skills_inventory.cpython-314.pyc +0 -0
  396. package/tools/arxiv_fetch.py +311 -0
  397. package/tools/capture_filter.py +126 -0
  398. package/tools/check_skills_inventory.py +273 -0
  399. package/tools/convert_skills_to_llm_chat.py +282 -0
  400. package/tools/copilot_native_evidence.py +818 -0
  401. package/tools/deepxiv_fetch.py +213 -0
  402. package/tools/evidence_check.py +212 -0
  403. package/tools/exa_search.py +425 -0
  404. package/tools/experiment_queue/README.md +118 -0
  405. package/tools/experiment_queue/build_manifest.py +44 -0
  406. package/tools/experiment_queue/queue_manager.py +44 -0
  407. package/tools/extract_paper_style.py +560 -0
  408. package/tools/figure_renderer.py +69 -0
  409. package/tools/forensics_gate.py +669 -0
  410. package/tools/generate_codex_claude_review_overrides.py +299 -0
  411. package/tools/idea_discovery_gate.py +256 -0
  412. package/tools/install_aris.ps1 +1372 -0
  413. package/tools/install_aris.sh +1370 -0
  414. package/tools/install_aris_codex.sh +1023 -0
  415. package/tools/install_aris_copilot.sh +1052 -0
  416. package/tools/iteration_log.py +143 -0
  417. package/tools/lint_skills_helpers.sh +84 -0
  418. package/tools/meta_opt/check_ready.sh +80 -0
  419. package/tools/meta_opt/log_event.sh +91 -0
  420. package/tools/meta_opt/trigger_eval.py +280 -0
  421. package/tools/meta_opt/trigger_evals.sample.json +28 -0
  422. package/tools/openalex_fetch.py +326 -0
  423. package/tools/overleaf_audit.sh +104 -0
  424. package/tools/overleaf_setup.sh +150 -0
  425. package/tools/paper_illustration_image2.py +62 -0
  426. package/tools/provenance.py +294 -0
  427. package/tools/research_wiki.py +1720 -0
  428. package/tools/review_gate.py +502 -0
  429. package/tools/run_state.py +399 -0
  430. package/tools/save_trace.sh +477 -0
  431. package/tools/semantic_scholar_fetch.py +438 -0
  432. package/tools/skill-groups.tsv +116 -0
  433. package/tools/skill_picker.py +238 -0
  434. package/tools/smart_update.ps1 +521 -0
  435. package/tools/smart_update.sh +591 -0
  436. package/tools/smart_update_codex.sh +419 -0
  437. package/tools/smart_update_copilot.sh +605 -0
  438. package/tools/threat_scan.py +222 -0
  439. package/tools/verify_paper_audits.sh +487 -0
  440. package/tools/verify_papers.py +613 -0
  441. package/tools/verify_wiki_coverage.sh +176 -0
  442. package/tools/watchdog.py +485 -0
@@ -0,0 +1,259 @@
1
+ ---
2
+ name: auto-review-loop-llm
3
+ description: Autonomous research review loop using any OpenAI-compatible LLM API. Configure via llm-chat MCP server or environment variables. Trigger with "auto review loop llm" or "llm review".
4
+ argument-hint: "[topic-or-scope]"
5
+ allowed-tools: Bash(*), Read, Grep, Glob, Write, Edit, Skill
6
+ ---
7
+
8
+ # Auto Review Loop (Generic LLM): Autonomous Research Improvement
9
+
10
+ > 🔒 **Do not wrap this skill in `/loop`, `/schedule`, or `CronCreate`.** Like
11
+ > `/auto-review-loop`, it already loops internally (review → fix → re-review),
12
+ > feeding each round's prior-round summary into the next review prompt (the
13
+ > backend is a stateless per-round API/MCP call, not a shared thread). An
14
+ > external timer re-enters from the top each tick, dropping that accumulated
15
+ > context and firing the verdict on wall-clock time instead of on artifact
16
+ > change — zero new signal, full token cost. Schedule the *external wait that
17
+ > precedes it*, not the verdict. See
18
+ > [`shared-references/external-cadence.md`](../shared-references/external-cadence.md).
19
+
20
+ Autonomously iterate: review → implement fixes → re-review, until the external reviewer gives a positive assessment or MAX_ROUNDS is reached.
21
+
22
+ ## Context: $ARGUMENTS
23
+
24
+ ## Constants
25
+
26
+ - MAX_ROUNDS = 4
27
+ - POSITIVE_THRESHOLD: score >= 6/10 **AND** verdict ∈ {"ready", "almost"} — **both** must hold, matching the operative STOP check below. Verdict vocabulary is {"ready", "almost", "not ready"}. (Earlier wording used `or` and a stale verdict set; the `AND` form is authoritative.)
28
+ - REVIEW_DOC: `review-stage/AUTO_REVIEW.md` (cumulative log) *(fall back to `./AUTO_REVIEW.md` for legacy projects)*
29
+
30
+ ## LLM Configuration
31
+
32
+ This skill uses **any OpenAI-compatible API** for external review via the `llm-chat` MCP server.
33
+
34
+ ### Configuration via MCP Server (Recommended)
35
+
36
+ Add to `~/.claude/settings.json`:
37
+
38
+ ```json
39
+ {
40
+ "mcpServers": {
41
+ "llm-chat": {
42
+ "command": "/usr/bin/python3",
43
+ "args": ["/Users/yourname/.claude/mcp-servers/llm-chat/server.py"],
44
+ "env": {
45
+ "LLM_API_KEY": "your-api-key",
46
+ "LLM_BASE_URL": "https://api.deepseek.com/v1",
47
+ "LLM_MODEL": "deepseek-chat"
48
+ }
49
+ }
50
+ }
51
+ }
52
+ ```
53
+
54
+ ### Supported Providers
55
+
56
+ | Provider | LLM_BASE_URL | LLM_MODEL |
57
+ |----------|--------------|-----------|
58
+ | **OpenAI** | `https://api.openai.com/v1` | `gpt-4o`, `o3` |
59
+ | **DeepSeek** | `https://api.deepseek.com/v1` | `deepseek-chat`, `deepseek-reasoner` |
60
+ | **MiniMax** | `https://api.minimax.io/v1` | `MiniMax-M3` |
61
+ | **Kimi (Moonshot)** | `https://api.moonshot.cn/v1` | `moonshot-v1-8k`, `moonshot-v1-32k` |
62
+ | **ZhiPu (GLM)** | `https://open.bigmodel.cn/api/paas/v4` | `glm-4`, `glm-4-plus` |
63
+ | **SiliconFlow** | `https://api.siliconflow.cn/v1` | `Qwen/Qwen2.5-72B-Instruct` |
64
+ | **阿里云百炼** | `https://dashscope.aliyuncs.com/compatible-mode/v1` | `qwen-max` |
65
+ | **零一万物** | `https://api.lingyiwanwu.com/v1` | `yi-large` |
66
+
67
+ ## API Call Method
68
+
69
+ **Primary: MCP Tool**
70
+
71
+ ```
72
+ mcp__llm-chat__chat:
73
+ prompt: |
74
+ [Review prompt content]
75
+ model: "deepseek-chat"
76
+ system: "You are a senior ML reviewer..."
77
+ ```
78
+
79
+ **Fallback: curl**
80
+
81
+ ```bash
82
+ curl -s "${LLM_BASE_URL}/chat/completions" \
83
+ -H "Content-Type: application/json" \
84
+ -H "Authorization: Bearer ${LLM_API_KEY}" \
85
+ -d '{
86
+ "model": "${LLM_MODEL}",
87
+ "messages": [
88
+ {"role": "system", "content": "You are a senior ML reviewer..."},
89
+ {"role": "user", "content": "[review prompt]"}
90
+ ],
91
+ "max_tokens": 4096
92
+ }'
93
+ ```
94
+
95
+ ## State Persistence (Compact Recovery)
96
+
97
+ Persist state to `review-stage/REVIEW_STATE.json` after each round:
98
+
99
+ ```json
100
+ {
101
+ "round": 2,
102
+ "status": "in_progress",
103
+ "last_score": 5.0,
104
+ "last_verdict": "not ready",
105
+ "pending_experiments": [],
106
+ "timestamp": "2026-03-15T10:00:00"
107
+ }
108
+ ```
109
+
110
+ **Write this file at the end of every Phase E** (after documenting the round).
111
+
112
+ **On completion**, set `"status": "completed"`.
113
+
114
+ ## Workflow
115
+
116
+ ### Initialization
117
+
118
+ 1. **Check `review-stage/REVIEW_STATE.json`** for recovery *(fall back to `./REVIEW_STATE.json` if not found — legacy path)*
119
+ 2. Read project context and prior reviews
120
+ 3. Initialize round counter
121
+
122
+ ### Loop (up to MAX_ROUNDS)
123
+
124
+ #### Phase A: Review
125
+
126
+ **If MCP available:**
127
+ ```
128
+ mcp__llm-chat__chat:
129
+ system: "You are a senior ML reviewer (NeurIPS/ICML level)."
130
+ prompt: |
131
+ [Round N/MAX_ROUNDS of autonomous review loop]
132
+
133
+ [Full research context: claims, methods, results, known weaknesses]
134
+ [Changes since last round, if any]
135
+
136
+ 1. Score this work 1-10 for a top venue
137
+ 2. List remaining critical weaknesses (ranked by severity)
138
+ 3. For each weakness, specify the MINIMUM fix
139
+ 4. State clearly: is this READY for submission? Yes/No/Almost
140
+
141
+ Be brutally honest. If the work is ready, say so clearly.
142
+ ```
143
+
144
+ **If MCP NOT available:**
145
+ ```bash
146
+ curl -s "${LLM_BASE_URL}/chat/completions" \
147
+ -H "Content-Type: application/json" \
148
+ -H "Authorization: Bearer ${LLM_API_KEY}" \
149
+ -d '{
150
+ "model": "${LLM_MODEL}",
151
+ "messages": [
152
+ {"role": "system", "content": "You are a senior ML reviewer (NeurIPS/ICML level)."},
153
+ {"role": "user", "content": "[Full review prompt]"}
154
+ ],
155
+ "max_tokens": 4096
156
+ }'
157
+ ```
158
+
159
+ #### Phase B: Parse Assessment
160
+
161
+ **CRITICAL: Save the FULL raw response** verbatim. Then extract:
162
+ - **Score** (numeric 1-10)
163
+ - **Verdict** ("ready" / "almost" / "not ready")
164
+ - **Action items** (ranked list of fixes)
165
+
166
+ **STOP**: If score >= 6 AND verdict ∈ {"ready", "almost"} (exact — "not ready" does NOT qualify)
167
+
168
+ #### Phase C: Implement Fixes
169
+
170
+ Priority: metric additions > reframing > new experiments
171
+
172
+ #### Phase D: Wait for Results
173
+
174
+ Monitor remote experiments
175
+
176
+ #### Phase E: Document Round
177
+
178
+ Append to `review-stage/AUTO_REVIEW.md`:
179
+
180
+ ```markdown
181
+ ## Round N (timestamp)
182
+
183
+ ### Assessment (Summary)
184
+ - Score: X/10
185
+ - Verdict: [ready/almost/not ready]
186
+ - Key criticisms: [bullet list]
187
+
188
+ ### Reviewer Raw Response
189
+
190
+ <details>
191
+ <summary>Click to expand full reviewer response</summary>
192
+
193
+ [Paste the COMPLETE raw response here — verbatim, unedited.]
194
+
195
+ </details>
196
+
197
+ ### Actions Taken
198
+ - [what was implemented/changed]
199
+
200
+ ### Results
201
+ - [experiment outcomes, if any]
202
+
203
+ ### Status
204
+ - [continuing to round N+1 / stopping]
205
+ ```
206
+
207
+ **Write `review-stage/REVIEW_STATE.json`** with current state.
208
+
209
+ ### Termination
210
+
211
+ 1. Set `review-stage/REVIEW_STATE.json` status to "completed"
212
+ 2. Write final summary
213
+
214
+ ## Key Rules
215
+
216
+ - **Large file handling**: If the Write tool fails due to file size, immediately retry using Bash (`cat << 'EOF' > file`) to write in chunks. Do NOT ask the user for permission — just do it silently.
217
+
218
+ - **Anti-hallucination citations**: When adding references, NEVER fabricate BibTeX. Use DBLP → CrossRef → `[VERIFY]` chain. Do NOT generate BibTeX from memory.
219
+ - Be honest about weaknesses
220
+ - Implement fixes BEFORE re-reviewing
221
+ - Document everything
222
+ - Include previous context in round 2+ prompts
223
+ - Prefer MCP tool over curl when available
224
+
225
+ ## Prompt Template for Round 2+
226
+
227
+ ```
228
+ mcp__llm-chat__chat:
229
+ system: "You are a senior ML reviewer (NeurIPS/ICML level)."
230
+ prompt: |
231
+ [Round N/MAX_ROUNDS of autonomous review loop]
232
+
233
+ ## Previous Review Summary (Round N-1)
234
+ - Previous Score: X/10
235
+ - Previous Verdict: [ready/almost/not ready]
236
+ - Previous Key Weaknesses: [list]
237
+
238
+ ## Changes Since Last Review
239
+ 1. [Action 1]: [result]
240
+ 2. [Action 2]: [result]
241
+
242
+ ## Updated Results
243
+ [paste updated metrics/tables]
244
+
245
+ Please re-score and re-assess:
246
+ 1. Score this work 1-10 for a top venue
247
+ 2. List remaining critical weaknesses (ranked by severity)
248
+ 3. For each weakness, specify the MINIMUM fix
249
+ 4. State clearly: is this READY for submission? Yes/No/Almost
250
+
251
+ Be brutally honest. If the work is ready, say so clearly.
252
+ ```
253
+
254
+ ## Output Protocols
255
+
256
+ > Follow these shared protocols for all output files:
257
+ > - **[Output Versioning Protocol](../shared-references/output-versioning.md)** — write timestamped file first, then copy to fixed name
258
+ > - **[Output Manifest Protocol](../shared-references/output-manifest.md)** — log every output to MANIFEST.md
259
+ > - **[Output Language Protocol](../shared-references/output-language.md)** — respect the project's language setting
@@ -0,0 +1,302 @@
1
+ ---
2
+ name: auto-review-loop-minimax
3
+ description: Autonomous multi-round research review loop using MiniMax API. Use when you want to use MiniMax instead of Codex MCP for external review. Trigger with "auto review loop minimax" or "minimax review".
4
+ argument-hint: "[topic-or-scope]"
5
+ allowed-tools: Bash(*), Read, Grep, Glob, Write, Edit, Skill
6
+ ---
7
+
8
+ # Auto Review Loop (MiniMax Version): Autonomous Research Improvement
9
+
10
+ > 🔒 **Do not wrap this skill in `/loop`, `/schedule`, or `CronCreate`.** Like
11
+ > `/auto-review-loop`, it already loops internally (review → fix → re-review),
12
+ > feeding each round's prior-round summary into the next review prompt (the
13
+ > backend is a stateless per-round API call, not a shared thread). An external
14
+ > timer re-enters from the top each tick, dropping that accumulated context and
15
+ > firing the verdict on wall-clock time instead of on artifact change — zero
16
+ > new signal, full token cost. Schedule the *external wait that precedes it*,
17
+ > not the verdict. See
18
+ > [`shared-references/external-cadence.md`](../shared-references/external-cadence.md).
19
+
20
+ Autonomously iterate: review → implement fixes → re-review, until the external reviewer gives a positive assessment or MAX_ROUNDS is reached.
21
+
22
+ ## Context: $ARGUMENTS
23
+
24
+ ## Constants
25
+
26
+ - MAX_ROUNDS = 4
27
+ - POSITIVE_THRESHOLD: score >= 6/10 **AND** verdict ∈ {"ready", "almost"} — **both** must hold, matching the operative STOP CONDITION below. Verdict vocabulary is {"ready", "almost", "not ready"}. (Earlier wording used `or` and a stale verdict set; the `AND` form is authoritative.)
28
+ - REVIEW_DOC: `review-stage/AUTO_REVIEW.md` (cumulative log) *(fall back to `./AUTO_REVIEW.md` for legacy projects)*
29
+ - REVIEWER_MODEL = `MiniMax-M3` — Model used via MiniMax API
30
+
31
+ ## API Configuration
32
+
33
+ This skill uses MiniMax API for external review. Two methods are supported:
34
+
35
+ ### Method 1: MCP Tool (Primary)
36
+
37
+ If `mcp__minimax-chat__minimax_chat` is available, use it:
38
+
39
+ ```
40
+ mcp__minimax-chat__minimax_chat:
41
+ prompt: |
42
+ [Review prompt content]
43
+ model: "MiniMax-M3"
44
+ system: "You are a senior machine learning researcher..."
45
+ ```
46
+
47
+ ### Method 2: curl (Fallback)
48
+
49
+ If MCP is not available, use curl directly:
50
+
51
+ ```bash
52
+ curl -s "https://api.minimax.io/v1/chat/completions" \
53
+ -H "Content-Type: application/json" \
54
+ -H "Authorization: Bearer $MINIMAX_API_KEY" \
55
+ -d '{
56
+ "model": "MiniMax-M3",
57
+ "messages": [
58
+ {"role": "system", "content": "You are a senior ML researcher..."},
59
+ {"role": "user", "content": "[Review prompt]"}
60
+ ],
61
+ "max_tokens": 4096
62
+ }'
63
+ ```
64
+
65
+ **API Key**: Read from `~/.claude/settings.json` under `env.MINIMAX_API_KEY`, or from environment variable.
66
+
67
+ **Why MiniMax instead of Codex MCP?** Codex CLI uses OpenAI's Responses API (`/v1/responses`) which is not supported by third-party providers. See: https://github.com/openai/codex/discussions/7782
68
+
69
+ ## State Persistence (Compact Recovery)
70
+
71
+ Long-running loops may hit the context window limit, triggering automatic compaction. To survive this, persist state to `review-stage/REVIEW_STATE.json` after each round:
72
+
73
+ ```json
74
+ {
75
+ "round": 2,
76
+ "status": "in_progress",
77
+ "last_score": 5.0,
78
+ "last_verdict": "not ready",
79
+ "pending_experiments": ["screen_name_1"],
80
+ "timestamp": "2026-03-13T21:00:00"
81
+ }
82
+ ```
83
+
84
+ **Write this file at the end of every Phase E** (after documenting the round). Overwrite each time — only the latest state matters.
85
+
86
+ **On completion** (positive assessment or max rounds), set `"status": "completed"` so future invocations don't accidentally resume a finished loop.
87
+
88
+ ## Workflow
89
+
90
+ ### Initialization
91
+
92
+ 1. **Check for `review-stage/REVIEW_STATE.json`** *(fall back to `./REVIEW_STATE.json` if not found — legacy path)*:
93
+ - If neither path exists: **fresh start** (normal case)
94
+ - If it exists AND `status` is `"completed"`: **fresh start** (previous loop finished normally)
95
+ - If it exists AND `status` is `"in_progress"` AND `timestamp` is older than 24 hours: **fresh start** (stale state from a killed/abandoned run — delete the file and start over)
96
+ - If it exists AND `status` is `"in_progress"` AND `timestamp` is within 24 hours: **resume**
97
+ - Read the state file to recover `round`, `last_score`, `pending_experiments`
98
+ - Read `review-stage/AUTO_REVIEW.md` to restore full context of prior rounds *(fall back to `./AUTO_REVIEW.md`)*
99
+ - If `pending_experiments` is non-empty, check if they have completed (e.g., check screen sessions)
100
+ - Resume from the next round (round = saved round + 1)
101
+ - Log: "Recovered from context compaction. Resuming at Round N."
102
+ 2. Read project narrative documents, memory files, and any prior review documents
103
+ 3. Read recent experiment results (check output directories, logs)
104
+ 4. Identify current weaknesses and open TODOs from prior reviews
105
+ 5. Initialize round counter = 1 (unless recovered from state file)
106
+ 6. Create/update `review-stage/AUTO_REVIEW.md` with header and timestamp
107
+
108
+ ### Loop (repeat up to MAX_ROUNDS)
109
+
110
+ #### Phase A: Review
111
+
112
+ Send comprehensive context to the external reviewer.
113
+
114
+ **Check MCP availability first**, then use appropriate method:
115
+
116
+ **If MCP available (Primary):**
117
+ ```
118
+ Use mcp__minimax-chat__minimax_chat tool with:
119
+ - system: "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
120
+ - prompt: [Full review prompt with context]
121
+ - model: "MiniMax-M3"
122
+ ```
123
+
124
+ **If MCP NOT available (Fallback):**
125
+ ```bash
126
+ curl -s "https://api.minimax.io/v1/chat/completions" \
127
+ -H "Content-Type: application/json" \
128
+ -H "Authorization: Bearer $MINIMAX_API_KEY" \
129
+ -d '{
130
+ "model": "MiniMax-M3",
131
+ "messages": [
132
+ {
133
+ "role": "system",
134
+ "content": "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
135
+ },
136
+ {
137
+ "role": "user",
138
+ "content": "[Round N/MAX_ROUNDS of autonomous review loop]\n\n[Full research context: claims, methods, results, known weaknesses]\n[Changes since last round, if any]\n[For round 2+: Summary of previous review feedback and what was addressed]\n\nPlease act as a senior ML reviewer (NeurIPS/ICML level).\n\n1. Score this work 1-10 for a top venue\n2. List remaining critical weaknesses (ranked by severity)\n3. For each weakness, specify the MINIMUM fix (experiment, analysis, or reframing)\n4. State clearly: is this READY for submission? Yes/No/Almost\n\nBe brutally honest. If the work is ready, say so clearly."
139
+ }
140
+ ],
141
+ "max_tokens": 4096
142
+ }'
143
+ ```
144
+
145
+ **Note**: Each round is a standalone API call. For round 2+, include the summary of previous reviews and changes in the prompt itself.
146
+
147
+ #### Phase B: Parse Assessment
148
+
149
+ **CRITICAL: Save the FULL raw response** from the external reviewer verbatim (store in a variable for Phase E). Do NOT discard or summarize — the raw text is the primary record.
150
+
151
+ Then extract structured fields:
152
+ - **Score** (numeric 1-10)
153
+ - **Verdict** ("ready" / "almost" / "not ready")
154
+ - **Action items** (ranked list of fixes)
155
+
156
+ **STOP CONDITION**: If score >= 6 AND verdict ∈ {"ready", "almost"} (exact match — "not ready" does NOT qualify) → stop loop, document final state.
157
+
158
+ #### Phase C: Implement Fixes (if not stopping)
159
+
160
+ For each action item (highest priority first):
161
+
162
+ 1. **Code changes**: Write/modify experiment scripts, model code, analysis scripts
163
+ 2. **Run experiments**: Deploy to GPU server via SSH + screen/tmux
164
+ 3. **Analysis**: Run evaluation, collect results, update figures/tables
165
+ 4. **Documentation**: Update project notes and review document
166
+
167
+ Prioritization rules:
168
+ - Skip fixes requiring excessive compute (flag for manual follow-up)
169
+ - Skip fixes requiring external data/models not available
170
+ - Prefer reframing/analysis over new experiments when both address the concern
171
+ - Always implement metric additions (cheap, high impact)
172
+
173
+ #### Phase D: Wait for Results
174
+
175
+ If experiments were launched:
176
+ - Monitor remote sessions for completion
177
+ - Collect results from output files and logs
178
+
179
+ #### Phase E: Document Round
180
+
181
+ Append to `review-stage/AUTO_REVIEW.md`:
182
+
183
+ ```markdown
184
+ ## Round N (timestamp)
185
+
186
+ ### Assessment (Summary)
187
+ - Score: X/10
188
+ - Verdict: [ready/almost/not ready]
189
+ - Key criticisms: [bullet list]
190
+
191
+ ### Reviewer Raw Response
192
+
193
+ <details>
194
+ <summary>Click to expand full reviewer response</summary>
195
+
196
+ [Paste the COMPLETE raw response from the external reviewer here — verbatim, unedited.
197
+ This is the authoritative record. Do NOT truncate or paraphrase.]
198
+
199
+ </details>
200
+
201
+ ### Actions Taken
202
+ - [what was implemented/changed]
203
+
204
+ ### Results
205
+ - [experiment outcomes, if any]
206
+
207
+ ### Status
208
+ - [continuing to round N+1 / stopping]
209
+ ```
210
+
211
+ **Write `review-stage/REVIEW_STATE.json`** with current round, score, verdict, and any pending experiments.
212
+
213
+ Increment round counter → back to Phase A.
214
+
215
+ ### Termination
216
+
217
+ When loop ends (positive assessment or max rounds):
218
+
219
+ 1. Update `review-stage/REVIEW_STATE.json` with `"status": "completed"`
220
+ 2. Write final summary to `review-stage/AUTO_REVIEW.md`
221
+ 3. Update project notes with conclusions
222
+ 4. If stopped at max rounds without positive assessment:
223
+ - List remaining blockers
224
+ - Estimate effort needed for each
225
+ - Suggest whether to continue manually or pivot
226
+
227
+ ## Key Rules
228
+
229
+ - **Large file handling**: If the Write tool fails due to file size, immediately retry using Bash (`cat << 'EOF' > file`) to write in chunks. Do NOT ask the user for permission — just do it silently.
230
+
231
+ - **Anti-hallucination citations**: When adding references, NEVER fabricate BibTeX. Use DBLP → CrossRef → `[VERIFY]` chain. Do NOT generate BibTeX from memory.
232
+ - Be honest — include negative results and failed experiments
233
+ - Do NOT hide weaknesses to game a positive score
234
+ - Implement fixes BEFORE re-reviewing (don't just promise to fix)
235
+ - If an experiment takes > 30 minutes, launch it and continue with other fixes while waiting
236
+ - Document EVERYTHING — the review log should be self-contained
237
+ - Update project notes after each round, not just at the end
238
+ - For round 2+, always include previous review context in the prompt
239
+ - Prefer MCP tool over curl when available (more reliable)
240
+
241
+ ## Prompt Template for Round 2+
242
+
243
+ **MCP Method (Primary):**
244
+ ```
245
+ mcp__minimax-chat__minimax_chat:
246
+ model: "MiniMax-M3"
247
+ system: "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
248
+ prompt: |
249
+ [Round N/MAX_ROUNDS of autonomous review loop]
250
+
251
+ ## Previous Review Summary (Round N-1)
252
+ - Previous Score: X/10
253
+ - Previous Verdict: [ready/almost/not ready]
254
+ - Previous Key Weaknesses: [list]
255
+
256
+ ## Changes Since Last Review
257
+ 1. [Action 1]: [result]
258
+ 2. [Action 2]: [result]
259
+ 3. [Action 3]: [result]
260
+
261
+ ## Updated Results
262
+ [paste updated metrics/tables]
263
+
264
+ ## Current Research Context
265
+ [brief summary of claims, methods, current state]
266
+
267
+ Please re-score and re-assess:
268
+ 1. Score this work 1-10 for a top venue
269
+ 2. List remaining critical weaknesses (ranked by severity)
270
+ 3. For each weakness, specify the MINIMUM fix
271
+ 4. State clearly: is this READY for submission? Yes/No/Almost
272
+
273
+ Be brutally honest. If the work is ready, say so clearly.
274
+ ```
275
+
276
+ **curl Fallback:**
277
+ ```bash
278
+ curl -s "https://api.minimax.io/v1/chat/completions" \
279
+ -H "Content-Type: application/json" \
280
+ -H "Authorization: Bearer $MINIMAX_API_KEY" \
281
+ -d '{
282
+ "model": "MiniMax-M3",
283
+ "messages": [
284
+ {
285
+ "role": "system",
286
+ "content": "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
287
+ },
288
+ {
289
+ "role": "user",
290
+ "content": "[Round N/MAX_ROUNDS of autonomous review loop]\n\n## Previous Review Summary (Round N-1)\n- Previous Score: X/10\n- Previous Verdict: [ready/almost/not ready]\n- Previous Key Weaknesses: [list]\n\n## Changes Since Last Review\n1. [Action 1]: [result]\n2. [Action 2]: [result]\n3. [Action 3]: [result]\n\n## Updated Results\n[paste updated metrics/tables]\n\n## Current Research Context\n[brief summary of claims, methods, current state]\n\nPlease re-score and re-assess:\n1. Score this work 1-10 for a top venue\n2. List remaining critical weaknesses (ranked by severity)\n3. For each weakness, specify the MINIMUM fix\n4. State clearly: is this READY for submission? Yes/No/Almost\n\nBe brutally honest. If the work is ready, say so clearly."
291
+ }
292
+ ],
293
+ "max_tokens": 4096
294
+ }'
295
+ ```
296
+
297
+ ## Output Protocols
298
+
299
+ > Follow these shared protocols for all output files:
300
+ > - **[Output Versioning Protocol](../shared-references/output-versioning.md)** — write timestamped file first, then copy to fixed name
301
+ > - **[Output Manifest Protocol](../shared-references/output-manifest.md)** — log every output to MANIFEST.md
302
+ > - **[Output Language Protocol](../shared-references/output-language.md)** — respect the project's language setting