@claude-flow/cli 3.28.0 → 3.30.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (465) hide show
  1. package/.claude/agents/analysis/analyze-code-quality.md +178 -178
  2. package/.claude/agents/analysis/code-analyzer.md +209 -209
  3. package/.claude/agents/analysis/code-review/analyze-code-quality.md +178 -178
  4. package/.claude/agents/architecture/arch-system-design.md +156 -156
  5. package/.claude/agents/architecture/system-design/arch-system-design.md +154 -154
  6. package/.claude/agents/browser/browser-agent.yaml +182 -182
  7. package/.claude/agents/consensus/byzantine-coordinator.md +62 -62
  8. package/.claude/agents/consensus/crdt-synchronizer.md +996 -996
  9. package/.claude/agents/consensus/gossip-coordinator.md +62 -62
  10. package/.claude/agents/consensus/performance-benchmarker.md +850 -850
  11. package/.claude/agents/consensus/quorum-manager.md +822 -822
  12. package/.claude/agents/consensus/raft-manager.md +62 -62
  13. package/.claude/agents/consensus/security-manager.md +621 -621
  14. package/.claude/agents/core/planner.md +374 -374
  15. package/.claude/agents/custom/test-long-runner.md +44 -44
  16. package/.claude/agents/data/data-ml-model.md +444 -444
  17. package/.claude/agents/data/ml/data-ml-model.md +192 -192
  18. package/.claude/agents/development/backend/dev-backend-api.md +141 -141
  19. package/.claude/agents/development/dev-backend-api.md +344 -344
  20. package/.claude/agents/devops/ci-cd/ops-cicd-github.md +163 -163
  21. package/.claude/agents/devops/ops-cicd-github.md +164 -164
  22. package/.claude/agents/documentation/api-docs/docs-api-openapi.md +173 -173
  23. package/.claude/agents/documentation/docs-api-openapi.md +354 -354
  24. package/.claude/agents/flow-nexus/app-store.md +87 -87
  25. package/.claude/agents/flow-nexus/authentication.md +68 -68
  26. package/.claude/agents/flow-nexus/challenges.md +80 -80
  27. package/.claude/agents/flow-nexus/neural-network.md +87 -87
  28. package/.claude/agents/flow-nexus/payments.md +82 -82
  29. package/.claude/agents/flow-nexus/sandbox.md +75 -75
  30. package/.claude/agents/flow-nexus/swarm.md +75 -75
  31. package/.claude/agents/flow-nexus/user-tools.md +95 -95
  32. package/.claude/agents/flow-nexus/workflow.md +83 -83
  33. package/.claude/agents/github/code-review-swarm.md +377 -377
  34. package/.claude/agents/github/github-modes.md +172 -172
  35. package/.claude/agents/github/issue-tracker.md +575 -575
  36. package/.claude/agents/github/multi-repo-swarm.md +552 -552
  37. package/.claude/agents/github/pr-manager.md +437 -437
  38. package/.claude/agents/github/project-board-sync.md +508 -508
  39. package/.claude/agents/github/release-manager.md +604 -604
  40. package/.claude/agents/github/release-swarm.md +582 -582
  41. package/.claude/agents/github/repo-architect.md +397 -397
  42. package/.claude/agents/github/swarm-issue.md +572 -572
  43. package/.claude/agents/github/swarm-pr.md +427 -427
  44. package/.claude/agents/github/sync-coordinator.md +451 -451
  45. package/.claude/agents/github/workflow-automation.md +902 -902
  46. package/.claude/agents/goal/agent.md +815 -815
  47. package/.claude/agents/optimization/benchmark-suite.md +664 -664
  48. package/.claude/agents/optimization/load-balancer.md +430 -430
  49. package/.claude/agents/optimization/performance-monitor.md +671 -671
  50. package/.claude/agents/optimization/resource-allocator.md +673 -673
  51. package/.claude/agents/optimization/topology-optimizer.md +807 -807
  52. package/.claude/agents/payments/agentic-payments.md +126 -126
  53. package/.claude/agents/sona/sona-learning-optimizer.md +74 -74
  54. package/.claude/agents/sparc/architecture.md +698 -698
  55. package/.claude/agents/sparc/pseudocode.md +519 -519
  56. package/.claude/agents/sparc/refinement.md +801 -801
  57. package/.claude/agents/sparc/specification.md +477 -477
  58. package/.claude/agents/specialized/mobile/spec-mobile-react-native.md +224 -224
  59. package/.claude/agents/specialized/spec-mobile-react-native.md +226 -226
  60. package/.claude/agents/sublinear/consensus-coordinator.md +337 -337
  61. package/.claude/agents/sublinear/matrix-optimizer.md +184 -184
  62. package/.claude/agents/sublinear/pagerank-analyzer.md +298 -298
  63. package/.claude/agents/sublinear/performance-optimizer.md +367 -367
  64. package/.claude/agents/sublinear/trading-predictor.md +245 -245
  65. package/.claude/agents/swarm/adaptive-coordinator.md +1126 -1126
  66. package/.claude/agents/swarm/hierarchical-coordinator.md +709 -709
  67. package/.claude/agents/swarm/mesh-coordinator.md +962 -962
  68. package/.claude/agents/templates/automation-smart-agent.md +204 -204
  69. package/.claude/agents/templates/base-template-generator.md +289 -289
  70. package/.claude/agents/templates/coordinator-swarm-init.md +89 -89
  71. package/.claude/agents/templates/github-pr-manager.md +176 -176
  72. package/.claude/agents/templates/implementer-sparc-coder.md +258 -258
  73. package/.claude/agents/templates/memory-coordinator.md +186 -186
  74. package/.claude/agents/templates/orchestrator-task.md +138 -138
  75. package/.claude/agents/templates/performance-analyzer.md +198 -198
  76. package/.claude/agents/templates/sparc-coordinator.md +513 -513
  77. package/.claude/agents/testing/production-validator.md +394 -394
  78. package/.claude/agents/testing/tdd-london-swarm.md +243 -243
  79. package/.claude/agents/v3/aidefence-guardian.md +282 -282
  80. package/.claude/agents/v3/claims-authorizer.md +208 -208
  81. package/.claude/agents/v3/collective-intelligence-coordinator.md +993 -993
  82. package/.claude/agents/v3/ddd-domain-expert.md +220 -220
  83. package/.claude/agents/v3/injection-analyst.md +236 -236
  84. package/.claude/agents/v3/performance-engineer.md +1233 -1233
  85. package/.claude/agents/v3/pii-detector.md +151 -151
  86. package/.claude/agents/v3/reasoningbank-learner.md +213 -213
  87. package/.claude/agents/v3/security-architect-aidefence.md +410 -410
  88. package/.claude/agents/v3/security-architect.md +867 -867
  89. package/.claude/agents/v3/swarm-memory-manager.md +157 -157
  90. package/.claude/agents/v3/v3-integration-architect.md +205 -205
  91. package/.claude/commands/agents/README.md +50 -50
  92. package/.claude/commands/agents/agent-capabilities.md +140 -140
  93. package/.claude/commands/agents/agent-coordination.md +28 -28
  94. package/.claude/commands/agents/agent-spawning.md +28 -28
  95. package/.claude/commands/agents/agent-types.md +216 -216
  96. package/.claude/commands/agents/health.md +139 -139
  97. package/.claude/commands/agents/list.md +100 -100
  98. package/.claude/commands/agents/logs.md +130 -130
  99. package/.claude/commands/agents/metrics.md +122 -122
  100. package/.claude/commands/agents/pool.md +127 -127
  101. package/.claude/commands/agents/spawn.md +140 -140
  102. package/.claude/commands/agents/status.md +115 -115
  103. package/.claude/commands/agents/stop.md +102 -102
  104. package/.claude/commands/analysis/COMMAND_COMPLIANCE_REPORT.md +53 -53
  105. package/.claude/commands/analysis/README.md +9 -9
  106. package/.claude/commands/analysis/bottleneck-detect.md +162 -162
  107. package/.claude/commands/analysis/performance-bottlenecks.md +58 -58
  108. package/.claude/commands/analysis/performance-report.md +25 -25
  109. package/.claude/commands/analysis/token-efficiency.md +44 -44
  110. package/.claude/commands/analysis/token-usage.md +25 -25
  111. package/.claude/commands/automation/README.md +9 -9
  112. package/.claude/commands/automation/auto-agent.md +122 -122
  113. package/.claude/commands/automation/self-healing.md +105 -105
  114. package/.claude/commands/automation/session-memory.md +89 -89
  115. package/.claude/commands/automation/smart-agents.md +72 -72
  116. package/.claude/commands/automation/smart-spawn.md +25 -25
  117. package/.claude/commands/automation/workflow-select.md +25 -25
  118. package/.claude/commands/claude-flow-help.md +103 -103
  119. package/.claude/commands/claude-flow-memory.md +107 -107
  120. package/.claude/commands/claude-flow-swarm.md +205 -205
  121. package/.claude/commands/coordination/README.md +9 -9
  122. package/.claude/commands/coordination/agent-spawn.md +25 -25
  123. package/.claude/commands/coordination/init.md +44 -44
  124. package/.claude/commands/coordination/orchestrate.md +43 -43
  125. package/.claude/commands/coordination/spawn.md +45 -45
  126. package/.claude/commands/coordination/swarm-init.md +85 -85
  127. package/.claude/commands/coordination/task-orchestrate.md +25 -25
  128. package/.claude/commands/github/README.md +11 -11
  129. package/.claude/commands/github/code-review-swarm.md +513 -513
  130. package/.claude/commands/github/code-review.md +25 -25
  131. package/.claude/commands/github/github-modes.md +146 -146
  132. package/.claude/commands/github/github-swarm.md +121 -121
  133. package/.claude/commands/github/issue-tracker.md +291 -291
  134. package/.claude/commands/github/issue-triage.md +25 -25
  135. package/.claude/commands/github/multi-repo-swarm.md +518 -518
  136. package/.claude/commands/github/pr-enhance.md +26 -26
  137. package/.claude/commands/github/pr-manager.md +169 -169
  138. package/.claude/commands/github/project-board-sync.md +470 -470
  139. package/.claude/commands/github/release-manager.md +339 -339
  140. package/.claude/commands/github/release-swarm.md +543 -543
  141. package/.claude/commands/github/repo-analyze.md +25 -25
  142. package/.claude/commands/github/repo-architect.md +366 -366
  143. package/.claude/commands/github/swarm-issue.md +484 -484
  144. package/.claude/commands/github/swarm-pr.md +287 -287
  145. package/.claude/commands/github/sync-coordinator.md +302 -302
  146. package/.claude/commands/github/workflow-automation.md +441 -441
  147. package/.claude/commands/hive-mind/README.md +17 -17
  148. package/.claude/commands/hive-mind/hive-mind-consensus.md +8 -8
  149. package/.claude/commands/hive-mind/hive-mind-init.md +18 -18
  150. package/.claude/commands/hive-mind/hive-mind-memory.md +8 -8
  151. package/.claude/commands/hive-mind/hive-mind-metrics.md +8 -8
  152. package/.claude/commands/hive-mind/hive-mind-resume.md +8 -8
  153. package/.claude/commands/hive-mind/hive-mind-sessions.md +8 -8
  154. package/.claude/commands/hive-mind/hive-mind-spawn.md +21 -21
  155. package/.claude/commands/hive-mind/hive-mind-status.md +8 -8
  156. package/.claude/commands/hive-mind/hive-mind-stop.md +8 -8
  157. package/.claude/commands/hive-mind/hive-mind-wizard.md +8 -8
  158. package/.claude/commands/hive-mind/hive-mind.md +27 -27
  159. package/.claude/commands/hooks/README.md +11 -11
  160. package/.claude/commands/hooks/overview.md +57 -57
  161. package/.claude/commands/hooks/post-edit.md +117 -117
  162. package/.claude/commands/hooks/post-task.md +112 -112
  163. package/.claude/commands/hooks/pre-edit.md +113 -113
  164. package/.claude/commands/hooks/pre-task.md +111 -111
  165. package/.claude/commands/hooks/session-end.md +118 -118
  166. package/.claude/commands/hooks/setup.md +102 -102
  167. package/.claude/commands/memory/README.md +9 -9
  168. package/.claude/commands/memory/memory-persist.md +25 -25
  169. package/.claude/commands/memory/memory-search.md +25 -25
  170. package/.claude/commands/memory/memory-usage.md +25 -25
  171. package/.claude/commands/memory/neural.md +47 -47
  172. package/.claude/commands/monitoring/README.md +9 -9
  173. package/.claude/commands/monitoring/agent-metrics.md +25 -25
  174. package/.claude/commands/monitoring/agents.md +44 -44
  175. package/.claude/commands/monitoring/real-time-view.md +25 -25
  176. package/.claude/commands/monitoring/status.md +46 -46
  177. package/.claude/commands/monitoring/swarm-monitor.md +25 -25
  178. package/.claude/commands/optimization/README.md +9 -9
  179. package/.claude/commands/optimization/auto-topology.md +61 -61
  180. package/.claude/commands/optimization/cache-manage.md +25 -25
  181. package/.claude/commands/optimization/parallel-execute.md +25 -25
  182. package/.claude/commands/optimization/parallel-execution.md +49 -49
  183. package/.claude/commands/optimization/topology-optimize.md +25 -25
  184. package/.claude/commands/pair/README.md +260 -260
  185. package/.claude/commands/pair/commands.md +545 -545
  186. package/.claude/commands/pair/config.md +509 -509
  187. package/.claude/commands/pair/examples.md +511 -511
  188. package/.claude/commands/pair/modes.md +347 -347
  189. package/.claude/commands/pair/session.md +406 -406
  190. package/.claude/commands/pair/start.md +208 -208
  191. package/.claude/commands/sparc/analyzer.md +51 -51
  192. package/.claude/commands/sparc/architect.md +53 -53
  193. package/.claude/commands/sparc/ask.md +97 -97
  194. package/.claude/commands/sparc/batch-executor.md +54 -54
  195. package/.claude/commands/sparc/code.md +89 -89
  196. package/.claude/commands/sparc/coder.md +54 -54
  197. package/.claude/commands/sparc/debug.md +83 -83
  198. package/.claude/commands/sparc/debugger.md +54 -54
  199. package/.claude/commands/sparc/designer.md +53 -53
  200. package/.claude/commands/sparc/devops.md +109 -109
  201. package/.claude/commands/sparc/docs-writer.md +80 -80
  202. package/.claude/commands/sparc/documenter.md +54 -54
  203. package/.claude/commands/sparc/innovator.md +54 -54
  204. package/.claude/commands/sparc/integration.md +83 -83
  205. package/.claude/commands/sparc/mcp.md +117 -117
  206. package/.claude/commands/sparc/memory-manager.md +54 -54
  207. package/.claude/commands/sparc/optimizer.md +54 -54
  208. package/.claude/commands/sparc/orchestrator.md +131 -131
  209. package/.claude/commands/sparc/post-deployment-monitoring-mode.md +83 -83
  210. package/.claude/commands/sparc/refinement-optimization-mode.md +83 -83
  211. package/.claude/commands/sparc/researcher.md +54 -54
  212. package/.claude/commands/sparc/reviewer.md +54 -54
  213. package/.claude/commands/sparc/security-review.md +80 -80
  214. package/.claude/commands/sparc/sparc-modes.md +174 -174
  215. package/.claude/commands/sparc/sparc.md +111 -111
  216. package/.claude/commands/sparc/spec-pseudocode.md +80 -80
  217. package/.claude/commands/sparc/supabase-admin.md +348 -348
  218. package/.claude/commands/sparc/swarm-coordinator.md +54 -54
  219. package/.claude/commands/sparc/tdd.md +54 -54
  220. package/.claude/commands/sparc/tester.md +54 -54
  221. package/.claude/commands/sparc/tutorial.md +79 -79
  222. package/.claude/commands/sparc/workflow-manager.md +54 -54
  223. package/.claude/commands/sparc.md +166 -166
  224. package/.claude/commands/stream-chain/pipeline.md +120 -120
  225. package/.claude/commands/stream-chain/run.md +69 -69
  226. package/.claude/commands/swarm/README.md +15 -15
  227. package/.claude/commands/swarm/analysis.md +95 -95
  228. package/.claude/commands/swarm/development.md +96 -96
  229. package/.claude/commands/swarm/examples.md +168 -168
  230. package/.claude/commands/swarm/maintenance.md +102 -102
  231. package/.claude/commands/swarm/optimization.md +117 -117
  232. package/.claude/commands/swarm/research.md +136 -136
  233. package/.claude/commands/swarm/swarm-analysis.md +8 -8
  234. package/.claude/commands/swarm/swarm-background.md +8 -8
  235. package/.claude/commands/swarm/swarm-init.md +19 -19
  236. package/.claude/commands/swarm/swarm-modes.md +8 -8
  237. package/.claude/commands/swarm/swarm-monitor.md +8 -8
  238. package/.claude/commands/swarm/swarm-spawn.md +19 -19
  239. package/.claude/commands/swarm/swarm-status.md +8 -8
  240. package/.claude/commands/swarm/swarm-strategies.md +8 -8
  241. package/.claude/commands/swarm/swarm.md +87 -87
  242. package/.claude/commands/swarm/testing.md +131 -131
  243. package/.claude/commands/training/README.md +9 -9
  244. package/.claude/commands/training/model-update.md +25 -25
  245. package/.claude/commands/training/neural-patterns.md +107 -107
  246. package/.claude/commands/training/neural-train.md +75 -75
  247. package/.claude/commands/training/pattern-learn.md +25 -25
  248. package/.claude/commands/training/specialization.md +62 -62
  249. package/.claude/commands/truth/start.md +142 -142
  250. package/.claude/commands/verify/check.md +49 -49
  251. package/.claude/commands/verify/start.md +127 -127
  252. package/.claude/commands/workflows/README.md +9 -9
  253. package/.claude/commands/workflows/development.md +77 -77
  254. package/.claude/commands/workflows/research.md +62 -62
  255. package/.claude/commands/workflows/workflow-create.md +25 -25
  256. package/.claude/commands/workflows/workflow-execute.md +25 -25
  257. package/.claude/commands/workflows/workflow-export.md +25 -25
  258. package/.claude/eval/human-relevance-frozen-v1.json +17 -17
  259. package/.claude/evolve-proof/generation-0.json +211 -211
  260. package/.claude/evolve-proof/real-generation-0.json +406 -406
  261. package/.claude/evolve-proof/real-generation-1.json +406 -406
  262. package/.claude/helpers/.helpers-version +1 -1
  263. package/.claude/helpers/README.md +96 -96
  264. package/.claude/helpers/adr-compliance.sh +186 -186
  265. package/.claude/helpers/auto-commit.sh +178 -178
  266. package/.claude/helpers/auto-memory-hook.mjs +430 -430
  267. package/.claude/helpers/checkpoint-manager.sh +251 -251
  268. package/.claude/helpers/daemon-manager.sh +252 -252
  269. package/.claude/helpers/ddd-tracker.sh +144 -144
  270. package/.claude/helpers/github-safe.js +156 -156
  271. package/.claude/helpers/github-setup.sh +45 -45
  272. package/.claude/helpers/guidance-hook.sh +13 -13
  273. package/.claude/helpers/guidance-hooks.sh +102 -102
  274. package/.claude/helpers/health-monitor.sh +108 -108
  275. package/.claude/helpers/helpers.manifest.json +6 -6
  276. package/.claude/helpers/hook-handler.cjs +564 -460
  277. package/.claude/helpers/intelligence.cjs +1058 -1058
  278. package/.claude/helpers/learning-hooks.sh +329 -329
  279. package/.claude/helpers/learning-optimizer.sh +127 -127
  280. package/.claude/helpers/learning-service.mjs +1144 -1144
  281. package/.claude/helpers/memory.js +83 -83
  282. package/.claude/helpers/metrics-db.mjs +503 -503
  283. package/.claude/helpers/pattern-consolidator.sh +86 -86
  284. package/.claude/helpers/perf-worker.sh +160 -160
  285. package/.claude/helpers/post-commit +16 -16
  286. package/.claude/helpers/pre-commit +26 -26
  287. package/.claude/helpers/quick-start.sh +19 -19
  288. package/.claude/helpers/router.js +105 -105
  289. package/.claude/helpers/security-scanner.sh +127 -127
  290. package/.claude/helpers/session.js +157 -157
  291. package/.claude/helpers/setup-mcp.sh +18 -18
  292. package/.claude/helpers/standard-checkpoint-hooks.sh +189 -189
  293. package/.claude/helpers/statusline-hook.sh +21 -21
  294. package/.claude/helpers/statusline.cjs +925 -925
  295. package/.claude/helpers/statusline.js +352 -352
  296. package/.claude/helpers/swarm-comms.sh +353 -353
  297. package/.claude/helpers/swarm-hooks.sh +761 -761
  298. package/.claude/helpers/swarm-monitor.sh +210 -210
  299. package/.claude/helpers/sync-v3-metrics.sh +245 -245
  300. package/.claude/helpers/update-v3-progress.sh +165 -165
  301. package/.claude/helpers/v3-quick-status.sh +57 -57
  302. package/.claude/helpers/v3.sh +110 -110
  303. package/.claude/helpers/validate-v3-config.sh +215 -215
  304. package/.claude/helpers/worker-manager.sh +170 -170
  305. package/.claude/proven-config.json +1 -1
  306. package/.claude/proven-config.manifest.json +37 -37
  307. package/.claude/proven-config.signed.json +41 -41
  308. package/.claude/settings.json +182 -182
  309. package/.claude/skills/agentdb-advanced/SKILL.md +550 -550
  310. package/.claude/skills/agentdb-learning/SKILL.md +545 -545
  311. package/.claude/skills/agentdb-memory-patterns/SKILL.md +339 -339
  312. package/.claude/skills/agentdb-optimization/SKILL.md +509 -509
  313. package/.claude/skills/agentdb-vector-search/SKILL.md +339 -339
  314. package/.claude/skills/browser/SKILL.md +204 -204
  315. package/.claude/skills/dual-mode/README.md +71 -71
  316. package/.claude/skills/dual-mode/dual-collect.md +103 -103
  317. package/.claude/skills/dual-mode/dual-coordinate.md +85 -85
  318. package/.claude/skills/dual-mode/dual-spawn.md +81 -81
  319. package/.claude/skills/flow-nexus-neural/SKILL.md +727 -727
  320. package/.claude/skills/flow-nexus-platform/SKILL.md +1154 -1154
  321. package/.claude/skills/flow-nexus-swarm/SKILL.md +604 -604
  322. package/.claude/skills/github-code-review/SKILL.md +1125 -1125
  323. package/.claude/skills/github-multi-repo/SKILL.md +862 -862
  324. package/.claude/skills/github-project-management/SKILL.md +1262 -1262
  325. package/.claude/skills/github-release-management/SKILL.md +1064 -1064
  326. package/.claude/skills/github-workflow-automation/SKILL.md +1047 -1047
  327. package/.claude/skills/hooks-automation/SKILL.md +1201 -1201
  328. package/.claude/skills/pair-programming/SKILL.md +1202 -1202
  329. package/.claude/skills/reasoningbank-agentdb/SKILL.md +446 -446
  330. package/.claude/skills/reasoningbank-intelligence/SKILL.md +201 -201
  331. package/.claude/skills/skill-builder/SKILL.md +910 -910
  332. package/.claude/skills/sparc-methodology/SKILL.md +1106 -1106
  333. package/.claude/skills/stream-chain/SKILL.md +560 -560
  334. package/.claude/skills/swarm-advanced/SKILL.md +970 -970
  335. package/.claude/skills/swarm-orchestration/SKILL.md +179 -179
  336. package/.claude/skills/v3-cli-modernization/SKILL.md +871 -871
  337. package/.claude/skills/v3-core-implementation/SKILL.md +796 -796
  338. package/.claude/skills/v3-ddd-architecture/SKILL.md +441 -441
  339. package/.claude/skills/v3-integration-deep/SKILL.md +240 -240
  340. package/.claude/skills/v3-mcp-optimization/SKILL.md +776 -776
  341. package/.claude/skills/v3-memory-unification/SKILL.md +173 -173
  342. package/.claude/skills/v3-performance-optimization/SKILL.md +389 -389
  343. package/.claude/skills/v3-security-overhaul/SKILL.md +81 -81
  344. package/.claude/skills/v3-swarm-coordination/SKILL.md +339 -339
  345. package/.claude/skills/verification-quality/SKILL.md +691 -691
  346. package/README.md +419 -419
  347. package/bin/cli.js +314 -314
  348. package/bin/mcp-server.js +224 -224
  349. package/bin/preinstall.cjs +2 -2
  350. package/catalog-manifest.json +2 -2
  351. package/dist/src/benchmarks/gaia-critic.js +24 -24
  352. package/dist/src/business-pods/bbs-budget-tracker.js +53 -53
  353. package/dist/src/commands/announcements.d.ts +17 -0
  354. package/dist/src/commands/announcements.js +0 -0
  355. package/dist/src/commands/completions.js +409 -409
  356. package/dist/src/commands/daemon.js +44 -44
  357. package/dist/src/commands/embeddings.js +26 -26
  358. package/dist/src/commands/funnel.d.ts +10 -5
  359. package/dist/src/commands/funnel.js +206 -7
  360. package/dist/src/commands/hive-mind.js +97 -97
  361. package/dist/src/commands/hooks.js +9 -9
  362. package/dist/src/commands/index.js +4 -0
  363. package/dist/src/commands/init.js +94 -38
  364. package/dist/src/commands/memory.js +85 -1
  365. package/dist/src/commands/ruvector/backup.js +23 -23
  366. package/dist/src/commands/ruvector/benchmark.js +31 -31
  367. package/dist/src/commands/ruvector/import.js +14 -14
  368. package/dist/src/commands/ruvector/init.js +115 -115
  369. package/dist/src/commands/ruvector/migrate.js +99 -99
  370. package/dist/src/commands/ruvector/optimize.js +51 -51
  371. package/dist/src/commands/ruvector/setup.js +624 -624
  372. package/dist/src/commands/ruvector/status.js +38 -38
  373. package/dist/src/commands/spinner.d.ts +16 -0
  374. package/dist/src/commands/spinner.js +329 -0
  375. package/dist/src/config/proven-config.js +2 -2
  376. package/dist/src/funnel/consent.js +3 -0
  377. package/dist/src/funnel/disclosure.d.ts +1 -0
  378. package/dist/src/funnel/disclosure.js +12 -0
  379. package/dist/src/funnel/index.d.ts +2 -1
  380. package/dist/src/funnel/index.js +2 -1
  381. package/dist/src/funnel/payout.d.ts +40 -0
  382. package/dist/src/funnel/payout.js +60 -0
  383. package/dist/src/funnel/types.d.ts +13 -1
  384. package/dist/src/init/claudemd-generator.js +231 -231
  385. package/dist/src/init/executor.js +453 -453
  386. package/dist/src/init/helper-refresh.js +17 -0
  387. package/dist/src/init/helper-signing.d.ts +8 -1
  388. package/dist/src/init/helper-signing.js +9 -2
  389. package/dist/src/init/helpers-generator.js +751 -751
  390. package/dist/src/init/statusline-generator.js +955 -949
  391. package/dist/src/mcp-tools/agentdb-tools.js +15 -15
  392. package/dist/src/mcp-tools/browser-intent-tools.js +19 -19
  393. package/dist/src/memory/graph-edge-writer.js +22 -22
  394. package/dist/src/memory/memory-bridge.d.ts +10 -0
  395. package/dist/src/memory/memory-bridge.js +139 -88
  396. package/dist/src/memory/memory-initializer.d.ts +18 -0
  397. package/dist/src/memory/memory-initializer.js +539 -407
  398. package/dist/src/memory/rabitq-index.js +5 -5
  399. package/dist/src/runtime/headless.js +28 -28
  400. package/dist/src/ruvector/flash-attention.d.ts +195 -0
  401. package/dist/src/ruvector/flash-attention.js +643 -0
  402. package/dist/src/ruvector/moe-router.d.ts +206 -0
  403. package/dist/src/ruvector/moe-router.js +626 -0
  404. package/dist/src/services/distill-tuning.js +7 -7
  405. package/dist/src/services/event-stream.d.ts +25 -0
  406. package/dist/src/services/event-stream.js +27 -0
  407. package/dist/src/services/headless-worker-executor.js +84 -84
  408. package/dist/src/services/loop-worker-runner.d.ts +16 -0
  409. package/dist/src/services/loop-worker-runner.js +34 -0
  410. package/dist/src/services/memory-distillation.js +4 -4
  411. package/dist/src/services/runtime-capabilities.d.ts +22 -0
  412. package/dist/src/services/runtime-capabilities.js +45 -0
  413. package/dist/src/transfer/deploy-seraphine.js +23 -23
  414. package/package.json +134 -133
  415. package/plugins/ruflo-metaharness/.claude-plugin/plugin.json +32 -32
  416. package/plugins/ruflo-metaharness/README.md +72 -72
  417. package/plugins/ruflo-metaharness/agents/metaharness-architect.md +58 -58
  418. package/plugins/ruflo-metaharness/commands/ruflo-metaharness.md +48 -48
  419. package/plugins/ruflo-metaharness/scripts/_darwin.mjs +210 -210
  420. package/plugins/ruflo-metaharness/scripts/_harness.mjs +330 -330
  421. package/plugins/ruflo-metaharness/scripts/_invoke.mjs +231 -231
  422. package/plugins/ruflo-metaharness/scripts/_redblue.mjs +143 -143
  423. package/plugins/ruflo-metaharness/scripts/_similarity.mjs +161 -161
  424. package/plugins/ruflo-metaharness/scripts/_spike-similarity.mjs +223 -223
  425. package/plugins/ruflo-metaharness/scripts/audit-list.mjs +158 -158
  426. package/plugins/ruflo-metaharness/scripts/audit-trend.mjs +272 -272
  427. package/plugins/ruflo-metaharness/scripts/bench-parse-mcp-scan.mjs +146 -146
  428. package/plugins/ruflo-metaharness/scripts/bench-recordpair-overhead.mjs +186 -186
  429. package/plugins/ruflo-metaharness/scripts/bench-similarity.mjs +177 -177
  430. package/plugins/ruflo-metaharness/scripts/bench.mjs +95 -95
  431. package/plugins/ruflo-metaharness/scripts/drift-from-history.mjs +363 -363
  432. package/plugins/ruflo-metaharness/scripts/evolve.mjs +404 -404
  433. package/plugins/ruflo-metaharness/scripts/genome.mjs +80 -80
  434. package/plugins/ruflo-metaharness/scripts/gepa.mjs +153 -153
  435. package/plugins/ruflo-metaharness/scripts/learn.mjs +127 -127
  436. package/plugins/ruflo-metaharness/scripts/mcp-scan.mjs +111 -111
  437. package/plugins/ruflo-metaharness/scripts/mint.mjs +126 -126
  438. package/plugins/ruflo-metaharness/scripts/oia-audit.mjs +228 -228
  439. package/plugins/ruflo-metaharness/scripts/redblue.mjs +286 -286
  440. package/plugins/ruflo-metaharness/scripts/router-parallel-analyze.mjs +250 -250
  441. package/plugins/ruflo-metaharness/scripts/score.mjs +92 -92
  442. package/plugins/ruflo-metaharness/scripts/security-bench.mjs +174 -174
  443. package/plugins/ruflo-metaharness/scripts/similarity.mjs +158 -158
  444. package/plugins/ruflo-metaharness/scripts/smoke.sh +2353 -2353
  445. package/plugins/ruflo-metaharness/scripts/test-graceful-degradation.mjs +165 -165
  446. package/plugins/ruflo-metaharness/scripts/test-mcp-tools.mjs +472 -472
  447. package/plugins/ruflo-metaharness/scripts/test-parallel-pipeline.mjs +204 -204
  448. package/plugins/ruflo-metaharness/scripts/test-pipeline-roundtrip.mjs +586 -586
  449. package/plugins/ruflo-metaharness/scripts/test-similarity.mjs +334 -334
  450. package/plugins/ruflo-metaharness/scripts/test-with-openrouter.mjs +229 -229
  451. package/plugins/ruflo-metaharness/scripts/threat-model.mjs +59 -59
  452. package/plugins/ruflo-metaharness/skills/harness-bench/SKILL.md +64 -64
  453. package/plugins/ruflo-metaharness/skills/harness-drift-from-history/SKILL.md +65 -65
  454. package/plugins/ruflo-metaharness/skills/harness-evolve/SKILL.md +131 -131
  455. package/plugins/ruflo-metaharness/skills/harness-genome/SKILL.md +54 -54
  456. package/plugins/ruflo-metaharness/skills/harness-gepa/SKILL.md +65 -65
  457. package/plugins/ruflo-metaharness/skills/harness-learn/SKILL.md +65 -65
  458. package/plugins/ruflo-metaharness/skills/harness-mcp-scan/SKILL.md +49 -49
  459. package/plugins/ruflo-metaharness/skills/harness-mint/SKILL.md +72 -72
  460. package/plugins/ruflo-metaharness/skills/harness-oia-audit/SKILL.md +79 -79
  461. package/plugins/ruflo-metaharness/skills/harness-score/SKILL.md +66 -66
  462. package/plugins/ruflo-metaharness/skills/harness-security-bench/SKILL.md +101 -101
  463. package/plugins/ruflo-metaharness/skills/harness-similarity/SKILL.md +67 -67
  464. package/plugins/ruflo-metaharness/skills/harness-threat-model/SKILL.md +41 -41
  465. package/scripts/postinstall.cjs +153 -153
@@ -1,665 +1,665 @@
1
- ---
2
- name: Benchmark Suite
3
- type: agent
4
- category: optimization
5
- description: Comprehensive performance benchmarking, regression detection and performance validation
6
- ---
7
-
8
- # Benchmark Suite Agent
9
-
10
- ## Agent Profile
11
- - **Name**: Benchmark Suite
12
- - **Type**: Performance Optimization Agent
13
- - **Specialization**: Comprehensive performance benchmarking and testing
14
- - **Performance Focus**: Automated benchmarking, regression detection, and performance validation
15
-
16
- ## Core Capabilities
17
-
18
- ### 1. Comprehensive Benchmarking Framework
19
- ```javascript
20
- // Advanced benchmarking system
21
- class ComprehensiveBenchmarkSuite {
22
- constructor() {
23
- this.benchmarks = {
24
- // Core performance benchmarks
25
- throughput: new ThroughputBenchmark(),
26
- latency: new LatencyBenchmark(),
27
- scalability: new ScalabilityBenchmark(),
28
- resource_usage: new ResourceUsageBenchmark(),
29
-
30
- // Swarm-specific benchmarks
31
- coordination: new CoordinationBenchmark(),
32
- load_balancing: new LoadBalancingBenchmark(),
33
- topology: new TopologyBenchmark(),
34
- fault_tolerance: new FaultToleranceBenchmark(),
35
-
36
- // Custom benchmarks
37
- custom: new CustomBenchmarkManager()
38
- };
39
-
40
- this.reporter = new BenchmarkReporter();
41
- this.comparator = new PerformanceComparator();
42
- this.analyzer = new BenchmarkAnalyzer();
43
- }
44
-
45
- // Execute comprehensive benchmark suite
46
- async runBenchmarkSuite(config = {}) {
47
- const suiteConfig = {
48
- duration: config.duration || 300000, // 5 minutes default
49
- iterations: config.iterations || 10,
50
- warmupTime: config.warmupTime || 30000, // 30 seconds
51
- cooldownTime: config.cooldownTime || 10000, // 10 seconds
52
- parallel: config.parallel || false,
53
- baseline: config.baseline || null
54
- };
55
-
56
- const results = {
57
- summary: {},
58
- detailed: new Map(),
59
- baseline_comparison: null,
60
- recommendations: []
61
- };
62
-
63
- // Warmup phase
64
- await this.warmup(suiteConfig.warmupTime);
65
-
66
- // Execute benchmarks
67
- if (suiteConfig.parallel) {
68
- results.detailed = await this.runBenchmarksParallel(suiteConfig);
69
- } else {
70
- results.detailed = await this.runBenchmarksSequential(suiteConfig);
71
- }
72
-
73
- // Generate summary
74
- results.summary = this.generateSummary(results.detailed);
75
-
76
- // Compare with baseline if provided
77
- if (suiteConfig.baseline) {
78
- results.baseline_comparison = await this.compareWithBaseline(
79
- results.detailed,
80
- suiteConfig.baseline
81
- );
82
- }
83
-
84
- // Generate recommendations
85
- results.recommendations = await this.generateRecommendations(results);
86
-
87
- // Cooldown phase
88
- await this.cooldown(suiteConfig.cooldownTime);
89
-
90
- return results;
91
- }
92
-
93
- // Parallel benchmark execution
94
- async runBenchmarksParallel(config) {
95
- const benchmarkPromises = Object.entries(this.benchmarks).map(
96
- async ([name, benchmark]) => {
97
- const result = await this.executeBenchmark(benchmark, name, config);
98
- return [name, result];
99
- }
100
- );
101
-
102
- const results = await Promise.all(benchmarkPromises);
103
- return new Map(results);
104
- }
105
-
106
- // Sequential benchmark execution
107
- async runBenchmarksSequential(config) {
108
- const results = new Map();
109
-
110
- for (const [name, benchmark] of Object.entries(this.benchmarks)) {
111
- const result = await this.executeBenchmark(benchmark, name, config);
112
- results.set(name, result);
113
-
114
- // Brief pause between benchmarks
115
- await this.sleep(1000);
116
- }
117
-
118
- return results;
119
- }
120
- }
121
- ```
122
-
123
- ### 2. Performance Regression Detection
124
- ```javascript
125
- // Advanced regression detection system
126
- class RegressionDetector {
127
- constructor() {
128
- this.detectors = {
129
- statistical: new StatisticalRegressionDetector(),
130
- machine_learning: new MLRegressionDetector(),
131
- threshold: new ThresholdRegressionDetector(),
132
- trend: new TrendRegressionDetector()
133
- };
134
-
135
- this.analyzer = new RegressionAnalyzer();
136
- this.alerting = new RegressionAlerting();
137
- }
138
-
139
- // Detect performance regressions
140
- async detectRegressions(currentResults, historicalData, config = {}) {
141
- const regressions = {
142
- detected: [],
143
- severity: 'none',
144
- confidence: 0,
145
- analysis: {}
146
- };
147
-
148
- // Run multiple detection algorithms
149
- const detectionPromises = Object.entries(this.detectors).map(
150
- async ([method, detector]) => {
151
- const detection = await detector.detect(currentResults, historicalData, config);
152
- return [method, detection];
153
- }
154
- );
155
-
156
- const detectionResults = await Promise.all(detectionPromises);
157
-
158
- // Aggregate detection results
159
- for (const [method, detection] of detectionResults) {
160
- if (detection.regression_detected) {
161
- regressions.detected.push({
162
- method,
163
- ...detection
164
- });
165
- }
166
- }
167
-
168
- // Calculate overall confidence and severity
169
- if (regressions.detected.length > 0) {
170
- regressions.confidence = this.calculateAggregateConfidence(regressions.detected);
171
- regressions.severity = this.calculateSeverity(regressions.detected);
172
- regressions.analysis = await this.analyzer.analyze(regressions.detected);
173
- }
174
-
175
- return regressions;
176
- }
177
-
178
- // Statistical regression detection using change point analysis
179
- async detectStatisticalRegression(metric, historicalData, sensitivity = 0.95) {
180
- // Use CUSUM (Cumulative Sum) algorithm for change point detection
181
- const cusum = this.calculateCUSUM(metric, historicalData);
182
-
183
- // Detect change points
184
- const changePoints = this.detectChangePoints(cusum, sensitivity);
185
-
186
- // Analyze significance of changes
187
- const analysis = changePoints.map(point => ({
188
- timestamp: point.timestamp,
189
- magnitude: point.magnitude,
190
- direction: point.direction,
191
- significance: point.significance,
192
- confidence: point.confidence
193
- }));
194
-
195
- return {
196
- regression_detected: changePoints.length > 0,
197
- change_points: analysis,
198
- cusum_statistics: cusum.statistics,
199
- sensitivity: sensitivity
200
- };
201
- }
202
-
203
- // Machine learning-based regression detection
204
- async detectMLRegression(metrics, historicalData) {
205
- // Train anomaly detection model on historical data
206
- const model = await this.trainAnomalyModel(historicalData);
207
-
208
- // Predict anomaly scores for current metrics
209
- const anomalyScores = await model.predict(metrics);
210
-
211
- // Identify regressions based on anomaly scores
212
- const threshold = this.calculateDynamicThreshold(anomalyScores);
213
- const regressions = anomalyScores.filter(score => score.anomaly > threshold);
214
-
215
- return {
216
- regression_detected: regressions.length > 0,
217
- anomaly_scores: anomalyScores,
218
- threshold: threshold,
219
- regressions: regressions,
220
- model_confidence: model.confidence
221
- };
222
- }
223
- }
224
- ```
225
-
226
- ### 3. Automated Performance Testing
227
- ```javascript
228
- // Comprehensive automated performance testing
229
- class AutomatedPerformanceTester {
230
- constructor() {
231
- this.testSuites = {
232
- load: new LoadTestSuite(),
233
- stress: new StressTestSuite(),
234
- volume: new VolumeTestSuite(),
235
- endurance: new EnduranceTestSuite(),
236
- spike: new SpikeTestSuite(),
237
- configuration: new ConfigurationTestSuite()
238
- };
239
-
240
- this.scheduler = new TestScheduler();
241
- this.orchestrator = new TestOrchestrator();
242
- this.validator = new ResultValidator();
243
- }
244
-
245
- // Execute automated performance test campaign
246
- async runTestCampaign(config) {
247
- const campaign = {
248
- id: this.generateCampaignId(),
249
- config,
250
- startTime: Date.now(),
251
- tests: [],
252
- results: new Map(),
253
- summary: null
254
- };
255
-
256
- // Schedule test execution
257
- const schedule = await this.scheduler.schedule(config.tests, config.constraints);
258
-
259
- // Execute tests according to schedule
260
- for (const scheduledTest of schedule) {
261
- const testResult = await this.executeScheduledTest(scheduledTest);
262
- campaign.tests.push(scheduledTest);
263
- campaign.results.set(scheduledTest.id, testResult);
264
-
265
- // Validate results in real-time
266
- const validation = await this.validator.validate(testResult);
267
- if (!validation.valid) {
268
- campaign.summary = {
269
- status: 'failed',
270
- reason: validation.reason,
271
- failedAt: scheduledTest.name
272
- };
273
- break;
274
- }
275
- }
276
-
277
- // Generate campaign summary
278
- if (!campaign.summary) {
279
- campaign.summary = await this.generateCampaignSummary(campaign);
280
- }
281
-
282
- campaign.endTime = Date.now();
283
- campaign.duration = campaign.endTime - campaign.startTime;
284
-
285
- return campaign;
286
- }
287
-
288
- // Load testing with gradual ramp-up
289
- async executeLoadTest(config) {
290
- const loadTest = {
291
- type: 'load',
292
- config,
293
- phases: [],
294
- metrics: new Map(),
295
- results: {}
296
- };
297
-
298
- // Ramp-up phase
299
- const rampUpResult = await this.executeRampUp(config.rampUp);
300
- loadTest.phases.push({ phase: 'ramp-up', result: rampUpResult });
301
-
302
- // Sustained load phase
303
- const sustainedResult = await this.executeSustainedLoad(config.sustained);
304
- loadTest.phases.push({ phase: 'sustained', result: sustainedResult });
305
-
306
- // Ramp-down phase
307
- const rampDownResult = await this.executeRampDown(config.rampDown);
308
- loadTest.phases.push({ phase: 'ramp-down', result: rampDownResult });
309
-
310
- // Analyze results
311
- loadTest.results = await this.analyzeLoadTestResults(loadTest.phases);
312
-
313
- return loadTest;
314
- }
315
-
316
- // Stress testing to find breaking points
317
- async executeStressTest(config) {
318
- const stressTest = {
319
- type: 'stress',
320
- config,
321
- breakingPoint: null,
322
- degradationCurve: [],
323
- results: {}
324
- };
325
-
326
- let currentLoad = config.startLoad;
327
- let systemBroken = false;
328
-
329
- while (!systemBroken && currentLoad <= config.maxLoad) {
330
- const testResult = await this.applyLoad(currentLoad, config.duration);
331
-
332
- stressTest.degradationCurve.push({
333
- load: currentLoad,
334
- performance: testResult.performance,
335
- stability: testResult.stability,
336
- errors: testResult.errors
337
- });
338
-
339
- // Check if system is breaking
340
- if (this.isSystemBreaking(testResult, config.breakingCriteria)) {
341
- stressTest.breakingPoint = {
342
- load: currentLoad,
343
- performance: testResult.performance,
344
- reason: this.identifyBreakingReason(testResult)
345
- };
346
- systemBroken = true;
347
- }
348
-
349
- currentLoad += config.loadIncrement;
350
- }
351
-
352
- stressTest.results = await this.analyzeStressTestResults(stressTest);
353
-
354
- return stressTest;
355
- }
356
- }
357
- ```
358
-
359
- ### 4. Performance Validation Framework
360
- ```javascript
361
- // Comprehensive performance validation
362
- class PerformanceValidator {
363
- constructor() {
364
- this.validators = {
365
- sla: new SLAValidator(),
366
- regression: new RegressionValidator(),
367
- scalability: new ScalabilityValidator(),
368
- reliability: new ReliabilityValidator(),
369
- efficiency: new EfficiencyValidator()
370
- };
371
-
372
- this.thresholds = new ThresholdManager();
373
- this.rules = new ValidationRuleEngine();
374
- }
375
-
376
- // Validate performance against defined criteria
377
- async validatePerformance(results, criteria) {
378
- const validation = {
379
- overall: {
380
- passed: true,
381
- score: 0,
382
- violations: []
383
- },
384
- detailed: new Map(),
385
- recommendations: []
386
- };
387
-
388
- // Run all validators
389
- const validationPromises = Object.entries(this.validators).map(
390
- async ([type, validator]) => {
391
- const result = await validator.validate(results, criteria[type]);
392
- return [type, result];
393
- }
394
- );
395
-
396
- const validationResults = await Promise.all(validationPromises);
397
-
398
- // Aggregate validation results
399
- for (const [type, result] of validationResults) {
400
- validation.detailed.set(type, result);
401
-
402
- if (!result.passed) {
403
- validation.overall.passed = false;
404
- validation.overall.violations.push(...result.violations);
405
- }
406
-
407
- validation.overall.score += result.score * (criteria[type]?.weight || 1);
408
- }
409
-
410
- // Normalize overall score
411
- const totalWeight = Object.values(criteria).reduce((sum, c) => sum + (c.weight || 1), 0);
412
- validation.overall.score /= totalWeight;
413
-
414
- // Generate recommendations
415
- validation.recommendations = await this.generateValidationRecommendations(validation);
416
-
417
- return validation;
418
- }
419
-
420
- // SLA validation
421
- async validateSLA(results, slaConfig) {
422
- const slaValidation = {
423
- passed: true,
424
- violations: [],
425
- score: 1.0,
426
- metrics: {}
427
- };
428
-
429
- // Validate each SLA metric
430
- for (const [metric, threshold] of Object.entries(slaConfig.thresholds)) {
431
- const actualValue = this.extractMetricValue(results, metric);
432
- const validation = this.validateThreshold(actualValue, threshold);
433
-
434
- slaValidation.metrics[metric] = {
435
- actual: actualValue,
436
- threshold: threshold.value,
437
- operator: threshold.operator,
438
- passed: validation.passed,
439
- deviation: validation.deviation
440
- };
441
-
442
- if (!validation.passed) {
443
- slaValidation.passed = false;
444
- slaValidation.violations.push({
445
- metric,
446
- actual: actualValue,
447
- expected: threshold.value,
448
- severity: threshold.severity || 'medium'
449
- });
450
-
451
- // Reduce score based on violation severity
452
- const severityMultiplier = this.getSeverityMultiplier(threshold.severity);
453
- slaValidation.score -= (validation.deviation * severityMultiplier);
454
- }
455
- }
456
-
457
- slaValidation.score = Math.max(0, slaValidation.score);
458
-
459
- return slaValidation;
460
- }
461
-
462
- // Scalability validation
463
- async validateScalability(results, scalabilityConfig) {
464
- const scalabilityValidation = {
465
- passed: true,
466
- violations: [],
467
- score: 1.0,
468
- analysis: {}
469
- };
470
-
471
- // Linear scalability analysis
472
- if (scalabilityConfig.linear) {
473
- const linearityAnalysis = this.analyzeLinearScalability(results);
474
- scalabilityValidation.analysis.linearity = linearityAnalysis;
475
-
476
- if (linearityAnalysis.coefficient < scalabilityConfig.linear.minCoefficient) {
477
- scalabilityValidation.passed = false;
478
- scalabilityValidation.violations.push({
479
- type: 'linearity',
480
- actual: linearityAnalysis.coefficient,
481
- expected: scalabilityConfig.linear.minCoefficient
482
- });
483
- }
484
- }
485
-
486
- // Efficiency retention analysis
487
- if (scalabilityConfig.efficiency) {
488
- const efficiencyAnalysis = this.analyzeEfficiencyRetention(results);
489
- scalabilityValidation.analysis.efficiency = efficiencyAnalysis;
490
-
491
- if (efficiencyAnalysis.retention < scalabilityConfig.efficiency.minRetention) {
492
- scalabilityValidation.passed = false;
493
- scalabilityValidation.violations.push({
494
- type: 'efficiency_retention',
495
- actual: efficiencyAnalysis.retention,
496
- expected: scalabilityConfig.efficiency.minRetention
497
- });
498
- }
499
- }
500
-
501
- return scalabilityValidation;
502
- }
503
- }
504
- ```
505
-
506
- ## MCP Integration Hooks
507
-
508
- ### Benchmark Execution Integration
509
- ```javascript
510
- // Comprehensive MCP benchmark integration
511
- const benchmarkIntegration = {
512
- // Execute performance benchmarks
513
- async runBenchmarks(config = {}) {
514
- // Run benchmark suite
515
- const benchmarkResult = await mcp.benchmark_run({
516
- suite: config.suite || 'comprehensive'
517
- });
518
-
519
- // Collect detailed metrics during benchmarking
520
- const metrics = await mcp.metrics_collect({
521
- components: ['system', 'agents', 'coordination', 'memory']
522
- });
523
-
524
- // Analyze performance trends
525
- const trends = await mcp.trend_analysis({
526
- metric: 'performance',
527
- period: '24h'
528
- });
529
-
530
- // Cost analysis
531
- const costAnalysis = await mcp.cost_analysis({
532
- timeframe: '24h'
533
- });
534
-
535
- return {
536
- benchmark: benchmarkResult,
537
- metrics,
538
- trends,
539
- costAnalysis,
540
- timestamp: Date.now()
541
- };
542
- },
543
-
544
- // Quality assessment
545
- async assessQuality(criteria) {
546
- const qualityAssessment = await mcp.quality_assess({
547
- target: 'swarm-performance',
548
- criteria: criteria || [
549
- 'throughput',
550
- 'latency',
551
- 'reliability',
552
- 'scalability',
553
- 'efficiency'
554
- ]
555
- });
556
-
557
- return qualityAssessment;
558
- },
559
-
560
- // Error pattern analysis
561
- async analyzeErrorPatterns() {
562
- // Collect system logs
563
- const logs = await this.collectSystemLogs();
564
-
565
- // Analyze error patterns
566
- const errorAnalysis = await mcp.error_analysis({
567
- logs: logs
568
- });
569
-
570
- return errorAnalysis;
571
- }
572
- };
573
- ```
574
-
575
- ## Operational Commands
576
-
577
- ### Benchmarking Commands
578
- ```bash
579
- # Run comprehensive benchmark suite
580
- npx claude-flow benchmark-run --suite comprehensive --duration 300
581
-
582
- # Execute specific benchmark
583
- npx claude-flow benchmark-run --suite throughput --iterations 10
584
-
585
- # Compare with baseline
586
- npx claude-flow benchmark-compare --current <results> --baseline <baseline>
587
-
588
- # Quality assessment
589
- npx claude-flow quality-assess --target swarm-performance --criteria throughput,latency
590
-
591
- # Performance validation
592
- npx claude-flow validate-performance --results <file> --criteria <file>
593
- ```
594
-
595
- ### Regression Detection Commands
596
- ```bash
597
- # Detect performance regressions
598
- npx claude-flow detect-regression --current <results> --historical <data>
599
-
600
- # Set up automated regression monitoring
601
- npx claude-flow regression-monitor --enable --sensitivity 0.95
602
-
603
- # Analyze error patterns
604
- npx claude-flow error-analysis --logs <log-files>
605
- ```
606
-
607
- ## Integration Points
608
-
609
- ### With Other Optimization Agents
610
- - **Performance Monitor**: Provides continuous monitoring data for benchmarking
611
- - **Load Balancer**: Validates load balancing effectiveness through benchmarks
612
- - **Topology Optimizer**: Tests topology configurations for optimal performance
613
-
614
- ### With CI/CD Pipeline
615
- - **Automated Testing**: Integrates with CI/CD for continuous performance validation
616
- - **Quality Gates**: Provides pass/fail criteria for deployment decisions
617
- - **Regression Prevention**: Catches performance regressions before production
618
-
619
- ## Performance Benchmarks
620
-
621
- ### Standard Benchmark Suite
622
- ```javascript
623
- // Comprehensive benchmark definitions
624
- const standardBenchmarks = {
625
- // Throughput benchmarks
626
- throughput: {
627
- name: 'Throughput Benchmark',
628
- metrics: ['requests_per_second', 'tasks_per_second', 'messages_per_second'],
629
- duration: 300000, // 5 minutes
630
- warmup: 30000, // 30 seconds
631
- targets: {
632
- requests_per_second: { min: 1000, optimal: 5000 },
633
- tasks_per_second: { min: 100, optimal: 500 },
634
- messages_per_second: { min: 10000, optimal: 50000 }
635
- }
636
- },
637
-
638
- // Latency benchmarks
639
- latency: {
640
- name: 'Latency Benchmark',
641
- metrics: ['p50', 'p90', 'p95', 'p99', 'max'],
642
- duration: 300000,
643
- targets: {
644
- p50: { max: 100 }, // 100ms
645
- p90: { max: 200 }, // 200ms
646
- p95: { max: 500 }, // 500ms
647
- p99: { max: 1000 }, // 1s
648
- max: { max: 5000 } // 5s
649
- }
650
- },
651
-
652
- // Scalability benchmarks
653
- scalability: {
654
- name: 'Scalability Benchmark',
655
- metrics: ['linear_coefficient', 'efficiency_retention'],
656
- load_points: [1, 2, 4, 8, 16, 32, 64],
657
- targets: {
658
- linear_coefficient: { min: 0.8 },
659
- efficiency_retention: { min: 0.7 }
660
- }
661
- }
662
- };
663
- ```
664
-
1
+ ---
2
+ name: Benchmark Suite
3
+ type: agent
4
+ category: optimization
5
+ description: Comprehensive performance benchmarking, regression detection and performance validation
6
+ ---
7
+
8
+ # Benchmark Suite Agent
9
+
10
+ ## Agent Profile
11
+ - **Name**: Benchmark Suite
12
+ - **Type**: Performance Optimization Agent
13
+ - **Specialization**: Comprehensive performance benchmarking and testing
14
+ - **Performance Focus**: Automated benchmarking, regression detection, and performance validation
15
+
16
+ ## Core Capabilities
17
+
18
+ ### 1. Comprehensive Benchmarking Framework
19
+ ```javascript
20
+ // Advanced benchmarking system
21
+ class ComprehensiveBenchmarkSuite {
22
+ constructor() {
23
+ this.benchmarks = {
24
+ // Core performance benchmarks
25
+ throughput: new ThroughputBenchmark(),
26
+ latency: new LatencyBenchmark(),
27
+ scalability: new ScalabilityBenchmark(),
28
+ resource_usage: new ResourceUsageBenchmark(),
29
+
30
+ // Swarm-specific benchmarks
31
+ coordination: new CoordinationBenchmark(),
32
+ load_balancing: new LoadBalancingBenchmark(),
33
+ topology: new TopologyBenchmark(),
34
+ fault_tolerance: new FaultToleranceBenchmark(),
35
+
36
+ // Custom benchmarks
37
+ custom: new CustomBenchmarkManager()
38
+ };
39
+
40
+ this.reporter = new BenchmarkReporter();
41
+ this.comparator = new PerformanceComparator();
42
+ this.analyzer = new BenchmarkAnalyzer();
43
+ }
44
+
45
+ // Execute comprehensive benchmark suite
46
+ async runBenchmarkSuite(config = {}) {
47
+ const suiteConfig = {
48
+ duration: config.duration || 300000, // 5 minutes default
49
+ iterations: config.iterations || 10,
50
+ warmupTime: config.warmupTime || 30000, // 30 seconds
51
+ cooldownTime: config.cooldownTime || 10000, // 10 seconds
52
+ parallel: config.parallel || false,
53
+ baseline: config.baseline || null
54
+ };
55
+
56
+ const results = {
57
+ summary: {},
58
+ detailed: new Map(),
59
+ baseline_comparison: null,
60
+ recommendations: []
61
+ };
62
+
63
+ // Warmup phase
64
+ await this.warmup(suiteConfig.warmupTime);
65
+
66
+ // Execute benchmarks
67
+ if (suiteConfig.parallel) {
68
+ results.detailed = await this.runBenchmarksParallel(suiteConfig);
69
+ } else {
70
+ results.detailed = await this.runBenchmarksSequential(suiteConfig);
71
+ }
72
+
73
+ // Generate summary
74
+ results.summary = this.generateSummary(results.detailed);
75
+
76
+ // Compare with baseline if provided
77
+ if (suiteConfig.baseline) {
78
+ results.baseline_comparison = await this.compareWithBaseline(
79
+ results.detailed,
80
+ suiteConfig.baseline
81
+ );
82
+ }
83
+
84
+ // Generate recommendations
85
+ results.recommendations = await this.generateRecommendations(results);
86
+
87
+ // Cooldown phase
88
+ await this.cooldown(suiteConfig.cooldownTime);
89
+
90
+ return results;
91
+ }
92
+
93
+ // Parallel benchmark execution
94
+ async runBenchmarksParallel(config) {
95
+ const benchmarkPromises = Object.entries(this.benchmarks).map(
96
+ async ([name, benchmark]) => {
97
+ const result = await this.executeBenchmark(benchmark, name, config);
98
+ return [name, result];
99
+ }
100
+ );
101
+
102
+ const results = await Promise.all(benchmarkPromises);
103
+ return new Map(results);
104
+ }
105
+
106
+ // Sequential benchmark execution
107
+ async runBenchmarksSequential(config) {
108
+ const results = new Map();
109
+
110
+ for (const [name, benchmark] of Object.entries(this.benchmarks)) {
111
+ const result = await this.executeBenchmark(benchmark, name, config);
112
+ results.set(name, result);
113
+
114
+ // Brief pause between benchmarks
115
+ await this.sleep(1000);
116
+ }
117
+
118
+ return results;
119
+ }
120
+ }
121
+ ```
122
+
123
+ ### 2. Performance Regression Detection
124
+ ```javascript
125
+ // Advanced regression detection system
126
+ class RegressionDetector {
127
+ constructor() {
128
+ this.detectors = {
129
+ statistical: new StatisticalRegressionDetector(),
130
+ machine_learning: new MLRegressionDetector(),
131
+ threshold: new ThresholdRegressionDetector(),
132
+ trend: new TrendRegressionDetector()
133
+ };
134
+
135
+ this.analyzer = new RegressionAnalyzer();
136
+ this.alerting = new RegressionAlerting();
137
+ }
138
+
139
+ // Detect performance regressions
140
+ async detectRegressions(currentResults, historicalData, config = {}) {
141
+ const regressions = {
142
+ detected: [],
143
+ severity: 'none',
144
+ confidence: 0,
145
+ analysis: {}
146
+ };
147
+
148
+ // Run multiple detection algorithms
149
+ const detectionPromises = Object.entries(this.detectors).map(
150
+ async ([method, detector]) => {
151
+ const detection = await detector.detect(currentResults, historicalData, config);
152
+ return [method, detection];
153
+ }
154
+ );
155
+
156
+ const detectionResults = await Promise.all(detectionPromises);
157
+
158
+ // Aggregate detection results
159
+ for (const [method, detection] of detectionResults) {
160
+ if (detection.regression_detected) {
161
+ regressions.detected.push({
162
+ method,
163
+ ...detection
164
+ });
165
+ }
166
+ }
167
+
168
+ // Calculate overall confidence and severity
169
+ if (regressions.detected.length > 0) {
170
+ regressions.confidence = this.calculateAggregateConfidence(regressions.detected);
171
+ regressions.severity = this.calculateSeverity(regressions.detected);
172
+ regressions.analysis = await this.analyzer.analyze(regressions.detected);
173
+ }
174
+
175
+ return regressions;
176
+ }
177
+
178
+ // Statistical regression detection using change point analysis
179
+ async detectStatisticalRegression(metric, historicalData, sensitivity = 0.95) {
180
+ // Use CUSUM (Cumulative Sum) algorithm for change point detection
181
+ const cusum = this.calculateCUSUM(metric, historicalData);
182
+
183
+ // Detect change points
184
+ const changePoints = this.detectChangePoints(cusum, sensitivity);
185
+
186
+ // Analyze significance of changes
187
+ const analysis = changePoints.map(point => ({
188
+ timestamp: point.timestamp,
189
+ magnitude: point.magnitude,
190
+ direction: point.direction,
191
+ significance: point.significance,
192
+ confidence: point.confidence
193
+ }));
194
+
195
+ return {
196
+ regression_detected: changePoints.length > 0,
197
+ change_points: analysis,
198
+ cusum_statistics: cusum.statistics,
199
+ sensitivity: sensitivity
200
+ };
201
+ }
202
+
203
+ // Machine learning-based regression detection
204
+ async detectMLRegression(metrics, historicalData) {
205
+ // Train anomaly detection model on historical data
206
+ const model = await this.trainAnomalyModel(historicalData);
207
+
208
+ // Predict anomaly scores for current metrics
209
+ const anomalyScores = await model.predict(metrics);
210
+
211
+ // Identify regressions based on anomaly scores
212
+ const threshold = this.calculateDynamicThreshold(anomalyScores);
213
+ const regressions = anomalyScores.filter(score => score.anomaly > threshold);
214
+
215
+ return {
216
+ regression_detected: regressions.length > 0,
217
+ anomaly_scores: anomalyScores,
218
+ threshold: threshold,
219
+ regressions: regressions,
220
+ model_confidence: model.confidence
221
+ };
222
+ }
223
+ }
224
+ ```
225
+
226
+ ### 3. Automated Performance Testing
227
+ ```javascript
228
+ // Comprehensive automated performance testing
229
+ class AutomatedPerformanceTester {
230
+ constructor() {
231
+ this.testSuites = {
232
+ load: new LoadTestSuite(),
233
+ stress: new StressTestSuite(),
234
+ volume: new VolumeTestSuite(),
235
+ endurance: new EnduranceTestSuite(),
236
+ spike: new SpikeTestSuite(),
237
+ configuration: new ConfigurationTestSuite()
238
+ };
239
+
240
+ this.scheduler = new TestScheduler();
241
+ this.orchestrator = new TestOrchestrator();
242
+ this.validator = new ResultValidator();
243
+ }
244
+
245
+ // Execute automated performance test campaign
246
+ async runTestCampaign(config) {
247
+ const campaign = {
248
+ id: this.generateCampaignId(),
249
+ config,
250
+ startTime: Date.now(),
251
+ tests: [],
252
+ results: new Map(),
253
+ summary: null
254
+ };
255
+
256
+ // Schedule test execution
257
+ const schedule = await this.scheduler.schedule(config.tests, config.constraints);
258
+
259
+ // Execute tests according to schedule
260
+ for (const scheduledTest of schedule) {
261
+ const testResult = await this.executeScheduledTest(scheduledTest);
262
+ campaign.tests.push(scheduledTest);
263
+ campaign.results.set(scheduledTest.id, testResult);
264
+
265
+ // Validate results in real-time
266
+ const validation = await this.validator.validate(testResult);
267
+ if (!validation.valid) {
268
+ campaign.summary = {
269
+ status: 'failed',
270
+ reason: validation.reason,
271
+ failedAt: scheduledTest.name
272
+ };
273
+ break;
274
+ }
275
+ }
276
+
277
+ // Generate campaign summary
278
+ if (!campaign.summary) {
279
+ campaign.summary = await this.generateCampaignSummary(campaign);
280
+ }
281
+
282
+ campaign.endTime = Date.now();
283
+ campaign.duration = campaign.endTime - campaign.startTime;
284
+
285
+ return campaign;
286
+ }
287
+
288
+ // Load testing with gradual ramp-up
289
+ async executeLoadTest(config) {
290
+ const loadTest = {
291
+ type: 'load',
292
+ config,
293
+ phases: [],
294
+ metrics: new Map(),
295
+ results: {}
296
+ };
297
+
298
+ // Ramp-up phase
299
+ const rampUpResult = await this.executeRampUp(config.rampUp);
300
+ loadTest.phases.push({ phase: 'ramp-up', result: rampUpResult });
301
+
302
+ // Sustained load phase
303
+ const sustainedResult = await this.executeSustainedLoad(config.sustained);
304
+ loadTest.phases.push({ phase: 'sustained', result: sustainedResult });
305
+
306
+ // Ramp-down phase
307
+ const rampDownResult = await this.executeRampDown(config.rampDown);
308
+ loadTest.phases.push({ phase: 'ramp-down', result: rampDownResult });
309
+
310
+ // Analyze results
311
+ loadTest.results = await this.analyzeLoadTestResults(loadTest.phases);
312
+
313
+ return loadTest;
314
+ }
315
+
316
+ // Stress testing to find breaking points
317
+ async executeStressTest(config) {
318
+ const stressTest = {
319
+ type: 'stress',
320
+ config,
321
+ breakingPoint: null,
322
+ degradationCurve: [],
323
+ results: {}
324
+ };
325
+
326
+ let currentLoad = config.startLoad;
327
+ let systemBroken = false;
328
+
329
+ while (!systemBroken && currentLoad <= config.maxLoad) {
330
+ const testResult = await this.applyLoad(currentLoad, config.duration);
331
+
332
+ stressTest.degradationCurve.push({
333
+ load: currentLoad,
334
+ performance: testResult.performance,
335
+ stability: testResult.stability,
336
+ errors: testResult.errors
337
+ });
338
+
339
+ // Check if system is breaking
340
+ if (this.isSystemBreaking(testResult, config.breakingCriteria)) {
341
+ stressTest.breakingPoint = {
342
+ load: currentLoad,
343
+ performance: testResult.performance,
344
+ reason: this.identifyBreakingReason(testResult)
345
+ };
346
+ systemBroken = true;
347
+ }
348
+
349
+ currentLoad += config.loadIncrement;
350
+ }
351
+
352
+ stressTest.results = await this.analyzeStressTestResults(stressTest);
353
+
354
+ return stressTest;
355
+ }
356
+ }
357
+ ```
358
+
359
+ ### 4. Performance Validation Framework
360
+ ```javascript
361
+ // Comprehensive performance validation
362
+ class PerformanceValidator {
363
+ constructor() {
364
+ this.validators = {
365
+ sla: new SLAValidator(),
366
+ regression: new RegressionValidator(),
367
+ scalability: new ScalabilityValidator(),
368
+ reliability: new ReliabilityValidator(),
369
+ efficiency: new EfficiencyValidator()
370
+ };
371
+
372
+ this.thresholds = new ThresholdManager();
373
+ this.rules = new ValidationRuleEngine();
374
+ }
375
+
376
+ // Validate performance against defined criteria
377
+ async validatePerformance(results, criteria) {
378
+ const validation = {
379
+ overall: {
380
+ passed: true,
381
+ score: 0,
382
+ violations: []
383
+ },
384
+ detailed: new Map(),
385
+ recommendations: []
386
+ };
387
+
388
+ // Run all validators
389
+ const validationPromises = Object.entries(this.validators).map(
390
+ async ([type, validator]) => {
391
+ const result = await validator.validate(results, criteria[type]);
392
+ return [type, result];
393
+ }
394
+ );
395
+
396
+ const validationResults = await Promise.all(validationPromises);
397
+
398
+ // Aggregate validation results
399
+ for (const [type, result] of validationResults) {
400
+ validation.detailed.set(type, result);
401
+
402
+ if (!result.passed) {
403
+ validation.overall.passed = false;
404
+ validation.overall.violations.push(...result.violations);
405
+ }
406
+
407
+ validation.overall.score += result.score * (criteria[type]?.weight || 1);
408
+ }
409
+
410
+ // Normalize overall score
411
+ const totalWeight = Object.values(criteria).reduce((sum, c) => sum + (c.weight || 1), 0);
412
+ validation.overall.score /= totalWeight;
413
+
414
+ // Generate recommendations
415
+ validation.recommendations = await this.generateValidationRecommendations(validation);
416
+
417
+ return validation;
418
+ }
419
+
420
+ // SLA validation
421
+ async validateSLA(results, slaConfig) {
422
+ const slaValidation = {
423
+ passed: true,
424
+ violations: [],
425
+ score: 1.0,
426
+ metrics: {}
427
+ };
428
+
429
+ // Validate each SLA metric
430
+ for (const [metric, threshold] of Object.entries(slaConfig.thresholds)) {
431
+ const actualValue = this.extractMetricValue(results, metric);
432
+ const validation = this.validateThreshold(actualValue, threshold);
433
+
434
+ slaValidation.metrics[metric] = {
435
+ actual: actualValue,
436
+ threshold: threshold.value,
437
+ operator: threshold.operator,
438
+ passed: validation.passed,
439
+ deviation: validation.deviation
440
+ };
441
+
442
+ if (!validation.passed) {
443
+ slaValidation.passed = false;
444
+ slaValidation.violations.push({
445
+ metric,
446
+ actual: actualValue,
447
+ expected: threshold.value,
448
+ severity: threshold.severity || 'medium'
449
+ });
450
+
451
+ // Reduce score based on violation severity
452
+ const severityMultiplier = this.getSeverityMultiplier(threshold.severity);
453
+ slaValidation.score -= (validation.deviation * severityMultiplier);
454
+ }
455
+ }
456
+
457
+ slaValidation.score = Math.max(0, slaValidation.score);
458
+
459
+ return slaValidation;
460
+ }
461
+
462
+ // Scalability validation
463
+ async validateScalability(results, scalabilityConfig) {
464
+ const scalabilityValidation = {
465
+ passed: true,
466
+ violations: [],
467
+ score: 1.0,
468
+ analysis: {}
469
+ };
470
+
471
+ // Linear scalability analysis
472
+ if (scalabilityConfig.linear) {
473
+ const linearityAnalysis = this.analyzeLinearScalability(results);
474
+ scalabilityValidation.analysis.linearity = linearityAnalysis;
475
+
476
+ if (linearityAnalysis.coefficient < scalabilityConfig.linear.minCoefficient) {
477
+ scalabilityValidation.passed = false;
478
+ scalabilityValidation.violations.push({
479
+ type: 'linearity',
480
+ actual: linearityAnalysis.coefficient,
481
+ expected: scalabilityConfig.linear.minCoefficient
482
+ });
483
+ }
484
+ }
485
+
486
+ // Efficiency retention analysis
487
+ if (scalabilityConfig.efficiency) {
488
+ const efficiencyAnalysis = this.analyzeEfficiencyRetention(results);
489
+ scalabilityValidation.analysis.efficiency = efficiencyAnalysis;
490
+
491
+ if (efficiencyAnalysis.retention < scalabilityConfig.efficiency.minRetention) {
492
+ scalabilityValidation.passed = false;
493
+ scalabilityValidation.violations.push({
494
+ type: 'efficiency_retention',
495
+ actual: efficiencyAnalysis.retention,
496
+ expected: scalabilityConfig.efficiency.minRetention
497
+ });
498
+ }
499
+ }
500
+
501
+ return scalabilityValidation;
502
+ }
503
+ }
504
+ ```
505
+
506
+ ## MCP Integration Hooks
507
+
508
+ ### Benchmark Execution Integration
509
+ ```javascript
510
+ // Comprehensive MCP benchmark integration
511
+ const benchmarkIntegration = {
512
+ // Execute performance benchmarks
513
+ async runBenchmarks(config = {}) {
514
+ // Run benchmark suite
515
+ const benchmarkResult = await mcp.benchmark_run({
516
+ suite: config.suite || 'comprehensive'
517
+ });
518
+
519
+ // Collect detailed metrics during benchmarking
520
+ const metrics = await mcp.metrics_collect({
521
+ components: ['system', 'agents', 'coordination', 'memory']
522
+ });
523
+
524
+ // Analyze performance trends
525
+ const trends = await mcp.trend_analysis({
526
+ metric: 'performance',
527
+ period: '24h'
528
+ });
529
+
530
+ // Cost analysis
531
+ const costAnalysis = await mcp.cost_analysis({
532
+ timeframe: '24h'
533
+ });
534
+
535
+ return {
536
+ benchmark: benchmarkResult,
537
+ metrics,
538
+ trends,
539
+ costAnalysis,
540
+ timestamp: Date.now()
541
+ };
542
+ },
543
+
544
+ // Quality assessment
545
+ async assessQuality(criteria) {
546
+ const qualityAssessment = await mcp.quality_assess({
547
+ target: 'swarm-performance',
548
+ criteria: criteria || [
549
+ 'throughput',
550
+ 'latency',
551
+ 'reliability',
552
+ 'scalability',
553
+ 'efficiency'
554
+ ]
555
+ });
556
+
557
+ return qualityAssessment;
558
+ },
559
+
560
+ // Error pattern analysis
561
+ async analyzeErrorPatterns() {
562
+ // Collect system logs
563
+ const logs = await this.collectSystemLogs();
564
+
565
+ // Analyze error patterns
566
+ const errorAnalysis = await mcp.error_analysis({
567
+ logs: logs
568
+ });
569
+
570
+ return errorAnalysis;
571
+ }
572
+ };
573
+ ```
574
+
575
+ ## Operational Commands
576
+
577
+ ### Benchmarking Commands
578
+ ```bash
579
+ # Run comprehensive benchmark suite
580
+ npx claude-flow benchmark-run --suite comprehensive --duration 300
581
+
582
+ # Execute specific benchmark
583
+ npx claude-flow benchmark-run --suite throughput --iterations 10
584
+
585
+ # Compare with baseline
586
+ npx claude-flow benchmark-compare --current <results> --baseline <baseline>
587
+
588
+ # Quality assessment
589
+ npx claude-flow quality-assess --target swarm-performance --criteria throughput,latency
590
+
591
+ # Performance validation
592
+ npx claude-flow validate-performance --results <file> --criteria <file>
593
+ ```
594
+
595
+ ### Regression Detection Commands
596
+ ```bash
597
+ # Detect performance regressions
598
+ npx claude-flow detect-regression --current <results> --historical <data>
599
+
600
+ # Set up automated regression monitoring
601
+ npx claude-flow regression-monitor --enable --sensitivity 0.95
602
+
603
+ # Analyze error patterns
604
+ npx claude-flow error-analysis --logs <log-files>
605
+ ```
606
+
607
+ ## Integration Points
608
+
609
+ ### With Other Optimization Agents
610
+ - **Performance Monitor**: Provides continuous monitoring data for benchmarking
611
+ - **Load Balancer**: Validates load balancing effectiveness through benchmarks
612
+ - **Topology Optimizer**: Tests topology configurations for optimal performance
613
+
614
+ ### With CI/CD Pipeline
615
+ - **Automated Testing**: Integrates with CI/CD for continuous performance validation
616
+ - **Quality Gates**: Provides pass/fail criteria for deployment decisions
617
+ - **Regression Prevention**: Catches performance regressions before production
618
+
619
+ ## Performance Benchmarks
620
+
621
+ ### Standard Benchmark Suite
622
+ ```javascript
623
+ // Comprehensive benchmark definitions
624
+ const standardBenchmarks = {
625
+ // Throughput benchmarks
626
+ throughput: {
627
+ name: 'Throughput Benchmark',
628
+ metrics: ['requests_per_second', 'tasks_per_second', 'messages_per_second'],
629
+ duration: 300000, // 5 minutes
630
+ warmup: 30000, // 30 seconds
631
+ targets: {
632
+ requests_per_second: { min: 1000, optimal: 5000 },
633
+ tasks_per_second: { min: 100, optimal: 500 },
634
+ messages_per_second: { min: 10000, optimal: 50000 }
635
+ }
636
+ },
637
+
638
+ // Latency benchmarks
639
+ latency: {
640
+ name: 'Latency Benchmark',
641
+ metrics: ['p50', 'p90', 'p95', 'p99', 'max'],
642
+ duration: 300000,
643
+ targets: {
644
+ p50: { max: 100 }, // 100ms
645
+ p90: { max: 200 }, // 200ms
646
+ p95: { max: 500 }, // 500ms
647
+ p99: { max: 1000 }, // 1s
648
+ max: { max: 5000 } // 5s
649
+ }
650
+ },
651
+
652
+ // Scalability benchmarks
653
+ scalability: {
654
+ name: 'Scalability Benchmark',
655
+ metrics: ['linear_coefficient', 'efficiency_retention'],
656
+ load_points: [1, 2, 4, 8, 16, 32, 64],
657
+ targets: {
658
+ linear_coefficient: { min: 0.8 },
659
+ efficiency_retention: { min: 0.7 }
660
+ }
661
+ }
662
+ };
663
+ ```
664
+
665
665
  This Benchmark Suite agent provides comprehensive automated performance testing, regression detection, and validation capabilities to ensure optimal swarm performance and prevent performance degradation.