claude-flow 3.32.8 → 3.32.10

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (455) hide show
  1. package/.claude/.proven-config-version +1 -0
  2. package/.claude/agents/MIGRATION_SUMMARY.md +221 -221
  3. package/.claude/agents/analysis/analyze-code-quality.md +57 -57
  4. package/.claude/agents/analysis/code-analyzer.md +188 -188
  5. package/.claude/agents/analysis/code-review/analyze-code-quality.md +57 -57
  6. package/.claude/agents/architecture/system-design/arch-system-design.md +35 -35
  7. package/.claude/agents/base-template-generator.md +41 -41
  8. package/.claude/agents/consensus/byzantine-coordinator.md +42 -42
  9. package/.claude/agents/consensus/crdt-synchronizer.md +976 -976
  10. package/.claude/agents/consensus/gossip-coordinator.md +42 -42
  11. package/.claude/agents/consensus/performance-benchmarker.md +830 -830
  12. package/.claude/agents/consensus/quorum-manager.md +802 -802
  13. package/.claude/agents/consensus/raft-manager.md +42 -42
  14. package/.claude/agents/consensus/security-manager.md +601 -601
  15. package/.claude/agents/core/coder.md +254 -254
  16. package/.claude/agents/core/planner.md +151 -151
  17. package/.claude/agents/core/researcher.md +173 -173
  18. package/.claude/agents/core/reviewer.md +308 -308
  19. package/.claude/agents/core/tester.md +299 -299
  20. package/.claude/agents/custom/test-long-runner.md +43 -43
  21. package/.claude/agents/data/ml/data-ml-model.md +75 -75
  22. package/.claude/agents/database-specialist.md +9 -9
  23. package/.claude/agents/development/backend/dev-backend-api.md +28 -28
  24. package/.claude/agents/development/dev-backend-api.md +177 -177
  25. package/.claude/agents/devops/ci-cd/ops-cicd-github.md +51 -51
  26. package/.claude/agents/documentation/api-docs/docs-api-openapi.md +62 -62
  27. package/.claude/agents/dual-mode/codex-coordinator.md +206 -206
  28. package/.claude/agents/dual-mode/codex-worker.md +190 -190
  29. package/.claude/agents/dual-mode/dual-orchestrator.md +253 -253
  30. package/.claude/agents/flow-nexus/app-store.md +87 -87
  31. package/.claude/agents/flow-nexus/authentication.md +68 -68
  32. package/.claude/agents/flow-nexus/challenges.md +80 -80
  33. package/.claude/agents/flow-nexus/neural-network.md +87 -87
  34. package/.claude/agents/flow-nexus/payments.md +82 -82
  35. package/.claude/agents/flow-nexus/sandbox.md +75 -75
  36. package/.claude/agents/flow-nexus/swarm.md +75 -75
  37. package/.claude/agents/flow-nexus/user-tools.md +95 -95
  38. package/.claude/agents/flow-nexus/workflow.md +83 -83
  39. package/.claude/agents/github/code-review-swarm.md +520 -520
  40. package/.claude/agents/github/github-modes.md +153 -153
  41. package/.claude/agents/github/issue-tracker.md +298 -298
  42. package/.claude/agents/github/multi-repo-swarm.md +524 -524
  43. package/.claude/agents/github/pr-manager.md +162 -162
  44. package/.claude/agents/github/project-board-sync.md +477 -477
  45. package/.claude/agents/github/release-manager.md +337 -337
  46. package/.claude/agents/github/release-swarm.md +550 -550
  47. package/.claude/agents/github/repo-architect.md +364 -364
  48. package/.claude/agents/github/swarm-issue.md +550 -550
  49. package/.claude/agents/github/swarm-pr.md +401 -401
  50. package/.claude/agents/github/sync-coordinator.md +424 -424
  51. package/.claude/agents/github/workflow-automation.md +604 -604
  52. package/.claude/agents/goal/agent.md +816 -816
  53. package/.claude/agents/goal/code-goal-planner.md +444 -444
  54. package/.claude/agents/goal/goal-planner.md +167 -167
  55. package/.claude/agents/hive-mind/collective-intelligence-coordinator.md +128 -128
  56. package/.claude/agents/hive-mind/queen-coordinator.md +201 -201
  57. package/.claude/agents/hive-mind/scout-explorer.md +240 -240
  58. package/.claude/agents/hive-mind/swarm-memory-manager.md +191 -191
  59. package/.claude/agents/hive-mind/worker-specialist.md +215 -215
  60. package/.claude/agents/neural/safla-neural.md +73 -73
  61. package/.claude/agents/optimization/benchmark-suite.md +662 -662
  62. package/.claude/agents/optimization/load-balancer.md +428 -428
  63. package/.claude/agents/optimization/performance-monitor.md +669 -669
  64. package/.claude/agents/optimization/resource-allocator.md +671 -671
  65. package/.claude/agents/optimization/topology-optimizer.md +805 -805
  66. package/.claude/agents/payments/agentic-payments.md +126 -126
  67. package/.claude/agents/project-coordinator.md +8 -8
  68. package/.claude/agents/python-specialist.md +9 -9
  69. package/.claude/agents/reasoning/agent.md +816 -816
  70. package/.claude/agents/reasoning/goal-planner.md +72 -72
  71. package/.claude/agents/security-auditor.md +9 -9
  72. package/.claude/agents/sona/sona-learning-optimizer.md +65 -65
  73. package/.claude/agents/sparc/architecture.md +452 -452
  74. package/.claude/agents/sparc/pseudocode.md +298 -298
  75. package/.claude/agents/sparc/refinement.md +503 -503
  76. package/.claude/agents/sparc/specification.md +257 -257
  77. package/.claude/agents/specialized/mobile/spec-mobile-react-native.md +87 -87
  78. package/.claude/agents/sublinear/consensus-coordinator.md +337 -337
  79. package/.claude/agents/sublinear/matrix-optimizer.md +184 -184
  80. package/.claude/agents/sublinear/pagerank-analyzer.md +298 -298
  81. package/.claude/agents/sublinear/performance-optimizer.md +367 -367
  82. package/.claude/agents/sublinear/trading-predictor.md +245 -245
  83. package/.claude/agents/swarm/adaptive-coordinator.md +363 -363
  84. package/.claude/agents/swarm/hierarchical-coordinator.md +299 -299
  85. package/.claude/agents/swarm/mesh-coordinator.md +362 -362
  86. package/.claude/agents/templates/automation-smart-agent.md +184 -184
  87. package/.claude/agents/templates/coordinator-swarm-init.md +82 -82
  88. package/.claude/agents/templates/github-pr-manager.md +154 -154
  89. package/.claude/agents/templates/implementer-sparc-coder.md +242 -242
  90. package/.claude/agents/templates/memory-coordinator.md +162 -162
  91. package/.claude/agents/templates/migration-plan.md +723 -723
  92. package/.claude/agents/templates/orchestrator-task.md +119 -119
  93. package/.claude/agents/templates/performance-analyzer.md +178 -178
  94. package/.claude/agents/templates/sparc-coordinator.md +162 -162
  95. package/.claude/agents/testing/production-validator.md +372 -372
  96. package/.claude/agents/testing/tdd-london-swarm.md +221 -221
  97. package/.claude/agents/testing/unit/tdd-london-swarm.md +221 -221
  98. package/.claude/agents/testing/validation/production-validator.md +372 -372
  99. package/.claude/agents/typescript-specialist.md +9 -9
  100. package/.claude/agents/v3/database-specialist.md +9 -9
  101. package/.claude/agents/v3/project-coordinator.md +8 -8
  102. package/.claude/agents/v3/python-specialist.md +9 -9
  103. package/.claude/agents/v3/test-architect.md +9 -9
  104. package/.claude/agents/v3/typescript-specialist.md +9 -9
  105. package/.claude/agents/v3/v3-integration-architect.md +311 -311
  106. package/.claude/agents/v3/v3-memory-specialist.md +280 -280
  107. package/.claude/agents/v3/v3-performance-engineer.md +362 -362
  108. package/.claude/agents/v3/v3-queen-coordinator.md +62 -62
  109. package/.claude/agents/v3/v3-security-architect.md +139 -139
  110. package/.claude/checkpoints/1767754460.json +8 -8
  111. package/.claude/commands/agents/README.md +10 -10
  112. package/.claude/commands/agents/agent-capabilities.md +21 -21
  113. package/.claude/commands/agents/agent-coordination.md +28 -28
  114. package/.claude/commands/agents/agent-spawning.md +28 -28
  115. package/.claude/commands/agents/agent-types.md +26 -26
  116. package/.claude/commands/analysis/COMMAND_COMPLIANCE_REPORT.md +53 -53
  117. package/.claude/commands/analysis/README.md +9 -9
  118. package/.claude/commands/analysis/bottleneck-detect.md +162 -162
  119. package/.claude/commands/analysis/performance-bottlenecks.md +58 -58
  120. package/.claude/commands/analysis/performance-report.md +25 -25
  121. package/.claude/commands/analysis/token-efficiency.md +44 -44
  122. package/.claude/commands/analysis/token-usage.md +25 -25
  123. package/.claude/commands/automation/README.md +9 -9
  124. package/.claude/commands/automation/auto-agent.md +122 -122
  125. package/.claude/commands/automation/self-healing.md +105 -105
  126. package/.claude/commands/automation/session-memory.md +89 -89
  127. package/.claude/commands/automation/smart-agents.md +72 -72
  128. package/.claude/commands/automation/smart-spawn.md +25 -25
  129. package/.claude/commands/automation/workflow-select.md +25 -25
  130. package/.claude/commands/claude-flow-help.md +103 -103
  131. package/.claude/commands/claude-flow-memory.md +107 -107
  132. package/.claude/commands/claude-flow-swarm.md +205 -205
  133. package/.claude/commands/coordination/README.md +9 -9
  134. package/.claude/commands/coordination/agent-spawn.md +25 -25
  135. package/.claude/commands/coordination/init.md +44 -44
  136. package/.claude/commands/coordination/orchestrate.md +43 -43
  137. package/.claude/commands/coordination/spawn.md +45 -45
  138. package/.claude/commands/coordination/swarm-init.md +85 -85
  139. package/.claude/commands/coordination/task-orchestrate.md +25 -25
  140. package/.claude/commands/flow-nexus/app-store.md +123 -123
  141. package/.claude/commands/flow-nexus/challenges.md +119 -119
  142. package/.claude/commands/flow-nexus/login-registration.md +64 -64
  143. package/.claude/commands/flow-nexus/neural-network.md +133 -133
  144. package/.claude/commands/flow-nexus/payments.md +115 -115
  145. package/.claude/commands/flow-nexus/sandbox.md +82 -82
  146. package/.claude/commands/flow-nexus/swarm.md +86 -86
  147. package/.claude/commands/flow-nexus/user-tools.md +151 -151
  148. package/.claude/commands/flow-nexus/workflow.md +114 -114
  149. package/.claude/commands/github/README.md +11 -11
  150. package/.claude/commands/github/code-review-swarm.md +513 -513
  151. package/.claude/commands/github/code-review.md +25 -25
  152. package/.claude/commands/github/github-modes.md +146 -146
  153. package/.claude/commands/github/github-swarm.md +121 -121
  154. package/.claude/commands/github/issue-tracker.md +291 -291
  155. package/.claude/commands/github/issue-triage.md +25 -25
  156. package/.claude/commands/github/multi-repo-swarm.md +518 -518
  157. package/.claude/commands/github/pr-enhance.md +26 -26
  158. package/.claude/commands/github/pr-manager.md +169 -169
  159. package/.claude/commands/github/project-board-sync.md +470 -470
  160. package/.claude/commands/github/release-manager.md +337 -337
  161. package/.claude/commands/github/release-swarm.md +543 -543
  162. package/.claude/commands/github/repo-analyze.md +25 -25
  163. package/.claude/commands/github/repo-architect.md +366 -366
  164. package/.claude/commands/github/swarm-issue.md +481 -481
  165. package/.claude/commands/github/swarm-pr.md +284 -284
  166. package/.claude/commands/github/sync-coordinator.md +300 -300
  167. package/.claude/commands/github/workflow-automation.md +441 -441
  168. package/.claude/commands/hive-mind/README.md +17 -17
  169. package/.claude/commands/hive-mind/hive-mind-consensus.md +8 -8
  170. package/.claude/commands/hive-mind/hive-mind-init.md +18 -18
  171. package/.claude/commands/hive-mind/hive-mind-memory.md +8 -8
  172. package/.claude/commands/hive-mind/hive-mind-metrics.md +8 -8
  173. package/.claude/commands/hive-mind/hive-mind-resume.md +8 -8
  174. package/.claude/commands/hive-mind/hive-mind-sessions.md +8 -8
  175. package/.claude/commands/hive-mind/hive-mind-spawn.md +21 -21
  176. package/.claude/commands/hive-mind/hive-mind-status.md +8 -8
  177. package/.claude/commands/hive-mind/hive-mind-stop.md +8 -8
  178. package/.claude/commands/hive-mind/hive-mind-wizard.md +8 -8
  179. package/.claude/commands/hive-mind/hive-mind.md +27 -27
  180. package/.claude/commands/hooks/README.md +11 -11
  181. package/.claude/commands/hooks/overview.md +57 -57
  182. package/.claude/commands/hooks/post-edit.md +117 -117
  183. package/.claude/commands/hooks/post-task.md +112 -112
  184. package/.claude/commands/hooks/pre-edit.md +113 -113
  185. package/.claude/commands/hooks/pre-task.md +111 -111
  186. package/.claude/commands/hooks/session-end.md +118 -118
  187. package/.claude/commands/hooks/setup.md +102 -102
  188. package/.claude/commands/memory/README.md +9 -9
  189. package/.claude/commands/memory/memory-persist.md +25 -25
  190. package/.claude/commands/memory/memory-search.md +25 -25
  191. package/.claude/commands/memory/memory-usage.md +25 -25
  192. package/.claude/commands/memory/neural.md +47 -47
  193. package/.claude/commands/monitoring/README.md +9 -9
  194. package/.claude/commands/monitoring/agent-metrics.md +25 -25
  195. package/.claude/commands/monitoring/agents.md +44 -44
  196. package/.claude/commands/monitoring/real-time-view.md +25 -25
  197. package/.claude/commands/monitoring/status.md +46 -46
  198. package/.claude/commands/monitoring/swarm-monitor.md +25 -25
  199. package/.claude/commands/optimization/README.md +9 -9
  200. package/.claude/commands/optimization/auto-topology.md +61 -61
  201. package/.claude/commands/optimization/cache-manage.md +25 -25
  202. package/.claude/commands/optimization/parallel-execute.md +25 -25
  203. package/.claude/commands/optimization/parallel-execution.md +49 -49
  204. package/.claude/commands/optimization/topology-optimize.md +25 -25
  205. package/.claude/commands/pair/README.md +260 -260
  206. package/.claude/commands/pair/commands.md +545 -545
  207. package/.claude/commands/pair/config.md +509 -509
  208. package/.claude/commands/pair/examples.md +511 -511
  209. package/.claude/commands/pair/modes.md +347 -347
  210. package/.claude/commands/pair/session.md +406 -406
  211. package/.claude/commands/pair/start.md +208 -208
  212. package/.claude/commands/sparc/analyzer.md +51 -51
  213. package/.claude/commands/sparc/architect.md +53 -53
  214. package/.claude/commands/sparc/ask.md +97 -97
  215. package/.claude/commands/sparc/batch-executor.md +54 -54
  216. package/.claude/commands/sparc/code.md +89 -89
  217. package/.claude/commands/sparc/coder.md +54 -54
  218. package/.claude/commands/sparc/debug.md +83 -83
  219. package/.claude/commands/sparc/debugger.md +54 -54
  220. package/.claude/commands/sparc/designer.md +53 -53
  221. package/.claude/commands/sparc/devops.md +109 -109
  222. package/.claude/commands/sparc/docs-writer.md +80 -80
  223. package/.claude/commands/sparc/documenter.md +54 -54
  224. package/.claude/commands/sparc/innovator.md +54 -54
  225. package/.claude/commands/sparc/integration.md +83 -83
  226. package/.claude/commands/sparc/mcp.md +117 -117
  227. package/.claude/commands/sparc/memory-manager.md +54 -54
  228. package/.claude/commands/sparc/optimizer.md +54 -54
  229. package/.claude/commands/sparc/orchestrator.md +131 -131
  230. package/.claude/commands/sparc/post-deployment-monitoring-mode.md +83 -83
  231. package/.claude/commands/sparc/refinement-optimization-mode.md +83 -83
  232. package/.claude/commands/sparc/researcher.md +54 -54
  233. package/.claude/commands/sparc/reviewer.md +54 -54
  234. package/.claude/commands/sparc/security-review.md +80 -80
  235. package/.claude/commands/sparc/sparc-modes.md +174 -174
  236. package/.claude/commands/sparc/sparc.md +111 -111
  237. package/.claude/commands/sparc/spec-pseudocode.md +80 -80
  238. package/.claude/commands/sparc/supabase-admin.md +348 -348
  239. package/.claude/commands/sparc/swarm-coordinator.md +54 -54
  240. package/.claude/commands/sparc/tdd.md +54 -54
  241. package/.claude/commands/sparc/tester.md +54 -54
  242. package/.claude/commands/sparc/tutorial.md +79 -79
  243. package/.claude/commands/sparc/workflow-manager.md +54 -54
  244. package/.claude/commands/sparc.md +166 -166
  245. package/.claude/commands/stream-chain/pipeline.md +120 -120
  246. package/.claude/commands/stream-chain/run.md +69 -69
  247. package/.claude/commands/swarm/README.md +15 -15
  248. package/.claude/commands/swarm/analysis.md +95 -95
  249. package/.claude/commands/swarm/development.md +96 -96
  250. package/.claude/commands/swarm/examples.md +168 -168
  251. package/.claude/commands/swarm/maintenance.md +102 -102
  252. package/.claude/commands/swarm/optimization.md +117 -117
  253. package/.claude/commands/swarm/research.md +136 -136
  254. package/.claude/commands/swarm/swarm-analysis.md +8 -8
  255. package/.claude/commands/swarm/swarm-background.md +8 -8
  256. package/.claude/commands/swarm/swarm-init.md +19 -19
  257. package/.claude/commands/swarm/swarm-modes.md +8 -8
  258. package/.claude/commands/swarm/swarm-monitor.md +8 -8
  259. package/.claude/commands/swarm/swarm-spawn.md +19 -19
  260. package/.claude/commands/swarm/swarm-status.md +8 -8
  261. package/.claude/commands/swarm/swarm-strategies.md +8 -8
  262. package/.claude/commands/swarm/swarm.md +27 -27
  263. package/.claude/commands/swarm/testing.md +131 -131
  264. package/.claude/commands/training/README.md +9 -9
  265. package/.claude/commands/training/model-update.md +25 -25
  266. package/.claude/commands/training/neural-patterns.md +73 -73
  267. package/.claude/commands/training/neural-train.md +25 -25
  268. package/.claude/commands/training/pattern-learn.md +25 -25
  269. package/.claude/commands/training/specialization.md +62 -62
  270. package/.claude/commands/truth/start.md +142 -142
  271. package/.claude/commands/verify/check.md +49 -49
  272. package/.claude/commands/verify/start.md +127 -127
  273. package/.claude/commands/workflows/README.md +9 -9
  274. package/.claude/commands/workflows/development.md +77 -77
  275. package/.claude/commands/workflows/research.md +62 -62
  276. package/.claude/commands/workflows/workflow-create.md +25 -25
  277. package/.claude/commands/workflows/workflow-execute.md +25 -25
  278. package/.claude/commands/workflows/workflow-export.md +25 -25
  279. package/.claude/config/v3-dependency-optimization.json +265 -265
  280. package/.claude/config/v3-performance-targets.json +250 -250
  281. package/.claude/helpers/.LOCKED +2 -2
  282. package/.claude/helpers/.helpers-version +1 -0
  283. package/.claude/helpers/README.md +96 -96
  284. package/.claude/helpers/adr-compliance.sh +186 -186
  285. package/.claude/helpers/aggressive-microcompact.mjs +36 -36
  286. package/.claude/helpers/auto-commit.sh +178 -178
  287. package/.claude/helpers/auto-memory-hook.mjs +430 -430
  288. package/.claude/helpers/checkpoint-manager.sh +251 -251
  289. package/.claude/helpers/context-persistence-hook.mjs +2001 -2001
  290. package/.claude/helpers/daemon-manager.sh +252 -252
  291. package/.claude/helpers/ddd-tracker.sh +144 -144
  292. package/.claude/helpers/github-safe.js +156 -156
  293. package/.claude/helpers/github-setup.sh +45 -45
  294. package/.claude/helpers/guidance-hook.sh +13 -13
  295. package/.claude/helpers/guidance-hooks.sh +102 -102
  296. package/.claude/helpers/health-monitor.sh +108 -108
  297. package/.claude/helpers/helpers.manifest.json +13 -0
  298. package/.claude/helpers/hook-handler.cjs +464 -464
  299. package/.claude/helpers/intelligence.cjs +1058 -1058
  300. package/.claude/helpers/learning-hooks.sh +329 -329
  301. package/.claude/helpers/learning-optimizer.sh +127 -127
  302. package/.claude/helpers/learning-service.mjs +1144 -1144
  303. package/.claude/helpers/memory.cjs +84 -84
  304. package/.claude/helpers/metrics-db.mjs +503 -503
  305. package/.claude/helpers/patch-aggressive-prune.mjs +184 -184
  306. package/.claude/helpers/pattern-consolidator.sh +86 -86
  307. package/.claude/helpers/perf-worker.sh +160 -160
  308. package/.claude/helpers/quick-start.sh +19 -19
  309. package/.claude/helpers/router.cjs +62 -62
  310. package/.claude/helpers/security-scanner.sh +127 -127
  311. package/.claude/helpers/session.cjs +125 -125
  312. package/.claude/helpers/setup-mcp.sh +18 -18
  313. package/.claude/helpers/standard-checkpoint-hooks.sh +189 -189
  314. package/.claude/helpers/statusline.cjs +211 -2
  315. package/.claude/helpers/swarm-comms.sh +353 -353
  316. package/.claude/helpers/swarm-hooks.sh +761 -761
  317. package/.claude/helpers/swarm-monitor.sh +210 -210
  318. package/.claude/helpers/sync-v3-metrics.sh +245 -245
  319. package/.claude/helpers/update-v3-progress.sh +165 -165
  320. package/.claude/helpers/v3-quick-status.sh +57 -57
  321. package/.claude/helpers/v3.sh +110 -110
  322. package/.claude/helpers/validate-v3-config.sh +215 -215
  323. package/.claude/helpers/worker-manager.sh +170 -170
  324. package/.claude/mcp.json +12 -12
  325. package/.claude/proven-config.json +42 -0
  326. package/.claude/settings.json +284 -284
  327. package/.claude/settings.json.bak +526 -526
  328. package/.claude/skills/agentdb-advanced/SKILL.md +550 -550
  329. package/.claude/skills/agentdb-learning/SKILL.md +545 -545
  330. package/.claude/skills/agentdb-memory-patterns/SKILL.md +339 -339
  331. package/.claude/skills/agentdb-optimization/SKILL.md +509 -509
  332. package/.claude/skills/agentdb-vector-search/SKILL.md +339 -339
  333. package/.claude/skills/agentic-jujutsu/SKILL.md +645 -645
  334. package/.claude/skills/browser/SKILL.md +204 -204
  335. package/.claude/skills/dual-mode/README.md +71 -71
  336. package/.claude/skills/dual-mode/dual-collect.md +103 -103
  337. package/.claude/skills/dual-mode/dual-coordinate.md +85 -85
  338. package/.claude/skills/dual-mode/dual-spawn.md +81 -81
  339. package/.claude/skills/flow-nexus-neural/SKILL.md +727 -727
  340. package/.claude/skills/flow-nexus-platform/SKILL.md +1154 -1154
  341. package/.claude/skills/flow-nexus-swarm/SKILL.md +604 -604
  342. package/.claude/skills/github-code-review/SKILL.md +1125 -1125
  343. package/.claude/skills/github-multi-repo/SKILL.md +862 -862
  344. package/.claude/skills/github-project-management/SKILL.md +1262 -1262
  345. package/.claude/skills/github-release-management/SKILL.md +1064 -1064
  346. package/.claude/skills/github-workflow-automation/SKILL.md +1047 -1047
  347. package/.claude/skills/hive-mind-advanced/SKILL.md +709 -709
  348. package/.claude/skills/hooks-automation/SKILL.md +1201 -1201
  349. package/.claude/skills/pair-programming/SKILL.md +1202 -1202
  350. package/.claude/skills/performance-analysis/SKILL.md +560 -560
  351. package/.claude/skills/reasoningbank-agentdb/SKILL.md +446 -446
  352. package/.claude/skills/reasoningbank-intelligence/SKILL.md +201 -201
  353. package/.claude/skills/skill-builder/SKILL.md +910 -910
  354. package/.claude/skills/sparc-methodology/SKILL.md +1106 -1106
  355. package/.claude/skills/stream-chain/SKILL.md +560 -560
  356. package/.claude/skills/swarm-advanced/SKILL.md +970 -970
  357. package/.claude/skills/swarm-orchestration/SKILL.md +179 -179
  358. package/.claude/skills/v3-cli-modernization/SKILL.md +871 -871
  359. package/.claude/skills/v3-core-implementation/SKILL.md +796 -796
  360. package/.claude/skills/v3-ddd-architecture/SKILL.md +441 -441
  361. package/.claude/skills/v3-integration-deep/SKILL.md +240 -240
  362. package/.claude/skills/v3-mcp-optimization/SKILL.md +776 -776
  363. package/.claude/skills/v3-memory-unification/SKILL.md +173 -173
  364. package/.claude/skills/v3-performance-optimization/SKILL.md +389 -389
  365. package/.claude/skills/v3-security-overhaul/SKILL.md +81 -81
  366. package/.claude/skills/v3-swarm-coordination/SKILL.md +339 -339
  367. package/.claude/skills/verification-quality/SKILL.md +691 -691
  368. package/.claude/skills/worker-benchmarks/SKILL.md +129 -129
  369. package/.claude/skills/worker-integration/SKILL.md +147 -147
  370. package/.claude/statusline-command.sh +176 -176
  371. package/.claude/statusline.mjs +109 -109
  372. package/.claude/statusline.sh +431 -431
  373. package/.claude/workflows/full-system-test.js +65 -65
  374. package/.claude/workflows/intelligence-system-hardening.js +120 -120
  375. package/.claude/workflows/plugin-contract-audit.js +91 -91
  376. package/.claude-plugin/README.md +720 -720
  377. package/.claude-plugin/docs/INSTALLATION.md +261 -261
  378. package/.claude-plugin/docs/PLUGIN_SUMMARY.md +361 -361
  379. package/.claude-plugin/docs/QUICKSTART.md +361 -361
  380. package/.claude-plugin/docs/STRUCTURE.md +128 -128
  381. package/.claude-plugin/hooks/hooks.json +79 -77
  382. package/.claude-plugin/marketplace.json +185 -185
  383. package/.claude-plugin/plugin.json +71 -71
  384. package/.claude-plugin/scripts/install.sh +234 -234
  385. package/.claude-plugin/scripts/ruflo-hook.cjs +166 -166
  386. package/.claude-plugin/scripts/ruflo-hook.sh +52 -52
  387. package/.claude-plugin/scripts/uninstall.sh +36 -36
  388. package/.claude-plugin/scripts/verify.sh +108 -108
  389. package/LICENSE +21 -21
  390. package/README.md +419 -419
  391. package/bin/cli.js +11 -11
  392. package/bin/npx-repair.js +7 -7
  393. package/bin/npx-safe-launch.js +9 -9
  394. package/package.json +192 -186
  395. package/v3/@claude-flow/cli/README.md +419 -419
  396. package/v3/@claude-flow/cli/bin/cli.js +314 -314
  397. package/v3/@claude-flow/cli/bin/mcp-server.js +224 -224
  398. package/v3/@claude-flow/cli/bin/preinstall.cjs +2 -2
  399. package/v3/@claude-flow/cli/catalog-manifest.json +2 -2
  400. package/v3/@claude-flow/cli/dist/src/autopilot-state.js +24 -7
  401. package/v3/@claude-flow/cli/dist/src/benchmarks/gaia-critic.js +24 -24
  402. package/v3/@claude-flow/cli/dist/src/business-pods/bbs-budget-tracker.js +53 -53
  403. package/v3/@claude-flow/cli/dist/src/commands/completions.js +409 -409
  404. package/v3/@claude-flow/cli/dist/src/commands/daemon.js +44 -44
  405. package/v3/@claude-flow/cli/dist/src/commands/doctor.js +267 -20
  406. package/v3/@claude-flow/cli/dist/src/commands/embeddings.js +26 -26
  407. package/v3/@claude-flow/cli/dist/src/commands/hive-mind.js +97 -97
  408. package/v3/@claude-flow/cli/dist/src/commands/hooks.js +74 -11
  409. package/v3/@claude-flow/cli/dist/src/commands/init.js +202 -34
  410. package/v3/@claude-flow/cli/dist/src/commands/memory.js +12 -1
  411. package/v3/@claude-flow/cli/dist/src/commands/ruvector/backup.js +23 -23
  412. package/v3/@claude-flow/cli/dist/src/commands/ruvector/benchmark.js +31 -31
  413. package/v3/@claude-flow/cli/dist/src/commands/ruvector/import.js +14 -14
  414. package/v3/@claude-flow/cli/dist/src/commands/ruvector/init.js +115 -115
  415. package/v3/@claude-flow/cli/dist/src/commands/ruvector/migrate.js +99 -99
  416. package/v3/@claude-flow/cli/dist/src/commands/ruvector/optimize.js +51 -51
  417. package/v3/@claude-flow/cli/dist/src/commands/ruvector/setup.js +624 -624
  418. package/v3/@claude-flow/cli/dist/src/commands/ruvector/status.js +38 -38
  419. package/v3/@claude-flow/cli/dist/src/config/proven-config.js +2 -2
  420. package/v3/@claude-flow/cli/dist/src/funnel/disclosure.js +13 -2
  421. package/v3/@claude-flow/cli/dist/src/funnel/message-transport.d.ts +11 -4
  422. package/v3/@claude-flow/cli/dist/src/funnel/message-transport.js +11 -4
  423. package/v3/@claude-flow/cli/dist/src/funnel/messages.d.ts +12 -10
  424. package/v3/@claude-flow/cli/dist/src/funnel/messages.js +83 -11
  425. package/v3/@claude-flow/cli/dist/src/init/claudemd-generator.js +231 -231
  426. package/v3/@claude-flow/cli/dist/src/init/executor.js +453 -453
  427. package/v3/@claude-flow/cli/dist/src/init/helper-signing.js +2 -2
  428. package/v3/@claude-flow/cli/dist/src/init/helpers-generator.js +751 -751
  429. package/v3/@claude-flow/cli/dist/src/init/statusline-generator.js +24 -24
  430. package/v3/@claude-flow/cli/dist/src/mcp-tools/agentdb-tools.js +15 -15
  431. package/v3/@claude-flow/cli/dist/src/mcp-tools/browser-intent-tools.js +19 -19
  432. package/v3/@claude-flow/cli/dist/src/mcp-tools/browser-tools.js +8 -0
  433. package/v3/@claude-flow/cli/dist/src/mcp-tools/hooks-tools.js +21 -0
  434. package/v3/@claude-flow/cli/dist/src/mcp-tools/memory-tools.js +4 -3
  435. package/v3/@claude-flow/cli/dist/src/memory/graph-edge-writer.d.ts +9 -0
  436. package/v3/@claude-flow/cli/dist/src/memory/graph-edge-writer.js +35 -22
  437. package/v3/@claude-flow/cli/dist/src/memory/memory-bridge.js +189 -93
  438. package/v3/@claude-flow/cli/dist/src/memory/memory-initializer.js +479 -407
  439. package/v3/@claude-flow/cli/dist/src/memory/rabitq-index.js +5 -5
  440. package/v3/@claude-flow/cli/dist/src/parser.js +25 -9
  441. package/v3/@claude-flow/cli/dist/src/proxy/verify.js +2 -2
  442. package/v3/@claude-flow/cli/dist/src/runtime/headless.js +28 -28
  443. package/v3/@claude-flow/cli/dist/src/services/distill-tuning.js +7 -7
  444. package/v3/@claude-flow/cli/dist/src/services/headless-worker-executor.js +84 -84
  445. package/v3/@claude-flow/cli/dist/src/services/memory-distillation.js +4 -4
  446. package/v3/@claude-flow/cli/dist/src/services/worker-daemon.js +7 -4
  447. package/v3/@claude-flow/cli/dist/src/transfer/deploy-seraphine.js +23 -23
  448. package/v3/@claude-flow/cli/package.json +137 -135
  449. package/v3/@claude-flow/guidance/README.md +1195 -1195
  450. package/v3/@claude-flow/guidance/package.json +198 -198
  451. package/v3/@claude-flow/shared/README.md +323 -323
  452. package/v3/@claude-flow/shared/dist/events/event-store.js +31 -31
  453. package/v3/@claude-flow/shared/dist/hooks/safety/git-commit.js +3 -3
  454. package/v3/@claude-flow/shared/package.json +43 -43
  455. package/v3/README.md +493 -493
@@ -1,663 +1,663 @@
1
- ---
2
- name: Benchmark Suite
3
- description: Comprehensive performance benchmarking, regression detection and performance validation
4
- ---
5
-
6
- # Benchmark Suite Agent
7
-
8
- ## Agent Profile
9
- - **Name**: Benchmark Suite
10
- - **Type**: Performance Optimization Agent
11
- - **Specialization**: Comprehensive performance benchmarking and testing
12
- - **Performance Focus**: Automated benchmarking, regression detection, and performance validation
13
-
14
- ## Core Capabilities
15
-
16
- ### 1. Comprehensive Benchmarking Framework
17
- ```javascript
18
- // Advanced benchmarking system
19
- class ComprehensiveBenchmarkSuite {
20
- constructor() {
21
- this.benchmarks = {
22
- // Core performance benchmarks
23
- throughput: new ThroughputBenchmark(),
24
- latency: new LatencyBenchmark(),
25
- scalability: new ScalabilityBenchmark(),
26
- resource_usage: new ResourceUsageBenchmark(),
27
-
28
- // Swarm-specific benchmarks
29
- coordination: new CoordinationBenchmark(),
30
- load_balancing: new LoadBalancingBenchmark(),
31
- topology: new TopologyBenchmark(),
32
- fault_tolerance: new FaultToleranceBenchmark(),
33
-
34
- // Custom benchmarks
35
- custom: new CustomBenchmarkManager()
36
- };
37
-
38
- this.reporter = new BenchmarkReporter();
39
- this.comparator = new PerformanceComparator();
40
- this.analyzer = new BenchmarkAnalyzer();
41
- }
42
-
43
- // Execute comprehensive benchmark suite
44
- async runBenchmarkSuite(config = {}) {
45
- const suiteConfig = {
46
- duration: config.duration || 300000, // 5 minutes default
47
- iterations: config.iterations || 10,
48
- warmupTime: config.warmupTime || 30000, // 30 seconds
49
- cooldownTime: config.cooldownTime || 10000, // 10 seconds
50
- parallel: config.parallel || false,
51
- baseline: config.baseline || null
52
- };
53
-
54
- const results = {
55
- summary: {},
56
- detailed: new Map(),
57
- baseline_comparison: null,
58
- recommendations: []
59
- };
60
-
61
- // Warmup phase
62
- await this.warmup(suiteConfig.warmupTime);
63
-
64
- // Execute benchmarks
65
- if (suiteConfig.parallel) {
66
- results.detailed = await this.runBenchmarksParallel(suiteConfig);
67
- } else {
68
- results.detailed = await this.runBenchmarksSequential(suiteConfig);
69
- }
70
-
71
- // Generate summary
72
- results.summary = this.generateSummary(results.detailed);
73
-
74
- // Compare with baseline if provided
75
- if (suiteConfig.baseline) {
76
- results.baseline_comparison = await this.compareWithBaseline(
77
- results.detailed,
78
- suiteConfig.baseline
79
- );
80
- }
81
-
82
- // Generate recommendations
83
- results.recommendations = await this.generateRecommendations(results);
84
-
85
- // Cooldown phase
86
- await this.cooldown(suiteConfig.cooldownTime);
87
-
88
- return results;
89
- }
90
-
91
- // Parallel benchmark execution
92
- async runBenchmarksParallel(config) {
93
- const benchmarkPromises = Object.entries(this.benchmarks).map(
94
- async ([name, benchmark]) => {
95
- const result = await this.executeBenchmark(benchmark, name, config);
96
- return [name, result];
97
- }
98
- );
99
-
100
- const results = await Promise.all(benchmarkPromises);
101
- return new Map(results);
102
- }
103
-
104
- // Sequential benchmark execution
105
- async runBenchmarksSequential(config) {
106
- const results = new Map();
107
-
108
- for (const [name, benchmark] of Object.entries(this.benchmarks)) {
109
- const result = await this.executeBenchmark(benchmark, name, config);
110
- results.set(name, result);
111
-
112
- // Brief pause between benchmarks
113
- await this.sleep(1000);
114
- }
115
-
116
- return results;
117
- }
118
- }
119
- ```
120
-
121
- ### 2. Performance Regression Detection
122
- ```javascript
123
- // Advanced regression detection system
124
- class RegressionDetector {
125
- constructor() {
126
- this.detectors = {
127
- statistical: new StatisticalRegressionDetector(),
128
- machine_learning: new MLRegressionDetector(),
129
- threshold: new ThresholdRegressionDetector(),
130
- trend: new TrendRegressionDetector()
131
- };
132
-
133
- this.analyzer = new RegressionAnalyzer();
134
- this.alerting = new RegressionAlerting();
135
- }
136
-
137
- // Detect performance regressions
138
- async detectRegressions(currentResults, historicalData, config = {}) {
139
- const regressions = {
140
- detected: [],
141
- severity: 'none',
142
- confidence: 0,
143
- analysis: {}
144
- };
145
-
146
- // Run multiple detection algorithms
147
- const detectionPromises = Object.entries(this.detectors).map(
148
- async ([method, detector]) => {
149
- const detection = await detector.detect(currentResults, historicalData, config);
150
- return [method, detection];
151
- }
152
- );
153
-
154
- const detectionResults = await Promise.all(detectionPromises);
155
-
156
- // Aggregate detection results
157
- for (const [method, detection] of detectionResults) {
158
- if (detection.regression_detected) {
159
- regressions.detected.push({
160
- method,
161
- ...detection
162
- });
163
- }
164
- }
165
-
166
- // Calculate overall confidence and severity
167
- if (regressions.detected.length > 0) {
168
- regressions.confidence = this.calculateAggregateConfidence(regressions.detected);
169
- regressions.severity = this.calculateSeverity(regressions.detected);
170
- regressions.analysis = await this.analyzer.analyze(regressions.detected);
171
- }
172
-
173
- return regressions;
174
- }
175
-
176
- // Statistical regression detection using change point analysis
177
- async detectStatisticalRegression(metric, historicalData, sensitivity = 0.95) {
178
- // Use CUSUM (Cumulative Sum) algorithm for change point detection
179
- const cusum = this.calculateCUSUM(metric, historicalData);
180
-
181
- // Detect change points
182
- const changePoints = this.detectChangePoints(cusum, sensitivity);
183
-
184
- // Analyze significance of changes
185
- const analysis = changePoints.map(point => ({
186
- timestamp: point.timestamp,
187
- magnitude: point.magnitude,
188
- direction: point.direction,
189
- significance: point.significance,
190
- confidence: point.confidence
191
- }));
192
-
193
- return {
194
- regression_detected: changePoints.length > 0,
195
- change_points: analysis,
196
- cusum_statistics: cusum.statistics,
197
- sensitivity: sensitivity
198
- };
199
- }
200
-
201
- // Machine learning-based regression detection
202
- async detectMLRegression(metrics, historicalData) {
203
- // Train anomaly detection model on historical data
204
- const model = await this.trainAnomalyModel(historicalData);
205
-
206
- // Predict anomaly scores for current metrics
207
- const anomalyScores = await model.predict(metrics);
208
-
209
- // Identify regressions based on anomaly scores
210
- const threshold = this.calculateDynamicThreshold(anomalyScores);
211
- const regressions = anomalyScores.filter(score => score.anomaly > threshold);
212
-
213
- return {
214
- regression_detected: regressions.length > 0,
215
- anomaly_scores: anomalyScores,
216
- threshold: threshold,
217
- regressions: regressions,
218
- model_confidence: model.confidence
219
- };
220
- }
221
- }
222
- ```
223
-
224
- ### 3. Automated Performance Testing
225
- ```javascript
226
- // Comprehensive automated performance testing
227
- class AutomatedPerformanceTester {
228
- constructor() {
229
- this.testSuites = {
230
- load: new LoadTestSuite(),
231
- stress: new StressTestSuite(),
232
- volume: new VolumeTestSuite(),
233
- endurance: new EnduranceTestSuite(),
234
- spike: new SpikeTestSuite(),
235
- configuration: new ConfigurationTestSuite()
236
- };
237
-
238
- this.scheduler = new TestScheduler();
239
- this.orchestrator = new TestOrchestrator();
240
- this.validator = new ResultValidator();
241
- }
242
-
243
- // Execute automated performance test campaign
244
- async runTestCampaign(config) {
245
- const campaign = {
246
- id: this.generateCampaignId(),
247
- config,
248
- startTime: Date.now(),
249
- tests: [],
250
- results: new Map(),
251
- summary: null
252
- };
253
-
254
- // Schedule test execution
255
- const schedule = await this.scheduler.schedule(config.tests, config.constraints);
256
-
257
- // Execute tests according to schedule
258
- for (const scheduledTest of schedule) {
259
- const testResult = await this.executeScheduledTest(scheduledTest);
260
- campaign.tests.push(scheduledTest);
261
- campaign.results.set(scheduledTest.id, testResult);
262
-
263
- // Validate results in real-time
264
- const validation = await this.validator.validate(testResult);
265
- if (!validation.valid) {
266
- campaign.summary = {
267
- status: 'failed',
268
- reason: validation.reason,
269
- failedAt: scheduledTest.name
270
- };
271
- break;
272
- }
273
- }
274
-
275
- // Generate campaign summary
276
- if (!campaign.summary) {
277
- campaign.summary = await this.generateCampaignSummary(campaign);
278
- }
279
-
280
- campaign.endTime = Date.now();
281
- campaign.duration = campaign.endTime - campaign.startTime;
282
-
283
- return campaign;
284
- }
285
-
286
- // Load testing with gradual ramp-up
287
- async executeLoadTest(config) {
288
- const loadTest = {
289
- type: 'load',
290
- config,
291
- phases: [],
292
- metrics: new Map(),
293
- results: {}
294
- };
295
-
296
- // Ramp-up phase
297
- const rampUpResult = await this.executeRampUp(config.rampUp);
298
- loadTest.phases.push({ phase: 'ramp-up', result: rampUpResult });
299
-
300
- // Sustained load phase
301
- const sustainedResult = await this.executeSustainedLoad(config.sustained);
302
- loadTest.phases.push({ phase: 'sustained', result: sustainedResult });
303
-
304
- // Ramp-down phase
305
- const rampDownResult = await this.executeRampDown(config.rampDown);
306
- loadTest.phases.push({ phase: 'ramp-down', result: rampDownResult });
307
-
308
- // Analyze results
309
- loadTest.results = await this.analyzeLoadTestResults(loadTest.phases);
310
-
311
- return loadTest;
312
- }
313
-
314
- // Stress testing to find breaking points
315
- async executeStressTest(config) {
316
- const stressTest = {
317
- type: 'stress',
318
- config,
319
- breakingPoint: null,
320
- degradationCurve: [],
321
- results: {}
322
- };
323
-
324
- let currentLoad = config.startLoad;
325
- let systemBroken = false;
326
-
327
- while (!systemBroken && currentLoad <= config.maxLoad) {
328
- const testResult = await this.applyLoad(currentLoad, config.duration);
329
-
330
- stressTest.degradationCurve.push({
331
- load: currentLoad,
332
- performance: testResult.performance,
333
- stability: testResult.stability,
334
- errors: testResult.errors
335
- });
336
-
337
- // Check if system is breaking
338
- if (this.isSystemBreaking(testResult, config.breakingCriteria)) {
339
- stressTest.breakingPoint = {
340
- load: currentLoad,
341
- performance: testResult.performance,
342
- reason: this.identifyBreakingReason(testResult)
343
- };
344
- systemBroken = true;
345
- }
346
-
347
- currentLoad += config.loadIncrement;
348
- }
349
-
350
- stressTest.results = await this.analyzeStressTestResults(stressTest);
351
-
352
- return stressTest;
353
- }
354
- }
355
- ```
356
-
357
- ### 4. Performance Validation Framework
358
- ```javascript
359
- // Comprehensive performance validation
360
- class PerformanceValidator {
361
- constructor() {
362
- this.validators = {
363
- sla: new SLAValidator(),
364
- regression: new RegressionValidator(),
365
- scalability: new ScalabilityValidator(),
366
- reliability: new ReliabilityValidator(),
367
- efficiency: new EfficiencyValidator()
368
- };
369
-
370
- this.thresholds = new ThresholdManager();
371
- this.rules = new ValidationRuleEngine();
372
- }
373
-
374
- // Validate performance against defined criteria
375
- async validatePerformance(results, criteria) {
376
- const validation = {
377
- overall: {
378
- passed: true,
379
- score: 0,
380
- violations: []
381
- },
382
- detailed: new Map(),
383
- recommendations: []
384
- };
385
-
386
- // Run all validators
387
- const validationPromises = Object.entries(this.validators).map(
388
- async ([type, validator]) => {
389
- const result = await validator.validate(results, criteria[type]);
390
- return [type, result];
391
- }
392
- );
393
-
394
- const validationResults = await Promise.all(validationPromises);
395
-
396
- // Aggregate validation results
397
- for (const [type, result] of validationResults) {
398
- validation.detailed.set(type, result);
399
-
400
- if (!result.passed) {
401
- validation.overall.passed = false;
402
- validation.overall.violations.push(...result.violations);
403
- }
404
-
405
- validation.overall.score += result.score * (criteria[type]?.weight || 1);
406
- }
407
-
408
- // Normalize overall score
409
- const totalWeight = Object.values(criteria).reduce((sum, c) => sum + (c.weight || 1), 0);
410
- validation.overall.score /= totalWeight;
411
-
412
- // Generate recommendations
413
- validation.recommendations = await this.generateValidationRecommendations(validation);
414
-
415
- return validation;
416
- }
417
-
418
- // SLA validation
419
- async validateSLA(results, slaConfig) {
420
- const slaValidation = {
421
- passed: true,
422
- violations: [],
423
- score: 1.0,
424
- metrics: {}
425
- };
426
-
427
- // Validate each SLA metric
428
- for (const [metric, threshold] of Object.entries(slaConfig.thresholds)) {
429
- const actualValue = this.extractMetricValue(results, metric);
430
- const validation = this.validateThreshold(actualValue, threshold);
431
-
432
- slaValidation.metrics[metric] = {
433
- actual: actualValue,
434
- threshold: threshold.value,
435
- operator: threshold.operator,
436
- passed: validation.passed,
437
- deviation: validation.deviation
438
- };
439
-
440
- if (!validation.passed) {
441
- slaValidation.passed = false;
442
- slaValidation.violations.push({
443
- metric,
444
- actual: actualValue,
445
- expected: threshold.value,
446
- severity: threshold.severity || 'medium'
447
- });
448
-
449
- // Reduce score based on violation severity
450
- const severityMultiplier = this.getSeverityMultiplier(threshold.severity);
451
- slaValidation.score -= (validation.deviation * severityMultiplier);
452
- }
453
- }
454
-
455
- slaValidation.score = Math.max(0, slaValidation.score);
456
-
457
- return slaValidation;
458
- }
459
-
460
- // Scalability validation
461
- async validateScalability(results, scalabilityConfig) {
462
- const scalabilityValidation = {
463
- passed: true,
464
- violations: [],
465
- score: 1.0,
466
- analysis: {}
467
- };
468
-
469
- // Linear scalability analysis
470
- if (scalabilityConfig.linear) {
471
- const linearityAnalysis = this.analyzeLinearScalability(results);
472
- scalabilityValidation.analysis.linearity = linearityAnalysis;
473
-
474
- if (linearityAnalysis.coefficient < scalabilityConfig.linear.minCoefficient) {
475
- scalabilityValidation.passed = false;
476
- scalabilityValidation.violations.push({
477
- type: 'linearity',
478
- actual: linearityAnalysis.coefficient,
479
- expected: scalabilityConfig.linear.minCoefficient
480
- });
481
- }
482
- }
483
-
484
- // Efficiency retention analysis
485
- if (scalabilityConfig.efficiency) {
486
- const efficiencyAnalysis = this.analyzeEfficiencyRetention(results);
487
- scalabilityValidation.analysis.efficiency = efficiencyAnalysis;
488
-
489
- if (efficiencyAnalysis.retention < scalabilityConfig.efficiency.minRetention) {
490
- scalabilityValidation.passed = false;
491
- scalabilityValidation.violations.push({
492
- type: 'efficiency_retention',
493
- actual: efficiencyAnalysis.retention,
494
- expected: scalabilityConfig.efficiency.minRetention
495
- });
496
- }
497
- }
498
-
499
- return scalabilityValidation;
500
- }
501
- }
502
- ```
503
-
504
- ## MCP Integration Hooks
505
-
506
- ### Benchmark Execution Integration
507
- ```javascript
508
- // Comprehensive MCP benchmark integration
509
- const benchmarkIntegration = {
510
- // Execute performance benchmarks
511
- async runBenchmarks(config = {}) {
512
- // Run benchmark suite
513
- const benchmarkResult = await mcp.benchmark_run({
514
- suite: config.suite || 'comprehensive'
515
- });
516
-
517
- // Collect detailed metrics during benchmarking
518
- const metrics = await mcp.metrics_collect({
519
- components: ['system', 'agents', 'coordination', 'memory']
520
- });
521
-
522
- // Analyze performance trends
523
- const trends = await mcp.trend_analysis({
524
- metric: 'performance',
525
- period: '24h'
526
- });
527
-
528
- // Cost analysis
529
- const costAnalysis = await mcp.cost_analysis({
530
- timeframe: '24h'
531
- });
532
-
533
- return {
534
- benchmark: benchmarkResult,
535
- metrics,
536
- trends,
537
- costAnalysis,
538
- timestamp: Date.now()
539
- };
540
- },
541
-
542
- // Quality assessment
543
- async assessQuality(criteria) {
544
- const qualityAssessment = await mcp.quality_assess({
545
- target: 'swarm-performance',
546
- criteria: criteria || [
547
- 'throughput',
548
- 'latency',
549
- 'reliability',
550
- 'scalability',
551
- 'efficiency'
552
- ]
553
- });
554
-
555
- return qualityAssessment;
556
- },
557
-
558
- // Error pattern analysis
559
- async analyzeErrorPatterns() {
560
- // Collect system logs
561
- const logs = await this.collectSystemLogs();
562
-
563
- // Analyze error patterns
564
- const errorAnalysis = await mcp.error_analysis({
565
- logs: logs
566
- });
567
-
568
- return errorAnalysis;
569
- }
570
- };
571
- ```
572
-
573
- ## Operational Commands
574
-
575
- ### Benchmarking Commands
576
- ```bash
577
- # Run comprehensive benchmark suite
578
- npx claude-flow benchmark-run --suite comprehensive --duration 300
579
-
580
- # Execute specific benchmark
581
- npx claude-flow benchmark-run --suite throughput --iterations 10
582
-
583
- # Compare with baseline
584
- npx claude-flow benchmark-compare --current <results> --baseline <baseline>
585
-
586
- # Quality assessment
587
- npx claude-flow quality-assess --target swarm-performance --criteria throughput,latency
588
-
589
- # Performance validation
590
- npx claude-flow validate-performance --results <file> --criteria <file>
591
- ```
592
-
593
- ### Regression Detection Commands
594
- ```bash
595
- # Detect performance regressions
596
- npx claude-flow detect-regression --current <results> --historical <data>
597
-
598
- # Set up automated regression monitoring
599
- npx claude-flow regression-monitor --enable --sensitivity 0.95
600
-
601
- # Analyze error patterns
602
- npx claude-flow error-analysis --logs <log-files>
603
- ```
604
-
605
- ## Integration Points
606
-
607
- ### With Other Optimization Agents
608
- - **Performance Monitor**: Provides continuous monitoring data for benchmarking
609
- - **Load Balancer**: Validates load balancing effectiveness through benchmarks
610
- - **Topology Optimizer**: Tests topology configurations for optimal performance
611
-
612
- ### With CI/CD Pipeline
613
- - **Automated Testing**: Integrates with CI/CD for continuous performance validation
614
- - **Quality Gates**: Provides pass/fail criteria for deployment decisions
615
- - **Regression Prevention**: Catches performance regressions before production
616
-
617
- ## Performance Benchmarks
618
-
619
- ### Standard Benchmark Suite
620
- ```javascript
621
- // Comprehensive benchmark definitions
622
- const standardBenchmarks = {
623
- // Throughput benchmarks
624
- throughput: {
625
- name: 'Throughput Benchmark',
626
- metrics: ['requests_per_second', 'tasks_per_second', 'messages_per_second'],
627
- duration: 300000, // 5 minutes
628
- warmup: 30000, // 30 seconds
629
- targets: {
630
- requests_per_second: { min: 1000, optimal: 5000 },
631
- tasks_per_second: { min: 100, optimal: 500 },
632
- messages_per_second: { min: 10000, optimal: 50000 }
633
- }
634
- },
635
-
636
- // Latency benchmarks
637
- latency: {
638
- name: 'Latency Benchmark',
639
- metrics: ['p50', 'p90', 'p95', 'p99', 'max'],
640
- duration: 300000,
641
- targets: {
642
- p50: { max: 100 }, // 100ms
643
- p90: { max: 200 }, // 200ms
644
- p95: { max: 500 }, // 500ms
645
- p99: { max: 1000 }, // 1s
646
- max: { max: 5000 } // 5s
647
- }
648
- },
649
-
650
- // Scalability benchmarks
651
- scalability: {
652
- name: 'Scalability Benchmark',
653
- metrics: ['linear_coefficient', 'efficiency_retention'],
654
- load_points: [1, 2, 4, 8, 16, 32, 64],
655
- targets: {
656
- linear_coefficient: { min: 0.8 },
657
- efficiency_retention: { min: 0.7 }
658
- }
659
- }
660
- };
661
- ```
662
-
1
+ ---
2
+ name: Benchmark Suite
3
+ description: Comprehensive performance benchmarking, regression detection and performance validation
4
+ ---
5
+
6
+ # Benchmark Suite Agent
7
+
8
+ ## Agent Profile
9
+ - **Name**: Benchmark Suite
10
+ - **Type**: Performance Optimization Agent
11
+ - **Specialization**: Comprehensive performance benchmarking and testing
12
+ - **Performance Focus**: Automated benchmarking, regression detection, and performance validation
13
+
14
+ ## Core Capabilities
15
+
16
+ ### 1. Comprehensive Benchmarking Framework
17
+ ```javascript
18
+ // Advanced benchmarking system
19
+ class ComprehensiveBenchmarkSuite {
20
+ constructor() {
21
+ this.benchmarks = {
22
+ // Core performance benchmarks
23
+ throughput: new ThroughputBenchmark(),
24
+ latency: new LatencyBenchmark(),
25
+ scalability: new ScalabilityBenchmark(),
26
+ resource_usage: new ResourceUsageBenchmark(),
27
+
28
+ // Swarm-specific benchmarks
29
+ coordination: new CoordinationBenchmark(),
30
+ load_balancing: new LoadBalancingBenchmark(),
31
+ topology: new TopologyBenchmark(),
32
+ fault_tolerance: new FaultToleranceBenchmark(),
33
+
34
+ // Custom benchmarks
35
+ custom: new CustomBenchmarkManager()
36
+ };
37
+
38
+ this.reporter = new BenchmarkReporter();
39
+ this.comparator = new PerformanceComparator();
40
+ this.analyzer = new BenchmarkAnalyzer();
41
+ }
42
+
43
+ // Execute comprehensive benchmark suite
44
+ async runBenchmarkSuite(config = {}) {
45
+ const suiteConfig = {
46
+ duration: config.duration || 300000, // 5 minutes default
47
+ iterations: config.iterations || 10,
48
+ warmupTime: config.warmupTime || 30000, // 30 seconds
49
+ cooldownTime: config.cooldownTime || 10000, // 10 seconds
50
+ parallel: config.parallel || false,
51
+ baseline: config.baseline || null
52
+ };
53
+
54
+ const results = {
55
+ summary: {},
56
+ detailed: new Map(),
57
+ baseline_comparison: null,
58
+ recommendations: []
59
+ };
60
+
61
+ // Warmup phase
62
+ await this.warmup(suiteConfig.warmupTime);
63
+
64
+ // Execute benchmarks
65
+ if (suiteConfig.parallel) {
66
+ results.detailed = await this.runBenchmarksParallel(suiteConfig);
67
+ } else {
68
+ results.detailed = await this.runBenchmarksSequential(suiteConfig);
69
+ }
70
+
71
+ // Generate summary
72
+ results.summary = this.generateSummary(results.detailed);
73
+
74
+ // Compare with baseline if provided
75
+ if (suiteConfig.baseline) {
76
+ results.baseline_comparison = await this.compareWithBaseline(
77
+ results.detailed,
78
+ suiteConfig.baseline
79
+ );
80
+ }
81
+
82
+ // Generate recommendations
83
+ results.recommendations = await this.generateRecommendations(results);
84
+
85
+ // Cooldown phase
86
+ await this.cooldown(suiteConfig.cooldownTime);
87
+
88
+ return results;
89
+ }
90
+
91
+ // Parallel benchmark execution
92
+ async runBenchmarksParallel(config) {
93
+ const benchmarkPromises = Object.entries(this.benchmarks).map(
94
+ async ([name, benchmark]) => {
95
+ const result = await this.executeBenchmark(benchmark, name, config);
96
+ return [name, result];
97
+ }
98
+ );
99
+
100
+ const results = await Promise.all(benchmarkPromises);
101
+ return new Map(results);
102
+ }
103
+
104
+ // Sequential benchmark execution
105
+ async runBenchmarksSequential(config) {
106
+ const results = new Map();
107
+
108
+ for (const [name, benchmark] of Object.entries(this.benchmarks)) {
109
+ const result = await this.executeBenchmark(benchmark, name, config);
110
+ results.set(name, result);
111
+
112
+ // Brief pause between benchmarks
113
+ await this.sleep(1000);
114
+ }
115
+
116
+ return results;
117
+ }
118
+ }
119
+ ```
120
+
121
+ ### 2. Performance Regression Detection
122
+ ```javascript
123
+ // Advanced regression detection system
124
+ class RegressionDetector {
125
+ constructor() {
126
+ this.detectors = {
127
+ statistical: new StatisticalRegressionDetector(),
128
+ machine_learning: new MLRegressionDetector(),
129
+ threshold: new ThresholdRegressionDetector(),
130
+ trend: new TrendRegressionDetector()
131
+ };
132
+
133
+ this.analyzer = new RegressionAnalyzer();
134
+ this.alerting = new RegressionAlerting();
135
+ }
136
+
137
+ // Detect performance regressions
138
+ async detectRegressions(currentResults, historicalData, config = {}) {
139
+ const regressions = {
140
+ detected: [],
141
+ severity: 'none',
142
+ confidence: 0,
143
+ analysis: {}
144
+ };
145
+
146
+ // Run multiple detection algorithms
147
+ const detectionPromises = Object.entries(this.detectors).map(
148
+ async ([method, detector]) => {
149
+ const detection = await detector.detect(currentResults, historicalData, config);
150
+ return [method, detection];
151
+ }
152
+ );
153
+
154
+ const detectionResults = await Promise.all(detectionPromises);
155
+
156
+ // Aggregate detection results
157
+ for (const [method, detection] of detectionResults) {
158
+ if (detection.regression_detected) {
159
+ regressions.detected.push({
160
+ method,
161
+ ...detection
162
+ });
163
+ }
164
+ }
165
+
166
+ // Calculate overall confidence and severity
167
+ if (regressions.detected.length > 0) {
168
+ regressions.confidence = this.calculateAggregateConfidence(regressions.detected);
169
+ regressions.severity = this.calculateSeverity(regressions.detected);
170
+ regressions.analysis = await this.analyzer.analyze(regressions.detected);
171
+ }
172
+
173
+ return regressions;
174
+ }
175
+
176
+ // Statistical regression detection using change point analysis
177
+ async detectStatisticalRegression(metric, historicalData, sensitivity = 0.95) {
178
+ // Use CUSUM (Cumulative Sum) algorithm for change point detection
179
+ const cusum = this.calculateCUSUM(metric, historicalData);
180
+
181
+ // Detect change points
182
+ const changePoints = this.detectChangePoints(cusum, sensitivity);
183
+
184
+ // Analyze significance of changes
185
+ const analysis = changePoints.map(point => ({
186
+ timestamp: point.timestamp,
187
+ magnitude: point.magnitude,
188
+ direction: point.direction,
189
+ significance: point.significance,
190
+ confidence: point.confidence
191
+ }));
192
+
193
+ return {
194
+ regression_detected: changePoints.length > 0,
195
+ change_points: analysis,
196
+ cusum_statistics: cusum.statistics,
197
+ sensitivity: sensitivity
198
+ };
199
+ }
200
+
201
+ // Machine learning-based regression detection
202
+ async detectMLRegression(metrics, historicalData) {
203
+ // Train anomaly detection model on historical data
204
+ const model = await this.trainAnomalyModel(historicalData);
205
+
206
+ // Predict anomaly scores for current metrics
207
+ const anomalyScores = await model.predict(metrics);
208
+
209
+ // Identify regressions based on anomaly scores
210
+ const threshold = this.calculateDynamicThreshold(anomalyScores);
211
+ const regressions = anomalyScores.filter(score => score.anomaly > threshold);
212
+
213
+ return {
214
+ regression_detected: regressions.length > 0,
215
+ anomaly_scores: anomalyScores,
216
+ threshold: threshold,
217
+ regressions: regressions,
218
+ model_confidence: model.confidence
219
+ };
220
+ }
221
+ }
222
+ ```
223
+
224
+ ### 3. Automated Performance Testing
225
+ ```javascript
226
+ // Comprehensive automated performance testing
227
+ class AutomatedPerformanceTester {
228
+ constructor() {
229
+ this.testSuites = {
230
+ load: new LoadTestSuite(),
231
+ stress: new StressTestSuite(),
232
+ volume: new VolumeTestSuite(),
233
+ endurance: new EnduranceTestSuite(),
234
+ spike: new SpikeTestSuite(),
235
+ configuration: new ConfigurationTestSuite()
236
+ };
237
+
238
+ this.scheduler = new TestScheduler();
239
+ this.orchestrator = new TestOrchestrator();
240
+ this.validator = new ResultValidator();
241
+ }
242
+
243
+ // Execute automated performance test campaign
244
+ async runTestCampaign(config) {
245
+ const campaign = {
246
+ id: this.generateCampaignId(),
247
+ config,
248
+ startTime: Date.now(),
249
+ tests: [],
250
+ results: new Map(),
251
+ summary: null
252
+ };
253
+
254
+ // Schedule test execution
255
+ const schedule = await this.scheduler.schedule(config.tests, config.constraints);
256
+
257
+ // Execute tests according to schedule
258
+ for (const scheduledTest of schedule) {
259
+ const testResult = await this.executeScheduledTest(scheduledTest);
260
+ campaign.tests.push(scheduledTest);
261
+ campaign.results.set(scheduledTest.id, testResult);
262
+
263
+ // Validate results in real-time
264
+ const validation = await this.validator.validate(testResult);
265
+ if (!validation.valid) {
266
+ campaign.summary = {
267
+ status: 'failed',
268
+ reason: validation.reason,
269
+ failedAt: scheduledTest.name
270
+ };
271
+ break;
272
+ }
273
+ }
274
+
275
+ // Generate campaign summary
276
+ if (!campaign.summary) {
277
+ campaign.summary = await this.generateCampaignSummary(campaign);
278
+ }
279
+
280
+ campaign.endTime = Date.now();
281
+ campaign.duration = campaign.endTime - campaign.startTime;
282
+
283
+ return campaign;
284
+ }
285
+
286
+ // Load testing with gradual ramp-up
287
+ async executeLoadTest(config) {
288
+ const loadTest = {
289
+ type: 'load',
290
+ config,
291
+ phases: [],
292
+ metrics: new Map(),
293
+ results: {}
294
+ };
295
+
296
+ // Ramp-up phase
297
+ const rampUpResult = await this.executeRampUp(config.rampUp);
298
+ loadTest.phases.push({ phase: 'ramp-up', result: rampUpResult });
299
+
300
+ // Sustained load phase
301
+ const sustainedResult = await this.executeSustainedLoad(config.sustained);
302
+ loadTest.phases.push({ phase: 'sustained', result: sustainedResult });
303
+
304
+ // Ramp-down phase
305
+ const rampDownResult = await this.executeRampDown(config.rampDown);
306
+ loadTest.phases.push({ phase: 'ramp-down', result: rampDownResult });
307
+
308
+ // Analyze results
309
+ loadTest.results = await this.analyzeLoadTestResults(loadTest.phases);
310
+
311
+ return loadTest;
312
+ }
313
+
314
+ // Stress testing to find breaking points
315
+ async executeStressTest(config) {
316
+ const stressTest = {
317
+ type: 'stress',
318
+ config,
319
+ breakingPoint: null,
320
+ degradationCurve: [],
321
+ results: {}
322
+ };
323
+
324
+ let currentLoad = config.startLoad;
325
+ let systemBroken = false;
326
+
327
+ while (!systemBroken && currentLoad <= config.maxLoad) {
328
+ const testResult = await this.applyLoad(currentLoad, config.duration);
329
+
330
+ stressTest.degradationCurve.push({
331
+ load: currentLoad,
332
+ performance: testResult.performance,
333
+ stability: testResult.stability,
334
+ errors: testResult.errors
335
+ });
336
+
337
+ // Check if system is breaking
338
+ if (this.isSystemBreaking(testResult, config.breakingCriteria)) {
339
+ stressTest.breakingPoint = {
340
+ load: currentLoad,
341
+ performance: testResult.performance,
342
+ reason: this.identifyBreakingReason(testResult)
343
+ };
344
+ systemBroken = true;
345
+ }
346
+
347
+ currentLoad += config.loadIncrement;
348
+ }
349
+
350
+ stressTest.results = await this.analyzeStressTestResults(stressTest);
351
+
352
+ return stressTest;
353
+ }
354
+ }
355
+ ```
356
+
357
+ ### 4. Performance Validation Framework
358
+ ```javascript
359
+ // Comprehensive performance validation
360
+ class PerformanceValidator {
361
+ constructor() {
362
+ this.validators = {
363
+ sla: new SLAValidator(),
364
+ regression: new RegressionValidator(),
365
+ scalability: new ScalabilityValidator(),
366
+ reliability: new ReliabilityValidator(),
367
+ efficiency: new EfficiencyValidator()
368
+ };
369
+
370
+ this.thresholds = new ThresholdManager();
371
+ this.rules = new ValidationRuleEngine();
372
+ }
373
+
374
+ // Validate performance against defined criteria
375
+ async validatePerformance(results, criteria) {
376
+ const validation = {
377
+ overall: {
378
+ passed: true,
379
+ score: 0,
380
+ violations: []
381
+ },
382
+ detailed: new Map(),
383
+ recommendations: []
384
+ };
385
+
386
+ // Run all validators
387
+ const validationPromises = Object.entries(this.validators).map(
388
+ async ([type, validator]) => {
389
+ const result = await validator.validate(results, criteria[type]);
390
+ return [type, result];
391
+ }
392
+ );
393
+
394
+ const validationResults = await Promise.all(validationPromises);
395
+
396
+ // Aggregate validation results
397
+ for (const [type, result] of validationResults) {
398
+ validation.detailed.set(type, result);
399
+
400
+ if (!result.passed) {
401
+ validation.overall.passed = false;
402
+ validation.overall.violations.push(...result.violations);
403
+ }
404
+
405
+ validation.overall.score += result.score * (criteria[type]?.weight || 1);
406
+ }
407
+
408
+ // Normalize overall score
409
+ const totalWeight = Object.values(criteria).reduce((sum, c) => sum + (c.weight || 1), 0);
410
+ validation.overall.score /= totalWeight;
411
+
412
+ // Generate recommendations
413
+ validation.recommendations = await this.generateValidationRecommendations(validation);
414
+
415
+ return validation;
416
+ }
417
+
418
+ // SLA validation
419
+ async validateSLA(results, slaConfig) {
420
+ const slaValidation = {
421
+ passed: true,
422
+ violations: [],
423
+ score: 1.0,
424
+ metrics: {}
425
+ };
426
+
427
+ // Validate each SLA metric
428
+ for (const [metric, threshold] of Object.entries(slaConfig.thresholds)) {
429
+ const actualValue = this.extractMetricValue(results, metric);
430
+ const validation = this.validateThreshold(actualValue, threshold);
431
+
432
+ slaValidation.metrics[metric] = {
433
+ actual: actualValue,
434
+ threshold: threshold.value,
435
+ operator: threshold.operator,
436
+ passed: validation.passed,
437
+ deviation: validation.deviation
438
+ };
439
+
440
+ if (!validation.passed) {
441
+ slaValidation.passed = false;
442
+ slaValidation.violations.push({
443
+ metric,
444
+ actual: actualValue,
445
+ expected: threshold.value,
446
+ severity: threshold.severity || 'medium'
447
+ });
448
+
449
+ // Reduce score based on violation severity
450
+ const severityMultiplier = this.getSeverityMultiplier(threshold.severity);
451
+ slaValidation.score -= (validation.deviation * severityMultiplier);
452
+ }
453
+ }
454
+
455
+ slaValidation.score = Math.max(0, slaValidation.score);
456
+
457
+ return slaValidation;
458
+ }
459
+
460
+ // Scalability validation
461
+ async validateScalability(results, scalabilityConfig) {
462
+ const scalabilityValidation = {
463
+ passed: true,
464
+ violations: [],
465
+ score: 1.0,
466
+ analysis: {}
467
+ };
468
+
469
+ // Linear scalability analysis
470
+ if (scalabilityConfig.linear) {
471
+ const linearityAnalysis = this.analyzeLinearScalability(results);
472
+ scalabilityValidation.analysis.linearity = linearityAnalysis;
473
+
474
+ if (linearityAnalysis.coefficient < scalabilityConfig.linear.minCoefficient) {
475
+ scalabilityValidation.passed = false;
476
+ scalabilityValidation.violations.push({
477
+ type: 'linearity',
478
+ actual: linearityAnalysis.coefficient,
479
+ expected: scalabilityConfig.linear.minCoefficient
480
+ });
481
+ }
482
+ }
483
+
484
+ // Efficiency retention analysis
485
+ if (scalabilityConfig.efficiency) {
486
+ const efficiencyAnalysis = this.analyzeEfficiencyRetention(results);
487
+ scalabilityValidation.analysis.efficiency = efficiencyAnalysis;
488
+
489
+ if (efficiencyAnalysis.retention < scalabilityConfig.efficiency.minRetention) {
490
+ scalabilityValidation.passed = false;
491
+ scalabilityValidation.violations.push({
492
+ type: 'efficiency_retention',
493
+ actual: efficiencyAnalysis.retention,
494
+ expected: scalabilityConfig.efficiency.minRetention
495
+ });
496
+ }
497
+ }
498
+
499
+ return scalabilityValidation;
500
+ }
501
+ }
502
+ ```
503
+
504
+ ## MCP Integration Hooks
505
+
506
+ ### Benchmark Execution Integration
507
+ ```javascript
508
+ // Comprehensive MCP benchmark integration
509
+ const benchmarkIntegration = {
510
+ // Execute performance benchmarks
511
+ async runBenchmarks(config = {}) {
512
+ // Run benchmark suite
513
+ const benchmarkResult = await mcp.benchmark_run({
514
+ suite: config.suite || 'comprehensive'
515
+ });
516
+
517
+ // Collect detailed metrics during benchmarking
518
+ const metrics = await mcp.metrics_collect({
519
+ components: ['system', 'agents', 'coordination', 'memory']
520
+ });
521
+
522
+ // Analyze performance trends
523
+ const trends = await mcp.trend_analysis({
524
+ metric: 'performance',
525
+ period: '24h'
526
+ });
527
+
528
+ // Cost analysis
529
+ const costAnalysis = await mcp.cost_analysis({
530
+ timeframe: '24h'
531
+ });
532
+
533
+ return {
534
+ benchmark: benchmarkResult,
535
+ metrics,
536
+ trends,
537
+ costAnalysis,
538
+ timestamp: Date.now()
539
+ };
540
+ },
541
+
542
+ // Quality assessment
543
+ async assessQuality(criteria) {
544
+ const qualityAssessment = await mcp.quality_assess({
545
+ target: 'swarm-performance',
546
+ criteria: criteria || [
547
+ 'throughput',
548
+ 'latency',
549
+ 'reliability',
550
+ 'scalability',
551
+ 'efficiency'
552
+ ]
553
+ });
554
+
555
+ return qualityAssessment;
556
+ },
557
+
558
+ // Error pattern analysis
559
+ async analyzeErrorPatterns() {
560
+ // Collect system logs
561
+ const logs = await this.collectSystemLogs();
562
+
563
+ // Analyze error patterns
564
+ const errorAnalysis = await mcp.error_analysis({
565
+ logs: logs
566
+ });
567
+
568
+ return errorAnalysis;
569
+ }
570
+ };
571
+ ```
572
+
573
+ ## Operational Commands
574
+
575
+ ### Benchmarking Commands
576
+ ```bash
577
+ # Run comprehensive benchmark suite
578
+ npx claude-flow benchmark-run --suite comprehensive --duration 300
579
+
580
+ # Execute specific benchmark
581
+ npx claude-flow benchmark-run --suite throughput --iterations 10
582
+
583
+ # Compare with baseline
584
+ npx claude-flow benchmark-compare --current <results> --baseline <baseline>
585
+
586
+ # Quality assessment
587
+ npx claude-flow quality-assess --target swarm-performance --criteria throughput,latency
588
+
589
+ # Performance validation
590
+ npx claude-flow validate-performance --results <file> --criteria <file>
591
+ ```
592
+
593
+ ### Regression Detection Commands
594
+ ```bash
595
+ # Detect performance regressions
596
+ npx claude-flow detect-regression --current <results> --historical <data>
597
+
598
+ # Set up automated regression monitoring
599
+ npx claude-flow regression-monitor --enable --sensitivity 0.95
600
+
601
+ # Analyze error patterns
602
+ npx claude-flow error-analysis --logs <log-files>
603
+ ```
604
+
605
+ ## Integration Points
606
+
607
+ ### With Other Optimization Agents
608
+ - **Performance Monitor**: Provides continuous monitoring data for benchmarking
609
+ - **Load Balancer**: Validates load balancing effectiveness through benchmarks
610
+ - **Topology Optimizer**: Tests topology configurations for optimal performance
611
+
612
+ ### With CI/CD Pipeline
613
+ - **Automated Testing**: Integrates with CI/CD for continuous performance validation
614
+ - **Quality Gates**: Provides pass/fail criteria for deployment decisions
615
+ - **Regression Prevention**: Catches performance regressions before production
616
+
617
+ ## Performance Benchmarks
618
+
619
+ ### Standard Benchmark Suite
620
+ ```javascript
621
+ // Comprehensive benchmark definitions
622
+ const standardBenchmarks = {
623
+ // Throughput benchmarks
624
+ throughput: {
625
+ name: 'Throughput Benchmark',
626
+ metrics: ['requests_per_second', 'tasks_per_second', 'messages_per_second'],
627
+ duration: 300000, // 5 minutes
628
+ warmup: 30000, // 30 seconds
629
+ targets: {
630
+ requests_per_second: { min: 1000, optimal: 5000 },
631
+ tasks_per_second: { min: 100, optimal: 500 },
632
+ messages_per_second: { min: 10000, optimal: 50000 }
633
+ }
634
+ },
635
+
636
+ // Latency benchmarks
637
+ latency: {
638
+ name: 'Latency Benchmark',
639
+ metrics: ['p50', 'p90', 'p95', 'p99', 'max'],
640
+ duration: 300000,
641
+ targets: {
642
+ p50: { max: 100 }, // 100ms
643
+ p90: { max: 200 }, // 200ms
644
+ p95: { max: 500 }, // 500ms
645
+ p99: { max: 1000 }, // 1s
646
+ max: { max: 5000 } // 5s
647
+ }
648
+ },
649
+
650
+ // Scalability benchmarks
651
+ scalability: {
652
+ name: 'Scalability Benchmark',
653
+ metrics: ['linear_coefficient', 'efficiency_retention'],
654
+ load_points: [1, 2, 4, 8, 16, 32, 64],
655
+ targets: {
656
+ linear_coefficient: { min: 0.8 },
657
+ efficiency_retention: { min: 0.7 }
658
+ }
659
+ }
660
+ };
661
+ ```
662
+
663
663
  This Benchmark Suite agent provides comprehensive automated performance testing, regression detection, and validation capabilities to ensure optimal swarm performance and prevent performance degradation.