loopx 0.4.8__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (811) hide show
  1. loopx/__init__.py +5 -0
  2. loopx/agent_onboarding.py +654 -0
  3. loopx/agent_registry.py +112 -0
  4. loopx/ark_managed_agent_host.py +59 -0
  5. loopx/authority.py +805 -0
  6. loopx/benchmark.py +2875 -0
  7. loopx/benchmark_adapters/__init__.py +1 -0
  8. loopx/benchmark_adapters/agentissue.py +2644 -0
  9. loopx/benchmark_adapters/agents_last_exam.py +3998 -0
  10. loopx/benchmark_adapters/edgebench.py +322 -0
  11. loopx/benchmark_adapters/skillsbench.py +5978 -0
  12. loopx/benchmark_adapters/skillsbench_acp_failure_policy.py +143 -0
  13. loopx/benchmark_adapters/skillsbench_acp_process.py +31 -0
  14. loopx/benchmark_adapters/skillsbench_acp_relay.py +4832 -0
  15. loopx/benchmark_adapters/skillsbench_batch.py +124 -0
  16. loopx/benchmark_adapters/skillsbench_bridge_guard.py +209 -0
  17. loopx/benchmark_adapters/skillsbench_bridge_summary.py +203 -0
  18. loopx/benchmark_adapters/skillsbench_codex_goal_recovery.py +271 -0
  19. loopx/benchmark_adapters/skillsbench_codex_goal_trace.py +81 -0
  20. loopx/benchmark_adapters/skillsbench_codex_runtime.py +339 -0
  21. loopx/benchmark_adapters/skillsbench_dockerfile_runtime.py +467 -0
  22. loopx/benchmark_adapters/skillsbench_failure_signals.py +652 -0
  23. loopx/benchmark_adapters/skillsbench_proxy_runtime.py +327 -0
  24. loopx/benchmark_adapters/skillsbench_remote_bridge.py +402 -0
  25. loopx/benchmark_adapters/skillsbench_result_discovery.py +143 -0
  26. loopx/benchmark_adapters/skillsbench_runner_profile.py +436 -0
  27. loopx/benchmark_adapters/skillsbench_runner_source.py +99 -0
  28. loopx/benchmark_adapters/skillsbench_setup_preflight.py +771 -0
  29. loopx/benchmark_adapters/skillsbench_signals.py +15 -0
  30. loopx/benchmark_adapters/skillsbench_task_source.py +141 -0
  31. loopx/benchmark_adapters/skillsbench_turn_route.py +723 -0
  32. loopx/benchmark_adapters/skillsbench_turn_runtime.py +1069 -0
  33. loopx/benchmark_adapters/skillsbench_typed_repair.py +689 -0
  34. loopx/benchmark_adapters/skillsbench_uv_cache.py +111 -0
  35. loopx/benchmark_adapters/skillsbench_verifier_bootstrap.py +227 -0
  36. loopx/benchmark_adapters/skillsbench_verifier_cache.py +138 -0
  37. loopx/benchmark_adapters/terminal_bench.py +10078 -0
  38. loopx/benchmark_case_analysis.py +1276 -0
  39. loopx/benchmark_case_state.py +1079 -0
  40. loopx/benchmark_core/__init__.py +239 -0
  41. loopx/benchmark_core/adapter.py +84 -0
  42. loopx/benchmark_core/artifacts.py +517 -0
  43. loopx/benchmark_core/attempts.py +199 -0
  44. loopx/benchmark_core/container_exec.py +216 -0
  45. loopx/benchmark_core/io.py +68 -0
  46. loopx/benchmark_core/lifecycle.py +211 -0
  47. loopx/benchmark_core/loop_protocol.py +689 -0
  48. loopx/benchmark_core/observable_handles.py +348 -0
  49. loopx/benchmark_core/parity.py +256 -0
  50. loopx/benchmark_core/remote_closeout.py +482 -0
  51. loopx/benchmark_core/rounds.py +215 -0
  52. loopx/benchmark_core/route_profile.py +509 -0
  53. loopx/benchmark_core/run_permissions.py +206 -0
  54. loopx/benchmark_core/split_control.py +925 -0
  55. loopx/benchmark_core/turn_fidelity.py +326 -0
  56. loopx/benchmark_ledger.py +3793 -0
  57. loopx/benchmark_ledger_countability.py +372 -0
  58. loopx/benchmark_ledger_current.py +724 -0
  59. loopx/benchmark_trajectory.py +405 -0
  60. loopx/benchmarks/__init__.py +1 -0
  61. loopx/benchmarks/qualification/__init__.py +1 -0
  62. loopx/benchmarks/qualification/release_outcome_baseline.py +360 -0
  63. loopx/benchmarks/read_models/__init__.py +1 -0
  64. loopx/benchmarks/read_models/benchmark_attempt_accounting.py +53 -0
  65. loopx/benchmarks/read_models/benchmark_comparison.py +414 -0
  66. loopx/benchmarks/read_models/benchmark_event_timeline.py +113 -0
  67. loopx/benchmarks/read_models/benchmark_experiment_report.py +475 -0
  68. loopx/benchmarks/read_models/benchmark_learning_ledger.py +137 -0
  69. loopx/benchmarks/read_models/benchmark_lifecycle_contracts.py +228 -0
  70. loopx/benchmarks/read_models/benchmark_projection.py +723 -0
  71. loopx/benchmarks/read_models/benchmark_result.py +146 -0
  72. loopx/benchmarks/read_models/benchmark_run_execution_contract.py +116 -0
  73. loopx/benchmarks/read_models/benchmark_run_failure.py +157 -0
  74. loopx/benchmarks/read_models/benchmark_run_metrics.py +213 -0
  75. loopx/benchmarks/read_models/benchmark_run_post_execution.py +635 -0
  76. loopx/benchmarks/read_models/benchmark_run_pre_execution.py +541 -0
  77. loopx/benchmarks/read_models/benchmark_status_compaction.py +1255 -0
  78. loopx/benchmarks/read_models/benchmark_status_runner.py +780 -0
  79. loopx/benchmarks/read_models/goal_start_control_score.py +857 -0
  80. loopx/benchmarks/read_models/skillsbench_post_run_debug.py +746 -0
  81. loopx/benchmarks/read_models/skillsbench_verifier_attribution.py +269 -0
  82. loopx/bootstrap.py +1116 -0
  83. loopx/bootstrap_command_pack.py +2167 -0
  84. loopx/boundary_authority.py +199 -0
  85. loopx/canary/__init__.py +1 -0
  86. loopx/canary/maintainability_ratchet.py +800 -0
  87. loopx/canary/planner.py +1984 -0
  88. loopx/canary/premerge.py +1130 -0
  89. loopx/canary/qualification_profiles.py +309 -0
  90. loopx/canary/quality_surface_catalog.py +838 -0
  91. loopx/canary/release_profiles.py +51 -0
  92. loopx/canary/runner.py +1107 -0
  93. loopx/canary/smoke_health.py +581 -0
  94. loopx/canary/smoke_profiles.py +212 -0
  95. loopx/capabilities/__init__.py +0 -0
  96. loopx/capabilities/agent_turn_recall/__init__.py +17 -0
  97. loopx/capabilities/agent_turn_recall/cli.py +369 -0
  98. loopx/capabilities/agent_turn_recall/core.py +296 -0
  99. loopx/capabilities/auto_research/__init__.py +16 -0
  100. loopx/capabilities/auto_research/bootstrap_contract.py +157 -0
  101. loopx/capabilities/auto_research/cli.py +1468 -0
  102. loopx/capabilities/auto_research/core.py +11 -0
  103. loopx/capabilities/auto_research/defaults.py +79 -0
  104. loopx/capabilities/auto_research/demo_e2e.py +1848 -0
  105. loopx/capabilities/auto_research/demo_supervisor.py +186 -0
  106. loopx/capabilities/auto_research/evidence_packet.py +767 -0
  107. loopx/capabilities/auto_research/human_view.py +794 -0
  108. loopx/capabilities/auto_research/kernel.py +191 -0
  109. loopx/capabilities/auto_research/knn_demo_workspace.py +322 -0
  110. loopx/capabilities/auto_research/live_evidence.py +248 -0
  111. loopx/capabilities/auto_research/preset.py +176 -0
  112. loopx/capabilities/auto_research/research_state.py +1085 -0
  113. loopx/capabilities/auto_research/role_profiles.py +394 -0
  114. loopx/capabilities/auto_research/rollout_append.py +97 -0
  115. loopx/capabilities/auto_research/terminal_result_contract.py +422 -0
  116. loopx/capabilities/auto_research/terminal_result_projection.py +171 -0
  117. loopx/capabilities/auto_research/terminal_result_query.py +233 -0
  118. loopx/capabilities/auto_research/terminal_results.py +349 -0
  119. loopx/capabilities/auto_research/user_contract.py +190 -0
  120. loopx/capabilities/auto_research/worker_loop.py +163 -0
  121. loopx/capabilities/auto_research/worker_runtime.py +777 -0
  122. loopx/capabilities/auto_research/worker_skill/SKILL.md +343 -0
  123. loopx/capabilities/benchmark_toolkit/__init__.py +19 -0
  124. loopx/capabilities/benchmark_toolkit/integrity.py +387 -0
  125. loopx/capabilities/catalog.py +1875 -0
  126. loopx/capabilities/change_quality/__init__.py +19 -0
  127. loopx/capabilities/change_quality/cli.py +171 -0
  128. loopx/capabilities/change_quality/context.py +156 -0
  129. loopx/capabilities/change_quality/oracles.py +269 -0
  130. loopx/capabilities/change_quality/policy.py +34 -0
  131. loopx/capabilities/change_quality/receipt.py +482 -0
  132. loopx/capabilities/change_quality/result.py +493 -0
  133. loopx/capabilities/change_quality/scope.py +171 -0
  134. loopx/capabilities/change_quality/shadow.py +680 -0
  135. loopx/capabilities/content_ops/__init__.py +0 -0
  136. loopx/capabilities/content_ops/cli.py +649 -0
  137. loopx/capabilities/content_ops/connector_packets.py +164 -0
  138. loopx/capabilities/content_ops/item_lifecycle.py +1000 -0
  139. loopx/capabilities/content_ops/layout.py +451 -0
  140. loopx/capabilities/content_ops/markdown.py +456 -0
  141. loopx/capabilities/content_ops/schemas.py +51 -0
  142. loopx/capabilities/content_ops/social_browser_x.py +107 -0
  143. loopx/capabilities/content_ops/surface.py +1956 -0
  144. loopx/capabilities/content_ops/templates/layout-catalog-v0.json +72 -0
  145. loopx/capabilities/context_providers/__init__.py +36 -0
  146. loopx/capabilities/context_providers/base.py +189 -0
  147. loopx/capabilities/context_providers/factory.py +32 -0
  148. loopx/capabilities/context_providers/openviking.py +702 -0
  149. loopx/capabilities/context_providers/service_ownership.py +185 -0
  150. loopx/capabilities/decision_context/__init__.py +129 -0
  151. loopx/capabilities/decision_context/architecture.py +83 -0
  152. loopx/capabilities/decision_context/assembler.py +849 -0
  153. loopx/capabilities/decision_context/catalog_entry.py +195 -0
  154. loopx/capabilities/decision_context/cli.py +310 -0
  155. loopx/capabilities/decision_context/cursor_commit.py +535 -0
  156. loopx/capabilities/decision_context/outcome_feedback.py +352 -0
  157. loopx/capabilities/decision_context/packets.py +654 -0
  158. loopx/capabilities/decision_context/private_state.py +189 -0
  159. loopx/capabilities/decision_context/profile.py +453 -0
  160. loopx/capabilities/decision_context/providers.py +228 -0
  161. loopx/capabilities/decision_context/review_settlement.py +136 -0
  162. loopx/capabilities/decision_context/runtime.py +273 -0
  163. loopx/capabilities/decision_context/sources.py +415 -0
  164. loopx/capabilities/explore/__init__.py +1 -0
  165. loopx/capabilities/explore/activation.py +198 -0
  166. loopx/capabilities/explore/adaptive_replay_planner.py +221 -0
  167. loopx/capabilities/explore/child_replay_runtime.py +463 -0
  168. loopx/capabilities/explore/composition_frontier.py +291 -0
  169. loopx/capabilities/explore/counterfactual_runtime.py +578 -0
  170. loopx/capabilities/explore/episode_runtime.py +647 -0
  171. loopx/capabilities/explore/harness_checkpoint.py +171 -0
  172. loopx/capabilities/explore/harness_gate.py +115 -0
  173. loopx/capabilities/explore/harness_runtime.py +1124 -0
  174. loopx/capabilities/explore/replay_metrics.py +206 -0
  175. loopx/capabilities/explore/replay_runtime.py +1271 -0
  176. loopx/capabilities/explore/resource_portfolio.py +173 -0
  177. loopx/capabilities/explore/result_log.py +974 -0
  178. loopx/capabilities/explore/router_state.py +432 -0
  179. loopx/capabilities/explore/source_history_reconcile.py +255 -0
  180. loopx/capabilities/explore/speculative_scheduler.py +498 -0
  181. loopx/capabilities/explore/todo_branch_plan.py +650 -0
  182. loopx/capabilities/explore/todo_evidence.py +141 -0
  183. loopx/capabilities/explore/trace_runtime.py +284 -0
  184. loopx/capabilities/explore/worker_branch_plan.py +1257 -0
  185. loopx/capabilities/integration_branch/__init__.py +13 -0
  186. loopx/capabilities/integration_branch/cli.py +148 -0
  187. loopx/capabilities/integration_branch/core.py +916 -0
  188. loopx/capabilities/issue_fix/__init__.py +19 -0
  189. loopx/capabilities/issue_fix/acceptance_loop.py +1050 -0
  190. loopx/capabilities/issue_fix/candidate_evidence.py +503 -0
  191. loopx/capabilities/issue_fix/candidate_preflight.py +676 -0
  192. loopx/capabilities/issue_fix/cli.py +1822 -0
  193. loopx/capabilities/issue_fix/cli_input.py +87 -0
  194. loopx/capabilities/issue_fix/content_ops_cli.py +148 -0
  195. loopx/capabilities/issue_fix/discovered_issue_promotion.py +947 -0
  196. loopx/capabilities/issue_fix/explore_projection.py +710 -0
  197. loopx/capabilities/issue_fix/feasibility.py +542 -0
  198. loopx/capabilities/issue_fix/github_public.py +661 -0
  199. loopx/capabilities/issue_fix/intake_surface.py +832 -0
  200. loopx/capabilities/issue_fix/metadata_preview.py +218 -0
  201. loopx/capabilities/issue_fix/metrics_projection.py +1340 -0
  202. loopx/capabilities/issue_fix/metrics_supplement.py +634 -0
  203. loopx/capabilities/issue_fix/metrics_supplement_cli.py +127 -0
  204. loopx/capabilities/issue_fix/outcome_projection.py +1235 -0
  205. loopx/capabilities/issue_fix/periodic_report.py +189 -0
  206. loopx/capabilities/issue_fix/pr_description.py +418 -0
  207. loopx/capabilities/issue_fix/pr_gate_reconcile.py +496 -0
  208. loopx/capabilities/issue_fix/pr_gate_reconcile_cli.py +464 -0
  209. loopx/capabilities/issue_fix/pr_lifecycle.py +1327 -0
  210. loopx/capabilities/issue_fix/pr_lifecycle_rollout.py +85 -0
  211. loopx/capabilities/issue_fix/pr_monitor_materialization.py +257 -0
  212. loopx/capabilities/issue_fix/pr_review_ack.py +439 -0
  213. loopx/capabilities/issue_fix/provider_hooks.py +24 -0
  214. loopx/capabilities/issue_fix/repository_commit_evidence.py +186 -0
  215. loopx/capabilities/issue_fix/repository_context.py +457 -0
  216. loopx/capabilities/issue_fix/repository_memory.py +459 -0
  217. loopx/capabilities/issue_fix/repository_memory_provider.py +1454 -0
  218. loopx/capabilities/issue_fix/repository_snapshot.py +454 -0
  219. loopx/capabilities/issue_fix/reviewer_cli.py +917 -0
  220. loopx/capabilities/issue_fix/reviewer_notification.py +882 -0
  221. loopx/capabilities/issue_fix/reviewer_notification_drain.py +942 -0
  222. loopx/capabilities/issue_fix/reviewer_recommendation.py +1057 -0
  223. loopx/capabilities/issue_fix/reviewer_request.py +1282 -0
  224. loopx/capabilities/issue_fix/reward_memory.py +879 -0
  225. loopx/capabilities/issue_fix/workflow_plan.py +1286 -0
  226. loopx/capabilities/material_lifecycle/__init__.py +161 -0
  227. loopx/capabilities/material_lifecycle/_validation.py +183 -0
  228. loopx/capabilities/material_lifecycle/apply.py +672 -0
  229. loopx/capabilities/material_lifecycle/architecture.py +122 -0
  230. loopx/capabilities/material_lifecycle/cli.py +161 -0
  231. loopx/capabilities/material_lifecycle/decision_planning.py +470 -0
  232. loopx/capabilities/material_lifecycle/explore_execution.py +306 -0
  233. loopx/capabilities/material_lifecycle/intake.py +869 -0
  234. loopx/capabilities/material_lifecycle/inventory.py +147 -0
  235. loopx/capabilities/material_lifecycle/lifecycle.py +98 -0
  236. loopx/capabilities/material_lifecycle/preparation.py +147 -0
  237. loopx/capabilities/material_lifecycle/project_skill.py +83 -0
  238. loopx/capabilities/material_lifecycle/ranking.py +267 -0
  239. loopx/capabilities/material_lifecycle/readable_projection.py +500 -0
  240. loopx/capabilities/material_lifecycle/rebuild.py +480 -0
  241. loopx/capabilities/material_lifecycle/settlement.py +238 -0
  242. loopx/capabilities/periodic_report/__init__.py +71 -0
  243. loopx/capabilities/periodic_report/adapters.py +939 -0
  244. loopx/capabilities/periodic_report/archive.py +422 -0
  245. loopx/capabilities/periodic_report/bindings.py +705 -0
  246. loopx/capabilities/periodic_report/cli.py +277 -0
  247. loopx/capabilities/periodic_report/core.py +691 -0
  248. loopx/capabilities/periodic_report/extension_envelope.py +66 -0
  249. loopx/capabilities/periodic_report/presets.py +103 -0
  250. loopx/capabilities/periodic_report/profile.py +235 -0
  251. loopx/capabilities/periodic_report/project_progress.py +179 -0
  252. loopx/capabilities/periodic_report/triggers.py +452 -0
  253. loopx/capabilities/pr_review_queue/__init__.py +17 -0
  254. loopx/capabilities/pr_review_queue/core.py +506 -0
  255. loopx/capabilities/pr_review_queue/review_contract.py +506 -0
  256. loopx/capabilities/registry.py +192 -0
  257. loopx/capabilities/reward_memory/__init__.py +75 -0
  258. loopx/capabilities/reward_memory/application.py +819 -0
  259. loopx/capabilities/reward_memory/architecture.py +572 -0
  260. loopx/capabilities/reward_memory/candidate_review.py +511 -0
  261. loopx/capabilities/reward_memory/cli.py +469 -0
  262. loopx/capabilities/reward_memory/dogfood.py +574 -0
  263. loopx/capabilities/reward_memory/evaluation.py +296 -0
  264. loopx/capabilities/reward_memory/evaluation_fixtures.py +362 -0
  265. loopx/capabilities/reward_memory/experiment.py +567 -0
  266. loopx/capabilities/reward_memory/health.py +222 -0
  267. loopx/capabilities/reward_memory/ingestion.py +519 -0
  268. loopx/capabilities/reward_memory/registry.py +600 -0
  269. loopx/capabilities/reward_memory/runtime_hooks.py +312 -0
  270. loopx/capabilities/reward_memory/scoped_feedback.py +173 -0
  271. loopx/capabilities/semantic_preference/__init__.py +12 -0
  272. loopx/capabilities/semantic_preference/cli.py +189 -0
  273. loopx/capabilities/semantic_preference/contract.py +592 -0
  274. loopx/capabilities/semantic_preference/reward_memory.py +62 -0
  275. loopx/capabilities/value_connectors/__init__.py +1 -0
  276. loopx/capabilities/value_connectors/cli.py +401 -0
  277. loopx/capabilities/value_connectors/finance_extension_migration.py +108 -0
  278. loopx/capabilities/value_connectors/install_check.py +147 -0
  279. loopx/capabilities/value_connectors/planner.py +733 -0
  280. loopx/capabilities/value_connectors/source_map.py +446 -0
  281. loopx/claude_goal_baseline.py +138 -0
  282. loopx/claude_goal_mode/__init__.py +23 -0
  283. loopx/claude_goal_mode/hooks/goal_policy.py +212 -0
  284. loopx/claude_goal_mode/hooks/goal_state.py +139 -0
  285. loopx/claude_goal_mode/mcp/loopx_mcp.py +167 -0
  286. loopx/claude_goal_mode/scripts/connect.py +103 -0
  287. loopx/claude_goal_mode/scripts/goalmode_cmd.py +241 -0
  288. loopx/claude_goal_mode/scripts/install.py +328 -0
  289. loopx/claude_goal_mode/statusline/goal_status.py +97 -0
  290. loopx/cli.py +836 -0
  291. loopx/cli_commands/__init__.py +334 -0
  292. loopx/cli_commands/_host_thread.py +13 -0
  293. loopx/cli_commands/agentissue_runner_flow.py +447 -0
  294. loopx/cli_commands/agents_last_exam.py +160 -0
  295. loopx/cli_commands/agents_last_exam_baked_input.py +302 -0
  296. loopx/cli_commands/agents_last_exam_host_codex.py +374 -0
  297. loopx/cli_commands/agents_last_exam_launch_dry_run.py +372 -0
  298. loopx/cli_commands/agents_last_exam_local_plan.py +322 -0
  299. loopx/cli_commands/agents_last_exam_runner_source.py +352 -0
  300. loopx/cli_commands/agents_last_exam_task_material.py +335 -0
  301. loopx/cli_commands/agents_last_exam_validation_gate.py +236 -0
  302. loopx/cli_commands/benchmark_boundary.py +499 -0
  303. loopx/cli_commands/benchmark_dispatch.py +161 -0
  304. loopx/cli_commands/benchmark_release_outcome.py +123 -0
  305. loopx/cli_commands/benchmark_review_lifecycle.py +1275 -0
  306. loopx/cli_commands/benchmark_run_ledger.py +763 -0
  307. loopx/cli_commands/benchmark_run_ledger_case_analysis.py +249 -0
  308. loopx/cli_commands/benchmark_run_ledger_classification.py +45 -0
  309. loopx/cli_commands/benchmark_run_ledger_maintenance.py +486 -0
  310. loopx/cli_commands/benchmark_run_ledger_maintenance_registration.py +342 -0
  311. loopx/cli_commands/benchmark_run_ledger_maintenance_rendering.py +233 -0
  312. loopx/cli_commands/benchmark_run_ledger_parity.py +92 -0
  313. loopx/cli_commands/bootstrap_connect.py +238 -0
  314. loopx/cli_commands/canary.py +707 -0
  315. loopx/cli_commands/canary_release_qualification.py +79 -0
  316. loopx/cli_commands/capability.py +96 -0
  317. loopx/cli_commands/doctor.py +43 -0
  318. loopx/cli_commands/dreaming.py +143 -0
  319. loopx/cli_commands/edgebench.py +205 -0
  320. loopx/cli_commands/evidence_log.py +275 -0
  321. loopx/cli_commands/explore.py +989 -0
  322. loopx/cli_commands/explore_planning_commands.py +157 -0
  323. loopx/cli_commands/extension.py +271 -0
  324. loopx/cli_commands/first_run_report.py +73 -0
  325. loopx/cli_commands/goal_channel.py +656 -0
  326. loopx/cli_commands/handoff_mode.py +158 -0
  327. loopx/cli_commands/history.py +622 -0
  328. loopx/cli_commands/host_mode_plan.py +113 -0
  329. loopx/cli_commands/lark_inbox.py +431 -0
  330. loopx/cli_commands/lark_kanban.py +629 -0
  331. loopx/cli_commands/ml_experiment.py +321 -0
  332. loopx/cli_commands/multi_agent.py +211 -0
  333. loopx/cli_commands/opencode2_goal_worker.py +217 -0
  334. loopx/cli_commands/pr_review.py +167 -0
  335. loopx/cli_commands/presentation.py +218 -0
  336. loopx/cli_commands/preset.py +96 -0
  337. loopx/cli_commands/project.py +150 -0
  338. loopx/cli_commands/project_lifecycle.py +915 -0
  339. loopx/cli_commands/quota.py +859 -0
  340. loopx/cli_commands/quota_registration.py +241 -0
  341. loopx/cli_commands/quota_request.py +113 -0
  342. loopx/cli_commands/ready_score.py +110 -0
  343. loopx/cli_commands/registry_admin.py +975 -0
  344. loopx/cli_commands/registry_admin_configure.py +344 -0
  345. loopx/cli_commands/registry_admin_peer.py +84 -0
  346. loopx/cli_commands/registry_authority.py +218 -0
  347. loopx/cli_commands/review_batch.py +146 -0
  348. loopx/cli_commands/slash_commands.py +145 -0
  349. loopx/cli_commands/start_goal.py +251 -0
  350. loopx/cli_commands/starter.py +175 -0
  351. loopx/cli_commands/starter_bootstrap.py +179 -0
  352. loopx/cli_commands/starter_bootstrap_registration.py +198 -0
  353. loopx/cli_commands/starter_runtime_idle.py +107 -0
  354. loopx/cli_commands/starter_scheduler.py +207 -0
  355. loopx/cli_commands/starter_session_runtime.py +152 -0
  356. loopx/cli_commands/starter_visible_common.py +54 -0
  357. loopx/cli_commands/starter_visible_driver.py +161 -0
  358. loopx/cli_commands/starter_visible_pilot.py +278 -0
  359. loopx/cli_commands/status.py +867 -0
  360. loopx/cli_commands/status_registration.py +239 -0
  361. loopx/cli_commands/summary_all.py +222 -0
  362. loopx/cli_commands/support_control.py +809 -0
  363. loopx/cli_commands/support_control_registry.py +68 -0
  364. loopx/cli_commands/support_control_supervisor.py +289 -0
  365. loopx/cli_commands/task_lease.py +306 -0
  366. loopx/cli_commands/terminal_bench_adapter.py +717 -0
  367. loopx/cli_commands/terminal_bench_environment_result.py +1246 -0
  368. loopx/cli_commands/todo.py +940 -0
  369. loopx/cli_commands/todo_argument_validation.py +572 -0
  370. loopx/cli_commands/todo_event.py +114 -0
  371. loopx/cli_commands/turn.py +804 -0
  372. loopx/cli_commands/version.py +46 -0
  373. loopx/cli_commands/worker_bridge.py +659 -0
  374. loopx/cli_rollout.py +314 -0
  375. loopx/codex_cli_goal_tui.py +672 -0
  376. loopx/codex_cli_probe.py +1530 -0
  377. loopx/codex_cli_probe_markdown.py +935 -0
  378. loopx/codex_cli_runtime_probe.py +733 -0
  379. loopx/codex_cli_scheduler.py +564 -0
  380. loopx/codex_goal_baseline.py +620 -0
  381. loopx/configuration_catalog.py +617 -0
  382. loopx/configure_goal.py +1375 -0
  383. loopx/contract.py +996 -0
  384. loopx/control_plane/__init__.py +71 -0
  385. loopx/control_plane/agents/__init__.py +1 -0
  386. loopx/control_plane/agents/agent_lane_recommendation.py +516 -0
  387. loopx/control_plane/agents/agent_scope.py +1578 -0
  388. loopx/control_plane/agents/agent_scope_frontier.py +60 -0
  389. loopx/control_plane/agents/capability_gate.py +531 -0
  390. loopx/control_plane/agents/identity.py +140 -0
  391. loopx/control_plane/agents/legacy_migration.py +169 -0
  392. loopx/control_plane/agents/management_projection.py +658 -0
  393. loopx/control_plane/agents/material_frontier.py +608 -0
  394. loopx/control_plane/agents/material_handoff.py +156 -0
  395. loopx/control_plane/agents/multi_agent/__init__.py +1 -0
  396. loopx/control_plane/agents/multi_agent/codex_executable.py +207 -0
  397. loopx/control_plane/agents/multi_agent/collective_round_ledger.py +387 -0
  398. loopx/control_plane/agents/multi_agent/contract.py +474 -0
  399. loopx/control_plane/agents/multi_agent/recipe.py +110 -0
  400. loopx/control_plane/agents/multi_agent/role_successor.py +297 -0
  401. loopx/control_plane/agents/multi_agent/runtime_scripts.py +426 -0
  402. loopx/control_plane/agents/multi_agent/visible_launch_policy.py +149 -0
  403. loopx/control_plane/agents/multi_agent/visible_wake_scheduler.py +392 -0
  404. loopx/control_plane/agents/profile.py +216 -0
  405. loopx/control_plane/agents/runtime_model.py +73 -0
  406. loopx/control_plane/agents/subagent_activity.py +164 -0
  407. loopx/control_plane/agents/supervisor.py +544 -0
  408. loopx/control_plane/agents/supervisor_events.py +462 -0
  409. loopx/control_plane/agents/supervisor_inject.py +204 -0
  410. loopx/control_plane/agents/work_mode.py +56 -0
  411. loopx/control_plane/agents/workspace_guard.py +364 -0
  412. loopx/control_plane/effect_program.py +644 -0
  413. loopx/control_plane/goals/__init__.py +1 -0
  414. loopx/control_plane/goals/active_state_event_projection.py +103 -0
  415. loopx/control_plane/goals/active_state_metadata.py +47 -0
  416. loopx/control_plane/goals/active_state_sections.py +58 -0
  417. loopx/control_plane/goals/configure_goal_service.py +354 -0
  418. loopx/control_plane/goals/contract_health.py +132 -0
  419. loopx/control_plane/goals/dreaming.py +152 -0
  420. loopx/control_plane/goals/global_registry_health.py +199 -0
  421. loopx/control_plane/goals/global_registry_shadow.py +33 -0
  422. loopx/control_plane/goals/goal_channel.py +34 -0
  423. loopx/control_plane/goals/goal_channel_projection.py +560 -0
  424. loopx/control_plane/goals/goal_frontier/__init__.py +1917 -0
  425. loopx/control_plane/goals/goal_frontier/ack_policy.py +149 -0
  426. loopx/control_plane/goals/goal_frontier/outcome_continuity.py +437 -0
  427. loopx/control_plane/goals/goal_frontier/replan_rules.py +210 -0
  428. loopx/control_plane/goals/goal_frontier/semantic_history.py +314 -0
  429. loopx/control_plane/goals/goal_frontier/terminal.py +180 -0
  430. loopx/control_plane/goals/goal_vision.py +443 -0
  431. loopx/control_plane/goals/goal_vision_policy.py +36 -0
  432. loopx/control_plane/goals/goal_vision_state.py +62 -0
  433. loopx/control_plane/goals/goal_vision_wait.py +290 -0
  434. loopx/control_plane/goals/path_resolution.py +20 -0
  435. loopx/control_plane/goals/start_contract.py +206 -0
  436. loopx/control_plane/goals/vision_checkpoint.py +92 -0
  437. loopx/control_plane/handoff/__init__.py +1 -0
  438. loopx/control_plane/handoff/cross_runtime_impl_review.py +311 -0
  439. loopx/control_plane/handoff/delivery_contract.py +161 -0
  440. loopx/control_plane/handoff/handoff_runs.py +71 -0
  441. loopx/control_plane/handoff/project_handoff.py +155 -0
  442. loopx/control_plane/handoff/review_batch.py +463 -0
  443. loopx/control_plane/handoff/review_packet_context.py +216 -0
  444. loopx/control_plane/heartbeat/agent.py +173 -0
  445. loopx/control_plane/heartbeat/budget.py +66 -0
  446. loopx/control_plane/heartbeat/builder.py +501 -0
  447. loopx/control_plane/heartbeat/host.py +64 -0
  448. loopx/control_plane/heartbeat/rules.py +68 -0
  449. loopx/control_plane/heartbeat/task_body.py +759 -0
  450. loopx/control_plane/heartbeat/visible_goal.py +86 -0
  451. loopx/control_plane/projects/__init__.py +1 -0
  452. loopx/control_plane/projects/contract.py +25 -0
  453. loopx/control_plane/projects/registry.py +663 -0
  454. loopx/control_plane/quota/__init__.py +1 -0
  455. loopx/control_plane/quota/cli_projection.py +704 -0
  456. loopx/control_plane/quota/decision_summary.py +431 -0
  457. loopx/control_plane/quota/effect_program.py +152 -0
  458. loopx/control_plane/quota/error_codes.py +19 -0
  459. loopx/control_plane/quota/goal_boundary.py +464 -0
  460. loopx/control_plane/quota/heartbeat_receipt.py +277 -0
  461. loopx/control_plane/quota/heartbeat_recommendation.py +718 -0
  462. loopx/control_plane/quota/host_poll_receipts.py +162 -0
  463. loopx/control_plane/quota/live_decision.py +142 -0
  464. loopx/control_plane/quota/monitor_poll.py +786 -0
  465. loopx/control_plane/quota/policy_constants.py +40 -0
  466. loopx/control_plane/quota/projection_repair.py +262 -0
  467. loopx/control_plane/quota/recent_runs.py +210 -0
  468. loopx/control_plane/quota/scheduler_ack.py +490 -0
  469. loopx/control_plane/quota/selected_todo_projection.py +139 -0
  470. loopx/control_plane/quota/settlement.py +437 -0
  471. loopx/control_plane/quota/settlement_cli.py +246 -0
  472. loopx/control_plane/quota/settlement_validation.py +64 -0
  473. loopx/control_plane/quota/settlement_workspace_causality.py +180 -0
  474. loopx/control_plane/quota/should_run.py +249 -0
  475. loopx/control_plane/quota/should_run_packet.py +1165 -0
  476. loopx/control_plane/quota/should_run_prepare.py +675 -0
  477. loopx/control_plane/quota/slot_accounting.py +1123 -0
  478. loopx/control_plane/quota/spend_sources.py +11 -0
  479. loopx/control_plane/quota/stall_repair.py +397 -0
  480. loopx/control_plane/quota/states.py +29 -0
  481. loopx/control_plane/quota/task_orchestration.py +448 -0
  482. loopx/control_plane/quota/task_orchestration_admission.py +497 -0
  483. loopx/control_plane/quota/turn_envelope.py +889 -0
  484. loopx/control_plane/quota/usage_summary.py +140 -0
  485. loopx/control_plane/reward_memory.py +43 -0
  486. loopx/control_plane/runtime/__init__.py +2 -0
  487. loopx/control_plane/runtime/active_user_assisted_pilot.py +275 -0
  488. loopx/control_plane/runtime/agent_scoped_evidence_log.py +435 -0
  489. loopx/control_plane/runtime/decision_freshness.py +203 -0
  490. loopx/control_plane/runtime/event_ledger.py +197 -0
  491. loopx/control_plane/runtime/event_store_migration_bridge.py +196 -0
  492. loopx/control_plane/runtime/goal_project_route.py +70 -0
  493. loopx/control_plane/runtime/local_state_write_correctness.py +242 -0
  494. loopx/control_plane/runtime/promotion_readiness.py +152 -0
  495. loopx/control_plane/runtime/public_safety.py +120 -0
  496. loopx/control_plane/runtime/run_artifacts.py +78 -0
  497. loopx/control_plane/runtime/run_compaction.py +397 -0
  498. loopx/control_plane/runtime/run_context_retention.py +241 -0
  499. loopx/control_plane/runtime/run_history.py +132 -0
  500. loopx/control_plane/runtime/run_index_duplicates.py +205 -0
  501. loopx/control_plane/runtime/run_index_rebuild.py +263 -0
  502. loopx/control_plane/runtime/run_ingest_health.py +336 -0
  503. loopx/control_plane/runtime/runtime_projection_route.py +624 -0
  504. loopx/control_plane/runtime/runtime_projection_writer.py +98 -0
  505. loopx/control_plane/runtime/session_runtime.py +339 -0
  506. loopx/control_plane/runtime/shared_runtime_material_projection.py +332 -0
  507. loopx/control_plane/runtime/shared_runtime_refresh_projection.py +183 -0
  508. loopx/control_plane/runtime/stale_latest_run.py +90 -0
  509. loopx/control_plane/runtime/status_classifications.py +49 -0
  510. loopx/control_plane/runtime/status_projection_cache.py +235 -0
  511. loopx/control_plane/runtime/stride_observation.py +144 -0
  512. loopx/control_plane/runtime/time.py +39 -0
  513. loopx/control_plane/runtime/trajectory_hygiene.py +149 -0
  514. loopx/control_plane/runtime/validation_command.py +69 -0
  515. loopx/control_plane/scheduler/__init__.py +1 -0
  516. loopx/control_plane/scheduler/ack.py +329 -0
  517. loopx/control_plane/scheduler/arbitration.py +188 -0
  518. loopx/control_plane/scheduler/automation_liveness.py +183 -0
  519. loopx/control_plane/scheduler/execution_context.py +555 -0
  520. loopx/control_plane/scheduler/external_evidence_observation.py +428 -0
  521. loopx/control_plane/scheduler/monitor_display.py +143 -0
  522. loopx/control_plane/scheduler/monitor_poll_policy.py +161 -0
  523. loopx/control_plane/scheduler/monitor_poll_writeback.py +351 -0
  524. loopx/control_plane/scheduler/monitor_target.py +64 -0
  525. loopx/control_plane/scheduler/monitor_todo.py +146 -0
  526. loopx/control_plane/scheduler/monitor_wait.py +237 -0
  527. loopx/control_plane/scheduler/scheduler_hint.py +1284 -0
  528. loopx/control_plane/scheduler/state.py +354 -0
  529. loopx/control_plane/scheduler/state_transition_rules.py +179 -0
  530. loopx/control_plane/scheduler/time.py +10 -0
  531. loopx/control_plane/settlement_driver.py +293 -0
  532. loopx/control_plane/status/__init__.py +6 -0
  533. loopx/control_plane/status/active_state_projection.py +105 -0
  534. loopx/control_plane/status/agent_lane_projection.py +375 -0
  535. loopx/control_plane/status/attention_projection.py +74 -0
  536. loopx/control_plane/status/autonomous_replan_projection.py +103 -0
  537. loopx/control_plane/status/collection.py +140 -0
  538. loopx/control_plane/status/contract_projection.py +31 -0
  539. loopx/control_plane/status/dreaming_projection.py +52 -0
  540. loopx/control_plane/status/goal_attention_projection.py +157 -0
  541. loopx/control_plane/status/lifecycle_projection.py +110 -0
  542. loopx/control_plane/status/monitor_display_projection.py +69 -0
  543. loopx/control_plane/status/registry_health_projection.py +75 -0
  544. loopx/control_plane/status/run_projection.py +70 -0
  545. loopx/control_plane/status/runtime_summaries.py +161 -0
  546. loopx/control_plane/testing/__init__.py +1 -0
  547. loopx/control_plane/testing/actual_default_model_behavior_portfolio.py +1371 -0
  548. loopx/control_plane/testing/canary_harness.py +182 -0
  549. loopx/control_plane/testing/capability_monitor_repair_tool_behavior.py +674 -0
  550. loopx/control_plane/testing/cli_output_budget.py +807 -0
  551. loopx/control_plane/testing/cli_output_differential.py +250 -0
  552. loopx/control_plane/testing/cli_output_semantics.py +87 -0
  553. loopx/control_plane/testing/control_plane_composition_scenarios.py +225 -0
  554. loopx/control_plane/testing/decision_replay.py +268 -0
  555. loopx/control_plane/testing/doubao_model_behavior_actor.py +559 -0
  556. loopx/control_plane/testing/model_behavior_corpus.py +344 -0
  557. loopx/control_plane/testing/model_behavior_qualification.py +769 -0
  558. loopx/control_plane/testing/model_behavior_retained_cases.py +235 -0
  559. loopx/control_plane/testing/model_tool_behavior.py +536 -0
  560. loopx/control_plane/testing/onboarding_model_behavior_qualification.py +642 -0
  561. loopx/control_plane/testing/quota_fixtures.py +208 -0
  562. loopx/control_plane/testing/quota_should_run_parity.py +57 -0
  563. loopx/control_plane/testing/release_commit_qualification.py +671 -0
  564. loopx/control_plane/testing/replan_semantic_action_behavior.py +1302 -0
  565. loopx/control_plane/testing/scoped_gate_successor_tool_behavior.py +527 -0
  566. loopx/control_plane/testing/selected_todo_tool_behavior.py +1002 -0
  567. loopx/control_plane/testing/terminal_settlement_tool_behavior.py +656 -0
  568. loopx/control_plane/todos/__init__.py +1 -0
  569. loopx/control_plane/todos/active_state_editing.py +296 -0
  570. loopx/control_plane/todos/active_state_todo_parser.py +138 -0
  571. loopx/control_plane/todos/active_state_todos.py +175 -0
  572. loopx/control_plane/todos/addition.py +103 -0
  573. loopx/control_plane/todos/claim_visibility.py +253 -0
  574. loopx/control_plane/todos/completed_archive.py +139 -0
  575. loopx/control_plane/todos/completion_fence.py +49 -0
  576. loopx/control_plane/todos/completion_policy.py +153 -0
  577. loopx/control_plane/todos/completion_validation.py +248 -0
  578. loopx/control_plane/todos/completion_validation_accountability.py +27 -0
  579. loopx/control_plane/todos/completion_validation_projection.py +57 -0
  580. loopx/control_plane/todos/contract.py +1476 -0
  581. loopx/control_plane/todos/decision_scope.py +554 -0
  582. loopx/control_plane/todos/deferred_resume.py +546 -0
  583. loopx/control_plane/todos/durable_completion.py +201 -0
  584. loopx/control_plane/todos/event_writeback.py +484 -0
  585. loopx/control_plane/todos/frontier_deadline.py +132 -0
  586. loopx/control_plane/todos/handoff_gate.py +283 -0
  587. loopx/control_plane/todos/handoff_mode.py +444 -0
  588. loopx/control_plane/todos/handoff_note.py +202 -0
  589. loopx/control_plane/todos/line_update.py +361 -0
  590. loopx/control_plane/todos/list_projection.py +205 -0
  591. loopx/control_plane/todos/markdown.py +199 -0
  592. loopx/control_plane/todos/monitor_metadata.py +88 -0
  593. loopx/control_plane/todos/mutation_authority.py +299 -0
  594. loopx/control_plane/todos/projection.py +655 -0
  595. loopx/control_plane/todos/quota_summary.py +1138 -0
  596. loopx/control_plane/todos/route_continuation.py +267 -0
  597. loopx/control_plane/todos/succession_warning.py +174 -0
  598. loopx/control_plane/todos/summary_item.py +223 -0
  599. loopx/control_plane/todos/text.py +30 -0
  600. loopx/control_plane/todos/todo_index.py +226 -0
  601. loopx/control_plane/todos/todo_summary.py +1458 -0
  602. loopx/control_plane/todos/unblock_resume.py +326 -0
  603. loopx/control_plane/todos/user_gate.py +263 -0
  604. loopx/control_plane/todos/write_hint.py +63 -0
  605. loopx/control_plane/todos/write_policy.py +135 -0
  606. loopx/control_plane/turn_driver/__init__.py +85 -0
  607. loopx/control_plane/turn_driver/codex_cli.py +502 -0
  608. loopx/control_plane/turn_driver/driver.py +355 -0
  609. loopx/control_plane/turn_driver/executor.py +1468 -0
  610. loopx/control_plane/turn_driver/loop_controller.py +669 -0
  611. loopx/control_plane/turn_driver/settlement.py +318 -0
  612. loopx/control_plane/turn_driver/transaction.py +375 -0
  613. loopx/control_plane/work_items/__init__.py +1 -0
  614. loopx/control_plane/work_items/attention_fields.py +56 -0
  615. loopx/control_plane/work_items/attention_item.py +77 -0
  616. loopx/control_plane/work_items/attention_queue.py +322 -0
  617. loopx/control_plane/work_items/attention_routing.py +213 -0
  618. loopx/control_plane/work_items/autonomous_candidates.py +135 -0
  619. loopx/control_plane/work_items/autonomous_replan_ack.py +276 -0
  620. loopx/control_plane/work_items/autonomous_replan_obligation.py +786 -0
  621. loopx/control_plane/work_items/backlog_hygiene.py +59 -0
  622. loopx/control_plane/work_items/capability_monitor_fallback.py +221 -0
  623. loopx/control_plane/work_items/delivery_batch_scale.py +66 -0
  624. loopx/control_plane/work_items/delivery_outcome.py +152 -0
  625. loopx/control_plane/work_items/delivery_signals.py +113 -0
  626. loopx/control_plane/work_items/execution_obligation.py +235 -0
  627. loopx/control_plane/work_items/goal_route_hint.py +320 -0
  628. loopx/control_plane/work_items/interaction_contract.py +1540 -0
  629. loopx/control_plane/work_items/issue_meta_surface.py +159 -0
  630. loopx/control_plane/work_items/lifecycle.py +139 -0
  631. loopx/control_plane/work_items/operator_inbox.py +266 -0
  632. loopx/control_plane/work_items/outcome_followthrough.py +69 -0
  633. loopx/control_plane/work_items/primary_action.py +326 -0
  634. loopx/control_plane/work_items/progress_observation.py +630 -0
  635. loopx/control_plane/work_items/project_asset.py +675 -0
  636. loopx/control_plane/work_items/repair_delta.py +693 -0
  637. loopx/control_plane/work_items/runtime_capability_reentry.py +168 -0
  638. loopx/control_plane/work_items/semantic_replan_writeback.py +177 -0
  639. loopx/control_plane/work_items/status_contract.py +49 -0
  640. loopx/control_plane/work_items/task_graph.py +1046 -0
  641. loopx/control_plane/work_items/task_lease.py +1254 -0
  642. loopx/control_plane/work_items/task_lease_settlement.py +422 -0
  643. loopx/control_plane/work_items/work_lane.py +510 -0
  644. loopx/control_plane/work_items/work_lane_context.py +161 -0
  645. loopx/demo.py +247 -0
  646. loopx/diagnose.py +633 -0
  647. loopx/doctor.py +1251 -0
  648. loopx/domain_packs/__init__.py +1 -0
  649. loopx/domain_packs/issue_fix.py +571 -0
  650. loopx/domain_packs/ml_experiment.py +854 -0
  651. loopx/domain_state.py +137 -0
  652. loopx/dreaming.py +706 -0
  653. loopx/entrypoint.py +16 -0
  654. loopx/event_sourced_state.py +981 -0
  655. loopx/execution_profile.py +286 -0
  656. loopx/experiments/__init__.py +1 -0
  657. loopx/experiments/planner_worker/__init__.py +1 -0
  658. loopx/experiments/planner_worker/contract.py +523 -0
  659. loopx/experiments/planner_worker/runtime.py +391 -0
  660. loopx/experiments/planner_worker/traex.py +461 -0
  661. loopx/explore_graph.py +11 -0
  662. loopx/extensions/__init__.py +1 -0
  663. loopx/extensions/bundled.py +28 -0
  664. loopx/extensions/execution_envelope.py +126 -0
  665. loopx/extensions/lark/__init__.py +11 -0
  666. loopx/extensions/lark/event_collector.py +478 -0
  667. loopx/extensions/lark/event_collector_runtime.py +506 -0
  668. loopx/extensions/lark/event_inbox.py +454 -0
  669. loopx/extensions/lark/extension.toml +88 -0
  670. loopx/extensions/lark/goal_channel.py +44 -0
  671. loopx/extensions/lark/goal_channel_contracts.py +388 -0
  672. loopx/extensions/lark/goal_channel_lifecycle.py +218 -0
  673. loopx/extensions/lark/goal_channel_runtime.py +792 -0
  674. loopx/extensions/lark/goal_channel_setup.py +805 -0
  675. loopx/extensions/lark/goal_channel_targets.py +215 -0
  676. loopx/extensions/lark/goal_channel_transport.py +281 -0
  677. loopx/extensions/lark/inbox_reactions.py +650 -0
  678. loopx/extensions/lark/inbox_reply.py +430 -0
  679. loopx/extensions/lark/presentation/__init__.py +11 -0
  680. loopx/extensions/lark/presentation/explore_results.py +2276 -0
  681. loopx/extensions/lark/presentation/explore_singleflight.py +127 -0
  682. loopx/extensions/lark/presentation/explore_source_guard.py +121 -0
  683. loopx/extensions/lark/presentation/explore_stage_document.py +703 -0
  684. loopx/extensions/lark/presentation/explore_visual_integrity.py +122 -0
  685. loopx/extensions/lark/presentation/explore_visual_readback.py +452 -0
  686. loopx/extensions/lark/presentation/explore_visual_styles.py +156 -0
  687. loopx/extensions/lark/presentation/issue_fix_surface.py +612 -0
  688. loopx/extensions/lark/presentation/kanban.py +2791 -0
  689. loopx/extensions/lark/presentation/message_card.py +112 -0
  690. loopx/extensions/lark/presentation/periodic_report.py +261 -0
  691. loopx/extensions/lark/presentation/projection_rows.py +600 -0
  692. loopx/extensions/lark/presentation/record_io.py +95 -0
  693. loopx/extensions/lark/presentation/sync_receipt.py +145 -0
  694. loopx/extensions/lark/private_json.py +40 -0
  695. loopx/extensions/lark/provider.py +86 -0
  696. loopx/extensions/lark/reviewer_notification.py +604 -0
  697. loopx/extensions/manifest.py +385 -0
  698. loopx/extensions/openviking_periodic_report/__init__.py +17 -0
  699. loopx/extensions/openviking_periodic_report/activation.py +173 -0
  700. loopx/extensions/openviking_periodic_report/extension.toml +17 -0
  701. loopx/extensions/openviking_periodic_report/provider.py +355 -0
  702. loopx/extensions/openviking_periodic_report/sink.py +117 -0
  703. loopx/extensions/openviking_semantic_preference/__init__.py +5 -0
  704. loopx/extensions/openviking_semantic_preference/extension.toml +16 -0
  705. loopx/extensions/openviking_semantic_preference/history_export.py +484 -0
  706. loopx/extensions/openviking_semantic_preference/project_peer.py +68 -0
  707. loopx/extensions/openviking_semantic_preference/provider.py +312 -0
  708. loopx/extensions/presentation.py +979 -0
  709. loopx/extensions/process_runtime.py +204 -0
  710. loopx/extensions/readiness.py +168 -0
  711. loopx/extensions/runtime.py +931 -0
  712. loopx/extensions/scaffold.py +335 -0
  713. loopx/feedback.py +581 -0
  714. loopx/file_lock.py +382 -0
  715. loopx/global_registry.py +842 -0
  716. loopx/global_risks.py +970 -0
  717. loopx/global_todos.py +568 -0
  718. loopx/handoff_budget.py +28 -0
  719. loopx/heartbeat_prequota.py +80 -0
  720. loopx/heartbeat_prompt.py +159 -0
  721. loopx/help_surface.py +516 -0
  722. loopx/history.py +1507 -0
  723. loopx/host_loop_activation.py +1311 -0
  724. loopx/host_mode_planner.py +991 -0
  725. loopx/install_contract.py +1 -0
  726. loopx/interface_budget.py +196 -0
  727. loopx/long_task_cadence.py +208 -0
  728. loopx/materials.py +185 -0
  729. loopx/ml_experiment.py +3 -0
  730. loopx/onboarding.py +214 -0
  731. loopx/opencode2_goal_mode/README.md +81 -0
  732. loopx/opencode2_goal_mode/__init__.py +9 -0
  733. loopx/opencode2_goal_mode/opencode2-goal-worker.mjs +1018 -0
  734. loopx/opencode_goal_mode/README.md +99 -0
  735. loopx/opencode_goal_mode/__init__.py +13 -0
  736. loopx/opencode_goal_mode/goal-bridge-runtime.mjs +858 -0
  737. loopx/opencode_goal_mode/loopx-goal.js +8 -0
  738. loopx/operator_gate.py +420 -0
  739. loopx/orchestration.py +127 -0
  740. loopx/paths.py +59 -0
  741. loopx/pi_goal_mode/README.md +67 -0
  742. loopx/pi_goal_mode/__init__.py +13 -0
  743. loopx/pi_goal_mode/loopx-goal.ts +254 -0
  744. loopx/pi_goal_mode/pi-goal-loop-runtime.mjs +574 -0
  745. loopx/pr_review.py +1206 -0
  746. loopx/presentation/__init__.py +1 -0
  747. loopx/presentation/explore_views.py +1334 -0
  748. loopx/presentation/markdown.py +61 -0
  749. loopx/presentation/projection_source_reconcile.py +140 -0
  750. loopx/presentation/public_safety.py +42 -0
  751. loopx/presentation/renderers/__init__.py +17 -0
  752. loopx/presentation/renderers/goal_channel_html.py +269 -0
  753. loopx/presentation/renderers/periodic_report_html.py +786 -0
  754. loopx/presentation/renderers/periodic_report_markdown.py +184 -0
  755. loopx/presentation/renderers/quota_event_markdown.py +116 -0
  756. loopx/presentation/renderers/quota_markdown.py +1112 -0
  757. loopx/presentation/renderers/status_markdown.py +1570 -0
  758. loopx/presentation/renderers/trajectory_hygiene_markdown.py +39 -0
  759. loopx/presentation/renderers/turn_envelope_markdown.py +33 -0
  760. loopx/presentation/sinks/__init__.py +5 -0
  761. loopx/presentation/sinks/openviking_periodic_report.py +7 -0
  762. loopx/presentation/static_site.py +691 -0
  763. loopx/presets.py +369 -0
  764. loopx/project_alias.py +217 -0
  765. loopx/project_map.py +589 -0
  766. loopx/project_prompt.py +1153 -0
  767. loopx/project_skill_cli.py +125 -0
  768. loopx/project_skill_delivery.py +470 -0
  769. loopx/project_uninstall.py +462 -0
  770. loopx/promotion_gate.py +197 -0
  771. loopx/quota.py +1197 -0
  772. loopx/ready_score.py +413 -0
  773. loopx/registry.py +621 -0
  774. loopx/registry_writability.py +64 -0
  775. loopx/release_candidate.py +148 -0
  776. loopx/release_manifest.py +316 -0
  777. loopx/repository_identity.py +100 -0
  778. loopx/review_packet.py +1024 -0
  779. loopx/rollout_event_log.py +505 -0
  780. loopx/runtime.py +112 -0
  781. loopx/self_update.py +750 -0
  782. loopx/session_runtime.py +418 -0
  783. loopx/skill_install_readback.py +500 -0
  784. loopx/slash_command_install.py +1393 -0
  785. loopx/slash_commands.py +264 -0
  786. loopx/state_backup.py +573 -0
  787. loopx/state_migration.py +350 -0
  788. loopx/state_projection.py +809 -0
  789. loopx/state_refresh.py +1416 -0
  790. loopx/status.py +1383 -0
  791. loopx/status_server.py +935 -0
  792. loopx/summary_all.py +725 -0
  793. loopx/terminal_bench_agent.py +2056 -0
  794. loopx/thread_agent_binding.py +408 -0
  795. loopx/todo_followups.py +168 -0
  796. loopx/todo_suggestion_prompt.py +204 -0
  797. loopx/todos.py +2229 -0
  798. loopx/turn_identity.py +17 -0
  799. loopx/upgrade.py +1083 -0
  800. loopx/visible_governance.py +667 -0
  801. loopx/visible_multi_agent_launcher.py +1253 -0
  802. loopx/visible_multi_agent_tmux.py +429 -0
  803. loopx/worker_bridge.py +1574 -0
  804. loopx-0.4.8.dist-info/METADATA +708 -0
  805. loopx-0.4.8.dist-info/RECORD +811 -0
  806. loopx-0.4.8.dist-info/WHEEL +5 -0
  807. loopx-0.4.8.dist-info/entry_points.txt +5 -0
  808. loopx-0.4.8.dist-info/licenses/LICENSE +202 -0
  809. loopx-0.4.8.dist-info/licenses/LICENSE-MIT +21 -0
  810. loopx-0.4.8.dist-info/licenses/NOTICE +6 -0
  811. loopx-0.4.8.dist-info/top_level.txt +1 -0
@@ -0,0 +1,1371 @@
1
+ from __future__ import annotations
2
+
3
+ import json
4
+ from collections.abc import Callable, Mapping
5
+ from copy import deepcopy
6
+ from dataclasses import dataclass
7
+ from hashlib import sha256
8
+ from pathlib import Path
9
+ from typing import Any
10
+
11
+ from ...bootstrap_command_pack import build_start_goal_guided_packet
12
+ from ..quota.cli_projection import compact_quota_should_run_cli_payload
13
+ from ..quota.turn_envelope import quota_action_signature_document
14
+ from ..work_items.interaction_contract import build_interaction_contract
15
+ from .control_plane_composition_scenarios import (
16
+ build_control_plane_composition_scenario_sources,
17
+ )
18
+ from .model_behavior_qualification import (
19
+ MODEL_BEHAVIOR_HARD_INVARIANT_FIELDS,
20
+ ModelBehaviorActor,
21
+ _actor_failure_code,
22
+ build_model_behavior_actor_request,
23
+ model_behavior_semantic_contract_from_packet,
24
+ run_model_behavior_qualification_arm,
25
+ )
26
+ from .onboarding_model_behavior_qualification import (
27
+ OnboardingActualBehaviorValidationError,
28
+ OnboardingModelBehaviorActor,
29
+ _behavior_contract_violations,
30
+ _semantic_contract,
31
+ _validate_actual_default_projection,
32
+ build_onboarding_model_behavior_actor_request,
33
+ build_onboarding_postcondition_observation,
34
+ run_onboarding_model_behavior_phase,
35
+ )
36
+ from .selected_todo_tool_behavior import (
37
+ SELECTED_TODO_TOOL_FIXTURE_ACTION_TEXT,
38
+ SELECTED_TODO_TOOL_FIXTURE_TODO_ID,
39
+ )
40
+
41
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_PORTFOLIO_SCHEMA_VERSION = (
42
+ "actual_default_model_behavior_portfolio_v0"
43
+ )
44
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_CATALOG_SCHEMA_VERSION = (
45
+ "actual_default_model_behavior_scenario_catalog_v0"
46
+ )
47
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS = 2
48
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID = "portfolio-goal"
49
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_AGENT_ID = "codex-portfolio"
50
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_HOT_PATH_JSON_BUDGET = 40_000
51
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_CONTRAST_SCHEMA_VERSION = (
52
+ "actual_default_model_behavior_contrast_v0"
53
+ )
54
+ _TOOL_ACTOR_KINDS = frozenset(
55
+ {
56
+ "turn_tool",
57
+ "replan_tool",
58
+ "scoped_gate_tool",
59
+ "capability_repair_tool",
60
+ "terminal_settlement_tool",
61
+ }
62
+ )
63
+ _TURN_ACTOR_KINDS = frozenset({"turn", *_TOOL_ACTOR_KINDS})
64
+
65
+
66
+ @dataclass(frozen=True)
67
+ class _ScenarioSpec:
68
+ scenario_id: str
69
+ actor_kind: str
70
+ phase: str | None
71
+ expected_route: str
72
+ scenario_family: str = "core_contract"
73
+ composition_dimensions: tuple[str, ...] = ()
74
+
75
+
76
+ @dataclass(frozen=True)
77
+ class _ContrastSpec:
78
+ contrast_id: str
79
+ contrast_kind: str
80
+ left_scenario_id: str
81
+ right_scenario_id: str
82
+ must_match_fields: tuple[str, ...]
83
+ must_differ_fields: tuple[str, ...] = ()
84
+
85
+
86
+ _SCENARIOS = (
87
+ _ScenarioSpec(
88
+ "onboarding_connect_default",
89
+ "onboarding",
90
+ "entry",
91
+ "connect_if_needed",
92
+ ),
93
+ _ScenarioSpec(
94
+ "onboarding_agent_identity_gate",
95
+ "onboarding",
96
+ "entry",
97
+ "select_agent_identity",
98
+ ),
99
+ _ScenarioSpec(
100
+ "onboarding_goal_selection_gate",
101
+ "onboarding",
102
+ "entry",
103
+ "select_goal",
104
+ ),
105
+ _ScenarioSpec(
106
+ "turn_selected_todo",
107
+ "turn_tool",
108
+ None,
109
+ "execute",
110
+ ),
111
+ _ScenarioSpec(
112
+ "turn_terminal_settlement",
113
+ "terminal_settlement_tool",
114
+ None,
115
+ "execute",
116
+ "effect_program_settlement",
117
+ ("writeback", "quota_spend", "terminal_closeout", "receipt_identity"),
118
+ ),
119
+ _ScenarioSpec(
120
+ "turn_peer_agent_identity",
121
+ "turn",
122
+ None,
123
+ "execute",
124
+ ),
125
+ _ScenarioSpec(
126
+ "turn_same_agent_continuation",
127
+ "turn",
128
+ None,
129
+ "execute",
130
+ ),
131
+ _ScenarioSpec(
132
+ "turn_human_gate",
133
+ "turn",
134
+ None,
135
+ "ask_user",
136
+ ),
137
+ _ScenarioSpec(
138
+ "turn_required_vision_replan",
139
+ "replan_tool",
140
+ None,
141
+ "execute",
142
+ "control_plane_composition",
143
+ ("vision", "monitor", "peer_ownership", "autonomous_replan"),
144
+ ),
145
+ _ScenarioSpec(
146
+ "turn_scoped_gate_successor_replan",
147
+ "scoped_gate_tool",
148
+ None,
149
+ "execute",
150
+ "control_plane_composition",
151
+ ("user_gate", "deferred_successor", "non_blocking_notice", "scheduler"),
152
+ ),
153
+ _ScenarioSpec(
154
+ "turn_capability_monitor_repair",
155
+ "capability_repair_tool",
156
+ None,
157
+ "execute",
158
+ "control_plane_composition",
159
+ ("capability_gate", "monitor_schedule", "bridge_repair", "selected_action"),
160
+ ),
161
+ _ScenarioSpec(
162
+ "turn_quota_hot_path_compaction_regression",
163
+ "turn",
164
+ None,
165
+ "execute",
166
+ "quota_cli_compaction_regression",
167
+ (
168
+ "json_budget",
169
+ "source_semantics",
170
+ "model_route",
171
+ ),
172
+ ),
173
+ _ScenarioSpec(
174
+ "turn_quota_hot_path_selected_todo_invariance",
175
+ "turn",
176
+ None,
177
+ "execute",
178
+ "quota_cli_compaction_contrast",
179
+ ("selected_todo", "omitted_diagnostics", "invariance"),
180
+ ),
181
+ _ScenarioSpec(
182
+ "turn_quota_hot_path_human_gate_invariance",
183
+ "turn",
184
+ None,
185
+ "ask_user",
186
+ "quota_cli_compaction_contrast",
187
+ ("blocking_user_gate", "omitted_diagnostics", "invariance"),
188
+ ),
189
+ _ScenarioSpec(
190
+ "onboarding_healthy_continue",
191
+ "onboarding",
192
+ "postcondition",
193
+ "continue_validation",
194
+ ),
195
+ _ScenarioSpec(
196
+ "onboarding_projection_repair",
197
+ "onboarding",
198
+ "postcondition",
199
+ "repair_projection",
200
+ ),
201
+ )
202
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_SCENARIO_COUNT = len(_SCENARIOS)
203
+
204
+ _HARD_INVARIANT_FIELDS = tuple(MODEL_BEHAVIOR_HARD_INVARIANT_FIELDS)
205
+ _CONTRASTS = (
206
+ _ContrastSpec(
207
+ "selected_todo_survives_omitted_diagnostics",
208
+ "invariance",
209
+ "turn_selected_todo",
210
+ "turn_quota_hot_path_selected_todo_invariance",
211
+ _HARD_INVARIANT_FIELDS,
212
+ ),
213
+ _ContrastSpec(
214
+ "blocking_gate_survives_omitted_diagnostics",
215
+ "invariance",
216
+ "turn_human_gate",
217
+ "turn_quota_hot_path_human_gate_invariance",
218
+ _HARD_INVARIANT_FIELDS,
219
+ ),
220
+ _ContrastSpec(
221
+ "blocking_gate_vs_non_blocking_notice",
222
+ "sensitivity",
223
+ "turn_human_gate",
224
+ "turn_scoped_gate_successor_replan",
225
+ (
226
+ "user_action_required",
227
+ "quiet_noop_allowed",
228
+ "external_write_requested",
229
+ ),
230
+ (
231
+ "decision",
232
+ "selected_todo_id",
233
+ "must_attempt_work",
234
+ "delivery_allowed",
235
+ ),
236
+ ),
237
+ _ContrastSpec(
238
+ "selected_work_vs_required_vision_replan",
239
+ "sensitivity",
240
+ "turn_selected_todo",
241
+ "turn_required_vision_replan",
242
+ (
243
+ "decision",
244
+ "user_action_required",
245
+ "must_attempt_work",
246
+ "delivery_allowed",
247
+ "quiet_noop_allowed",
248
+ "external_write_requested",
249
+ ),
250
+ ("selected_todo_id",),
251
+ ),
252
+ )
253
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_CONTRAST_COUNT = len(_CONTRASTS)
254
+
255
+
256
+ def _canonical_json(value: Any) -> str:
257
+ return json.dumps(value, ensure_ascii=False, sort_keys=True, separators=(",", ":"))
258
+
259
+
260
+ def _pretty_json_size(value: Any) -> int:
261
+ return len(
262
+ json.dumps(value, ensure_ascii=False, sort_keys=True, indent=2).encode("utf-8")
263
+ )
264
+
265
+
266
+ def _digest(value: Any) -> str:
267
+ return "sha256:" + sha256(_canonical_json(value).encode("utf-8")).hexdigest()
268
+
269
+
270
+ def actual_default_model_behavior_scenario_catalog() -> dict[str, Any]:
271
+ return {
272
+ "schema_version": ACTUAL_DEFAULT_MODEL_BEHAVIOR_CATALOG_SCHEMA_VERSION,
273
+ "topology": "actual_default_one_arm",
274
+ "scenarios": [
275
+ {
276
+ "scenario_id": spec.scenario_id,
277
+ "actor_kind": spec.actor_kind,
278
+ "phase": spec.phase,
279
+ "expected_route": spec.expected_route,
280
+ "scenario_family": spec.scenario_family,
281
+ "composition_dimensions": list(spec.composition_dimensions),
282
+ "packet_view": (
283
+ "production_heartbeat_tool_loop"
284
+ if spec.actor_kind in _TOOL_ACTOR_KINDS
285
+ else (
286
+ "quota_should_run_default"
287
+ if spec.actor_kind == "turn"
288
+ else "guided_onboarding_default"
289
+ )
290
+ ),
291
+ "repeat_policy": {
292
+ "attempts": ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS,
293
+ "pass_condition": "all_attempts_source_aligned",
294
+ "automatic_retry_on_actor_error": False,
295
+ },
296
+ }
297
+ for spec in _SCENARIOS
298
+ ],
299
+ "contrasts": [
300
+ {
301
+ "contrast_id": spec.contrast_id,
302
+ "contrast_kind": spec.contrast_kind,
303
+ "left_scenario_id": spec.left_scenario_id,
304
+ "right_scenario_id": spec.right_scenario_id,
305
+ "must_match_fields": list(spec.must_match_fields),
306
+ "must_differ_fields": list(spec.must_differ_fields),
307
+ }
308
+ for spec in _CONTRASTS
309
+ ],
310
+ }
311
+
312
+
313
+ def _write_scenario_project(root: Path) -> tuple[Path, Path]:
314
+ project = root / "project"
315
+ state_file = (
316
+ project
317
+ / ".codex"
318
+ / "goals"
319
+ / ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID
320
+ / "ACTIVE_GOAL_STATE.md"
321
+ )
322
+ state_file.parent.mkdir(parents=True)
323
+ state_file.write_text("# Active Goal State\n", encoding="utf-8")
324
+ registry_path = project / ".loopx" / "registry.json"
325
+ registry_path.parent.mkdir(parents=True)
326
+ registry_path.write_text(
327
+ json.dumps(
328
+ {
329
+ "schema_version": "0.1",
330
+ "goals": [
331
+ {
332
+ "id": ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID,
333
+ "status": "active",
334
+ "repo": str(project),
335
+ "state_file": str(state_file.relative_to(project)),
336
+ "coordination": {
337
+ "agent_model": "peer_v1",
338
+ "registered_agents": [
339
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_AGENT_ID
340
+ ],
341
+ },
342
+ }
343
+ ],
344
+ },
345
+ indent=2,
346
+ )
347
+ + "\n",
348
+ encoding="utf-8",
349
+ )
350
+ return project, registry_path
351
+
352
+
353
+ def _guided_scenario_packet(
354
+ project: Path,
355
+ *,
356
+ goal_id: str | None,
357
+ agent_id: str | None,
358
+ ) -> dict[str, Any]:
359
+ return build_start_goal_guided_packet(
360
+ project=project,
361
+ goal_id=goal_id,
362
+ agent_id=agent_id,
363
+ cli_bin="loopx",
364
+ host_surface="codex-app",
365
+ goal_text="Establish one public-safe quality contract.",
366
+ available_capabilities=["network"],
367
+ include_command_pack_detail=False,
368
+ )
369
+
370
+
371
+ def _entry_scenario_packets(root: Path) -> dict[str, dict[str, Any]]:
372
+ project, registry_path = _write_scenario_project(root)
373
+ connect = _guided_scenario_packet(
374
+ project,
375
+ goal_id=ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID,
376
+ agent_id=ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_AGENT_ID,
377
+ )
378
+
379
+ registry = json.loads(registry_path.read_text(encoding="utf-8"))
380
+ registry["goals"][0]["coordination"]["registered_agents"].append(
381
+ "codex-portfolio-reviewer"
382
+ )
383
+ registry_path.write_text(json.dumps(registry, indent=2) + "\n", encoding="utf-8")
384
+ identity = _guided_scenario_packet(
385
+ project,
386
+ goal_id=ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID,
387
+ agent_id=None,
388
+ )
389
+
390
+ second_goal = "portfolio-second-goal"
391
+ second_state = project / ".codex" / "goals" / second_goal / "ACTIVE_GOAL_STATE.md"
392
+ second_state.parent.mkdir(parents=True)
393
+ second_state.write_text("# Second Active Goal State\n", encoding="utf-8")
394
+ registry["goals"].append(
395
+ {
396
+ "id": second_goal,
397
+ "status": "active",
398
+ "repo": str(project),
399
+ "state_file": str(second_state.relative_to(project)),
400
+ "coordination": {
401
+ "agent_model": "peer_v1",
402
+ "registered_agents": ["codex-portfolio-second"],
403
+ },
404
+ }
405
+ )
406
+ registry_path.write_text(json.dumps(registry, indent=2) + "\n", encoding="utf-8")
407
+ goal_selection = _guided_scenario_packet(project, goal_id=None, agent_id=None)
408
+ return {
409
+ "onboarding_connect_default": connect,
410
+ "onboarding_agent_identity_gate": identity,
411
+ "onboarding_goal_selection_gate": goal_selection,
412
+ }
413
+
414
+
415
+ def _turn_scenario_source(
416
+ *,
417
+ human_gate: bool,
418
+ agent_id: str = ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_AGENT_ID,
419
+ continuation_policy: str | None = None,
420
+ ) -> dict[str, Any]:
421
+ selected_todo = None
422
+ if not human_gate:
423
+ selected_todo = {
424
+ "todo_id": "todo_portfolio001",
425
+ "status": "open",
426
+ "task_class": "advancement_task",
427
+ "claimed_by": agent_id,
428
+ "text": "Implement one bounded public-safe slice.",
429
+ }
430
+ if continuation_policy:
431
+ selected_todo["continuation_policy"] = continuation_policy
432
+ payload: dict[str, Any] = {
433
+ "ok": True,
434
+ "mode": "should-run",
435
+ "goal_id": ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID,
436
+ "decision": "skip" if human_gate else "run",
437
+ "should_run": not human_gate,
438
+ "effective_action": "operator_gate" if human_gate else "normal_run",
439
+ "state": "operator_gate" if human_gate else "eligible",
440
+ "requires_user_action": human_gate,
441
+ "gate_prompt": ("Approve the bounded public release." if human_gate else None),
442
+ "recommended_action": (
443
+ "Approve the bounded public release."
444
+ if human_gate
445
+ else "Implement one bounded public-safe slice."
446
+ ),
447
+ "selected_todo": selected_todo,
448
+ "agent_identity": {"agent_id": agent_id},
449
+ "execution_obligation": {
450
+ "must_attempt_work": not human_gate,
451
+ "delivery_allowed": not human_gate,
452
+ },
453
+ "normal_delivery_allowed": not human_gate,
454
+ "heartbeat_recommendation": {
455
+ "notify": "NOTIFY" if human_gate else "DONT_NOTIFY"
456
+ },
457
+ "goal_boundary": {
458
+ "write_scope": ["loopx/**", "tests/**"],
459
+ "guards": ["stop before external writes"],
460
+ },
461
+ }
462
+ payload["interaction_contract"] = build_interaction_contract(
463
+ payload,
464
+ available_capabilities=["network"],
465
+ )
466
+ payload["action_required"] = human_gate
467
+ payload["open_count"] = 1 if human_gate else 0
468
+ return payload
469
+
470
+
471
+ def _selected_todo_scenario_source() -> dict[str, Any]:
472
+ payload = _turn_scenario_source(human_gate=False)
473
+ payload["selected_todo"] = {
474
+ **dict(payload["selected_todo"]),
475
+ "text": SELECTED_TODO_TOOL_FIXTURE_ACTION_TEXT,
476
+ }
477
+ payload["recommended_action"] = SELECTED_TODO_TOOL_FIXTURE_ACTION_TEXT
478
+ payload["interaction_contract"] = build_interaction_contract(
479
+ payload,
480
+ available_capabilities=["shell", "filesystem_read"],
481
+ )
482
+ return payload
483
+
484
+
485
+ def _terminal_settlement_scenario_source() -> dict[str, Any]:
486
+ payload = _turn_scenario_source(human_gate=False)
487
+ payload["selected_todo"] = {
488
+ **dict(payload["selected_todo"]),
489
+ "todo_id": "todo_terminal001",
490
+ "text": (
491
+ "Read `fixture/settlement-proof.json`, verify the bounded delivery, "
492
+ "then settle this final Todo from the quota-projected plan. "
493
+ "It has no successor."
494
+ ),
495
+ }
496
+ payload["recommended_action"] = payload["selected_todo"]["text"]
497
+ payload["interaction_contract"] = build_interaction_contract(
498
+ payload,
499
+ available_capabilities=[
500
+ "shell",
501
+ "filesystem_read",
502
+ "filesystem_write",
503
+ ],
504
+ )
505
+ return payload
506
+
507
+
508
+ def build_quota_hot_path_compaction_regression_source() -> dict[str, Any]:
509
+ """Build a coherent over-budget turn whose cold diagnostics are removable."""
510
+
511
+ payload = _turn_scenario_source(human_gate=False)
512
+ selected_todo = dict(payload["selected_todo"])
513
+ selected_todo.update(
514
+ {
515
+ "todo_id": "todo_c0ffee123456",
516
+ "text": "Implement the bounded hot-path qualification slice.",
517
+ }
518
+ )
519
+ payload["selected_todo"] = selected_todo
520
+ payload["recommended_action"] = selected_todo["text"]
521
+ payload["interaction_contract"] = build_interaction_contract(
522
+ payload,
523
+ available_capabilities=["network", "shell", "filesystem_write"],
524
+ )
525
+
526
+ repeated_detail = "Bounded public-safe candidate diagnostic. " * 28
527
+
528
+ def candidate(kind: str, index: int) -> dict[str, Any]:
529
+ return {
530
+ "todo_id": f"todo_{index:012x}",
531
+ "status": "open",
532
+ "priority": "P1",
533
+ "task_class": "advancement_task",
534
+ "action_kind": f"regression_{kind}",
535
+ "claimed_by": ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_AGENT_ID,
536
+ "text": f"{kind} candidate {index}. {repeated_detail}",
537
+ }
538
+
539
+ runnable_candidates = [candidate("runnable", index) for index in range(1, 25)]
540
+ blocked_candidates = [candidate("blocked", index) for index in range(25, 41)]
541
+ resolution_bindings = [candidate("binding", index) for index in range(41, 57)]
542
+ payload["capability_gate"] = {
543
+ "schema_version": "capability_gate_v0",
544
+ "required": ["shell", "filesystem_write"],
545
+ "available": ["network", "shell", "filesystem_write"],
546
+ "missing": [],
547
+ "action": "run",
548
+ "decision_owner": "agent",
549
+ "candidate_order_policy": "claim_then_profile_then_priority",
550
+ "runnable_count": len(runnable_candidates),
551
+ "runnable_candidates": runnable_candidates,
552
+ "blocked_candidates": blocked_candidates,
553
+ "resolution_bindings": resolution_bindings,
554
+ }
555
+
556
+ lane_title = "Implement the bounded hot-path qualification slice."
557
+ payload["agent_lane_next_action"] = {
558
+ **selected_todo,
559
+ "title": lane_title,
560
+ "text": f"[P1] {lane_title}",
561
+ }
562
+ payload["active_state_next_action"] = (
563
+ "Preserve the durable quality route while the peer lane remains independent."
564
+ )
565
+ payload["latest_run_recommended_action"] = selected_todo["text"]
566
+ payload["next_action_projection_warning"] = {
567
+ "schema_version": "next_action_projection_warning_v0",
568
+ "kind": "next_action_projection_mismatch",
569
+ "severity": "info",
570
+ "requires_state_writeback": False,
571
+ "active_state_next_action": payload["active_state_next_action"],
572
+ "latest_run_recommended_action": payload["latest_run_recommended_action"],
573
+ "agent_lane_next_action": payload["agent_lane_next_action"]["text"],
574
+ "reason": "The agent lane is intentionally narrower than the durable route.",
575
+ "recommended_action": selected_todo["text"],
576
+ }
577
+ peer_actions = [candidate("peer", index) for index in range(57, 81)]
578
+ payload["goal_route_hint"] = {
579
+ "schema_version": "goal_route_hint_v0",
580
+ "route_decision": "run_current_agent_lane",
581
+ "counts": {"other_agent_claimed_advancement_count": len(peer_actions)},
582
+ "other_agent_next_actions": peer_actions,
583
+ }
584
+ return payload
585
+
586
+
587
+ def _build_quota_hot_path_selected_todo_invariance_source() -> dict[str, Any]:
588
+ payload = build_quota_hot_path_compaction_regression_source()
589
+ selected_todo = dict(payload["selected_todo"])
590
+ selected_todo.update(
591
+ {
592
+ "todo_id": "todo_portfolio001",
593
+ "text": "Implement one bounded public-safe slice.",
594
+ "continuation_hint": "Next run the restart canary.",
595
+ }
596
+ )
597
+ payload["selected_todo"] = selected_todo
598
+ payload["recommended_action"] = selected_todo["text"]
599
+ payload["active_state_next_action"] = "Inspect an obsolete branch."
600
+ payload["latest_run_recommended_action"] = "Continue the previous run."
601
+ payload["agent_lane_next_action"] = {
602
+ **selected_todo,
603
+ "title": selected_todo["text"],
604
+ "text": f"[P1] {selected_todo['text']}",
605
+ }
606
+ warning = dict(payload["next_action_projection_warning"])
607
+ warning["active_state_next_action"] = payload["active_state_next_action"]
608
+ warning["latest_run_recommended_action"] = payload["latest_run_recommended_action"]
609
+ warning["agent_lane_next_action"] = payload["agent_lane_next_action"]["text"]
610
+ warning["recommended_action"] = selected_todo["text"]
611
+ payload["next_action_projection_warning"] = warning
612
+ payload["goal_route_hint"] = {
613
+ **payload["goal_route_hint"],
614
+ "selected_action_differs_from_durable": True,
615
+ "current_agent_next_action": {"todo_id": selected_todo["todo_id"]},
616
+ }
617
+ payload["interaction_contract"] = build_interaction_contract(
618
+ payload,
619
+ available_capabilities=["network", "shell", "filesystem_write"],
620
+ )
621
+ return payload
622
+
623
+
624
+ def _build_quota_hot_path_human_gate_invariance_source() -> dict[str, Any]:
625
+ payload = _turn_scenario_source(human_gate=True)
626
+ noisy_source = build_quota_hot_path_compaction_regression_source()
627
+ capability_gate = deepcopy(noisy_source["capability_gate"])
628
+ capability_gate.update(
629
+ {
630
+ "required": [],
631
+ "available": ["network"],
632
+ "missing": [],
633
+ "action": "operator_gate",
634
+ "decision_owner": "user",
635
+ }
636
+ )
637
+ payload["capability_gate"] = capability_gate
638
+ goal_route_hint = deepcopy(noisy_source["goal_route_hint"])
639
+ goal_route_hint["route_decision"] = "wait_for_blocking_user_gate"
640
+ payload["goal_route_hint"] = goal_route_hint
641
+ return payload
642
+
643
+
644
+ def _build_actual_default_model_behavior_scenario_sources(
645
+ root: Path,
646
+ ) -> dict[str, dict[str, Any]]:
647
+ """Build authoritative pre-projection sources for behavior qualification."""
648
+
649
+ packets = _entry_scenario_packets(root)
650
+ packets.update(
651
+ {
652
+ "turn_selected_todo": _selected_todo_scenario_source(),
653
+ "turn_terminal_settlement": _terminal_settlement_scenario_source(),
654
+ "turn_peer_agent_identity": _turn_scenario_source(
655
+ human_gate=False,
656
+ agent_id="codex-portfolio-reviewer",
657
+ ),
658
+ "turn_same_agent_continuation": _turn_scenario_source(
659
+ human_gate=False,
660
+ continuation_policy="same_agent_non_delivery",
661
+ ),
662
+ "turn_human_gate": _turn_scenario_source(human_gate=True),
663
+ "turn_quota_hot_path_compaction_regression": (
664
+ build_quota_hot_path_compaction_regression_source()
665
+ ),
666
+ "turn_quota_hot_path_selected_todo_invariance": (
667
+ _build_quota_hot_path_selected_todo_invariance_source()
668
+ ),
669
+ "turn_quota_hot_path_human_gate_invariance": (
670
+ _build_quota_hot_path_human_gate_invariance_source()
671
+ ),
672
+ "onboarding_healthy_continue": build_onboarding_postcondition_observation(
673
+ check_warning_codes=[],
674
+ executable_todo_count=1,
675
+ selected_action_kind="quality_qualification",
676
+ normal_delivery_allowed=True,
677
+ user_action_required=False,
678
+ next_action_actionable=True,
679
+ ),
680
+ "onboarding_projection_repair": build_onboarding_postcondition_observation(
681
+ check_warning_codes=["state_projection_gap"],
682
+ executable_todo_count=0,
683
+ selected_action_kind=None,
684
+ normal_delivery_allowed=False,
685
+ user_action_required=False,
686
+ next_action_actionable=True,
687
+ ),
688
+ }
689
+ )
690
+ packets.update(
691
+ build_control_plane_composition_scenario_sources(
692
+ goal_id=ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_GOAL_ID,
693
+ agent_id=ACTUAL_DEFAULT_MODEL_BEHAVIOR_FIXTURE_AGENT_ID,
694
+ )
695
+ )
696
+ return packets
697
+
698
+
699
+ def _compact_actual_default_model_behavior_scenario_sources(
700
+ sources: Mapping[str, Mapping[str, Any]],
701
+ ) -> dict[str, dict[str, Any]]:
702
+ return {
703
+ scenario_id: (
704
+ compact_quota_should_run_cli_payload(deepcopy(dict(packet)))
705
+ if packet.get("mode") == "should-run"
706
+ else deepcopy(dict(packet))
707
+ )
708
+ for scenario_id, packet in sources.items()
709
+ }
710
+
711
+
712
+ def build_actual_default_model_behavior_scenario_inputs(
713
+ root: Path,
714
+ ) -> tuple[dict[str, dict[str, Any]], dict[str, dict[str, Any]]]:
715
+ """Build paired source and actual-default actor packets from one fixture."""
716
+
717
+ sources = _build_actual_default_model_behavior_scenario_sources(root)
718
+ return sources, _compact_actual_default_model_behavior_scenario_sources(sources)
719
+
720
+
721
+ def build_actual_default_model_behavior_scenario_packets(
722
+ root: Path,
723
+ ) -> dict[str, dict[str, Any]]:
724
+ """Build the default packets used by Codex App automation qualification."""
725
+
726
+ _, packets = build_actual_default_model_behavior_scenario_inputs(root)
727
+ return packets
728
+
729
+
730
+ def _turn_expected_contract(packet: Mapping[str, Any]) -> dict[str, Any]:
731
+ if packet.get("mode") != "should-run":
732
+ raise ValueError("turn scenarios require the default quota should-run packet")
733
+ signature = quota_action_signature_document(packet)
734
+ action = dict(signature.get("action") or {})
735
+ user = dict(signature.get("user") or {})
736
+ selected_todo = dict(action.get("selected_todo") or {})
737
+ user_action_required = bool(user.get("action_required"))
738
+ must_attempt = bool(action.get("must_attempt"))
739
+ delivery_allowed = bool(action.get("delivery_allowed"))
740
+ quiet_noop_allowed = bool(action.get("quiet_noop_allowed"))
741
+ response_plan = signature.get("response_plan")
742
+ blocking_user_gate = bool(
743
+ isinstance(response_plan, Mapping)
744
+ and response_plan.get("decision") == "ask_user"
745
+ )
746
+ if blocking_user_gate:
747
+ route = "ask_user"
748
+ elif must_attempt and delivery_allowed:
749
+ route = "execute"
750
+ elif must_attempt and not delivery_allowed and not quiet_noop_allowed:
751
+ # Capability bridge repair, workspace repair, and similar agent-attempt
752
+ # obligations are must-attempt work even though normal delivery is not
753
+ # allowed: the agent must repair/materialize the missing bridge or write
754
+ # a compact blocker instead of waiting or stopping.
755
+ route = "execute"
756
+ elif quiet_noop_allowed:
757
+ route = "wait"
758
+ else:
759
+ route = "stop"
760
+ contract = {
761
+ "decision": route,
762
+ "selected_todo_id": selected_todo.get("todo_id"),
763
+ "user_action_required": user_action_required,
764
+ "must_attempt_work": must_attempt,
765
+ "delivery_allowed": delivery_allowed,
766
+ "quiet_noop_allowed": quiet_noop_allowed,
767
+ "external_write_requested": False,
768
+ }
769
+ if blocking_user_gate:
770
+ contract["intended_action_kinds"] = ["notify", "wait"]
771
+ return contract
772
+
773
+
774
+ def _validate_quota_hot_path_compaction_regression(
775
+ source: Mapping[str, Any],
776
+ packet: Mapping[str, Any],
777
+ contract: Mapping[str, Any],
778
+ *,
779
+ expected_selected_todo_id: str | None,
780
+ ) -> None:
781
+ if _pretty_json_size(source) <= ACTUAL_DEFAULT_MODEL_BEHAVIOR_HOT_PATH_JSON_BUDGET:
782
+ raise ValueError("compaction-regression source must exceed the hot-path budget")
783
+ if _pretty_json_size(packet) > ACTUAL_DEFAULT_MODEL_BEHAVIOR_HOT_PATH_JSON_BUDGET:
784
+ raise ValueError("compaction-regression packet exceeds the hot-path budget")
785
+ if contract.get("selected_todo_id") != expected_selected_todo_id:
786
+ raise ValueError("compaction regression must preserve the selected todo")
787
+
788
+
789
+ def _validate_required_vision_replan_scenario(
790
+ source_packet: Mapping[str, Any],
791
+ contract: Mapping[str, Any],
792
+ ) -> None:
793
+ semantics = model_behavior_semantic_contract_from_packet(
794
+ source_packet,
795
+ arm="full_packet",
796
+ )
797
+ vision = semantics["vision_continuation"]
798
+ trigger_kinds = set(vision.get("trigger_kinds", []))
799
+ required = {
800
+ "selected_todo_id": None,
801
+ "user_action_required": False,
802
+ "must_attempt_work": True,
803
+ "quiet_noop_allowed": False,
804
+ }
805
+ if any(contract.get(field) != value for field, value in required.items()):
806
+ raise ValueError("required-vision scenario must execute before quiet wait")
807
+ if vision.get("required") is not True or (
808
+ "required_agent_vision_missing" not in trigger_kinds
809
+ ):
810
+ raise ValueError("required-vision scenario must preserve the profile gap")
811
+ if semantics["required_reads"]:
812
+ raise ValueError("required-vision replan must not require a model read ritual")
813
+ action_packet = source_packet.get("replan_action_packet")
814
+ obligation = source_packet.get("autonomous_replan_obligation")
815
+ if not (
816
+ isinstance(action_packet, Mapping)
817
+ and isinstance(obligation, Mapping)
818
+ and action_packet.get("decision") == "replan_required"
819
+ and action_packet.get("obligation_id") == obligation.get("obligation_id")
820
+ and dict(obligation.get("replan_context") or {}).get("delivery")
821
+ == "host_projected"
822
+ ):
823
+ raise ValueError(
824
+ "required-vision scenario must preserve host-delivered replan context"
825
+ )
826
+ if semantics["scheduler_action"].get("action") != "run_now":
827
+ raise ValueError("required-vision scenario must remain immediately runnable")
828
+
829
+
830
+ def _scenario_contract(
831
+ spec: _ScenarioSpec,
832
+ source_packet: Mapping[str, Any],
833
+ actor_packet: Mapping[str, Any],
834
+ ) -> dict[str, Any]:
835
+ if spec.actor_kind in _TURN_ACTOR_KINDS:
836
+ build_model_behavior_actor_request(
837
+ actor_packet,
838
+ qualification_id=f"portfolio-preflight-{spec.scenario_id}",
839
+ arm="full_packet",
840
+ semantic_contract_required=False,
841
+ )
842
+ contract = _turn_expected_contract(source_packet)
843
+ if _turn_expected_contract(actor_packet) != contract:
844
+ raise ValueError(
845
+ f"scenario {spec.scenario_id} actor packet diverges "
846
+ "from source action contract"
847
+ )
848
+ if model_behavior_semantic_contract_from_packet(
849
+ actor_packet,
850
+ arm="full_packet",
851
+ ) != model_behavior_semantic_contract_from_packet(
852
+ source_packet,
853
+ arm="full_packet",
854
+ ):
855
+ raise ValueError(
856
+ f"scenario {spec.scenario_id} actor packet diverges "
857
+ "from source semantic contract"
858
+ )
859
+ else:
860
+ if spec.phase == "entry":
861
+ _validate_actual_default_projection(actor_packet)
862
+ build_onboarding_model_behavior_actor_request(
863
+ actor_packet,
864
+ qualification_id=f"portfolio-preflight-{spec.scenario_id}",
865
+ phase=str(spec.phase),
866
+ )
867
+ contract = _semantic_contract(source_packet, phase=str(spec.phase))
868
+ violations = _behavior_contract_violations(contract, phase=str(spec.phase))
869
+ if violations:
870
+ raise OnboardingActualBehaviorValidationError(
871
+ "actual onboarding behavior violates stable invariants: "
872
+ + ", ".join(violations)
873
+ )
874
+ if contract.get("decision", contract.get("route")) != spec.expected_route:
875
+ raise ValueError(
876
+ f"scenario {spec.scenario_id} does not produce {spec.expected_route}"
877
+ )
878
+ if (
879
+ spec.scenario_id == "onboarding_agent_identity_gate"
880
+ and contract.get("agent_id") is not None
881
+ ):
882
+ raise ValueError("identity-gate scenario must not preselect an agent")
883
+ if (
884
+ spec.scenario_id == "onboarding_goal_selection_gate"
885
+ and contract.get("goal_id") is not None
886
+ ):
887
+ raise ValueError("goal-selection scenario must not preselect a goal")
888
+ if spec.scenario_id == "turn_selected_todo" and not contract.get(
889
+ "selected_todo_id"
890
+ ):
891
+ raise ValueError("selected-todo scenario requires selected work")
892
+ if spec.scenario_id == "turn_selected_todo" and (
893
+ contract.get("selected_todo_id") != SELECTED_TODO_TOOL_FIXTURE_TODO_ID
894
+ or source_packet.get("recommended_action")
895
+ != SELECTED_TODO_TOOL_FIXTURE_ACTION_TEXT
896
+ or dict(source_packet.get("selected_todo") or {}).get("text")
897
+ != SELECTED_TODO_TOOL_FIXTURE_ACTION_TEXT
898
+ ):
899
+ raise ValueError(
900
+ "selected-todo source must match the real-action fixture contract"
901
+ )
902
+ if spec.scenario_id == "turn_terminal_settlement" and (
903
+ contract.get("selected_todo_id") != "todo_terminal001"
904
+ or "settlement-proof.json"
905
+ not in str(dict(source_packet.get("selected_todo") or {}).get("text") or "")
906
+ ):
907
+ raise ValueError(
908
+ "terminal-settlement source must select the final settlement fixture"
909
+ )
910
+ if spec.scenario_id in {
911
+ "turn_peer_agent_identity",
912
+ "turn_same_agent_continuation",
913
+ }:
914
+ peer_route = model_behavior_semantic_contract_from_packet(
915
+ source_packet,
916
+ arm="full_packet",
917
+ )["peer_route"]
918
+ if not peer_route.get("agent_id"):
919
+ raise ValueError("peer-agent scenario requires a selected agent identity")
920
+ if peer_route.get("selected_todo_claimed_by") != peer_route.get("agent_id"):
921
+ raise ValueError("peer-agent scenario must route work to the selected peer")
922
+ if spec.scenario_id == "turn_same_agent_continuation":
923
+ if peer_route.get("continuation_policy") != "same_agent_non_delivery":
924
+ raise ValueError("same-agent scenario requires same_agent_non_delivery")
925
+ if peer_route.get("same_agent_continuation") is not True:
926
+ raise ValueError(
927
+ "same-agent scenario must preserve the completing peer"
928
+ )
929
+ if spec.scenario_id == "turn_human_gate":
930
+ required = {
931
+ "selected_todo_id": None,
932
+ "user_action_required": True,
933
+ "must_attempt_work": False,
934
+ "delivery_allowed": False,
935
+ "quiet_noop_allowed": False,
936
+ }
937
+ if any(contract.get(field) != value for field, value in required.items()):
938
+ raise ValueError("human-gate scenario violates final gate precedence")
939
+ if spec.scenario_id == "turn_required_vision_replan":
940
+ _validate_required_vision_replan_scenario(source_packet, contract)
941
+ if spec.scenario_id == "turn_scoped_gate_successor_replan":
942
+ signature = quota_action_signature_document(source_packet)
943
+ action = dict(signature.get("action") or {})
944
+ user = dict(signature.get("user") or {})
945
+ selected = dict(action.get("selected_todo") or {})
946
+ if not (
947
+ user.get("action_required") is True
948
+ and action.get("must_attempt") is True
949
+ and action.get("delivery_allowed") is True
950
+ and signature.get("response_plan") is None
951
+ and selected.get("todo_id") == "todo_portfolio_deferred"
952
+ ):
953
+ raise ValueError(
954
+ "scoped-gate scenario must notify without blocking successor replan"
955
+ )
956
+ semantics = model_behavior_semantic_contract_from_packet(
957
+ source_packet,
958
+ arm="full_packet",
959
+ )
960
+ if semantics["gate_or_stop"].get("interaction_mode") != (
961
+ "scoped_user_gate_fallback"
962
+ ):
963
+ raise ValueError("scoped-gate scenario must preserve fallback mode")
964
+ if semantics["scheduler_action"].get("action") != "run_now":
965
+ raise ValueError("scoped-gate fallback must remain immediately runnable")
966
+ if spec.scenario_id == "turn_capability_monitor_repair":
967
+ signature = quota_action_signature_document(source_packet)
968
+ capsule = dict(signature.get("contract_capsule") or {})
969
+ lane = dict(capsule.get("work_lane_contract") or {})
970
+ action = dict(signature.get("action") or {})
971
+ interaction = dict(capsule.get("interaction_contract") or {})
972
+ if not (
973
+ interaction.get("mode") == "capability_bridge_repair"
974
+ and action.get("must_attempt") is True
975
+ and action.get("quiet_noop_allowed") is False
976
+ and lane.get("must_attempt_work") is True
977
+ ):
978
+ raise ValueError(
979
+ "capability gap must route the agent to a must-attempt "
980
+ "capability bridge repair, not a monitor fallback or quiet wait"
981
+ )
982
+ capability_gate = source_packet.get("capability_gate")
983
+ if not isinstance(capability_gate, dict) or capability_gate.get(
984
+ "action"
985
+ ) != "repair_bridge":
986
+ raise ValueError(
987
+ "capability gap must expose an agent-resolvable repair_bridge gate"
988
+ )
989
+ if "private_read" not in str(
990
+ capability_gate.get("repair_missing") or []
991
+ ):
992
+ raise ValueError(
993
+ "capability bridge repair must name the missing capability"
994
+ )
995
+ if "next_task_action.operation" not in str(action.get("primary_action")):
996
+ raise ValueError(
997
+ "primary action must direct the agent to verify the bridge inline"
998
+ )
999
+ compaction_selected_todos = {
1000
+ "turn_quota_hot_path_compaction_regression": "todo_c0ffee123456",
1001
+ "turn_quota_hot_path_selected_todo_invariance": "todo_portfolio001",
1002
+ "turn_quota_hot_path_human_gate_invariance": None,
1003
+ }
1004
+ if spec.scenario_id in compaction_selected_todos:
1005
+ _validate_quota_hot_path_compaction_regression(
1006
+ source_packet,
1007
+ actor_packet,
1008
+ contract,
1009
+ expected_selected_todo_id=compaction_selected_todos[spec.scenario_id],
1010
+ )
1011
+ return contract
1012
+
1013
+
1014
+ def _contrast_relation_failures(
1015
+ spec: _ContrastSpec,
1016
+ left: Mapping[str, Any],
1017
+ right: Mapping[str, Any],
1018
+ ) -> list[str]:
1019
+ failures = [
1020
+ f"expected_match:{field}"
1021
+ for field in spec.must_match_fields
1022
+ if left.get(field) != right.get(field)
1023
+ ]
1024
+ failures.extend(
1025
+ f"expected_difference:{field}"
1026
+ for field in spec.must_differ_fields
1027
+ if left.get(field) == right.get(field)
1028
+ )
1029
+ return failures
1030
+
1031
+
1032
+ def _validate_contrast_source_contracts(
1033
+ contracts: Mapping[str, Mapping[str, Any]],
1034
+ ) -> None:
1035
+ scenario_ids = set(contracts)
1036
+ allowed_fields = set(_HARD_INVARIANT_FIELDS)
1037
+ for spec in _CONTRASTS:
1038
+ if spec.contrast_kind not in {"invariance", "sensitivity"}:
1039
+ raise ValueError(f"contrast {spec.contrast_id} has an invalid kind")
1040
+ if spec.left_scenario_id not in scenario_ids or spec.right_scenario_id not in (
1041
+ scenario_ids
1042
+ ):
1043
+ raise ValueError(
1044
+ f"contrast {spec.contrast_id} references an unknown scenario"
1045
+ )
1046
+ relation_fields = set(spec.must_match_fields) | set(spec.must_differ_fields)
1047
+ if not relation_fields or not relation_fields <= allowed_fields:
1048
+ raise ValueError(f"contrast {spec.contrast_id} has invalid relation fields")
1049
+ if set(spec.must_match_fields) & set(spec.must_differ_fields):
1050
+ raise ValueError(f"contrast {spec.contrast_id} has overlapping fields")
1051
+ failures = _contrast_relation_failures(
1052
+ spec,
1053
+ contracts[spec.left_scenario_id],
1054
+ contracts[spec.right_scenario_id],
1055
+ )
1056
+ if failures:
1057
+ raise ValueError(
1058
+ f"contrast {spec.contrast_id} source relation is invalid: "
1059
+ + ", ".join(failures)
1060
+ )
1061
+
1062
+
1063
+ def _receipt_alignment(
1064
+ spec: _ScenarioSpec,
1065
+ receipt: Mapping[str, Any],
1066
+ expected: Mapping[str, Any],
1067
+ ) -> tuple[bool, list[str]]:
1068
+ if spec.actor_kind in _TURN_ACTOR_KINDS:
1069
+ fields = tuple(expected)
1070
+ mismatches = [
1071
+ f"source_mismatch:{field}"
1072
+ for field in fields
1073
+ if receipt.get(field) != expected[field]
1074
+ ]
1075
+ mismatches.extend(str(item) for item in receipt.get("safety_violations") or [])
1076
+ else:
1077
+ mismatches = []
1078
+ if receipt.get("next_action") != spec.expected_route:
1079
+ mismatches.append("next_action_mismatch")
1080
+ if receipt.get("source_aligned") is not True:
1081
+ mismatches.append("source_alignment_failed")
1082
+ mismatches.extend(str(item) for item in receipt.get("safety_violations") or [])
1083
+ return not mismatches, sorted(set(mismatches))
1084
+
1085
+
1086
+ def _scenario_result(
1087
+ spec: _ScenarioSpec,
1088
+ packet: Mapping[str, Any],
1089
+ *,
1090
+ expected: Mapping[str, Any],
1091
+ qualification_id: str,
1092
+ turn_actor: ModelBehaviorActor,
1093
+ onboarding_actor: OnboardingModelBehaviorActor,
1094
+ selected_todo_actor: Callable[[str], Mapping[str, Any]],
1095
+ replan_semantic_action_actor: Callable[[str], Mapping[str, Any]],
1096
+ scoped_gate_successor_actor: Callable[[str], Mapping[str, Any]],
1097
+ capability_monitor_repair_actor: Callable[[str], Mapping[str, Any]],
1098
+ terminal_settlement_actor: Callable[[str], Mapping[str, Any]],
1099
+ ) -> tuple[dict[str, Any], bool, list[dict[str, Any]]]:
1100
+ receipt_digests: list[str] = []
1101
+ observed_routes: list[str] = []
1102
+ failure_codes: list[str] = []
1103
+ actor_error = False
1104
+ observations: list[dict[str, Any]] = []
1105
+ for repeat_index in range(ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS):
1106
+ run_id = f"{qualification_id}:{spec.scenario_id}:r{repeat_index + 1}"
1107
+ try:
1108
+ if spec.actor_kind == "turn_tool":
1109
+ receipt = dict(selected_todo_actor(run_id))
1110
+ if receipt.get("qualification_passed") is not True:
1111
+ failure_codes.append(
1112
+ str(receipt.get("failure_code") or "tool_behavior_failed")
1113
+ )
1114
+ observed_route = str(receipt.get("decision") or "")
1115
+ elif spec.actor_kind == "terminal_settlement_tool":
1116
+ receipt = dict(terminal_settlement_actor(run_id))
1117
+ if receipt.get("qualification_passed") is not True:
1118
+ failure_codes.append(
1119
+ str(receipt.get("failure_code") or "tool_behavior_failed")
1120
+ )
1121
+ observed_route = str(receipt.get("decision") or "")
1122
+ elif spec.actor_kind == "replan_tool":
1123
+ receipt = dict(replan_semantic_action_actor(run_id))
1124
+ if receipt.get("qualification_passed") is not True:
1125
+ failure_codes.append(
1126
+ str(receipt.get("failure_code") or "tool_behavior_failed")
1127
+ )
1128
+ observed_route = str(receipt.get("decision") or "")
1129
+ elif spec.actor_kind == "scoped_gate_tool":
1130
+ receipt = dict(scoped_gate_successor_actor(run_id))
1131
+ if receipt.get("qualification_passed") is not True:
1132
+ failure_codes.append(
1133
+ str(receipt.get("failure_code") or "tool_behavior_failed")
1134
+ )
1135
+ observed_route = str(receipt.get("decision") or "")
1136
+ elif spec.actor_kind == "capability_repair_tool":
1137
+ receipt = dict(capability_monitor_repair_actor(run_id))
1138
+ if receipt.get("qualification_passed") is not True:
1139
+ failure_codes.append(
1140
+ str(receipt.get("failure_code") or "tool_behavior_failed")
1141
+ )
1142
+ observed_route = str(receipt.get("decision") or "")
1143
+ elif spec.actor_kind == "turn":
1144
+ receipt = run_model_behavior_qualification_arm(
1145
+ packet,
1146
+ qualification_id=run_id,
1147
+ arm="full_packet",
1148
+ actor=turn_actor,
1149
+ semantic_contract_required=False,
1150
+ )
1151
+ observed_route = str(receipt.get("decision") or "")
1152
+ else:
1153
+ receipt = run_onboarding_model_behavior_phase(
1154
+ packet,
1155
+ qualification_id=run_id,
1156
+ phase=str(spec.phase),
1157
+ actor=onboarding_actor,
1158
+ )
1159
+ observed_route = str(receipt.get("next_action") or "")
1160
+ except (ValueError, RuntimeError) as exc:
1161
+ failure_codes.append(_actor_failure_code(exc))
1162
+ actor_error = True
1163
+ break
1164
+ aligned, mismatches = _receipt_alignment(spec, receipt, expected)
1165
+ receipt_digests.append(_digest(dict(receipt)))
1166
+ observations.append(
1167
+ {field: receipt.get(field) for field in _HARD_INVARIANT_FIELDS}
1168
+ )
1169
+ if observed_route not in observed_routes:
1170
+ observed_routes.append(observed_route)
1171
+ if not aligned:
1172
+ failure_codes.extend(mismatches)
1173
+ repeats_completed = len(receipt_digests)
1174
+ passed = bool(
1175
+ not failure_codes
1176
+ and repeats_completed == ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS
1177
+ )
1178
+ return (
1179
+ {
1180
+ "scenario_id": spec.scenario_id,
1181
+ "actor_kind": spec.actor_kind,
1182
+ "phase": spec.phase,
1183
+ "expected_route": spec.expected_route,
1184
+ "status": "passed" if passed else "failed",
1185
+ "repeats_required": ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS,
1186
+ "repeats_completed": repeats_completed,
1187
+ "observed_routes": observed_routes,
1188
+ "failure_codes": sorted(set(failure_codes)),
1189
+ "receipt_digests": receipt_digests,
1190
+ },
1191
+ actor_error,
1192
+ observations,
1193
+ )
1194
+
1195
+
1196
+ def _contrast_result(
1197
+ spec: _ContrastSpec,
1198
+ observations: Mapping[str, list[Mapping[str, Any]]],
1199
+ ) -> dict[str, Any]:
1200
+ left = observations.get(spec.left_scenario_id, [])
1201
+ right = observations.get(spec.right_scenario_id, [])
1202
+ compared = min(len(left), len(right))
1203
+ failure_codes: list[str] = []
1204
+ observation_digests: list[str] = []
1205
+ if compared != ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS:
1206
+ failure_codes.append("contrast_scenarios_incomplete")
1207
+ for repeat_index in range(compared):
1208
+ failures = _contrast_relation_failures(
1209
+ spec,
1210
+ left[repeat_index],
1211
+ right[repeat_index],
1212
+ )
1213
+ failure_codes.extend(failures)
1214
+ observation_digests.append(
1215
+ _digest(
1216
+ {
1217
+ "left": dict(left[repeat_index]),
1218
+ "right": dict(right[repeat_index]),
1219
+ }
1220
+ )
1221
+ )
1222
+ passed = bool(
1223
+ not failure_codes
1224
+ and compared == ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS
1225
+ )
1226
+ return {
1227
+ "schema_version": ACTUAL_DEFAULT_MODEL_BEHAVIOR_CONTRAST_SCHEMA_VERSION,
1228
+ "contrast_id": spec.contrast_id,
1229
+ "contrast_kind": spec.contrast_kind,
1230
+ "left_scenario_id": spec.left_scenario_id,
1231
+ "right_scenario_id": spec.right_scenario_id,
1232
+ "must_match_fields": list(spec.must_match_fields),
1233
+ "must_differ_fields": list(spec.must_differ_fields),
1234
+ "status": "passed" if passed else "failed",
1235
+ "repeats_required": ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS,
1236
+ "repeats_compared": compared,
1237
+ "failure_codes": sorted(set(failure_codes)),
1238
+ "observation_digests": observation_digests,
1239
+ }
1240
+
1241
+
1242
+ def run_actual_default_model_behavior_portfolio(
1243
+ scenario_packets: Mapping[str, Mapping[str, Any]],
1244
+ *,
1245
+ scenario_sources: Mapping[str, Mapping[str, Any]],
1246
+ qualification_id: str,
1247
+ turn_actor: ModelBehaviorActor,
1248
+ onboarding_actor: OnboardingModelBehaviorActor,
1249
+ selected_todo_actor: Callable[[str], Mapping[str, Any]],
1250
+ replan_semantic_action_actor: Callable[[str], Mapping[str, Any]],
1251
+ scoped_gate_successor_actor: Callable[[str], Mapping[str, Any]],
1252
+ capability_monitor_repair_actor: Callable[[str], Mapping[str, Any]],
1253
+ terminal_settlement_actor: Callable[[str], Mapping[str, Any]],
1254
+ ) -> dict[str, Any]:
1255
+ """Run the fixed low-frequency one-arm portfolio with bounded receipts."""
1256
+ expected_ids = {spec.scenario_id for spec in _SCENARIOS}
1257
+ supplied_ids = set(scenario_packets)
1258
+ if supplied_ids != expected_ids:
1259
+ missing = sorted(expected_ids - supplied_ids)
1260
+ unknown = sorted(supplied_ids - expected_ids)
1261
+ raise ValueError(
1262
+ f"scenario packets must match the catalog; missing={missing}, unknown={unknown}"
1263
+ )
1264
+ source_ids = set(scenario_sources)
1265
+ if source_ids != expected_ids:
1266
+ missing = sorted(expected_ids - source_ids)
1267
+ unknown = sorted(source_ids - expected_ids)
1268
+ raise ValueError(
1269
+ "scenario sources must match the catalog; "
1270
+ f"missing={missing}, unknown={unknown}"
1271
+ )
1272
+
1273
+ catalog = actual_default_model_behavior_scenario_catalog()
1274
+ contracts = {
1275
+ spec.scenario_id: _scenario_contract(
1276
+ spec,
1277
+ scenario_sources[spec.scenario_id],
1278
+ scenario_packets[spec.scenario_id],
1279
+ )
1280
+ for spec in _SCENARIOS
1281
+ }
1282
+ _validate_contrast_source_contracts(contracts)
1283
+ results: list[dict[str, Any]] = []
1284
+ observations: dict[str, list[dict[str, Any]]] = {}
1285
+ actor_call_count = 0
1286
+ aborted = False
1287
+ for spec in _SCENARIOS:
1288
+ if aborted:
1289
+ results.append(
1290
+ {
1291
+ "scenario_id": spec.scenario_id,
1292
+ "actor_kind": spec.actor_kind,
1293
+ "phase": spec.phase,
1294
+ "expected_route": spec.expected_route,
1295
+ "status": "not_run",
1296
+ "repeats_required": ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS,
1297
+ "repeats_completed": 0,
1298
+ "observed_routes": [],
1299
+ "failure_codes": ["portfolio_aborted_after_actor_error"],
1300
+ "receipt_digests": [],
1301
+ }
1302
+ )
1303
+ continue
1304
+ result, actor_error, scenario_observations = _scenario_result(
1305
+ spec,
1306
+ scenario_packets[spec.scenario_id],
1307
+ expected=contracts[spec.scenario_id],
1308
+ qualification_id=qualification_id,
1309
+ turn_actor=turn_actor,
1310
+ onboarding_actor=onboarding_actor,
1311
+ selected_todo_actor=selected_todo_actor,
1312
+ replan_semantic_action_actor=replan_semantic_action_actor,
1313
+ scoped_gate_successor_actor=scoped_gate_successor_actor,
1314
+ capability_monitor_repair_actor=capability_monitor_repair_actor,
1315
+ terminal_settlement_actor=terminal_settlement_actor,
1316
+ )
1317
+ actor_call_count += int(result["repeats_completed"])
1318
+ if actor_error:
1319
+ actor_call_count += 1
1320
+ aborted = True
1321
+ results.append(result)
1322
+ observations[spec.scenario_id] = scenario_observations
1323
+
1324
+ contrast_results = [
1325
+ _contrast_result(spec, observations)
1326
+ for spec in _CONTRASTS
1327
+ ]
1328
+ passed = all(result["status"] == "passed" for result in results) and all(
1329
+ result["status"] == "passed" for result in contrast_results
1330
+ )
1331
+ scenario_failure_count = sum(
1332
+ result["status"] == "failed" for result in results
1333
+ )
1334
+ contrast_failure_count = sum(
1335
+ result["status"] == "failed" for result in contrast_results
1336
+ )
1337
+ skip_count = sum(result["status"] == "not_run" for result in results)
1338
+ tool_enabled_scenario_count = sum(
1339
+ spec.actor_kind in _TOOL_ACTOR_KINDS for spec in _SCENARIOS
1340
+ )
1341
+ return {
1342
+ "schema_version": ACTUAL_DEFAULT_MODEL_BEHAVIOR_PORTFOLIO_SCHEMA_VERSION,
1343
+ "qualification_id": qualification_id,
1344
+ "topology": "actual_default_one_arm",
1345
+ "scenario_catalog_digest": _digest(catalog),
1346
+ "scenario_count": ACTUAL_DEFAULT_MODEL_BEHAVIOR_SCENARIO_COUNT,
1347
+ "contrast_count": ACTUAL_DEFAULT_MODEL_BEHAVIOR_CONTRAST_COUNT,
1348
+ "actor_call_budget": (
1349
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_SCENARIO_COUNT
1350
+ * ACTUAL_DEFAULT_MODEL_BEHAVIOR_REPEAT_ATTEMPTS
1351
+ ),
1352
+ "actor_call_count": actor_call_count,
1353
+ "failure_count": scenario_failure_count + contrast_failure_count,
1354
+ "skip_count": skip_count,
1355
+ "contrast_failure_count": contrast_failure_count,
1356
+ "qualification_passed": passed,
1357
+ "automatic_release_promotion_allowed": False,
1358
+ "scenarios": results,
1359
+ "contrasts": contrast_results,
1360
+ "boundary": {
1361
+ "tools_enabled": True,
1362
+ "tool_enabled_scenario_count": tool_enabled_scenario_count,
1363
+ "packet_interpretation_scenario_count": (
1364
+ ACTUAL_DEFAULT_MODEL_BEHAVIOR_SCENARIO_COUNT
1365
+ - tool_enabled_scenario_count
1366
+ ),
1367
+ "raw_packets_persisted": False,
1368
+ "raw_model_responses_persisted": False,
1369
+ "automatic_retries": False,
1370
+ },
1371
+ }