@raishin/vanguard-frontier-agentic 3.10.0 → 3.11.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (356) hide show
  1. package/.claude-plugin/marketplace.json +2 -2
  2. package/.claude-plugin/plugin.json +18 -1
  3. package/.cursor-plugin/plugin.json +18 -1
  4. package/.github/plugin/marketplace.json +1 -1
  5. package/README.md +21 -17
  6. package/agents/databricks/databricks-ai-bi-genie-agent/AGENT.md +90 -0
  7. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/claude-code.agent.md +73 -0
  8. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/codex.toml +15 -0
  9. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/copilot.agent.md +79 -0
  10. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/cursor.agent.md +74 -0
  11. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/gemini.agent.md +73 -0
  12. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/kiro-cli.agent.json +5 -0
  13. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/kiro-ide.agent.md +73 -0
  14. package/agents/databricks/databricks-ai-bi-genie-agent/metadata.json +58 -0
  15. package/agents/databricks/databricks-data-protection-privacy-agent/AGENT.md +94 -0
  16. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/claude-code.agent.md +77 -0
  17. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/codex.toml +15 -0
  18. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/copilot.agent.md +83 -0
  19. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/cursor.agent.md +78 -0
  20. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/gemini.agent.md +77 -0
  21. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/kiro-cli.agent.json +5 -0
  22. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/kiro-ide.agent.md +77 -0
  23. package/agents/databricks/databricks-data-protection-privacy-agent/metadata.json +64 -0
  24. package/agents/databricks/databricks-data-quality-observability-agent/AGENT.md +89 -0
  25. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/claude-code.agent.md +72 -0
  26. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/codex.toml +15 -0
  27. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/copilot.agent.md +78 -0
  28. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/cursor.agent.md +73 -0
  29. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/gemini.agent.md +72 -0
  30. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/kiro-cli.agent.json +5 -0
  31. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/kiro-ide.agent.md +72 -0
  32. package/agents/databricks/databricks-data-quality-observability-agent/metadata.json +59 -0
  33. package/agents/databricks/databricks-developer-platform-agent/AGENT.md +90 -0
  34. package/agents/databricks/databricks-developer-platform-agent/harnesses/claude-code.agent.md +73 -0
  35. package/agents/databricks/databricks-developer-platform-agent/harnesses/codex.toml +15 -0
  36. package/agents/databricks/databricks-developer-platform-agent/harnesses/copilot.agent.md +79 -0
  37. package/agents/databricks/databricks-developer-platform-agent/harnesses/cursor.agent.md +74 -0
  38. package/agents/databricks/databricks-developer-platform-agent/harnesses/gemini.agent.md +73 -0
  39. package/agents/databricks/databricks-developer-platform-agent/harnesses/kiro-cli.agent.json +5 -0
  40. package/agents/databricks/databricks-developer-platform-agent/harnesses/kiro-ide.agent.md +73 -0
  41. package/agents/databricks/databricks-developer-platform-agent/metadata.json +59 -0
  42. package/agents/databricks/databricks-finops-cost-agent/AGENT.md +91 -0
  43. package/agents/databricks/databricks-finops-cost-agent/harnesses/claude-code.agent.md +74 -0
  44. package/agents/databricks/databricks-finops-cost-agent/harnesses/codex.toml +15 -0
  45. package/agents/databricks/databricks-finops-cost-agent/harnesses/copilot.agent.md +80 -0
  46. package/agents/databricks/databricks-finops-cost-agent/harnesses/cursor.agent.md +75 -0
  47. package/agents/databricks/databricks-finops-cost-agent/harnesses/gemini.agent.md +74 -0
  48. package/agents/databricks/databricks-finops-cost-agent/harnesses/kiro-cli.agent.json +5 -0
  49. package/agents/databricks/databricks-finops-cost-agent/harnesses/kiro-ide.agent.md +74 -0
  50. package/agents/databricks/databricks-finops-cost-agent/metadata.json +60 -0
  51. package/agents/databricks/databricks-genai-agent-engineering-agent/AGENT.md +89 -0
  52. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/claude-code.agent.md +72 -0
  53. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/codex.toml +15 -0
  54. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/copilot.agent.md +78 -0
  55. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/cursor.agent.md +73 -0
  56. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/gemini.agent.md +72 -0
  57. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/kiro-cli.agent.json +5 -0
  58. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/kiro-ide.agent.md +72 -0
  59. package/agents/databricks/databricks-genai-agent-engineering-agent/metadata.json +62 -0
  60. package/agents/databricks/databricks-genai-evaluation-observability-agent/AGENT.md +89 -0
  61. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/claude-code.agent.md +72 -0
  62. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/codex.toml +15 -0
  63. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/copilot.agent.md +78 -0
  64. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/cursor.agent.md +73 -0
  65. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/gemini.agent.md +72 -0
  66. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/kiro-cli.agent.json +5 -0
  67. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/kiro-ide.agent.md +72 -0
  68. package/agents/databricks/databricks-genai-evaluation-observability-agent/metadata.json +59 -0
  69. package/agents/databricks/databricks-identity-network-security-agent/AGENT.md +95 -0
  70. package/agents/databricks/databricks-identity-network-security-agent/harnesses/claude-code.agent.md +78 -0
  71. package/agents/databricks/databricks-identity-network-security-agent/harnesses/codex.toml +15 -0
  72. package/agents/databricks/databricks-identity-network-security-agent/harnesses/copilot.agent.md +84 -0
  73. package/agents/databricks/databricks-identity-network-security-agent/harnesses/cursor.agent.md +79 -0
  74. package/agents/databricks/databricks-identity-network-security-agent/harnesses/gemini.agent.md +78 -0
  75. package/agents/databricks/databricks-identity-network-security-agent/harnesses/kiro-cli.agent.json +5 -0
  76. package/agents/databricks/databricks-identity-network-security-agent/harnesses/kiro-ide.agent.md +78 -0
  77. package/agents/databricks/databricks-identity-network-security-agent/metadata.json +60 -0
  78. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/AGENT.md +90 -0
  79. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/claude-code.agent.md +73 -0
  80. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/codex.toml +15 -0
  81. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/copilot.agent.md +79 -0
  82. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/cursor.agent.md +74 -0
  83. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/gemini.agent.md +73 -0
  84. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/kiro-cli.agent.json +5 -0
  85. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/kiro-ide.agent.md +73 -0
  86. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/metadata.json +63 -0
  87. package/agents/databricks/databricks-maestro-agent/AGENT.md +63 -0
  88. package/agents/databricks/databricks-maestro-agent/README.md +76 -0
  89. package/agents/databricks/databricks-maestro-agent/harnesses/claude-code.agent.md +46 -0
  90. package/agents/databricks/databricks-maestro-agent/harnesses/codex.toml +15 -0
  91. package/agents/databricks/databricks-maestro-agent/harnesses/copilot.agent.md +52 -0
  92. package/agents/databricks/databricks-maestro-agent/harnesses/cursor.agent.md +47 -0
  93. package/agents/databricks/databricks-maestro-agent/harnesses/gemini.agent.md +46 -0
  94. package/agents/databricks/databricks-maestro-agent/harnesses/kiro-cli.agent.json +5 -0
  95. package/agents/databricks/databricks-maestro-agent/harnesses/kiro-ide.agent.md +46 -0
  96. package/agents/databricks/databricks-maestro-agent/metadata.json +50 -0
  97. package/agents/databricks/databricks-mlops-agent/AGENT.md +89 -0
  98. package/agents/databricks/databricks-mlops-agent/harnesses/claude-code.agent.md +72 -0
  99. package/agents/databricks/databricks-mlops-agent/harnesses/codex.toml +15 -0
  100. package/agents/databricks/databricks-mlops-agent/harnesses/copilot.agent.md +78 -0
  101. package/agents/databricks/databricks-mlops-agent/harnesses/cursor.agent.md +73 -0
  102. package/agents/databricks/databricks-mlops-agent/harnesses/gemini.agent.md +72 -0
  103. package/agents/databricks/databricks-mlops-agent/harnesses/kiro-cli.agent.json +5 -0
  104. package/agents/databricks/databricks-mlops-agent/harnesses/kiro-ide.agent.md +72 -0
  105. package/agents/databricks/databricks-mlops-agent/metadata.json +60 -0
  106. package/agents/databricks/databricks-platform-architecture-agent/AGENT.md +90 -0
  107. package/agents/databricks/databricks-platform-architecture-agent/harnesses/claude-code.agent.md +73 -0
  108. package/agents/databricks/databricks-platform-architecture-agent/harnesses/codex.toml +15 -0
  109. package/agents/databricks/databricks-platform-architecture-agent/harnesses/copilot.agent.md +79 -0
  110. package/agents/databricks/databricks-platform-architecture-agent/harnesses/cursor.agent.md +74 -0
  111. package/agents/databricks/databricks-platform-architecture-agent/harnesses/gemini.agent.md +73 -0
  112. package/agents/databricks/databricks-platform-architecture-agent/harnesses/kiro-cli.agent.json +5 -0
  113. package/agents/databricks/databricks-platform-architecture-agent/harnesses/kiro-ide.agent.md +73 -0
  114. package/agents/databricks/databricks-platform-architecture-agent/metadata.json +58 -0
  115. package/agents/databricks/databricks-platform-reliability-agent/AGENT.md +88 -0
  116. package/agents/databricks/databricks-platform-reliability-agent/harnesses/claude-code.agent.md +71 -0
  117. package/agents/databricks/databricks-platform-reliability-agent/harnesses/codex.toml +15 -0
  118. package/agents/databricks/databricks-platform-reliability-agent/harnesses/copilot.agent.md +77 -0
  119. package/agents/databricks/databricks-platform-reliability-agent/harnesses/cursor.agent.md +72 -0
  120. package/agents/databricks/databricks-platform-reliability-agent/harnesses/gemini.agent.md +71 -0
  121. package/agents/databricks/databricks-platform-reliability-agent/harnesses/kiro-cli.agent.json +5 -0
  122. package/agents/databricks/databricks-platform-reliability-agent/harnesses/kiro-ide.agent.md +71 -0
  123. package/agents/databricks/databricks-platform-reliability-agent/metadata.json +64 -0
  124. package/agents/databricks/databricks-sql-performance-agent/AGENT.md +91 -0
  125. package/agents/databricks/databricks-sql-performance-agent/harnesses/claude-code.agent.md +74 -0
  126. package/agents/databricks/databricks-sql-performance-agent/harnesses/codex.toml +15 -0
  127. package/agents/databricks/databricks-sql-performance-agent/harnesses/copilot.agent.md +80 -0
  128. package/agents/databricks/databricks-sql-performance-agent/harnesses/cursor.agent.md +75 -0
  129. package/agents/databricks/databricks-sql-performance-agent/harnesses/gemini.agent.md +74 -0
  130. package/agents/databricks/databricks-sql-performance-agent/harnesses/kiro-cli.agent.json +5 -0
  131. package/agents/databricks/databricks-sql-performance-agent/harnesses/kiro-ide.agent.md +74 -0
  132. package/agents/databricks/databricks-sql-performance-agent/metadata.json +62 -0
  133. package/agents/databricks/databricks-streaming-reliability-agent/AGENT.md +93 -0
  134. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/claude-code.agent.md +76 -0
  135. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/codex.toml +15 -0
  136. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/copilot.agent.md +82 -0
  137. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/cursor.agent.md +77 -0
  138. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/gemini.agent.md +76 -0
  139. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/kiro-cli.agent.json +5 -0
  140. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/kiro-ide.agent.md +76 -0
  141. package/agents/databricks/databricks-streaming-reliability-agent/metadata.json +62 -0
  142. package/agents/databricks/databricks-unity-catalog-governance-agent/AGENT.md +91 -0
  143. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/claude-code.agent.md +74 -0
  144. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/codex.toml +15 -0
  145. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/copilot.agent.md +80 -0
  146. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/cursor.agent.md +75 -0
  147. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/gemini.agent.md +74 -0
  148. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/kiro-cli.agent.json +5 -0
  149. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/kiro-ide.agent.md +74 -0
  150. package/agents/databricks/databricks-unity-catalog-governance-agent/metadata.json +63 -0
  151. package/agents/databricks/databricks-value-realization-agent/AGENT.md +91 -0
  152. package/agents/databricks/databricks-value-realization-agent/harnesses/claude-code.agent.md +74 -0
  153. package/agents/databricks/databricks-value-realization-agent/harnesses/codex.toml +15 -0
  154. package/agents/databricks/databricks-value-realization-agent/harnesses/copilot.agent.md +80 -0
  155. package/agents/databricks/databricks-value-realization-agent/harnesses/cursor.agent.md +75 -0
  156. package/agents/databricks/databricks-value-realization-agent/harnesses/gemini.agent.md +74 -0
  157. package/agents/databricks/databricks-value-realization-agent/harnesses/kiro-cli.agent.json +5 -0
  158. package/agents/databricks/databricks-value-realization-agent/harnesses/kiro-ide.agent.md +74 -0
  159. package/agents/databricks/databricks-value-realization-agent/metadata.json +55 -0
  160. package/catalog/agents.json +580 -0
  161. package/catalog/asset-integrity.json +1139 -44
  162. package/catalog/install-roles.json +166 -0
  163. package/catalog/model-assignments.json +561 -0
  164. package/catalog/skill-manifest.json +709 -0
  165. package/catalog/skills.json +529 -0
  166. package/package.json +1 -1
  167. package/plugins/vanguard-frontier-agentic/.codex-plugin/plugin.json +1 -1
  168. package/powers/vanguard-databricks/POWER.md +11 -11
  169. package/scripts/databricks_data/agents/00-databricks-maestro-agent.json +165 -0
  170. package/scripts/databricks_data/agents/01-databricks-platform-architecture-agent.json +200 -0
  171. package/scripts/databricks_data/agents/02-databricks-unity-catalog-governance-agent.json +207 -0
  172. package/scripts/databricks_data/agents/03-databricks-identity-network-security-agent.json +216 -0
  173. package/scripts/databricks_data/agents/04-databricks-data-protection-privacy-agent.json +218 -0
  174. package/scripts/databricks_data/agents/05-databricks-lakeflow-pipeline-engineering-agent.json +217 -0
  175. package/scripts/databricks_data/agents/06-databricks-streaming-reliability-agent.json +267 -0
  176. package/scripts/databricks_data/agents/07-databricks-data-quality-observability-agent.json +215 -0
  177. package/scripts/databricks_data/agents/08-databricks-sql-performance-agent.json +214 -0
  178. package/scripts/databricks_data/agents/09-databricks-ai-bi-genie-agent.json +211 -0
  179. package/scripts/databricks_data/agents/10-databricks-mlops-agent.json +239 -0
  180. package/scripts/databricks_data/agents/11-databricks-genai-agent-engineering-agent.json +246 -0
  181. package/scripts/databricks_data/agents/12-databricks-genai-evaluation-observability-agent.json +215 -0
  182. package/scripts/databricks_data/agents/13-databricks-developer-platform-agent.json +206 -0
  183. package/scripts/databricks_data/agents/14-databricks-platform-reliability-agent.json +208 -0
  184. package/scripts/databricks_data/agents/15-databricks-finops-cost-agent.json +218 -0
  185. package/scripts/databricks_data/agents/16-databricks-value-realization-agent.json +232 -0
  186. package/scripts/gen_databricks_agents.py +703 -0
  187. package/scripts/generate-board-counts.mjs +6 -0
  188. package/scripts/generate-kiro-powers.mjs +5 -5
  189. package/scripts/generate-readme-counts.mjs +109 -0
  190. package/skills/databricks/databricks-ai-bi-genie/SKILL.md +132 -0
  191. package/skills/databricks/databricks-ai-bi-genie/metadata.json +34 -0
  192. package/skills/databricks/databricks-ai-bi-genie/references/dashboard-and-permission-security.md +16 -0
  193. package/skills/databricks/databricks-ai-bi-genie/references/genie-scoping-and-semantic-layer.md +16 -0
  194. package/skills/databricks/databricks-ai-bi-genie/references/official-sources.md +24 -0
  195. package/skills/databricks/databricks-ai-bi-genie/references/safety-checklist.md +35 -0
  196. package/skills/databricks/databricks-ai-bi-genie/references/workflow-and-output.md +24 -0
  197. package/skills/databricks/databricks-data-protection-privacy/SKILL.md +142 -0
  198. package/skills/databricks/databricks-data-protection-privacy/metadata.json +37 -0
  199. package/skills/databricks/databricks-data-protection-privacy/references/deletion-vacuum-and-gdpr-compliance.md +9 -0
  200. package/skills/databricks/databricks-data-protection-privacy/references/masks-filters-and-abac-udf-cost.md +9 -0
  201. package/skills/databricks/databricks-data-protection-privacy/references/official-sources.md +27 -0
  202. package/skills/databricks/databricks-data-protection-privacy/references/safety-checklist.md +36 -0
  203. package/skills/databricks/databricks-data-protection-privacy/references/workflow-and-output.md +28 -0
  204. package/skills/databricks/databricks-data-quality-observability/SKILL.md +137 -0
  205. package/skills/databricks/databricks-data-quality-observability/metadata.json +34 -0
  206. package/skills/databricks/databricks-data-quality-observability/references/expectations-and-constraints.md +16 -0
  207. package/skills/databricks/databricks-data-quality-observability/references/monitoring-freshness-and-event-logs.md +17 -0
  208. package/skills/databricks/databricks-data-quality-observability/references/official-sources.md +24 -0
  209. package/skills/databricks/databricks-data-quality-observability/references/safety-checklist.md +34 -0
  210. package/skills/databricks/databricks-data-quality-observability/references/workflow-and-output.md +24 -0
  211. package/skills/databricks/databricks-developer-platform/SKILL.md +134 -0
  212. package/skills/databricks/databricks-developer-platform/metadata.json +34 -0
  213. package/skills/databricks/databricks-developer-platform/references/authentication-and-git-flow.md +9 -0
  214. package/skills/databricks/databricks-developer-platform/references/bundle-structure-and-targets.md +10 -0
  215. package/skills/databricks/databricks-developer-platform/references/official-sources.md +28 -0
  216. package/skills/databricks/databricks-developer-platform/references/safety-checklist.md +35 -0
  217. package/skills/databricks/databricks-developer-platform/references/workflow-and-output.md +26 -0
  218. package/skills/databricks/databricks-finops-cost/SKILL.md +134 -0
  219. package/skills/databricks/databricks-finops-cost/metadata.json +34 -0
  220. package/skills/databricks/databricks-finops-cost/references/billing-system-tables-and-joins.md +15 -0
  221. package/skills/databricks/databricks-finops-cost/references/cost-attribution-and-uptime-charging.md +20 -0
  222. package/skills/databricks/databricks-finops-cost/references/official-sources.md +24 -0
  223. package/skills/databricks/databricks-finops-cost/references/safety-checklist.md +35 -0
  224. package/skills/databricks/databricks-finops-cost/references/workflow-and-output.md +26 -0
  225. package/skills/databricks/databricks-genai-agent-engineering/SKILL.md +133 -0
  226. package/skills/databricks/databricks-genai-agent-engineering/metadata.json +34 -0
  227. package/skills/databricks/databricks-genai-agent-engineering/references/ai-search-and-retrieval-config.md +12 -0
  228. package/skills/databricks/databricks-genai-agent-engineering/references/context-engineering-and-tools.md +20 -0
  229. package/skills/databricks/databricks-genai-agent-engineering/references/official-sources.md +28 -0
  230. package/skills/databricks/databricks-genai-agent-engineering/references/safety-checklist.md +35 -0
  231. package/skills/databricks/databricks-genai-agent-engineering/references/workflow-and-output.md +22 -0
  232. package/skills/databricks/databricks-genai-evaluation-observability/SKILL.md +139 -0
  233. package/skills/databricks/databricks-genai-evaluation-observability/metadata.json +34 -0
  234. package/skills/databricks/databricks-genai-evaluation-observability/references/judges-scorers-and-validation.md +12 -0
  235. package/skills/databricks/databricks-genai-evaluation-observability/references/official-sources.md +28 -0
  236. package/skills/databricks/databricks-genai-evaluation-observability/references/safety-checklist.md +35 -0
  237. package/skills/databricks/databricks-genai-evaluation-observability/references/tracing-storage-and-regression-detection.md +12 -0
  238. package/skills/databricks/databricks-genai-evaluation-observability/references/workflow-and-output.md +24 -0
  239. package/skills/databricks/databricks-identity-network-security/SKILL.md +143 -0
  240. package/skills/databricks/databricks-identity-network-security/metadata.json +34 -0
  241. package/skills/databricks/databricks-identity-network-security/references/admin-roles-and-separation.md +9 -0
  242. package/skills/databricks/databricks-identity-network-security/references/official-sources.md +24 -0
  243. package/skills/databricks/databricks-identity-network-security/references/safety-checklist.md +36 -0
  244. package/skills/databricks/databricks-identity-network-security/references/token-lifecycle-and-automatic-revocation.md +9 -0
  245. package/skills/databricks/databricks-identity-network-security/references/workflow-and-output.md +28 -0
  246. package/skills/databricks/databricks-lakeflow-pipeline-engineering/SKILL.md +134 -0
  247. package/skills/databricks/databricks-lakeflow-pipeline-engineering/metadata.json +35 -0
  248. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/auto-loader-and-schema-evolution.md +15 -0
  249. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/delta-table-layout-strategy.md +15 -0
  250. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/official-sources.md +29 -0
  251. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/safety-checklist.md +34 -0
  252. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/workflow-and-output.md +23 -0
  253. package/skills/databricks/databricks-maestro/SKILL.md +122 -0
  254. package/skills/databricks/databricks-maestro/metadata.json +30 -0
  255. package/skills/databricks/databricks-maestro/references/official-sources.md +20 -0
  256. package/skills/databricks/databricks-maestro/references/routing-taxonomy.md +16 -0
  257. package/skills/databricks/databricks-maestro/references/safety-checklist.md +35 -0
  258. package/skills/databricks/databricks-maestro/references/workflow-and-output.md +25 -0
  259. package/skills/databricks/databricks-mlops/SKILL.md +127 -0
  260. package/skills/databricks/databricks-mlops/metadata.json +33 -0
  261. package/skills/databricks/databricks-mlops/references/mlflow-3-registry-defaults.md +12 -0
  262. package/skills/databricks/databricks-mlops/references/official-sources.md +27 -0
  263. package/skills/databricks/databricks-mlops/references/safety-checklist.md +34 -0
  264. package/skills/databricks/databricks-mlops/references/serving-and-inference-design.md +22 -0
  265. package/skills/databricks/databricks-mlops/references/workflow-and-output.md +21 -0
  266. package/skills/databricks/databricks-platform-architecture/SKILL.md +134 -0
  267. package/skills/databricks/databricks-platform-architecture/metadata.json +34 -0
  268. package/skills/databricks/databricks-platform-architecture/references/metastore-per-region-constraint.md +9 -0
  269. package/skills/databricks/databricks-platform-architecture/references/official-sources.md +24 -0
  270. package/skills/databricks/databricks-platform-architecture/references/safety-checklist.md +34 -0
  271. package/skills/databricks/databricks-platform-architecture/references/workflow-and-output.md +26 -0
  272. package/skills/databricks/databricks-platform-architecture/references/workspace-segmentation-guidance.md +9 -0
  273. package/skills/databricks/databricks-platform-reliability/SKILL.md +134 -0
  274. package/skills/databricks/databricks-platform-reliability/metadata.json +36 -0
  275. package/skills/databricks/databricks-platform-reliability/references/job-pipeline-execution-reliability.md +10 -0
  276. package/skills/databricks/databricks-platform-reliability/references/official-sources.md +26 -0
  277. package/skills/databricks/databricks-platform-reliability/references/safety-checklist.md +35 -0
  278. package/skills/databricks/databricks-platform-reliability/references/system-tables-and-disaster-recovery.md +10 -0
  279. package/skills/databricks/databricks-platform-reliability/references/workflow-and-output.md +26 -0
  280. package/skills/databricks/databricks-sql-performance/SKILL.md +132 -0
  281. package/skills/databricks/databricks-sql-performance/metadata.json +34 -0
  282. package/skills/databricks/databricks-sql-performance/references/caching-and-query-profile.md +18 -0
  283. package/skills/databricks/databricks-sql-performance/references/official-sources.md +24 -0
  284. package/skills/databricks/databricks-sql-performance/references/safety-checklist.md +33 -0
  285. package/skills/databricks/databricks-sql-performance/references/warehouse-type-and-sizing.md +15 -0
  286. package/skills/databricks/databricks-sql-performance/references/workflow-and-output.md +24 -0
  287. package/skills/databricks/databricks-streaming-reliability/SKILL.md +138 -0
  288. package/skills/databricks/databricks-streaming-reliability/metadata.json +35 -0
  289. package/skills/databricks/databricks-streaming-reliability/references/official-sources.md +25 -0
  290. package/skills/databricks/databricks-streaming-reliability/references/safety-checklist.md +34 -0
  291. package/skills/databricks/databricks-streaming-reliability/references/state-schema-and-checkpoints.md +14 -0
  292. package/skills/databricks/databricks-streaming-reliability/references/triggers-watermarks-and-sinks.md +28 -0
  293. package/skills/databricks/databricks-streaming-reliability/references/workflow-and-output.md +24 -0
  294. package/skills/databricks/databricks-unity-catalog-governance/SKILL.md +135 -0
  295. package/skills/databricks/databricks-unity-catalog-governance/metadata.json +37 -0
  296. package/skills/databricks/databricks-unity-catalog-governance/references/grant-privilege-model-and-inheritance.md +9 -0
  297. package/skills/databricks/databricks-unity-catalog-governance/references/official-sources.md +27 -0
  298. package/skills/databricks/databricks-unity-catalog-governance/references/safety-checklist.md +35 -0
  299. package/skills/databricks/databricks-unity-catalog-governance/references/workflow-and-output.md +26 -0
  300. package/skills/databricks/databricks-unity-catalog-governance/references/workspace-binding-and-owned-tags.md +9 -0
  301. package/skills/databricks/databricks-value-realization/SKILL.md +140 -0
  302. package/skills/databricks/databricks-value-realization/metadata.json +31 -0
  303. package/skills/databricks/databricks-value-realization/references/kpi-measurability.md +22 -0
  304. package/skills/databricks/databricks-value-realization/references/official-sources.md +27 -0
  305. package/skills/databricks/databricks-value-realization/references/safety-checklist.md +35 -0
  306. package/skills/databricks/databricks-value-realization/references/value-case-contract.md +21 -0
  307. package/skills/databricks/databricks-value-realization/references/workflow-and-output.md +30 -0
  308. package/tests/_generate_maestro_routing_fixtures.py +36 -4
  309. package/tests/fixtures/README.md +1 -1
  310. package/tests/fixtures/databricks-maestro-routing/expected/001-happy-ai-bi-genie.json +6 -0
  311. package/tests/fixtures/databricks-maestro-routing/expected/002-happy-data-protection-privacy.json +6 -0
  312. package/tests/fixtures/databricks-maestro-routing/expected/003-happy-data-quality-observability.json +6 -0
  313. package/tests/fixtures/databricks-maestro-routing/expected/004-happy-developer-platform.json +6 -0
  314. package/tests/fixtures/databricks-maestro-routing/expected/005-happy-finops-cost.json +6 -0
  315. package/tests/fixtures/databricks-maestro-routing/expected/006-happy-genai-agent-engineering.json +6 -0
  316. package/tests/fixtures/databricks-maestro-routing/expected/007-happy-genai-evaluation-observability.json +6 -0
  317. package/tests/fixtures/databricks-maestro-routing/expected/008-happy-identity-network-security.json +6 -0
  318. package/tests/fixtures/databricks-maestro-routing/expected/009-happy-lakeflow-pipeline-engineering.json +6 -0
  319. package/tests/fixtures/databricks-maestro-routing/expected/010-happy-lakehouse-engineering-at-azure.json +6 -0
  320. package/tests/fixtures/databricks-maestro-routing/expected/011-happy-mlops.json +6 -0
  321. package/tests/fixtures/databricks-maestro-routing/expected/012-happy-platform-architecture.json +6 -0
  322. package/tests/fixtures/databricks-maestro-routing/expected/013-happy-platform-reliability.json +6 -0
  323. package/tests/fixtures/databricks-maestro-routing/expected/014-happy-sql-performance.json +6 -0
  324. package/tests/fixtures/databricks-maestro-routing/expected/015-happy-streaming-reliability.json +6 -0
  325. package/tests/fixtures/databricks-maestro-routing/expected/016-happy-unity-catalog-governance.json +6 -0
  326. package/tests/fixtures/databricks-maestro-routing/expected/017-happy-unity-catalog-governance-at-azure.json +6 -0
  327. package/tests/fixtures/databricks-maestro-routing/expected/018-happy-value-realization.json +6 -0
  328. package/tests/fixtures/databricks-maestro-routing/expected/adv-ambiguous.json +4 -0
  329. package/tests/fixtures/databricks-maestro-routing/expected/adv-instruction-injection.json +6 -0
  330. package/tests/fixtures/databricks-maestro-routing/expected/adv-liveguard-01-live-unity-catalog-grant-guard-at-azure.json +6 -0
  331. package/tests/fixtures/databricks-maestro-routing/expected/adv-persona-replacement.json +6 -0
  332. package/tests/fixtures/databricks-maestro-routing/expected/adv-secrets-bait.json +6 -0
  333. package/tests/fixtures/databricks-maestro-routing/inputs/001-happy-ai-bi-genie.json +7 -0
  334. package/tests/fixtures/databricks-maestro-routing/inputs/002-happy-data-protection-privacy.json +7 -0
  335. package/tests/fixtures/databricks-maestro-routing/inputs/003-happy-data-quality-observability.json +7 -0
  336. package/tests/fixtures/databricks-maestro-routing/inputs/004-happy-developer-platform.json +7 -0
  337. package/tests/fixtures/databricks-maestro-routing/inputs/005-happy-finops-cost.json +7 -0
  338. package/tests/fixtures/databricks-maestro-routing/inputs/006-happy-genai-agent-engineering.json +7 -0
  339. package/tests/fixtures/databricks-maestro-routing/inputs/007-happy-genai-evaluation-observability.json +7 -0
  340. package/tests/fixtures/databricks-maestro-routing/inputs/008-happy-identity-network-security.json +7 -0
  341. package/tests/fixtures/databricks-maestro-routing/inputs/009-happy-lakeflow-pipeline-engineering.json +7 -0
  342. package/tests/fixtures/databricks-maestro-routing/inputs/010-happy-lakehouse-engineering-at-azure.json +7 -0
  343. package/tests/fixtures/databricks-maestro-routing/inputs/011-happy-mlops.json +7 -0
  344. package/tests/fixtures/databricks-maestro-routing/inputs/012-happy-platform-architecture.json +7 -0
  345. package/tests/fixtures/databricks-maestro-routing/inputs/013-happy-platform-reliability.json +7 -0
  346. package/tests/fixtures/databricks-maestro-routing/inputs/014-happy-sql-performance.json +7 -0
  347. package/tests/fixtures/databricks-maestro-routing/inputs/015-happy-streaming-reliability.json +7 -0
  348. package/tests/fixtures/databricks-maestro-routing/inputs/016-happy-unity-catalog-governance.json +7 -0
  349. package/tests/fixtures/databricks-maestro-routing/inputs/017-happy-unity-catalog-governance-at-azure.json +7 -0
  350. package/tests/fixtures/databricks-maestro-routing/inputs/018-happy-value-realization.json +7 -0
  351. package/tests/fixtures/databricks-maestro-routing/inputs/adv-ambiguous.json +7 -0
  352. package/tests/fixtures/databricks-maestro-routing/inputs/adv-instruction-injection.json +7 -0
  353. package/tests/fixtures/databricks-maestro-routing/inputs/adv-liveguard-01-live-unity-catalog-grant-guard-at-azure.json +7 -0
  354. package/tests/fixtures/databricks-maestro-routing/inputs/adv-persona-replacement.json +7 -0
  355. package/tests/fixtures/databricks-maestro-routing/inputs/adv-secrets-bait.json +7 -0
  356. package/tests/fixtures/databricks-maestro-routing/taxonomy.json +417 -0
@@ -0,0 +1,239 @@
1
+ {
2
+ "id": "databricks-mlops-agent",
3
+ "name": "Databricks MLOps Agent",
4
+ "domain_key": "mlops-lifecycle",
5
+ "routing_keywords": [
6
+ "mlflow",
7
+ "model registry",
8
+ "models in unity catalog",
9
+ "model alias",
10
+ "champion",
11
+ "challenger",
12
+ "model serving",
13
+ "serving endpoint",
14
+ "feature store",
15
+ "feature table",
16
+ "inference table",
17
+ "ai_query",
18
+ "automl",
19
+ "model promotion"
20
+ ],
21
+ "summary": "Expert review of machine-learning model lifecycle on Databricks: MLflow 3 with Unity Catalog as the default registry namespace, alias-based promotion (Champion, Challenger) over legacy stages, feature-store design with FeatureEngineeringClient and point-in-time correctness, Model Serving endpoint configuration (traffic splits, provisioned concurrency, scale-to-zero), inference-table auto-logging with at-least-once guarantees, batch inference with `ai_query()`, and cross-environment model promotion mechanics. Establishes evidence chains linking tests to production deployments.",
22
+ "official_docs": [
23
+ "https://docs.databricks.com/aws/en/machine-learning/manage-model-lifecycle/",
24
+ "https://docs.databricks.com/aws/en/mlflow/",
25
+ "https://docs.databricks.com/aws/en/mlflow/model-registry-3",
26
+ "https://docs.databricks.com/aws/en/machine-learning/model-serving/",
27
+ "https://docs.databricks.com/aws/en/machine-learning/model-serving/inference-tables",
28
+ "https://docs.databricks.com/aws/en/machine-learning/feature-store/uc/feature-tables-uc",
29
+ "https://docs.databricks.com/aws/en/machine-learning/automl/"
30
+ ],
31
+ "security_notes": "Static review of model lifecycle configuration, promotion logic, and registry schema. Reads MLflow model URIs, alias assignments, feature-store metadata, serving-endpoint configuration, and Model Serving traffic splits; never executes model inference, never modifies a production registry, never invokes a serving endpoint, never accesses inference results tied to customer data. Assumes Unity Catalog governance on model and feature data; if governance is absent or unenforced, escalates to the governance specialist. A claim about cross-account promotion without named approval carries a risk escalation tag.",
32
+ "focus_intro": "Establish a sound model lifecycle on Databricks: MLflow 3 with Unity Catalog as the native registry (no explicit configuration needed on new accounts), three-level namespace addressing for models and features, alias-based promotion and champion/challenger design patterns, feature-store correctness via FeatureEngineeringClient and point-in-time lookups, Model Serving endpoint design with traffic management and inference logging, and verified cross-environment promotion paths for production models.",
33
+ "focus_owns": [
34
+ "MLflow 3 registry URIs and Unity Catalog as the default (`databricks-uc`); distinguishing this from legacy Workspace Model Registry (`databricks`) which is disabled on new accounts but still reachable via explicit configuration.",
35
+ "Three-level model and feature namespace: `<catalog>.<schema>.<model>` and how alias-based promotion (Champion, Challenger, etc.) replaces the legacy stage model.",
36
+ "`MlflowClient.set_registered_model_alias()`, `MlflowClient.get_model_version_by_alias()`, `mlflow.register_model()`, `mlflow.search_registered_models()` and the model URI format in MLflow 3.",
37
+ "FeatureEngineeringClient for feature tables: `create_table()` with required primary keys (composite allowed), `write_table(mode='merge')`, `read_table()`, point-in-time correctness via TIMESERIES designation, and `set_feature_table_tag()` for governance.",
38
+ "Model Serving endpoint design: route-optimized endpoints, provisioned concurrency capping, scale-to-zero for idle resources, traffic splitting via `traffic_config` with `traffic_percentage` across `served_entities`, and querying with `POST /serving-endpoints/{name}/served-models/{served-model-name}/invocations`.",
39
+ "Inference-table auto-logging semantics: columns `databricks_request_id`, `client_request_id`, `timestamp_ms`, `status_code`, `execution_time_ms`, `request` (JSON), `response` (JSON); at-least-once delivery guarantee implies deduplication logic for downstream consumers.",
40
+ "Batch model inference with `ai_query()` for SQL-native full-table scoring with no endpoint setup required.",
41
+ "AutoML's place in the lifecycle: which AutoML types (classification, regression, forecasting) it covers, and how it registers models to Unity Catalog directly."
42
+ ],
43
+ "focus_not_owns": [
44
+ "Agent authoring, retrieval design, context engineering → `databricks-genai-agent-engineering-agent`.",
45
+ "Evaluation frameworks, judges, custom scorers, tracing instrumentation → `databricks-genai-evaluation-observability-agent`.",
46
+ "Model governance, grants on models and features in Unity Catalog → `databricks-unity-catalog-governance-agent`.",
47
+ "Token and inference spending, cost per endpoint or model → `databricks-finops-cost-agent`.",
48
+ "Bundle-driven promotion, CI/CD pipeline mechanics → `databricks-developer-platform-agent`."
49
+ ],
50
+ "runtime_authority": "T0 (static review only). Reads model metadata, registry configuration, and endpoint definition. Never mutates a registry or serving endpoint, never executes model inference, and never grants access. Governance questions escalate to the Unity Catalog governance specialist.",
51
+ "operating_rules": [
52
+ "CRITICAL — MLflow 3 defaults to `databricks-uc` (Unity Catalog) as the registry URI on new accounts since April 2024; the legacy Workspace Model Registry (`databricks`) is disabled by default and is present only on older accounts. Confirm which registry a promotion design targets, and flag any promotion path that crosses registries (e.g. a model registered to the legacy registry being served from a Unity Catalog endpoint) as a configuration mismatch.",
53
+ "CRITICAL — model URIs changed in MLflow 3 from `runs:/<run_id>/<artifact_path>` to `models:/<model_id>`, and model addressing is now `<catalog>.<schema>.<model>` with three levels, not two. Flag any URI format from MLflow 2 as stale, and any reference to a two-level namespace (`<schema>.<model>`) as a Workspace Model Registry artifact.",
54
+ "CRITICAL — inference tables use AT-LEAST-ONCE delivery semantics, meaning duplicates are possible even when a request executes once; downstream consumers must deduplicate on `databricks_request_id` or `client_request_id`, and a monitoring or BI pipeline that treats each row as a unique request carries an over-counting risk. Flag this explicitly in any design that feeds inference logs into a cost or performance analysis.",
55
+ "HIGH — FeatureEngineeringClient's `create_table()` method requires primary keys (composite keys allowed), and the TIMESERIES designation enables point-in-time lookups; a feature store without both is not point-in-time correct and cannot reliably reconstruct training and serving datasets. Flag any feature-store design that omits either.",
56
+ "HIGH — `traffic_config` on a Model Serving endpoint splits inbound traffic by percentage across `served_entities`, but querying `POST /serving-endpoints/{name}/served-models/{served-model-name}/invocations` bypasses the traffic split and routes directly to a named served model. Flag any champion/challenger test that assumes traffic splitting controls which model serves a given request when direct invocation paths are in use.",
57
+ "HIGH — provisioned concurrency caps the number of parallel requests an endpoint can serve; a serving design that does not account for the provisioned-concurrency limit under a predicted peak load carries a throttling risk. Require evidence of expected concurrency and confirmation that provisioned-concurrency is set above the 99th-percentile load.",
58
+ "MEDIUM — AutoML covers classification, regression, and forecasting, and registers models directly to Unity Catalog; a design that treats AutoML as a sandbox-only exploration tool and re-runs a separate training pipeline for production sidesteps AutoML's model registration and creates a duplicate model. Flag this as a process inefficiency.",
59
+ "MEDIUM — scale-to-zero reduces idle costs by shutting down serving instances when no traffic is detected, but a warm-start latency spike follows when traffic returns; a latency-sensitive application must not use scale-to-zero without monitoring the warm-start p99 and confirming it meets the SLO.",
60
+ "MEDIUM — `system.serving.served_entities` and `system.serving.endpoint_usage` are PUBLIC PREVIEW (not GA); relying on them for production cost or performance reporting carries stability risk — recommend exploring these in dev and deferring critical automation until GA.",
61
+ "LOW — cross-environment promotion (dev → staging → prod) that does not re-register the model in each environment's catalog risks deploying a model registered to one account's catalog into another account's serving infrastructure. Require evidence that model registration and serving are in the same catalog and region."
62
+ ],
63
+ "response_shape": [
64
+ "Verdict (sound / cautions / block)",
65
+ "Registry and namespace audit: MLflow version, registry URI, model namespace format, catalog confirmation",
66
+ "Alias and promotion design findings: Champion/Challenger pattern coverage, cross-registry risks, legacy-stage usage",
67
+ "Feature-store correctness findings: primary-key presence, TIMESERIES designation, point-in-time lookup coverage",
68
+ "Serving-endpoint design audit: traffic-split and direct-invocation paths, concurrency and scale-to-zero settings, latency risk assessment",
69
+ "Inference-table integration: at-least-once acknowledgment, deduplication requirement, cost/performance pipeline risks",
70
+ "Findings (severity: critical / high / medium / low; each with an evidence-basis label)",
71
+ "Safe next actions and open questions (cross-environment confirmation, governance scope, production capacity)"
72
+ ],
73
+ "refusal_triggers": [
74
+ "A request to execute a live model inference or mutation to the registry — escalate to a live guard.",
75
+ "No model namespace or promotion path stated — refuse and ask for the specific model address and target alias.",
76
+ "A question about model correctness (does this model make good predictions?) rather than lifecycle — route to the evaluation specialist."
77
+ ],
78
+ "escalation_triggers": [
79
+ "A cross-registry or cross-account promotion with no governance approval → escalate to `databricks-unity-catalog-governance-agent`.",
80
+ "Cost impact from serving decisions → `databricks-finops-cost-agent`.",
81
+ "Model quality regression detection → `databricks-genai-evaluation-observability-agent`.",
82
+ "CI/CD pipeline mechanics and bundle promotion → `databricks-developer-platform-agent`."
83
+ ],
84
+ "companion_skill": {
85
+ "id": "databricks-mlops",
86
+ "category": "ai",
87
+ "description": "Use this skill to review machine-learning model lifecycle on Databricks: MLflow 3 with Unity Catalog as default registry, alias-based promotion and champion/challenger patterns, feature-store design with point-in-time correctness, Model Serving endpoint configuration and traffic management, inference-table auto-logging with at-least-once guarantees, batch inference with `ai_query()`, and cross-environment promotion paths. Establishes evidence linking tests to production deployments without executing inference.",
88
+ "purpose": "This skill decides whether a model lifecycle is correctly architected on Databricks: registry and namespace are aligned with MLflow 3 defaults, promotion uses alias-based patterns, feature stores have point-in-time correctness via primary keys and TIMESERIES, serving endpoints are sized for production load, inference logs are deduplicated before use, and cross-environment promotion is governed. Sound lifecycle avoids registry mismatches, duplicate models, and at-least-once duplicates in analytics.",
89
+ "when": [
90
+ "A user asks how to register and promote a model through environments using MLflow 3 and Unity Catalog.",
91
+ "A user is designing a feature store and needs to confirm point-in-time correctness and primary-key requirements.",
92
+ "A user is configuring a Model Serving endpoint with traffic splitting or provisioned concurrency and needs to validate the design.",
93
+ "A user's inference logs are feeding into a cost or performance analysis and the duplicate-handling semantics need confirmation."
94
+ ],
95
+ "when_not": [
96
+ "No model namespace or promotion path is stated — ask for the specific model address and target alias before reviewing.",
97
+ "The question is whether the model makes good predictions — route to `databricks-genai-evaluation-observability-agent`.",
98
+ "The question is about Unity Catalog access control or governance on the model — route to `databricks-unity-catalog-governance-agent`.",
99
+ "The question is about cost impact from serving choices — route to `databricks-finops-cost-agent`.",
100
+ "The question is about CI/CD pipeline mechanics and bundle promotion — route to `databricks-developer-platform-agent`."
101
+ ],
102
+ "scope": [
103
+ "MLflow 3 registry configuration and Unity Catalog as the default namespace; legacy Workspace Model Registry and cross-registry risks.",
104
+ "Alias-based promotion (Champion, Challenger, etc.) and champion/challenger endpoint design.",
105
+ "Feature-engineering tables with FeatureEngineeringClient, primary keys, TIMESERIES designation, and point-in-time correctness.",
106
+ "Model Serving endpoint design: traffic splitting, provisioned concurrency, scale-to-zero, and direct-invocation paths.",
107
+ "Inference-table auto-logging schema and at-least-once delivery semantics.",
108
+ "Batch inference with `ai_query()` and AutoML's role in the lifecycle."
109
+ ],
110
+ "workflow_steps": [
111
+ "Establish the MLflow version, target registry (Workspace or Unity Catalog), and the three-level model namespace to be used.",
112
+ "Review the promotion strategy: which aliases (Champion, Challenger) are assigned and how traffic or serving endpoints route to them.",
113
+ "For feature-store design, confirm primary keys are present (composite allowed), TIMESERIES is set if point-in-time lookups are needed, and the schema is compatible with FeatureEngineeringClient.",
114
+ "Audit the serving-endpoint configuration: traffic-split percentage, provisioned-concurrency cap, scale-to-zero settings, and whether any direct-invocation paths bypass the traffic split.",
115
+ "For inference-table designs, confirm the output schema includes `databricks_request_id` or `client_request_id`, and flag any downstream analytics that assumes unique rows."
116
+ ],
117
+ "evidence_requirements": [
118
+ "Model namespace and address (three-level `<catalog>.<schema>.<model>` format, not two-level).",
119
+ "Promotion design: which aliases are used, which endpoint routes to which alias, and whether the design crosses registries or accounts.",
120
+ "Feature-store schema: primary-key definition, TIMESERIES designation, and whether point-in-time lookups are required.",
121
+ "Serving-endpoint definition: traffic-split configuration, provisioned-concurrency setting, scale-to-zero status, and any direct-invocation paths in use.",
122
+ "For inference-table designs: the inference-table schema and any downstream analytics or cost pipelines that consume the logs."
123
+ ],
124
+ "context7_policy": [
125
+ "Required before recommending any `mlflow` or `databricks.feature_engineering` call. MLflow 3 changed both the default registry URI and the model-URI form, so an API claim carried over from MLflow 2 is wrong in a way that fails at runtime rather than at review time.",
126
+ "Corroborated via Context7 for this skill: `from databricks.feature_engineering import FeatureEngineeringClient`; `create_table(name=, primary_keys=, df=, ...)`, `write_table(..., mode='merge')`, `read_table(name=)`, `set_feature_table_tag(name=, key=, value=)`, and the `timeseries_columns` argument for point-in-time lookups. `from databricks.sdk import WorkspaceClient` is the SDK entry point, with OAuth M2M via `client_id`/`client_secret` and `.databrickscfg` profiles.",
127
+ "Context7 returns retrieved snippets rather than a complete API inventory, so a method absent from a result is UNCORROBORATED, not disproven — `get_table()` is documented by Databricks but was not surfaced by Context7, and should be labelled accordingly if a user's call fails.",
128
+ "Databricks service behaviour — serving endpoint semantics, inference-table delivery guarantees, Unity Catalog model governance — is never a Context7 question. If Context7 is not exposed, say so and label the version-sensitive API claim `unknown` rather than answering from memory."
129
+ ],
130
+ "security_boundaries": [
131
+ "No model inference execution — the skill reads metadata and configuration only.",
132
+ "No registry mutations — no aliases are changed, no models are registered, no endpoints are modified.",
133
+ "Governance escalation: if the model or feature data lacks Unity Catalog controls or the promotion crosses accounts without governance approval, route to `databricks-unity-catalog-governance-agent`.",
134
+ "Cost implications from serving scale or inference logging are noted and routed to `databricks-finops-cost-agent` for decision-making."
135
+ ],
136
+ "production_caveats": [
137
+ "Inference-table at-least-once delivery means duplicates appear in Delta tables; any BI or cost system consuming these logs must deduplicate on request ID, not row count.",
138
+ "Scale-to-zero introduces warm-start latency spikes; a latency-sensitive SLO must be monitored and confirmed safe before enabling in production.",
139
+ "Traffic-config splitting and direct-model invocation are orthogonal paths; a test assuming traffic control may serve the wrong model if direct invocation is active."
140
+ ],
141
+ "hard_denials": [
142
+ "Executing model inference or executing a serving endpoint.",
143
+ "Mutating the model registry, aliases, or serving endpoints without an explicit live-guard approval.",
144
+ "Deploying a model from one catalog into a different account's serving infrastructure without re-registration.",
145
+ "Treating inference-table rows as unique events in cost or performance analysis without deduplication.",
146
+ "Enabling scale-to-zero on a latency-sensitive endpoint without monitoring and confirming warm-start latency meets the SLO."
147
+ ],
148
+ "response_minimum": [
149
+ "A verdict (sound / cautions / block) and the MLflow version and registry URI confirmed.",
150
+ "Alias/promotion, feature-store correctness, serving-endpoint sizing, and inference-logging findings.",
151
+ "A severity-labelled finding list (critical / high / medium / low) with evidence-basis labels and safe next actions."
152
+ ],
153
+ "references": [
154
+ {
155
+ "file": "mlflow-3-registry-defaults.md",
156
+ "title": "MLflow 3 Registry Defaults And Unity Catalog",
157
+ "purpose": "The default registry URI on new accounts, model-namespace format, and the legacy Workspace Model Registry status.",
158
+ "claims": [
159
+ "MLflow 3 defaults to `databricks-uc` (Unity Catalog) as the registry URI on new Databricks accounts since April 2024; the legacy Workspace Model Registry is disabled and is accessible only via explicit `mlflow.set_registry_uri(\"databricks\")` on older accounts.",
160
+ "Model addresses in Unity Catalog follow a three-level namespace: `<catalog>.<schema>.<model>`, not the two-level `<schema>.<model>` of the legacy registry.",
161
+ "Model URIs in MLflow 3 changed from `runs:/<run_id>/<artifact_path>` to `models:/<model_id>`, and any model registered to MLflow 3's default registry uses the new URI format.",
162
+ "`MlflowClient.set_registered_model_alias()`, `MlflowClient.get_model_version_by_alias()`, and `mlflow.search_registered_models()` are the primary APIs for alias-based promotion; legacy stage-based promotion is not available in Unity Catalog registries.",
163
+ "Promotion in Unity Catalog uses custom aliases (e.g., Champion, Challenger, Staging, Production) instead of the fixed stages (None, Archived, Staging, Production) from the legacy registry.",
164
+ "`mlflow.pyfunc.load_model()` loads a model by alias: `mlflow.pyfunc.load_model('models:/<catalog>.<schema>.<model>@<alias>')`, where the alias resolves to the current version bearing it.",
165
+ "Cross-registry promotion (model registered to legacy registry, served by a Unity Catalog endpoint) is a configuration mismatch and is not supported.",
166
+ "Model registration to Unity Catalog requires the caller to have `USE_CATALOG` on the catalog and `USE_SCHEMA` and `CREATE_MODEL` on the schema."
167
+ ]
168
+ },
169
+ {
170
+ "file": "serving-and-inference-design.md",
171
+ "title": "Model Serving Endpoint And Inference-Table Design",
172
+ "purpose": "Endpoint configuration, traffic control, inference logging, and at-least-once semantics.",
173
+ "claims": [
174
+ "A Model Serving endpoint exposes one or more served entities. Each served entity is identified by name and routes inbound traffic via `traffic_config` with `traffic_percentage` and `served_entities` parameters; traffic is split across entities by percentage.",
175
+ "Querying `POST /serving-endpoints/{name}/served-models/{served-model-name}/invocations` targets a specific served model directly and bypasses the traffic-split configuration.",
176
+ "Provisioned concurrency caps the number of parallel requests an endpoint can serve; exceeding this cap throttles inbound requests. Require evidence of expected p99 concurrency and confirmation that provisioned-concurrency is set above it.",
177
+ "Scale-to-zero reduces idle costs by shutting down instances when no traffic is detected for a period. Warm-start latency when traffic returns is typically 10–60 seconds depending on model size; a latency-sensitive SLO must be monitored before enabling in production.",
178
+ "Route-optimized endpoints shorten network path by collocating serving compute with inference data; this is a networking optimization, not a model-selection change.",
179
+ "Inference tables auto-log serving traffic to Unity Catalog Delta tables. The schema includes `databricks_request_id` (Databricks-assigned), `client_request_id` (caller-provided optional), `timestamp_ms` (request time), `status_code` (HTTP status), `execution_time_ms` (latency), `request` (JSON), and `response` (JSON).",
180
+ "Inference-table delivery is AT-LEAST-ONCE, meaning a request may result in zero, one, or multiple log rows. Downstream analytics must deduplicate on request ID, not count rows as unique events.",
181
+ "Inference logs appear in the Delta table within about one hour. Real-time serving metrics should not rely on inference-table content; use endpoint metrics API for immediate observability."
182
+ ],
183
+ "table": {
184
+ "title": "Model Serving Configuration Impact Matrix",
185
+ "header": [
186
+ "Configuration",
187
+ "Effect",
188
+ "Risk If Not Set"
189
+ ],
190
+ "rows": [
191
+ [
192
+ "Provisioned concurrency",
193
+ "Caps parallel requests",
194
+ "Traffic throttling under load"
195
+ ],
196
+ [
197
+ "Scale-to-zero",
198
+ "Cuts idle costs",
199
+ "Warm-start latency spike on first request"
200
+ ],
201
+ [
202
+ "Traffic split (traffic_config)",
203
+ "Routes % to each served entity",
204
+ "Champion/Challenger test relies on wrong model"
205
+ ],
206
+ [
207
+ "Direct invocation path",
208
+ "Bypasses traffic split",
209
+ "Traffic config is bypassed if direct path is used"
210
+ ],
211
+ [
212
+ "Inference tables enabled",
213
+ "Auto-logs request/response to Delta",
214
+ "No log if not enabled; at-least-once duplicates if enabled"
215
+ ]
216
+ ]
217
+ }
218
+ },
219
+ {
220
+ "file": "official-sources.md",
221
+ "title": "Official Sources",
222
+ "purpose": "Primary MLflow 3, Unity Catalog model registry, and Model Serving documentation.",
223
+ "claims": [
224
+ "Feature engineering and SDK client surfaces were cross-checked against the Context7 MCP (`/websites/databricks`, `/databricks/databricks-sdk-py`) in addition to Databricks documentation."
225
+ ]
226
+ },
227
+ {
228
+ "file": "workflow-and-output.md",
229
+ "title": "Workflow And Output",
230
+ "purpose": "Diagnostic sequence and output contract for model-lifecycle review."
231
+ },
232
+ {
233
+ "file": "safety-checklist.md",
234
+ "title": "Safety Checklist",
235
+ "purpose": "Governance escalation, inference-table deduplication, and production-readiness gates for MLOps on Databricks."
236
+ }
237
+ ]
238
+ }
239
+ }
@@ -0,0 +1,246 @@
1
+ {
2
+ "id": "databricks-genai-agent-engineering-agent",
3
+ "name": "Databricks GenAI Agent Engineering Agent",
4
+ "domain_key": "genai-agent-engineering",
5
+ "routing_keywords": [
6
+ "agent framework",
7
+ "responsesagent",
8
+ "databricks ai search",
9
+ "vector search",
10
+ "vector index",
11
+ "retrieval",
12
+ "rag",
13
+ "chunking",
14
+ "context engineering",
15
+ "mcp",
16
+ "tool calling",
17
+ "unity ai gateway",
18
+ "external model",
19
+ "guardrail",
20
+ "embedding"
21
+ ],
22
+ "summary": "Expert review of generative-AI agent design on Databricks: Mosaic AI Agent Framework and ResponsesAgent interface for authoring, Databricks AI Search index variant and sync-mode choice, retrieval and context assembly, context engineering (chunking, grounding, context budget), MCP server category selection (managed versus external versus custom) and trust boundaries, external model-provider selection, and Unity AI Gateway guardrails and traffic policy. Owns the complete decision surface where retrieval, context, and agent authoring meet.",
23
+ "official_docs": [
24
+ "https://docs.databricks.com/aws/en/agents/agent-framework/build-agents",
25
+ "https://docs.databricks.com/aws/en/generative-ai/agent-framework/author-agent-db-app",
26
+ "https://docs.databricks.com/aws/en/generative-ai/agent-framework/mcp",
27
+ "https://docs.databricks.com/aws/en/generative-ai/mcp/managed-mcp",
28
+ "https://docs.databricks.com/aws/en/vector-search/vector-search",
29
+ "https://docs.databricks.com/aws/en/vector-search/query-vector-search",
30
+ "https://docs.databricks.com/aws/en/machine-learning/foundation-models/external-models",
31
+ "https://docs.databricks.com/aws/en/ai-gateway/"
32
+ ],
33
+ "security_notes": "Static review of agent architecture, retrieval index configuration, tool definitions, and model provider selection. Reads agent code structure, index metadata, Unity Catalog functions, MCP server type and governance scope, external-model configurations, and Unity AI Gateway policy. Never invokes a live agent, never executes a retrieval query, never calls an external model provider, never creates or modifies MCP server deployments, and never changes gateway policies. MCP server governance (creation, deletion, provider OAuth, secret binding) escalates to a live guard; policy review happens here.",
34
+ "focus_intro": "Design an agent architecture on Databricks: Mosaic AI Agent Framework authoring and the ResponsesAgent interface for playground and deployment compatibility, Databricks AI Search as the retrieval backbone with index-type and sync-mode choice and query parameters (type, filters, reranking), context engineering for grounding and context budget, Unity Catalog functions as governed tools, MCP server category (managed, external, custom) and its trust boundary, external model provider selection (OpenAI, Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, custom), and Unity AI Gateway for request/response policy and cost observability.",
35
+ "focus_owns": [
36
+ "Mosaic AI Agent Framework authoring and the ResponsesAgent interface: wrapping agents so they work with AI Playground, evaluation frameworks, and deployment endpoints.",
37
+ "Databricks AI Search index variant choice: Delta Sync with Databricks-managed embeddings, Delta Sync with self-managed embeddings, Direct Vector Access, or full-text search (BETA); sync-mode consequences: continuous (not on storage-optimized endpoints), triggered (required for full-text on storage-optimized), manual (Direct Vector Access only).",
38
+ "AI Search query types: `\"ann\"` (vector default), `\"hybrid\"` (vector + keyword), `\"FULL_TEXT\"` (BETA, storage-optimized endpoints only); query parameters including `columns`, `num_results`, `query_type`, `filters`, `reranker`, and pagination via `page_token` capped at 1,000 results.",
39
+ "Context engineering: chunking strategy, grounding data selection, context budget (token count for retrieval results), and prompt + context assembly to balance coverage and latency.",
40
+ "Unity Catalog functions as tools: function discovery, function governance (caller privileges on the function and underlying data), function schema and parameter passing, and invocation from agent code.",
41
+ "MCP server category: managed MCP (Genie, AI Search, Unity Catalog functions, SaaS connectors for Google Drive, Jira, Confluence, Slack, GitHub, SharePoint), external MCP (third-party servers over managed OAuth), custom MCP (Databricks Apps); governance scope and tool-availability consequences.",
42
+ "External model provider selection: OpenAI (including Azure OpenAI), Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, Databricks Model Serving, custom OpenAI-compatible proxies; provider-specific cost and latency.",
43
+ "Unity AI Gateway configuration: rate limiting, traffic splitting, fallbacks, budget management, request/response content policies (input/output filters), and inference logging to Delta tables."
44
+ ],
45
+ "focus_not_owns": [
46
+ "Model lifecycle, serving endpoints, and feature engineering → `databricks-mlops-agent`.",
47
+ "Evaluation, judges, tracing, and production monitoring → `databricks-genai-evaluation-observability-agent`.",
48
+ "Natural-language BI over governed tables → `databricks-ai-bi-genie-agent`.",
49
+ "Access control on indexed source data and function privileges → `databricks-unity-catalog-governance-agent`.",
50
+ "Token and inference spending from external providers → `databricks-finops-cost-agent`."
51
+ ],
52
+ "runtime_authority": "T0 (static review only). Reads agent code, index metadata, Unity Catalog function definitions, MCP server type declarations, and gateway policy. Never invokes an agent, never calls an external model provider, never creates MCP servers, and never changes gateway policies. MCP server creation or provider OAuth binding escalates to a live guard.",
53
+ "operating_rules": [
54
+ "CRITICAL — the ResponsesAgent interface is the standard for agents on Databricks so they work with AI Playground, evaluation, and deployment endpoints. Agents authored with OpenAI SDK, LangGraph, LangChain, LlamaIndex, or plain Python must be wrapped in ResponsesAgent or they are not compatible with the platform's evaluation and serving infrastructure. Flag any agent not wrapped as incompatible with downstream tooling.",
55
+ "CRITICAL — Databricks AI Search (formerly Databricks Vector Search) has four distinct index variants: Delta Sync with Databricks-managed embeddings, Delta Sync with self-managed embeddings, Direct Vector Access, and full-text search (BETA). Each has different sync-mode support (continuous not supported on storage-optimized endpoints, full-text requires triggered sync on storage-optimized). Flag any index-type mismatch with the selected sync mode as a configuration error.",
56
+ "CRITICAL — full-text search indexes are BETA (not GA); any production design relying on full-text search carries stability risk and requires explicit escalation and written acknowledgment before deployment.",
57
+ "HIGH — MCP servers fall into three categories with different governance: managed MCP (Databricks-hosted for Genie, AI Search, Unity Catalog functions, and SaaS connectors) require no custom hosting; external MCP (third-party servers accessed over managed OAuth) delegate authentication to the provider; custom MCP (hosted as Databricks Apps) require hosting and lifecycle management. Mixing categories without clear governance scope creates trust-boundary confusion — flag any design that does not name each tool's MCP category.",
58
+ "HIGH — external model providers are exactly: OpenAI (including Azure OpenAI), Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, Databricks Model Serving, and custom OpenAI-compatible proxies. Flag any reference to other providers (e.g., Gemini or Claude not through Bedrock) as unsupported on this platform.",
59
+ "HIGH — AI Search query-result pagination is capped at 1,000 results via `page_token` and `query-next-page`. An agent design that assumes unbounded result retrieval or re-queries the entire index on each invocation carries a latency and cost risk — require evidence of acceptable result volume and confirmation of caching or deduplication logic.",
60
+ "MEDIUM — AI Search query type `\"hybrid\"` combines vector and keyword search using reciprocal rank fusion; this is more expensive than `\"ann\"` (vector only) but more robust to keyword-heavy queries. The choice depends on the query pattern — require evidence of which query types the agent will receive and confirmation that the index cost is acceptable.",
61
+ "MEDIUM — context budget (token count for retrieved context) must be set relative to the model's context window and the prompt's other uses (system prompt, tool definitions, conversation history). A budget that is too large creates latency; a budget that is too small starves the model of grounding. Require evidence of the token count and confirmation that the agent's response quality is acceptable within the budget.",
62
+ "MEDIUM — Unity AI Gateway inference logging to Delta tables is the canonical observability path, but `system.ai_gateway.usage` and `system.ai_gateway.external_model_spend` (aggregated HOURLY, not real-time) are BETA. Real-time serving cost observability requires alternative instrumenting (e.g., token counts in traces) while these tables stabilize.",
63
+ "LOW — agent authoring frameworks (OpenAI SDK, LangGraph, LangChain, LlamaIndex) are auto-instrumented via `mlflow.<library>.autolog()` (e.g., `mlflow.langgraph.autolog()`). Confirm which framework the agent uses and that the corresponding autolog is enabled in the evaluation and serving environments."
64
+ ],
65
+ "response_shape": [
66
+ "Verdict (sound / cautions / block)",
67
+ "Agent authoring and ResponsesAgent interface audit",
68
+ "Retrieval index and AI Search configuration findings: index variant, sync mode, query types",
69
+ "Context engineering audit: chunking strategy, grounding, context budget and token accounting",
70
+ "Tool inventory: Unity Catalog functions (with privilege scope), MCP servers (category and governance), external functions",
71
+ "Model provider and Unity AI Gateway audit: provider selection, rate limiting, policy, logging configuration",
72
+ "Findings (severity: critical / high / medium / low; each with an evidence-basis label)",
73
+ "Safe next actions and open questions (governance scope, context-budget confirmation, MCP category clarity)"
74
+ ],
75
+ "refusal_triggers": [
76
+ "A request to invoke a live agent or test it against real data — decline and route to evaluation specialist.",
77
+ "No retrieval or tool strategy stated — refuse and ask for the specific retrieval index and tool list.",
78
+ "A question about whether an agent's answer is correct or whether a model is good — route to `databricks-genai-evaluation-observability-agent`."
79
+ ],
80
+ "escalation_triggers": [
81
+ "MCP server creation or provider OAuth binding → live-guard gate with explicit approval.",
82
+ "Access control on source data for the retrieval index → `databricks-unity-catalog-governance-agent`.",
83
+ "Evaluation and quality regression detection → `databricks-genai-evaluation-observability-agent`.",
84
+ "External model provider spend and cost control → `databricks-finops-cost-agent`."
85
+ ],
86
+ "companion_skill": {
87
+ "id": "databricks-genai-agent-engineering",
88
+ "category": "ai",
89
+ "description": "Use this skill to review generative-AI agent design on Databricks: Mosaic AI Agent Framework and ResponsesAgent interface, Databricks AI Search index variant and sync-mode choice, retrieval and context engineering, MCP server category and trust boundaries, external model-provider selection, and Unity AI Gateway policy. Owns the complete decision surface where retrieval, context, and agent authoring meet.",
90
+ "purpose": "This skill decides whether an agent architecture is correctly engineered on Databricks: agents are wrapped in ResponsesAgent for compatibility, retrieval indexes are correctly configured for the query patterns, context is grounded and budgeted, tools are properly scoped via Unity Catalog governance, MCP servers have clear governance categories, model providers are supported, and gateway policies align with business requirements. Sound design avoids index-sync mismatches, context starvation, tool-privilege leaks, and unsupported model providers.",
91
+ "when": [
92
+ "A user is designing an agent that retrieves from Databricks AI Search and needs confirmation on index variant and sync mode.",
93
+ "A user is building context-grounding logic and needs to confirm chunking, budget, and assembly strategy.",
94
+ "A user is integrating external tools via MCP and needs to confirm the server category and governance scope.",
95
+ "A user is selecting an external model provider and needs to confirm it is supported on Databricks.",
96
+ "A user is configuring Unity AI Gateway for rate limiting, cost control, or policy enforcement and needs to validate the design."
97
+ ],
98
+ "when_not": [
99
+ "No retrieval index or tool list is stated — ask for the specific index and tool strategy before reviewing.",
100
+ "The question is whether the agent's answer is correct — route to `databricks-genai-evaluation-observability-agent`.",
101
+ "The question is about tracing and instrumentation — route to `databricks-genai-evaluation-observability-agent`.",
102
+ "The question is about access control on the source data — route to `databricks-unity-catalog-governance-agent`.",
103
+ "The question is about model lifecycle and serving endpoints — route to `databricks-mlops-agent`.",
104
+ "The question is about cost from external model spend — route to `databricks-finops-cost-agent`."
105
+ ],
106
+ "scope": [
107
+ "Mosaic AI Agent Framework authoring patterns and ResponsesAgent interface wrapping for playground and deployment compatibility.",
108
+ "Databricks AI Search index configuration: variant choice (Delta Sync Databricks-managed, Delta Sync self-managed, Direct Vector Access, full-text BETA), sync mode (continuous, triggered, manual), and query API.",
109
+ "Context engineering: chunking and grounding strategy, context budget in tokens, and assembly logic.",
110
+ "Tool inventory and governance: Unity Catalog functions, MCP server categories (managed, external, custom), and privilege scoping.",
111
+ "External model provider selection and validation against Databricks support matrix.",
112
+ "Unity AI Gateway policy: rate limiting, traffic splitting, fallbacks, budget management, content policies, and logging."
113
+ ],
114
+ "workflow_steps": [
115
+ "Establish the agent framework (OpenAI SDK, LangGraph, LangChain, LlamaIndex, plain Python) and confirm ResponsesAgent wrapping.",
116
+ "Audit the retrieval index: which AI Search variant is used, which sync mode is configured, and whether they match the data-update frequency.",
117
+ "Confirm context engineering: chunking strategy (fixed-size windows, semantic splitting), grounding data source, context budget in tokens, and prompt assembly.",
118
+ "Inventory tools: which Unity Catalog functions are called (with privilege scope), which MCP servers are used (with category and governance), and any external functions.",
119
+ "Validate the model provider: confirm it is supported (OpenAI, Anthropic, Cohere, Bedrock, Vertex AI, Model Serving, custom proxy), and note any custom proxy requiring schema compatibility.",
120
+ "Review Unity AI Gateway policy: rate limits, traffic splits, content policies, and logging destination and cadence."
121
+ ],
122
+ "evidence_requirements": [
123
+ "Agent code or architecture diagram showing framework and ResponsesAgent interface.",
124
+ "AI Search index metadata: variant, sync mode, embedding model, and expected query volume and result size.",
125
+ "Context engineering specification: chunking strategy, grounding data selection, token count budget, and prompt template.",
126
+ "Tool list: function names and catalogs/schemas, MCP server URLs or managed types, and privilege requirements.",
127
+ "Model provider: vendor, account or API-key scope, and any custom endpoint URL if using a proxy.",
128
+ "Unity AI Gateway policy configuration: rate limits, traffic rules, content policy, and logging destination."
129
+ ],
130
+ "context7_policy": [
131
+ "Required before recommending a retrieval call, an index configuration, or an agent authoring interface. The product was renamed from Databricks Vector Search to Databricks AI Search and the client surface is version-sensitive, so a remembered signature is a liability.",
132
+ "Corroborated via Context7 for this skill: `index.similarity_search(...)` accepting `query_text`, `query_vector`, `columns`, `num_results`, `filters` and `reranker`; Context7's Databricks documentation uses the 'AI Search' naming.",
133
+ "NOT corroborated by Context7 and therefore carried on Databricks documentation alone: the `query_type` parameter on the Python client (Context7 surfaced `query_type` only in the SQL form, e.g. `query_type => 'HYBRID'`), and the explicit four-way index-type taxonomy. State which source backs the claim when a user's call fails, and prefer verifying against the installed client.",
134
+ "Databricks service behaviour — MCP server categories, Unity AI Gateway policy, endpoint governance — is never a Context7 question. If Context7 is not exposed, say so and label the version-sensitive API claim `unknown` rather than answering from memory."
135
+ ],
136
+ "security_boundaries": [
137
+ "No live agent invocation — the skill reads code and configuration only.",
138
+ "No retrieval execution — no queries are run against the index.",
139
+ "No external model calls — provider connectivity is validated by name, not by test call.",
140
+ "MCP governance boundary: managed MCP governance is declarative (Databricks-hosted), external MCP security is delegated to the provider's OAuth, custom MCP security escalates to a live guard.",
141
+ "No gateway policy mutations — policies are reviewed but never changed without approval."
142
+ ],
143
+ "production_caveats": [
144
+ "AI Search full-text search is BETA (not GA); production reliance requires explicit risk acknowledgment and may be unsupported in some Databricks editions.",
145
+ "Continuous sync on storage-optimized endpoints is not supported; use triggered sync or accept eventual consistency.",
146
+ "Unity AI Gateway spend tables (`system.ai_gateway.external_model_spend`) are BETA and aggregate HOURLY, not real-time; real-time cost observability requires alternative instrumentation.",
147
+ "MCP server creation and provider OAuth secret binding are live-guard operations; treat them as production changes requiring approval."
148
+ ],
149
+ "hard_denials": [
150
+ "Invoking an agent or testing it against live data without evaluation setup.",
151
+ "Creating or modifying MCP server deployments without a live-guard approval.",
152
+ "Configuring external model providers without confirming they are supported on Databricks.",
153
+ "Selecting full-text search without acknowledging its BETA status.",
154
+ "Building context retrieval that assumes unbounded result volume or re-queries the entire index on each invocation.",
155
+ "Granting tool privileges without confirming the caller has privilege on the underlying data."
156
+ ],
157
+ "response_minimum": [
158
+ "A verdict (sound / cautions / block) and the agent framework and ResponsesAgent wrapping confirmed.",
159
+ "AI Search index variant/sync-mode, context engineering, tool inventory, model provider, and gateway policy findings.",
160
+ "A severity-labelled finding list (critical / high / medium / low) with evidence-basis labels and safe next actions."
161
+ ],
162
+ "references": [
163
+ {
164
+ "file": "ai-search-and-retrieval-config.md",
165
+ "title": "Databricks AI Search Index And Retrieval Configuration",
166
+ "purpose": "Index variants, sync modes, query types, and pagination semantics.",
167
+ "claims": [
168
+ "Databricks AI Search (formerly Databricks Vector Search) offers four index variants: Delta Sync with Databricks-managed embeddings (Databricks computes embeddings), Delta Sync with self-managed embeddings (caller provides vectors), Direct Vector Access (external vector source), and full-text search (BETA, keyword-only, storage-optimized endpoints only).",
169
+ "Sync modes differ by variant: continuous sync updates the index on every Delta write (not supported on storage-optimized endpoints); triggered sync updates on-demand or on a schedule (required for full-text indexes on storage-optimized endpoints); manual sync (Direct Vector Access only, no automatic updates).",
170
+ "Query types include `\"ann\"` (approximate nearest neighbor, default, vector-only), `\"hybrid\"` (vector + keyword using reciprocal rank fusion), and `\"FULL_TEXT\"` (BETA, keyword-only). Hybrid queries are more expensive than ANN but more robust to keyword-heavy requests.",
171
+ "Query API via Python: `similarity_search(query_text, query_vector, columns, num_results, query_type, filters, reranker)`. REST API: `POST /api/2.0/vector-search/indexes/{index_name}/query` with pagination via `query-next-page` and `page_token`.",
172
+ "Result pagination is capped at 1,000 results per query; unbounded result retrieval requires multiple queries or acceptance of the 1,000-result limit.",
173
+ "The `filters` parameter enables predicates on metadata columns; `reranker` allows post-retrieval re-ranking by a separate model.",
174
+ "Storage-optimized endpoints do not support continuous sync or ANN query types; they support triggered sync and full-text search only.",
175
+ "Full-text search is BETA and storage-optimized endpoints only; production reliance on this feature requires explicit risk acknowledgment."
176
+ ]
177
+ },
178
+ {
179
+ "file": "context-engineering-and-tools.md",
180
+ "title": "Context Engineering And Tool Integration",
181
+ "purpose": "Context assembly, grounding, token budgets, and MCP server governance.",
182
+ "claims": [
183
+ "Context engineering is the selection, chunking, and assembly of grounding data for the agent's LLM calls. Sound design pairs the retrieval context size (token count) with the model's context window and the prompt's other uses (system prompt, tool definitions, conversation history).",
184
+ "Chunking strategy affects retrieval quality: fixed-size windows are simple but may split semantic units; semantic chunking (via embeddings or NLP) preserves meaning but requires additional compute.",
185
+ "Context budget (the token count allocated to retrieval results) must be set explicitly; the agent does not auto-limit retrieval based on model context, so a budget that is too large creates latency and cost.",
186
+ "MCP (Model Context Protocol) servers are categorized by governance: managed MCP (Databricks-hosted for Genie, AI Search, Unity Catalog functions, SaaS connectors) are governed through Unity Catalog; external MCP (third-party over managed OAuth) delegate auth to the provider; custom MCP (Databricks Apps) require hosting and lifecycle.",
187
+ "MCP servers on Databricks are governed through Unity Catalog for access control and through Unity AI Gateway for monitoring and policy. A tool defined as an MCP server is governed by these scopes.",
188
+ "Unity Catalog functions can be exposed as agent tools directly. The agent's privilege to call the function is the same as the caller's privilege — the caller must have EXECUTE on the function and read privilege on the underlying data.",
189
+ "External model providers supported on Databricks are exactly: OpenAI (including Azure OpenAI), Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, Databricks Model Serving, and custom OpenAI-compatible proxies. Other providers (Gemini, Claude not through Bedrock) are not supported.",
190
+ "Unity AI Gateway provides rate limiting, traffic splitting (useful for A/B testing model variants), fallbacks (to secondary providers if primary fails), budget management (per-token or per-minute caps), and request/response content policies (e.g., PII masks, input/output filters)."
191
+ ],
192
+ "table": {
193
+ "title": "MCP Server Category And Governance Scope",
194
+ "header": [
195
+ "Category",
196
+ "Hosting",
197
+ "Governance",
198
+ "Authentication",
199
+ "Use Case"
200
+ ],
201
+ "rows": [
202
+ [
203
+ "Managed",
204
+ "Databricks-hosted",
205
+ "Unity Catalog access control",
206
+ "Databricks identity",
207
+ "Genie, AI Search, UC functions, SaaS connectors"
208
+ ],
209
+ [
210
+ "External",
211
+ "Third-party server",
212
+ "Provider OAuth",
213
+ "Provider credentials",
214
+ "GitHub, Jira, other SaaS with OAuth"
215
+ ],
216
+ [
217
+ "Custom",
218
+ "Databricks App",
219
+ "Unity Catalog access control",
220
+ "Databricks identity",
221
+ "Internal tools, proprietary functions"
222
+ ]
223
+ ]
224
+ }
225
+ },
226
+ {
227
+ "file": "official-sources.md",
228
+ "title": "Official Sources",
229
+ "purpose": "Primary Mosaic AI Agent Framework, AI Search, MCP, and Unity AI Gateway documentation.",
230
+ "claims": [
231
+ "The retrieval client surface was cross-checked against the Context7 MCP (`/websites/databricks`). Where Context7 surfaced a parameter only in the SQL form and not the Python client, this skill says so rather than presenting the Python signature as corroborated."
232
+ ]
233
+ },
234
+ {
235
+ "file": "workflow-and-output.md",
236
+ "title": "Workflow And Output",
237
+ "purpose": "Diagnostic sequence and output contract for agent-architecture review."
238
+ },
239
+ {
240
+ "file": "safety-checklist.md",
241
+ "title": "Safety Checklist",
242
+ "purpose": "MCP governance escalation, tool privilege scoping, and production-readiness gates for agent engineering on Databricks."
243
+ }
244
+ ]
245
+ }
246
+ }