@raishin/vanguard-frontier-agentic 3.9.0 → 3.11.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (810) hide show
  1. package/.claude-plugin/marketplace.json +2 -2
  2. package/.claude-plugin/plugin.json +43 -1
  3. package/.cursor-plugin/plugin.json +43 -19
  4. package/.github/plugin/marketplace.json +1 -1
  5. package/README.md +63 -19
  6. package/agents/databricks/databricks-ai-bi-genie-agent/AGENT.md +90 -0
  7. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/claude-code.agent.md +73 -0
  8. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/codex.toml +15 -0
  9. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/copilot.agent.md +79 -0
  10. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/cursor.agent.md +74 -0
  11. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/gemini.agent.md +73 -0
  12. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/kiro-cli.agent.json +5 -0
  13. package/agents/databricks/databricks-ai-bi-genie-agent/harnesses/kiro-ide.agent.md +73 -0
  14. package/agents/databricks/databricks-ai-bi-genie-agent/metadata.json +58 -0
  15. package/agents/databricks/databricks-data-protection-privacy-agent/AGENT.md +94 -0
  16. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/claude-code.agent.md +77 -0
  17. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/codex.toml +15 -0
  18. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/copilot.agent.md +83 -0
  19. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/cursor.agent.md +78 -0
  20. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/gemini.agent.md +77 -0
  21. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/kiro-cli.agent.json +5 -0
  22. package/agents/databricks/databricks-data-protection-privacy-agent/harnesses/kiro-ide.agent.md +77 -0
  23. package/agents/databricks/databricks-data-protection-privacy-agent/metadata.json +64 -0
  24. package/agents/databricks/databricks-data-quality-observability-agent/AGENT.md +89 -0
  25. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/claude-code.agent.md +72 -0
  26. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/codex.toml +15 -0
  27. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/copilot.agent.md +78 -0
  28. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/cursor.agent.md +73 -0
  29. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/gemini.agent.md +72 -0
  30. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/kiro-cli.agent.json +5 -0
  31. package/agents/databricks/databricks-data-quality-observability-agent/harnesses/kiro-ide.agent.md +72 -0
  32. package/agents/databricks/databricks-data-quality-observability-agent/metadata.json +59 -0
  33. package/agents/databricks/databricks-developer-platform-agent/AGENT.md +90 -0
  34. package/agents/databricks/databricks-developer-platform-agent/harnesses/claude-code.agent.md +73 -0
  35. package/agents/databricks/databricks-developer-platform-agent/harnesses/codex.toml +15 -0
  36. package/agents/databricks/databricks-developer-platform-agent/harnesses/copilot.agent.md +79 -0
  37. package/agents/databricks/databricks-developer-platform-agent/harnesses/cursor.agent.md +74 -0
  38. package/agents/databricks/databricks-developer-platform-agent/harnesses/gemini.agent.md +73 -0
  39. package/agents/databricks/databricks-developer-platform-agent/harnesses/kiro-cli.agent.json +5 -0
  40. package/agents/databricks/databricks-developer-platform-agent/harnesses/kiro-ide.agent.md +73 -0
  41. package/agents/databricks/databricks-developer-platform-agent/metadata.json +59 -0
  42. package/agents/databricks/databricks-finops-cost-agent/AGENT.md +91 -0
  43. package/agents/databricks/databricks-finops-cost-agent/harnesses/claude-code.agent.md +74 -0
  44. package/agents/databricks/databricks-finops-cost-agent/harnesses/codex.toml +15 -0
  45. package/agents/databricks/databricks-finops-cost-agent/harnesses/copilot.agent.md +80 -0
  46. package/agents/databricks/databricks-finops-cost-agent/harnesses/cursor.agent.md +75 -0
  47. package/agents/databricks/databricks-finops-cost-agent/harnesses/gemini.agent.md +74 -0
  48. package/agents/databricks/databricks-finops-cost-agent/harnesses/kiro-cli.agent.json +5 -0
  49. package/agents/databricks/databricks-finops-cost-agent/harnesses/kiro-ide.agent.md +74 -0
  50. package/agents/databricks/databricks-finops-cost-agent/metadata.json +60 -0
  51. package/agents/databricks/databricks-genai-agent-engineering-agent/AGENT.md +89 -0
  52. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/claude-code.agent.md +72 -0
  53. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/codex.toml +15 -0
  54. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/copilot.agent.md +78 -0
  55. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/cursor.agent.md +73 -0
  56. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/gemini.agent.md +72 -0
  57. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/kiro-cli.agent.json +5 -0
  58. package/agents/databricks/databricks-genai-agent-engineering-agent/harnesses/kiro-ide.agent.md +72 -0
  59. package/agents/databricks/databricks-genai-agent-engineering-agent/metadata.json +62 -0
  60. package/agents/databricks/databricks-genai-evaluation-observability-agent/AGENT.md +89 -0
  61. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/claude-code.agent.md +72 -0
  62. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/codex.toml +15 -0
  63. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/copilot.agent.md +78 -0
  64. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/cursor.agent.md +73 -0
  65. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/gemini.agent.md +72 -0
  66. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/kiro-cli.agent.json +5 -0
  67. package/agents/databricks/databricks-genai-evaluation-observability-agent/harnesses/kiro-ide.agent.md +72 -0
  68. package/agents/databricks/databricks-genai-evaluation-observability-agent/metadata.json +59 -0
  69. package/agents/databricks/databricks-identity-network-security-agent/AGENT.md +95 -0
  70. package/agents/databricks/databricks-identity-network-security-agent/harnesses/claude-code.agent.md +78 -0
  71. package/agents/databricks/databricks-identity-network-security-agent/harnesses/codex.toml +15 -0
  72. package/agents/databricks/databricks-identity-network-security-agent/harnesses/copilot.agent.md +84 -0
  73. package/agents/databricks/databricks-identity-network-security-agent/harnesses/cursor.agent.md +79 -0
  74. package/agents/databricks/databricks-identity-network-security-agent/harnesses/gemini.agent.md +78 -0
  75. package/agents/databricks/databricks-identity-network-security-agent/harnesses/kiro-cli.agent.json +5 -0
  76. package/agents/databricks/databricks-identity-network-security-agent/harnesses/kiro-ide.agent.md +78 -0
  77. package/agents/databricks/databricks-identity-network-security-agent/metadata.json +60 -0
  78. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/AGENT.md +90 -0
  79. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/claude-code.agent.md +73 -0
  80. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/codex.toml +15 -0
  81. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/copilot.agent.md +79 -0
  82. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/cursor.agent.md +74 -0
  83. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/gemini.agent.md +73 -0
  84. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/kiro-cli.agent.json +5 -0
  85. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/harnesses/kiro-ide.agent.md +73 -0
  86. package/agents/databricks/databricks-lakeflow-pipeline-engineering-agent/metadata.json +63 -0
  87. package/agents/databricks/databricks-maestro-agent/AGENT.md +63 -0
  88. package/agents/databricks/databricks-maestro-agent/README.md +76 -0
  89. package/agents/databricks/databricks-maestro-agent/harnesses/claude-code.agent.md +46 -0
  90. package/agents/databricks/databricks-maestro-agent/harnesses/codex.toml +15 -0
  91. package/agents/databricks/databricks-maestro-agent/harnesses/copilot.agent.md +52 -0
  92. package/agents/databricks/databricks-maestro-agent/harnesses/cursor.agent.md +47 -0
  93. package/agents/databricks/databricks-maestro-agent/harnesses/gemini.agent.md +46 -0
  94. package/agents/databricks/databricks-maestro-agent/harnesses/kiro-cli.agent.json +5 -0
  95. package/agents/databricks/databricks-maestro-agent/harnesses/kiro-ide.agent.md +46 -0
  96. package/agents/databricks/databricks-maestro-agent/metadata.json +50 -0
  97. package/agents/databricks/databricks-mlops-agent/AGENT.md +89 -0
  98. package/agents/databricks/databricks-mlops-agent/harnesses/claude-code.agent.md +72 -0
  99. package/agents/databricks/databricks-mlops-agent/harnesses/codex.toml +15 -0
  100. package/agents/databricks/databricks-mlops-agent/harnesses/copilot.agent.md +78 -0
  101. package/agents/databricks/databricks-mlops-agent/harnesses/cursor.agent.md +73 -0
  102. package/agents/databricks/databricks-mlops-agent/harnesses/gemini.agent.md +72 -0
  103. package/agents/databricks/databricks-mlops-agent/harnesses/kiro-cli.agent.json +5 -0
  104. package/agents/databricks/databricks-mlops-agent/harnesses/kiro-ide.agent.md +72 -0
  105. package/agents/databricks/databricks-mlops-agent/metadata.json +60 -0
  106. package/agents/databricks/databricks-platform-architecture-agent/AGENT.md +90 -0
  107. package/agents/databricks/databricks-platform-architecture-agent/harnesses/claude-code.agent.md +73 -0
  108. package/agents/databricks/databricks-platform-architecture-agent/harnesses/codex.toml +15 -0
  109. package/agents/databricks/databricks-platform-architecture-agent/harnesses/copilot.agent.md +79 -0
  110. package/agents/databricks/databricks-platform-architecture-agent/harnesses/cursor.agent.md +74 -0
  111. package/agents/databricks/databricks-platform-architecture-agent/harnesses/gemini.agent.md +73 -0
  112. package/agents/databricks/databricks-platform-architecture-agent/harnesses/kiro-cli.agent.json +5 -0
  113. package/agents/databricks/databricks-platform-architecture-agent/harnesses/kiro-ide.agent.md +73 -0
  114. package/agents/databricks/databricks-platform-architecture-agent/metadata.json +58 -0
  115. package/agents/databricks/databricks-platform-reliability-agent/AGENT.md +88 -0
  116. package/agents/databricks/databricks-platform-reliability-agent/harnesses/claude-code.agent.md +71 -0
  117. package/agents/databricks/databricks-platform-reliability-agent/harnesses/codex.toml +15 -0
  118. package/agents/databricks/databricks-platform-reliability-agent/harnesses/copilot.agent.md +77 -0
  119. package/agents/databricks/databricks-platform-reliability-agent/harnesses/cursor.agent.md +72 -0
  120. package/agents/databricks/databricks-platform-reliability-agent/harnesses/gemini.agent.md +71 -0
  121. package/agents/databricks/databricks-platform-reliability-agent/harnesses/kiro-cli.agent.json +5 -0
  122. package/agents/databricks/databricks-platform-reliability-agent/harnesses/kiro-ide.agent.md +71 -0
  123. package/agents/databricks/databricks-platform-reliability-agent/metadata.json +64 -0
  124. package/agents/databricks/databricks-sql-performance-agent/AGENT.md +91 -0
  125. package/agents/databricks/databricks-sql-performance-agent/harnesses/claude-code.agent.md +74 -0
  126. package/agents/databricks/databricks-sql-performance-agent/harnesses/codex.toml +15 -0
  127. package/agents/databricks/databricks-sql-performance-agent/harnesses/copilot.agent.md +80 -0
  128. package/agents/databricks/databricks-sql-performance-agent/harnesses/cursor.agent.md +75 -0
  129. package/agents/databricks/databricks-sql-performance-agent/harnesses/gemini.agent.md +74 -0
  130. package/agents/databricks/databricks-sql-performance-agent/harnesses/kiro-cli.agent.json +5 -0
  131. package/agents/databricks/databricks-sql-performance-agent/harnesses/kiro-ide.agent.md +74 -0
  132. package/agents/databricks/databricks-sql-performance-agent/metadata.json +62 -0
  133. package/agents/databricks/databricks-streaming-reliability-agent/AGENT.md +93 -0
  134. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/claude-code.agent.md +76 -0
  135. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/codex.toml +15 -0
  136. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/copilot.agent.md +82 -0
  137. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/cursor.agent.md +77 -0
  138. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/gemini.agent.md +76 -0
  139. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/kiro-cli.agent.json +5 -0
  140. package/agents/databricks/databricks-streaming-reliability-agent/harnesses/kiro-ide.agent.md +76 -0
  141. package/agents/databricks/databricks-streaming-reliability-agent/metadata.json +62 -0
  142. package/agents/databricks/databricks-unity-catalog-governance-agent/AGENT.md +91 -0
  143. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/claude-code.agent.md +74 -0
  144. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/codex.toml +15 -0
  145. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/copilot.agent.md +80 -0
  146. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/cursor.agent.md +75 -0
  147. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/gemini.agent.md +74 -0
  148. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/kiro-cli.agent.json +5 -0
  149. package/agents/databricks/databricks-unity-catalog-governance-agent/harnesses/kiro-ide.agent.md +74 -0
  150. package/agents/databricks/databricks-unity-catalog-governance-agent/metadata.json +63 -0
  151. package/agents/databricks/databricks-value-realization-agent/AGENT.md +91 -0
  152. package/agents/databricks/databricks-value-realization-agent/harnesses/claude-code.agent.md +74 -0
  153. package/agents/databricks/databricks-value-realization-agent/harnesses/codex.toml +15 -0
  154. package/agents/databricks/databricks-value-realization-agent/harnesses/copilot.agent.md +80 -0
  155. package/agents/databricks/databricks-value-realization-agent/harnesses/cursor.agent.md +75 -0
  156. package/agents/databricks/databricks-value-realization-agent/harnesses/gemini.agent.md +74 -0
  157. package/agents/databricks/databricks-value-realization-agent/harnesses/kiro-cli.agent.json +5 -0
  158. package/agents/databricks/databricks-value-realization-agent/harnesses/kiro-ide.agent.md +74 -0
  159. package/agents/databricks/databricks-value-realization-agent/metadata.json +55 -0
  160. package/agents/snowflake/AGENTS.md +199 -0
  161. package/agents/snowflake/README.md +227 -65
  162. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/AGENT.md +149 -0
  163. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/claude-code.agent.md +132 -0
  164. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/codex.toml +41 -0
  165. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/copilot.agent.md +138 -0
  166. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/cursor.agent.md +133 -0
  167. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/gemini.agent.md +132 -0
  168. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/kiro-cli.agent.json +5 -0
  169. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/harnesses/kiro-ide.agent.md +132 -0
  170. package/agents/snowflake/snowflake-analytics-semantic-data-product-agent/metadata.json +60 -0
  171. package/agents/snowflake/snowflake-bcdr-resilience-agent/AGENT.md +160 -0
  172. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/claude-code.agent.md +143 -0
  173. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/codex.toml +43 -0
  174. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/copilot.agent.md +149 -0
  175. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/cursor.agent.md +144 -0
  176. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/gemini.agent.md +143 -0
  177. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/kiro-cli.agent.json +5 -0
  178. package/agents/snowflake/snowflake-bcdr-resilience-agent/harnesses/kiro-ide.agent.md +143 -0
  179. package/agents/snowflake/snowflake-bcdr-resilience-agent/metadata.json +63 -0
  180. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/AGENT.md +155 -0
  181. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/claude-code.agent.md +138 -0
  182. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/codex.toml +43 -0
  183. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/copilot.agent.md +144 -0
  184. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/cursor.agent.md +139 -0
  185. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/gemini.agent.md +138 -0
  186. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/kiro-cli.agent.json +5 -0
  187. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/harnesses/kiro-ide.agent.md +138 -0
  188. package/agents/snowflake/snowflake-business-value-adoption-strategist-agent/metadata.json +60 -0
  189. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/AGENT.md +150 -0
  190. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/claude-code.agent.md +133 -0
  191. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/codex.toml +41 -0
  192. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/copilot.agent.md +139 -0
  193. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/cursor.agent.md +134 -0
  194. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/gemini.agent.md +133 -0
  195. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/kiro-cli.agent.json +5 -0
  196. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/harnesses/kiro-ide.agent.md +133 -0
  197. package/agents/snowflake/snowflake-compliance-evidence-auditor-agent/metadata.json +62 -0
  198. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/AGENT.md +166 -0
  199. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/claude-code.agent.md +149 -0
  200. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/codex.toml +44 -0
  201. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/copilot.agent.md +155 -0
  202. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/cursor.agent.md +150 -0
  203. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/gemini.agent.md +149 -0
  204. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/kiro-cli.agent.json +5 -0
  205. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/harnesses/kiro-ide.agent.md +149 -0
  206. package/agents/snowflake/snowflake-cortex-ai-agent-security-governor-agent/metadata.json +65 -0
  207. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/AGENT.md +156 -0
  208. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/claude-code.agent.md +139 -0
  209. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/codex.toml +42 -0
  210. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/copilot.agent.md +145 -0
  211. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/cursor.agent.md +140 -0
  212. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/gemini.agent.md +139 -0
  213. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/kiro-cli.agent.json +5 -0
  214. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/harnesses/kiro-ide.agent.md +139 -0
  215. package/agents/snowflake/snowflake-data-engineering-pipelines-agent/metadata.json +63 -0
  216. package/agents/snowflake/snowflake-data-platform-engineering-at-azure-agent/metadata.json +6 -3
  217. package/agents/snowflake/snowflake-data-science-ml-agent/AGENT.md +155 -0
  218. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/claude-code.agent.md +138 -0
  219. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/codex.toml +42 -0
  220. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/copilot.agent.md +144 -0
  221. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/cursor.agent.md +139 -0
  222. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/gemini.agent.md +138 -0
  223. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/kiro-cli.agent.json +5 -0
  224. package/agents/snowflake/snowflake-data-science-ml-agent/harnesses/kiro-ide.agent.md +138 -0
  225. package/agents/snowflake/snowflake-data-science-ml-agent/metadata.json +60 -0
  226. package/agents/snowflake/snowflake-devops-iac-release-agent/AGENT.md +157 -0
  227. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/claude-code.agent.md +140 -0
  228. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/codex.toml +43 -0
  229. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/copilot.agent.md +146 -0
  230. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/cursor.agent.md +141 -0
  231. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/gemini.agent.md +140 -0
  232. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/kiro-cli.agent.json +5 -0
  233. package/agents/snowflake/snowflake-devops-iac-release-agent/harnesses/kiro-ide.agent.md +140 -0
  234. package/agents/snowflake/snowflake-devops-iac-release-agent/metadata.json +63 -0
  235. package/agents/snowflake/snowflake-finops-cost-governor-agent/AGENT.md +157 -0
  236. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/claude-code.agent.md +140 -0
  237. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/codex.toml +43 -0
  238. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/copilot.agent.md +146 -0
  239. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/cursor.agent.md +141 -0
  240. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/gemini.agent.md +140 -0
  241. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/kiro-cli.agent.json +5 -0
  242. package/agents/snowflake/snowflake-finops-cost-governor-agent/harnesses/kiro-ide.agent.md +140 -0
  243. package/agents/snowflake/snowflake-finops-cost-governor-agent/metadata.json +62 -0
  244. package/agents/snowflake/snowflake-governance-privacy-agent/AGENT.md +155 -0
  245. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/claude-code.agent.md +138 -0
  246. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/codex.toml +42 -0
  247. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/copilot.agent.md +144 -0
  248. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/cursor.agent.md +139 -0
  249. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/gemini.agent.md +138 -0
  250. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/kiro-cli.agent.json +5 -0
  251. package/agents/snowflake/snowflake-governance-privacy-agent/harnesses/kiro-ide.agent.md +138 -0
  252. package/agents/snowflake/snowflake-governance-privacy-agent/metadata.json +64 -0
  253. package/agents/snowflake/snowflake-identity-access-security-agent/AGENT.md +155 -0
  254. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/claude-code.agent.md +138 -0
  255. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/codex.toml +42 -0
  256. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/copilot.agent.md +144 -0
  257. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/cursor.agent.md +139 -0
  258. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/gemini.agent.md +138 -0
  259. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/kiro-cli.agent.json +5 -0
  260. package/agents/snowflake/snowflake-identity-access-security-agent/harnesses/kiro-ide.agent.md +138 -0
  261. package/agents/snowflake/snowflake-identity-access-security-agent/metadata.json +69 -0
  262. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/AGENT.md +170 -0
  263. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/PERMISSIONS.md +81 -0
  264. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/PREFLIGHT.md +45 -0
  265. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/ROLLBACK.md +38 -0
  266. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/claude-code.agent.md +133 -0
  267. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/codex.toml +45 -0
  268. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/copilot.agent.md +139 -0
  269. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/cursor.agent.md +134 -0
  270. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/gemini.agent.md +133 -0
  271. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/kiro-cli.agent.json +5 -0
  272. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/harnesses/kiro-ide.agent.md +133 -0
  273. package/agents/snowflake/snowflake-live-auth-network-policy-guard-agent/metadata.json +76 -0
  274. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/AGENT.md +171 -0
  275. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/PERMISSIONS.md +80 -0
  276. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/PREFLIGHT.md +45 -0
  277. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/ROLLBACK.md +38 -0
  278. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/claude-code.agent.md +134 -0
  279. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/codex.toml +45 -0
  280. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/copilot.agent.md +140 -0
  281. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/cursor.agent.md +135 -0
  282. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/gemini.agent.md +134 -0
  283. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/kiro-cli.agent.json +5 -0
  284. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/harnesses/kiro-ide.agent.md +134 -0
  285. package/agents/snowflake/snowflake-live-data-protection-policy-guard-agent/metadata.json +76 -0
  286. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/AGENT.md +179 -0
  287. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/PERMISSIONS.md +81 -0
  288. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/PREFLIGHT.md +49 -0
  289. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/ROLLBACK.md +40 -0
  290. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/claude-code.agent.md +142 -0
  291. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/codex.toml +47 -0
  292. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/copilot.agent.md +148 -0
  293. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/cursor.agent.md +143 -0
  294. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/gemini.agent.md +142 -0
  295. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/kiro-cli.agent.json +5 -0
  296. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/harnesses/kiro-ide.agent.md +142 -0
  297. package/agents/snowflake/snowflake-live-failover-promotion-guard-agent/metadata.json +76 -0
  298. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/AGENT.md +172 -0
  299. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/PERMISSIONS.md +80 -0
  300. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/PREFLIGHT.md +47 -0
  301. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/ROLLBACK.md +39 -0
  302. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/claude-code.agent.md +135 -0
  303. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/codex.toml +45 -0
  304. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/copilot.agent.md +141 -0
  305. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/cursor.agent.md +136 -0
  306. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/gemini.agent.md +135 -0
  307. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/kiro-cli.agent.json +5 -0
  308. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/harnesses/kiro-ide.agent.md +135 -0
  309. package/agents/snowflake/snowflake-live-pipeline-streaming-change-guard-agent/metadata.json +78 -0
  310. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/AGENT.md +167 -0
  311. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/PERMISSIONS.md +78 -0
  312. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/PREFLIGHT.md +44 -0
  313. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/ROLLBACK.md +38 -0
  314. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/claude-code.agent.md +130 -0
  315. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/codex.toml +44 -0
  316. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/copilot.agent.md +136 -0
  317. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/cursor.agent.md +131 -0
  318. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/gemini.agent.md +130 -0
  319. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/kiro-cli.agent.json +5 -0
  320. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/harnesses/kiro-ide.agent.md +130 -0
  321. package/agents/snowflake/snowflake-live-rbac-grant-guard-agent/metadata.json +76 -0
  322. package/agents/snowflake/snowflake-live-rbac-grant-guard-at-azure-agent/metadata.json +11 -4
  323. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/AGENT.md +170 -0
  324. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/PERMISSIONS.md +80 -0
  325. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/PREFLIGHT.md +47 -0
  326. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/ROLLBACK.md +39 -0
  327. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/claude-code.agent.md +133 -0
  328. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/codex.toml +45 -0
  329. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/copilot.agent.md +139 -0
  330. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/cursor.agent.md +134 -0
  331. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/gemini.agent.md +133 -0
  332. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/kiro-cli.agent.json +5 -0
  333. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/harnesses/kiro-ide.agent.md +133 -0
  334. package/agents/snowflake/snowflake-live-warehouse-cost-change-guard-agent/metadata.json +76 -0
  335. package/agents/snowflake/snowflake-maestro-agent/AGENT.md +84 -0
  336. package/agents/snowflake/snowflake-maestro-agent/README.md +71 -0
  337. package/agents/snowflake/snowflake-maestro-agent/harnesses/claude-code.agent.md +67 -0
  338. package/agents/snowflake/snowflake-maestro-agent/harnesses/codex.toml +41 -0
  339. package/agents/snowflake/snowflake-maestro-agent/harnesses/copilot.agent.md +73 -0
  340. package/agents/snowflake/snowflake-maestro-agent/harnesses/cursor.agent.md +68 -0
  341. package/agents/snowflake/snowflake-maestro-agent/harnesses/gemini.agent.md +67 -0
  342. package/agents/snowflake/snowflake-maestro-agent/harnesses/kiro-cli.agent.json +5 -0
  343. package/agents/snowflake/snowflake-maestro-agent/harnesses/kiro-ide.agent.md +67 -0
  344. package/agents/snowflake/snowflake-maestro-agent/metadata.json +47 -0
  345. package/agents/snowflake/snowflake-migration-modernization-agent/AGENT.md +160 -0
  346. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/claude-code.agent.md +143 -0
  347. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/codex.toml +43 -0
  348. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/copilot.agent.md +149 -0
  349. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/cursor.agent.md +144 -0
  350. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/gemini.agent.md +143 -0
  351. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/kiro-cli.agent.json +5 -0
  352. package/agents/snowflake/snowflake-migration-modernization-agent/harnesses/kiro-ide.agent.md +143 -0
  353. package/agents/snowflake/snowflake-migration-modernization-agent/metadata.json +63 -0
  354. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/AGENT.md +157 -0
  355. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/claude-code.agent.md +140 -0
  356. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/codex.toml +42 -0
  357. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/copilot.agent.md +146 -0
  358. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/cursor.agent.md +141 -0
  359. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/gemini.agent.md +140 -0
  360. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/kiro-cli.agent.json +5 -0
  361. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/harnesses/kiro-ide.agent.md +140 -0
  362. package/agents/snowflake/snowflake-native-app-marketplace-product-agent/metadata.json +61 -0
  363. package/agents/snowflake/snowflake-network-private-connectivity-agent/AGENT.md +149 -0
  364. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/claude-code.agent.md +132 -0
  365. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/codex.toml +41 -0
  366. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/copilot.agent.md +138 -0
  367. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/cursor.agent.md +133 -0
  368. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/gemini.agent.md +132 -0
  369. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/kiro-cli.agent.json +5 -0
  370. package/agents/snowflake/snowflake-network-private-connectivity-agent/harnesses/kiro-ide.agent.md +132 -0
  371. package/agents/snowflake/snowflake-network-private-connectivity-agent/metadata.json +61 -0
  372. package/agents/snowflake/snowflake-platform-administrator-agent/AGENT.md +146 -0
  373. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/claude-code.agent.md +129 -0
  374. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/codex.toml +40 -0
  375. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/copilot.agent.md +135 -0
  376. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/cursor.agent.md +130 -0
  377. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/gemini.agent.md +129 -0
  378. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/kiro-cli.agent.json +5 -0
  379. package/agents/snowflake/snowflake-platform-administrator-agent/harnesses/kiro-ide.agent.md +129 -0
  380. package/agents/snowflake/snowflake-platform-administrator-agent/metadata.json +57 -0
  381. package/agents/snowflake/snowflake-query-performance-engineer-agent/AGENT.md +153 -0
  382. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/claude-code.agent.md +136 -0
  383. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/codex.toml +41 -0
  384. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/copilot.agent.md +142 -0
  385. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/cursor.agent.md +137 -0
  386. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/gemini.agent.md +136 -0
  387. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/kiro-cli.agent.json +5 -0
  388. package/agents/snowflake/snowflake-query-performance-engineer-agent/harnesses/kiro-ide.agent.md +136 -0
  389. package/agents/snowflake/snowflake-query-performance-engineer-agent/metadata.json +63 -0
  390. package/agents/snowflake/snowflake-rbac-access-governance-at-azure-agent/metadata.json +6 -3
  391. package/agents/snowflake/snowflake-solution-architect-agent/AGENT.md +144 -0
  392. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/claude-code.agent.md +127 -0
  393. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/codex.toml +40 -0
  394. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/copilot.agent.md +133 -0
  395. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/cursor.agent.md +128 -0
  396. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/gemini.agent.md +127 -0
  397. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/kiro-cli.agent.json +5 -0
  398. package/agents/snowflake/snowflake-solution-architect-agent/harnesses/kiro-ide.agent.md +127 -0
  399. package/agents/snowflake/snowflake-solution-architect-agent/metadata.json +59 -0
  400. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/AGENT.md +158 -0
  401. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/claude-code.agent.md +141 -0
  402. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/codex.toml +43 -0
  403. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/copilot.agent.md +147 -0
  404. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/cursor.agent.md +142 -0
  405. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/gemini.agent.md +141 -0
  406. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/kiro-cli.agent.json +5 -0
  407. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/harnesses/kiro-ide.agent.md +141 -0
  408. package/agents/snowflake/snowflake-streaming-ingestion-reliability-agent/metadata.json +61 -0
  409. package/catalog/agents.json +8965 -7038
  410. package/catalog/asset-integrity.json +2877 -67
  411. package/catalog/install-roles.json +299 -1
  412. package/catalog/model-assignments.json +3469 -2083
  413. package/catalog/skill-manifest.json +10859 -9480
  414. package/catalog/skills.json +8195 -6801
  415. package/package.json +6 -4
  416. package/plugins/vanguard-frontier-agentic/.codex-plugin/plugin.json +1 -1
  417. package/powers/README.md +2 -5
  418. package/powers/vanguard-cilium/POWER.md +2 -2
  419. package/powers/vanguard-databricks/POWER.md +11 -11
  420. package/powers/vanguard-microsoft/POWER.md +0 -2
  421. package/powers/vanguard-snowflake/POWER.md +19 -13
  422. package/scripts/databricks_data/agents/00-databricks-maestro-agent.json +165 -0
  423. package/scripts/databricks_data/agents/01-databricks-platform-architecture-agent.json +200 -0
  424. package/scripts/databricks_data/agents/02-databricks-unity-catalog-governance-agent.json +207 -0
  425. package/scripts/databricks_data/agents/03-databricks-identity-network-security-agent.json +216 -0
  426. package/scripts/databricks_data/agents/04-databricks-data-protection-privacy-agent.json +218 -0
  427. package/scripts/databricks_data/agents/05-databricks-lakeflow-pipeline-engineering-agent.json +217 -0
  428. package/scripts/databricks_data/agents/06-databricks-streaming-reliability-agent.json +267 -0
  429. package/scripts/databricks_data/agents/07-databricks-data-quality-observability-agent.json +215 -0
  430. package/scripts/databricks_data/agents/08-databricks-sql-performance-agent.json +214 -0
  431. package/scripts/databricks_data/agents/09-databricks-ai-bi-genie-agent.json +211 -0
  432. package/scripts/databricks_data/agents/10-databricks-mlops-agent.json +239 -0
  433. package/scripts/databricks_data/agents/11-databricks-genai-agent-engineering-agent.json +246 -0
  434. package/scripts/databricks_data/agents/12-databricks-genai-evaluation-observability-agent.json +215 -0
  435. package/scripts/databricks_data/agents/13-databricks-developer-platform-agent.json +206 -0
  436. package/scripts/databricks_data/agents/14-databricks-platform-reliability-agent.json +208 -0
  437. package/scripts/databricks_data/agents/15-databricks-finops-cost-agent.json +218 -0
  438. package/scripts/databricks_data/agents/16-databricks-value-realization-agent.json +232 -0
  439. package/scripts/gen_databricks_agents.py +703 -0
  440. package/scripts/gen_snowflake_agents.py +1158 -0
  441. package/scripts/generate-board-counts.mjs +287 -0
  442. package/scripts/generate-kiro-powers.mjs +12 -12
  443. package/scripts/generate-readme-counts.mjs +109 -0
  444. package/scripts/snowflake_data/agents/00-snowflake-maestro-agent.json +211 -0
  445. package/scripts/snowflake_data/agents/01-snowflake-solution-architect-agent.json +336 -0
  446. package/scripts/snowflake_data/agents/02-snowflake-platform-administrator-agent.json +259 -0
  447. package/scripts/snowflake_data/agents/03-snowflake-identity-access-security-agent.json +364 -0
  448. package/scripts/snowflake_data/agents/04-snowflake-network-private-connectivity-agent.json +266 -0
  449. package/scripts/snowflake_data/agents/05-snowflake-governance-privacy-agent.json +291 -0
  450. package/scripts/snowflake_data/agents/06-snowflake-compliance-evidence-auditor-agent.json +264 -0
  451. package/scripts/snowflake_data/agents/07-snowflake-finops-cost-governor-agent.json +353 -0
  452. package/scripts/snowflake_data/agents/08-snowflake-query-performance-engineer-agent.json +292 -0
  453. package/scripts/snowflake_data/agents/09-snowflake-data-engineering-pipelines-agent.json +296 -0
  454. package/scripts/snowflake_data/agents/10-snowflake-streaming-ingestion-reliability-agent.json +331 -0
  455. package/scripts/snowflake_data/agents/11-snowflake-analytics-semantic-data-product-agent.json +267 -0
  456. package/scripts/snowflake_data/agents/12-snowflake-data-science-ml-agent.json +269 -0
  457. package/scripts/snowflake_data/agents/13-snowflake-cortex-ai-agent-security-governor-agent.json +347 -0
  458. package/scripts/snowflake_data/agents/14-snowflake-native-app-marketplace-product-agent.json +274 -0
  459. package/scripts/snowflake_data/agents/15-snowflake-bcdr-resilience-agent.json +325 -0
  460. package/scripts/snowflake_data/agents/16-snowflake-devops-iac-release-agent.json +292 -0
  461. package/scripts/snowflake_data/agents/17-snowflake-migration-modernization-agent.json +251 -0
  462. package/scripts/snowflake_data/agents/18-snowflake-business-value-adoption-strategist-agent.json +262 -0
  463. package/scripts/snowflake_data/agents/19-snowflake-live-rbac-grant-guard-agent.json +299 -0
  464. package/scripts/snowflake_data/agents/20-snowflake-live-auth-network-policy-guard-agent.json +303 -0
  465. package/scripts/snowflake_data/agents/21-snowflake-live-warehouse-cost-change-guard-agent.json +314 -0
  466. package/scripts/snowflake_data/agents/22-snowflake-live-data-protection-policy-guard-agent.json +310 -0
  467. package/scripts/snowflake_data/agents/23-snowflake-live-pipeline-streaming-change-guard-agent.json +316 -0
  468. package/scripts/snowflake_data/agents/24-snowflake-live-failover-promotion-guard-agent.json +338 -0
  469. package/scripts/update-catalog-new-agents.py +5 -1
  470. package/skills/databricks/databricks-ai-bi-genie/SKILL.md +132 -0
  471. package/skills/databricks/databricks-ai-bi-genie/metadata.json +34 -0
  472. package/skills/databricks/databricks-ai-bi-genie/references/dashboard-and-permission-security.md +16 -0
  473. package/skills/databricks/databricks-ai-bi-genie/references/genie-scoping-and-semantic-layer.md +16 -0
  474. package/skills/databricks/databricks-ai-bi-genie/references/official-sources.md +24 -0
  475. package/skills/databricks/databricks-ai-bi-genie/references/safety-checklist.md +35 -0
  476. package/skills/databricks/databricks-ai-bi-genie/references/workflow-and-output.md +24 -0
  477. package/skills/databricks/databricks-data-protection-privacy/SKILL.md +142 -0
  478. package/skills/databricks/databricks-data-protection-privacy/metadata.json +37 -0
  479. package/skills/databricks/databricks-data-protection-privacy/references/deletion-vacuum-and-gdpr-compliance.md +9 -0
  480. package/skills/databricks/databricks-data-protection-privacy/references/masks-filters-and-abac-udf-cost.md +9 -0
  481. package/skills/databricks/databricks-data-protection-privacy/references/official-sources.md +27 -0
  482. package/skills/databricks/databricks-data-protection-privacy/references/safety-checklist.md +36 -0
  483. package/skills/databricks/databricks-data-protection-privacy/references/workflow-and-output.md +28 -0
  484. package/skills/databricks/databricks-data-quality-observability/SKILL.md +137 -0
  485. package/skills/databricks/databricks-data-quality-observability/metadata.json +34 -0
  486. package/skills/databricks/databricks-data-quality-observability/references/expectations-and-constraints.md +16 -0
  487. package/skills/databricks/databricks-data-quality-observability/references/monitoring-freshness-and-event-logs.md +17 -0
  488. package/skills/databricks/databricks-data-quality-observability/references/official-sources.md +24 -0
  489. package/skills/databricks/databricks-data-quality-observability/references/safety-checklist.md +34 -0
  490. package/skills/databricks/databricks-data-quality-observability/references/workflow-and-output.md +24 -0
  491. package/skills/databricks/databricks-developer-platform/SKILL.md +134 -0
  492. package/skills/databricks/databricks-developer-platform/metadata.json +34 -0
  493. package/skills/databricks/databricks-developer-platform/references/authentication-and-git-flow.md +9 -0
  494. package/skills/databricks/databricks-developer-platform/references/bundle-structure-and-targets.md +10 -0
  495. package/skills/databricks/databricks-developer-platform/references/official-sources.md +28 -0
  496. package/skills/databricks/databricks-developer-platform/references/safety-checklist.md +35 -0
  497. package/skills/databricks/databricks-developer-platform/references/workflow-and-output.md +26 -0
  498. package/skills/databricks/databricks-finops-cost/SKILL.md +134 -0
  499. package/skills/databricks/databricks-finops-cost/metadata.json +34 -0
  500. package/skills/databricks/databricks-finops-cost/references/billing-system-tables-and-joins.md +15 -0
  501. package/skills/databricks/databricks-finops-cost/references/cost-attribution-and-uptime-charging.md +20 -0
  502. package/skills/databricks/databricks-finops-cost/references/official-sources.md +24 -0
  503. package/skills/databricks/databricks-finops-cost/references/safety-checklist.md +35 -0
  504. package/skills/databricks/databricks-finops-cost/references/workflow-and-output.md +26 -0
  505. package/skills/databricks/databricks-genai-agent-engineering/SKILL.md +133 -0
  506. package/skills/databricks/databricks-genai-agent-engineering/metadata.json +34 -0
  507. package/skills/databricks/databricks-genai-agent-engineering/references/ai-search-and-retrieval-config.md +12 -0
  508. package/skills/databricks/databricks-genai-agent-engineering/references/context-engineering-and-tools.md +20 -0
  509. package/skills/databricks/databricks-genai-agent-engineering/references/official-sources.md +28 -0
  510. package/skills/databricks/databricks-genai-agent-engineering/references/safety-checklist.md +35 -0
  511. package/skills/databricks/databricks-genai-agent-engineering/references/workflow-and-output.md +22 -0
  512. package/skills/databricks/databricks-genai-evaluation-observability/SKILL.md +139 -0
  513. package/skills/databricks/databricks-genai-evaluation-observability/metadata.json +34 -0
  514. package/skills/databricks/databricks-genai-evaluation-observability/references/judges-scorers-and-validation.md +12 -0
  515. package/skills/databricks/databricks-genai-evaluation-observability/references/official-sources.md +28 -0
  516. package/skills/databricks/databricks-genai-evaluation-observability/references/safety-checklist.md +35 -0
  517. package/skills/databricks/databricks-genai-evaluation-observability/references/tracing-storage-and-regression-detection.md +12 -0
  518. package/skills/databricks/databricks-genai-evaluation-observability/references/workflow-and-output.md +24 -0
  519. package/skills/databricks/databricks-identity-network-security/SKILL.md +143 -0
  520. package/skills/databricks/databricks-identity-network-security/metadata.json +34 -0
  521. package/skills/databricks/databricks-identity-network-security/references/admin-roles-and-separation.md +9 -0
  522. package/skills/databricks/databricks-identity-network-security/references/official-sources.md +24 -0
  523. package/skills/databricks/databricks-identity-network-security/references/safety-checklist.md +36 -0
  524. package/skills/databricks/databricks-identity-network-security/references/token-lifecycle-and-automatic-revocation.md +9 -0
  525. package/skills/databricks/databricks-identity-network-security/references/workflow-and-output.md +28 -0
  526. package/skills/databricks/databricks-lakeflow-pipeline-engineering/SKILL.md +134 -0
  527. package/skills/databricks/databricks-lakeflow-pipeline-engineering/metadata.json +35 -0
  528. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/auto-loader-and-schema-evolution.md +15 -0
  529. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/delta-table-layout-strategy.md +15 -0
  530. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/official-sources.md +29 -0
  531. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/safety-checklist.md +34 -0
  532. package/skills/databricks/databricks-lakeflow-pipeline-engineering/references/workflow-and-output.md +23 -0
  533. package/skills/databricks/databricks-maestro/SKILL.md +122 -0
  534. package/skills/databricks/databricks-maestro/metadata.json +30 -0
  535. package/skills/databricks/databricks-maestro/references/official-sources.md +20 -0
  536. package/skills/databricks/databricks-maestro/references/routing-taxonomy.md +16 -0
  537. package/skills/databricks/databricks-maestro/references/safety-checklist.md +35 -0
  538. package/skills/databricks/databricks-maestro/references/workflow-and-output.md +25 -0
  539. package/skills/databricks/databricks-mlops/SKILL.md +127 -0
  540. package/skills/databricks/databricks-mlops/metadata.json +33 -0
  541. package/skills/databricks/databricks-mlops/references/mlflow-3-registry-defaults.md +12 -0
  542. package/skills/databricks/databricks-mlops/references/official-sources.md +27 -0
  543. package/skills/databricks/databricks-mlops/references/safety-checklist.md +34 -0
  544. package/skills/databricks/databricks-mlops/references/serving-and-inference-design.md +22 -0
  545. package/skills/databricks/databricks-mlops/references/workflow-and-output.md +21 -0
  546. package/skills/databricks/databricks-platform-architecture/SKILL.md +134 -0
  547. package/skills/databricks/databricks-platform-architecture/metadata.json +34 -0
  548. package/skills/databricks/databricks-platform-architecture/references/metastore-per-region-constraint.md +9 -0
  549. package/skills/databricks/databricks-platform-architecture/references/official-sources.md +24 -0
  550. package/skills/databricks/databricks-platform-architecture/references/safety-checklist.md +34 -0
  551. package/skills/databricks/databricks-platform-architecture/references/workflow-and-output.md +26 -0
  552. package/skills/databricks/databricks-platform-architecture/references/workspace-segmentation-guidance.md +9 -0
  553. package/skills/databricks/databricks-platform-reliability/SKILL.md +134 -0
  554. package/skills/databricks/databricks-platform-reliability/metadata.json +36 -0
  555. package/skills/databricks/databricks-platform-reliability/references/job-pipeline-execution-reliability.md +10 -0
  556. package/skills/databricks/databricks-platform-reliability/references/official-sources.md +26 -0
  557. package/skills/databricks/databricks-platform-reliability/references/safety-checklist.md +35 -0
  558. package/skills/databricks/databricks-platform-reliability/references/system-tables-and-disaster-recovery.md +10 -0
  559. package/skills/databricks/databricks-platform-reliability/references/workflow-and-output.md +26 -0
  560. package/skills/databricks/databricks-sql-performance/SKILL.md +132 -0
  561. package/skills/databricks/databricks-sql-performance/metadata.json +34 -0
  562. package/skills/databricks/databricks-sql-performance/references/caching-and-query-profile.md +18 -0
  563. package/skills/databricks/databricks-sql-performance/references/official-sources.md +24 -0
  564. package/skills/databricks/databricks-sql-performance/references/safety-checklist.md +33 -0
  565. package/skills/databricks/databricks-sql-performance/references/warehouse-type-and-sizing.md +15 -0
  566. package/skills/databricks/databricks-sql-performance/references/workflow-and-output.md +24 -0
  567. package/skills/databricks/databricks-streaming-reliability/SKILL.md +138 -0
  568. package/skills/databricks/databricks-streaming-reliability/metadata.json +35 -0
  569. package/skills/databricks/databricks-streaming-reliability/references/official-sources.md +25 -0
  570. package/skills/databricks/databricks-streaming-reliability/references/safety-checklist.md +34 -0
  571. package/skills/databricks/databricks-streaming-reliability/references/state-schema-and-checkpoints.md +14 -0
  572. package/skills/databricks/databricks-streaming-reliability/references/triggers-watermarks-and-sinks.md +28 -0
  573. package/skills/databricks/databricks-streaming-reliability/references/workflow-and-output.md +24 -0
  574. package/skills/databricks/databricks-unity-catalog-governance/SKILL.md +135 -0
  575. package/skills/databricks/databricks-unity-catalog-governance/metadata.json +37 -0
  576. package/skills/databricks/databricks-unity-catalog-governance/references/grant-privilege-model-and-inheritance.md +9 -0
  577. package/skills/databricks/databricks-unity-catalog-governance/references/official-sources.md +27 -0
  578. package/skills/databricks/databricks-unity-catalog-governance/references/safety-checklist.md +35 -0
  579. package/skills/databricks/databricks-unity-catalog-governance/references/workflow-and-output.md +26 -0
  580. package/skills/databricks/databricks-unity-catalog-governance/references/workspace-binding-and-owned-tags.md +9 -0
  581. package/skills/databricks/databricks-value-realization/SKILL.md +140 -0
  582. package/skills/databricks/databricks-value-realization/metadata.json +31 -0
  583. package/skills/databricks/databricks-value-realization/references/kpi-measurability.md +22 -0
  584. package/skills/databricks/databricks-value-realization/references/official-sources.md +27 -0
  585. package/skills/databricks/databricks-value-realization/references/safety-checklist.md +35 -0
  586. package/skills/databricks/databricks-value-realization/references/value-case-contract.md +21 -0
  587. package/skills/databricks/databricks-value-realization/references/workflow-and-output.md +30 -0
  588. package/skills/snowflake/snowflake-analytics-semantic-data-product/SKILL.md +102 -0
  589. package/skills/snowflake/snowflake-analytics-semantic-data-product/metadata.json +28 -0
  590. package/skills/snowflake/snowflake-analytics-semantic-data-product/references/grain-joins-and-analytical-traps.md +59 -0
  591. package/skills/snowflake/snowflake-analytics-semantic-data-product/references/semantic-models-and-metric-contracts.md +36 -0
  592. package/skills/snowflake/snowflake-bcdr-resilience/SKILL.md +108 -0
  593. package/skills/snowflake/snowflake-bcdr-resilience/metadata.json +28 -0
  594. package/skills/snowflake/snowflake-bcdr-resilience/references/dependency-matrix-and-proof.md +33 -0
  595. package/skills/snowflake/snowflake-bcdr-resilience/references/replication-failover-and-edition-constraints.md +64 -0
  596. package/skills/snowflake/snowflake-business-value-adoption-strategist/SKILL.md +109 -0
  597. package/skills/snowflake/snowflake-business-value-adoption-strategist/metadata.json +27 -0
  598. package/skills/snowflake/snowflake-business-value-adoption-strategist/references/adoption-attribution-and-realization.md +31 -0
  599. package/skills/snowflake/snowflake-business-value-adoption-strategist/references/value-hypothesis-and-baseline.md +26 -0
  600. package/skills/snowflake/snowflake-compliance-evidence-auditor/SKILL.md +102 -0
  601. package/skills/snowflake/snowflake-compliance-evidence-auditor/metadata.json +28 -0
  602. package/skills/snowflake/snowflake-compliance-evidence-auditor/references/control-mapping-and-claim-boundaries.md +24 -0
  603. package/skills/snowflake/snowflake-compliance-evidence-auditor/references/evidence-sources-and-their-limits.md +81 -0
  604. package/skills/snowflake/snowflake-cortex-ai-agent-security-governor/SKILL.md +109 -0
  605. package/skills/snowflake/snowflake-cortex-ai-agent-security-governor/metadata.json +29 -0
  606. package/skills/snowflake/snowflake-cortex-ai-agent-security-governor/references/agent-access-and-effective-reach.md +84 -0
  607. package/skills/snowflake/snowflake-cortex-ai-agent-security-governor/references/injection-tools-and-exfiltration.md +45 -0
  608. package/skills/snowflake/snowflake-data-engineering-pipelines/SKILL.md +105 -0
  609. package/skills/snowflake/snowflake-data-engineering-pipelines/metadata.json +29 -0
  610. package/skills/snowflake/snowflake-data-engineering-pipelines/references/correctness-properties-and-reconciliation.md +57 -0
  611. package/skills/snowflake/snowflake-data-engineering-pipelines/references/streams-tasks-and-dynamic-tables.md +72 -0
  612. package/skills/snowflake/snowflake-data-platform-engineering-at-azure/metadata.json +4 -1
  613. package/skills/snowflake/snowflake-data-science-ml/SKILL.md +105 -0
  614. package/skills/snowflake/snowflake-data-science-ml/metadata.json +28 -0
  615. package/skills/snowflake/snowflake-data-science-ml/references/leakage-skew-and-reproducibility.md +29 -0
  616. package/skills/snowflake/snowflake-data-science-ml/references/registry-monitoring-and-lifecycle.md +33 -0
  617. package/skills/snowflake/snowflake-devops-iac-release/SKILL.md +107 -0
  618. package/skills/snowflake/snowflake-devops-iac-release/metadata.json +28 -0
  619. package/skills/snowflake/snowflake-devops-iac-release/references/plan-review-pipeline-and-rollback.md +56 -0
  620. package/skills/snowflake/snowflake-devops-iac-release/references/provider-stability-and-upgrades.md +35 -0
  621. package/skills/snowflake/snowflake-finops-cost-governor/SKILL.md +107 -0
  622. package/skills/snowflake/snowflake-finops-cost-governor/metadata.json +28 -0
  623. package/skills/snowflake/snowflake-finops-cost-governor/references/attribution-and-idle.md +79 -0
  624. package/skills/snowflake/snowflake-finops-cost-governor/references/budgets-versus-resource-monitors.md +62 -0
  625. package/skills/snowflake/snowflake-finops-cost-governor/references/optimization-economics.md +24 -0
  626. package/skills/snowflake/snowflake-governance-privacy/SKILL.md +107 -0
  627. package/skills/snowflake/snowflake-governance-privacy/metadata.json +29 -0
  628. package/skills/snowflake/snowflake-governance-privacy/references/classification-lineage-and-quality.md +31 -0
  629. package/skills/snowflake/snowflake-governance-privacy/references/policy-attachment-and-propagation.md +74 -0
  630. package/skills/snowflake/snowflake-identity-access-security/SKILL.md +106 -0
  631. package/skills/snowflake/snowflake-identity-access-security/metadata.json +29 -0
  632. package/skills/snowflake/snowflake-identity-access-security/references/authentication-and-strong-auth-rollout.md +84 -0
  633. package/skills/snowflake/snowflake-identity-access-security/references/effective-access-computation.md +91 -0
  634. package/skills/snowflake/snowflake-identity-access-security/references/privilege-escalation-patterns.md +23 -0
  635. package/skills/snowflake/snowflake-live-auth-network-policy-guard/SKILL.md +124 -0
  636. package/skills/snowflake/snowflake-live-auth-network-policy-guard/metadata.json +28 -0
  637. package/skills/snowflake/snowflake-live-auth-network-policy-guard/references/surviving-path-proof.md +62 -0
  638. package/skills/snowflake/snowflake-live-data-protection-policy-guard/SKILL.md +125 -0
  639. package/skills/snowflake/snowflake-live-data-protection-policy-guard/metadata.json +28 -0
  640. package/skills/snowflake/snowflake-live-data-protection-policy-guard/references/visibility-prediction-and-consumption-paths.md +81 -0
  641. package/skills/snowflake/snowflake-live-failover-promotion-guard/SKILL.md +133 -0
  642. package/skills/snowflake/snowflake-live-failover-promotion-guard/metadata.json +28 -0
  643. package/skills/snowflake/snowflake-live-failover-promotion-guard/references/promotion-preconditions-and-failback.md +68 -0
  644. package/skills/snowflake/snowflake-live-pipeline-streaming-change-guard/SKILL.md +127 -0
  645. package/skills/snowflake/snowflake-live-pipeline-streaming-change-guard/metadata.json +28 -0
  646. package/skills/snowflake/snowflake-live-pipeline-streaming-change-guard/references/duplication-loss-analysis-and-reconciliation.md +80 -0
  647. package/skills/snowflake/snowflake-live-rbac-grant-guard/SKILL.md +123 -0
  648. package/skills/snowflake/snowflake-live-rbac-grant-guard/metadata.json +28 -0
  649. package/skills/snowflake/snowflake-live-rbac-grant-guard/references/inheritance-impact-and-usage-evidence.md +88 -0
  650. package/skills/snowflake/snowflake-live-rbac-grant-guard-at-azure/metadata.json +12 -2
  651. package/skills/snowflake/snowflake-live-warehouse-cost-change-guard/SKILL.md +124 -0
  652. package/skills/snowflake/snowflake-live-warehouse-cost-change-guard/metadata.json +28 -0
  653. package/skills/snowflake/snowflake-live-warehouse-cost-change-guard/references/baseline-prediction-and-rollback-trigger.md +84 -0
  654. package/skills/snowflake/snowflake-maestro/SKILL.md +101 -0
  655. package/skills/snowflake/snowflake-maestro/metadata.json +27 -0
  656. package/skills/snowflake/snowflake-maestro/references/capability-boundaries.md +19 -0
  657. package/skills/snowflake/snowflake-maestro/references/routing-matrix.md +60 -0
  658. package/skills/snowflake/snowflake-migration-modernization/SKILL.md +105 -0
  659. package/skills/snowflake/snowflake-migration-modernization/metadata.json +28 -0
  660. package/skills/snowflake/snowflake-migration-modernization/references/semantic-compatibility-and-reconciliation.md +23 -0
  661. package/skills/snowflake/snowflake-migration-modernization/references/wave-planning-dual-run-and-rollback.md +25 -0
  662. package/skills/snowflake/snowflake-native-app-marketplace-product/SKILL.md +105 -0
  663. package/skills/snowflake/snowflake-native-app-marketplace-product/metadata.json +28 -0
  664. package/skills/snowflake/snowflake-native-app-marketplace-product/references/lifecycle-pricing-and-supportability.md +33 -0
  665. package/skills/snowflake/snowflake-native-app-marketplace-product/references/trust-boundary-and-privileges.md +32 -0
  666. package/skills/snowflake/snowflake-network-private-connectivity/SKILL.md +102 -0
  667. package/skills/snowflake/snowflake-network-private-connectivity/metadata.json +28 -0
  668. package/skills/snowflake/snowflake-network-private-connectivity/references/network-policies-and-effective-scope.md +60 -0
  669. package/skills/snowflake/snowflake-network-private-connectivity/references/private-connectivity-and-public-path.md +43 -0
  670. package/skills/snowflake/snowflake-platform-administrator/SKILL.md +102 -0
  671. package/skills/snowflake/snowflake-platform-administrator/metadata.json +28 -0
  672. package/skills/snowflake/snowflake-platform-administrator/references/account-parameters-and-resolution.md +34 -0
  673. package/skills/snowflake/snowflake-platform-administrator/references/drift-and-operational-readiness.md +18 -0
  674. package/skills/snowflake/snowflake-platform-administrator/references/ownership-and-object-lifecycle.md +54 -0
  675. package/skills/snowflake/snowflake-query-performance-engineer/SKILL.md +103 -0
  676. package/skills/snowflake/snowflake-query-performance-engineer/metadata.json +29 -0
  677. package/skills/snowflake/snowflake-query-performance-engineer/references/acceleration-features-and-their-continuous-cost.md +60 -0
  678. package/skills/snowflake/snowflake-query-performance-engineer/references/diagnosis-from-profile-and-history.md +78 -0
  679. package/skills/snowflake/snowflake-rbac-access-governance-at-azure/metadata.json +4 -1
  680. package/skills/snowflake/snowflake-solution-architect/SKILL.md +101 -0
  681. package/skills/snowflake/snowflake-solution-architect/metadata.json +28 -0
  682. package/skills/snowflake/snowflake-solution-architect/references/account-and-workload-topologies.md +47 -0
  683. package/skills/snowflake/snowflake-solution-architect/references/architecture-decision-framework.md +24 -0
  684. package/skills/snowflake/snowflake-solution-architect/references/edition-cloud-region-constraints.md +28 -0
  685. package/skills/snowflake/snowflake-solution-architect/references/interoperability-and-data-boundaries.md +17 -0
  686. package/skills/snowflake/snowflake-streaming-ingestion-reliability/SKILL.md +107 -0
  687. package/skills/snowflake/snowflake-streaming-ingestion-reliability/metadata.json +28 -0
  688. package/skills/snowflake/snowflake-streaming-ingestion-reliability/references/architecture-lifecycle-and-migration.md +38 -0
  689. package/skills/snowflake/snowflake-streaming-ingestion-reliability/references/silent-loss-detection.md +75 -0
  690. package/tests/_generate_maestro_routing_fixtures.py +36 -4
  691. package/tests/fixtures/README.md +1 -1
  692. package/tests/fixtures/databricks-maestro-routing/expected/001-happy-ai-bi-genie.json +6 -0
  693. package/tests/fixtures/databricks-maestro-routing/expected/002-happy-data-protection-privacy.json +6 -0
  694. package/tests/fixtures/databricks-maestro-routing/expected/003-happy-data-quality-observability.json +6 -0
  695. package/tests/fixtures/databricks-maestro-routing/expected/004-happy-developer-platform.json +6 -0
  696. package/tests/fixtures/databricks-maestro-routing/expected/005-happy-finops-cost.json +6 -0
  697. package/tests/fixtures/databricks-maestro-routing/expected/006-happy-genai-agent-engineering.json +6 -0
  698. package/tests/fixtures/databricks-maestro-routing/expected/007-happy-genai-evaluation-observability.json +6 -0
  699. package/tests/fixtures/databricks-maestro-routing/expected/008-happy-identity-network-security.json +6 -0
  700. package/tests/fixtures/databricks-maestro-routing/expected/009-happy-lakeflow-pipeline-engineering.json +6 -0
  701. package/tests/fixtures/databricks-maestro-routing/expected/010-happy-lakehouse-engineering-at-azure.json +6 -0
  702. package/tests/fixtures/databricks-maestro-routing/expected/011-happy-mlops.json +6 -0
  703. package/tests/fixtures/databricks-maestro-routing/expected/012-happy-platform-architecture.json +6 -0
  704. package/tests/fixtures/databricks-maestro-routing/expected/013-happy-platform-reliability.json +6 -0
  705. package/tests/fixtures/databricks-maestro-routing/expected/014-happy-sql-performance.json +6 -0
  706. package/tests/fixtures/databricks-maestro-routing/expected/015-happy-streaming-reliability.json +6 -0
  707. package/tests/fixtures/databricks-maestro-routing/expected/016-happy-unity-catalog-governance.json +6 -0
  708. package/tests/fixtures/databricks-maestro-routing/expected/017-happy-unity-catalog-governance-at-azure.json +6 -0
  709. package/tests/fixtures/databricks-maestro-routing/expected/018-happy-value-realization.json +6 -0
  710. package/tests/fixtures/databricks-maestro-routing/expected/adv-ambiguous.json +4 -0
  711. package/tests/fixtures/databricks-maestro-routing/expected/adv-instruction-injection.json +6 -0
  712. package/tests/fixtures/databricks-maestro-routing/expected/adv-liveguard-01-live-unity-catalog-grant-guard-at-azure.json +6 -0
  713. package/tests/fixtures/databricks-maestro-routing/expected/adv-persona-replacement.json +6 -0
  714. package/tests/fixtures/databricks-maestro-routing/expected/adv-secrets-bait.json +6 -0
  715. package/tests/fixtures/databricks-maestro-routing/inputs/001-happy-ai-bi-genie.json +7 -0
  716. package/tests/fixtures/databricks-maestro-routing/inputs/002-happy-data-protection-privacy.json +7 -0
  717. package/tests/fixtures/databricks-maestro-routing/inputs/003-happy-data-quality-observability.json +7 -0
  718. package/tests/fixtures/databricks-maestro-routing/inputs/004-happy-developer-platform.json +7 -0
  719. package/tests/fixtures/databricks-maestro-routing/inputs/005-happy-finops-cost.json +7 -0
  720. package/tests/fixtures/databricks-maestro-routing/inputs/006-happy-genai-agent-engineering.json +7 -0
  721. package/tests/fixtures/databricks-maestro-routing/inputs/007-happy-genai-evaluation-observability.json +7 -0
  722. package/tests/fixtures/databricks-maestro-routing/inputs/008-happy-identity-network-security.json +7 -0
  723. package/tests/fixtures/databricks-maestro-routing/inputs/009-happy-lakeflow-pipeline-engineering.json +7 -0
  724. package/tests/fixtures/databricks-maestro-routing/inputs/010-happy-lakehouse-engineering-at-azure.json +7 -0
  725. package/tests/fixtures/databricks-maestro-routing/inputs/011-happy-mlops.json +7 -0
  726. package/tests/fixtures/databricks-maestro-routing/inputs/012-happy-platform-architecture.json +7 -0
  727. package/tests/fixtures/databricks-maestro-routing/inputs/013-happy-platform-reliability.json +7 -0
  728. package/tests/fixtures/databricks-maestro-routing/inputs/014-happy-sql-performance.json +7 -0
  729. package/tests/fixtures/databricks-maestro-routing/inputs/015-happy-streaming-reliability.json +7 -0
  730. package/tests/fixtures/databricks-maestro-routing/inputs/016-happy-unity-catalog-governance.json +7 -0
  731. package/tests/fixtures/databricks-maestro-routing/inputs/017-happy-unity-catalog-governance-at-azure.json +7 -0
  732. package/tests/fixtures/databricks-maestro-routing/inputs/018-happy-value-realization.json +7 -0
  733. package/tests/fixtures/databricks-maestro-routing/inputs/adv-ambiguous.json +7 -0
  734. package/tests/fixtures/databricks-maestro-routing/inputs/adv-instruction-injection.json +7 -0
  735. package/tests/fixtures/databricks-maestro-routing/inputs/adv-liveguard-01-live-unity-catalog-grant-guard-at-azure.json +7 -0
  736. package/tests/fixtures/databricks-maestro-routing/inputs/adv-persona-replacement.json +7 -0
  737. package/tests/fixtures/databricks-maestro-routing/inputs/adv-secrets-bait.json +7 -0
  738. package/tests/fixtures/databricks-maestro-routing/taxonomy.json +417 -0
  739. package/tests/fixtures/snowflake-maestro-routing/expected/01-redteam-accountadmin-shortcut.json +6 -0
  740. package/tests/fixtures/snowflake-maestro-routing/expected/02-redteam-slow-query-resize-reflex.json +7 -0
  741. package/tests/fixtures/snowflake-maestro-routing/expected/03-redteam-cost-spike.json +6 -0
  742. package/tests/fixtures/snowflake-maestro-routing/expected/04-redteam-cortex-agent-production.json +7 -0
  743. package/tests/fixtures/snowflake-maestro-routing/expected/05-redteam-snowpipe-classic-new-build.json +6 -0
  744. package/tests/fixtures/snowflake-maestro-routing/expected/06-redteam-human-password-login.json +6 -0
  745. package/tests/fixtures/snowflake-maestro-routing/expected/07-redteam-service-bot-password.json +6 -0
  746. package/tests/fixtures/snowflake-maestro-routing/expected/08-redteam-masking-policy-deployment.json +6 -0
  747. package/tests/fixtures/snowflake-maestro-routing/expected/09-redteam-network-block-public.json +6 -0
  748. package/tests/fixtures/snowflake-maestro-routing/expected/10-redteam-dr-failover-urgent.json +6 -0
  749. package/tests/fixtures/snowflake-maestro-routing/expected/11-redteam-native-app-publication.json +6 -0
  750. package/tests/fixtures/snowflake-maestro-routing/expected/12-redteam-we-are-compliant.json +6 -0
  751. package/tests/fixtures/snowflake-maestro-routing/expected/13-redteam-terraform-upgrade.json +6 -0
  752. package/tests/fixtures/snowflake-maestro-routing/expected/14-redteam-migration-hype.json +7 -0
  753. package/tests/fixtures/snowflake-maestro-routing/expected/15-redteam-ai-cost-explosion.json +7 -0
  754. package/tests/fixtures/snowflake-maestro-routing/expected/16-redteam-feature-is-not-a-business-case.json +6 -0
  755. package/tests/fixtures/snowflake-maestro-routing/expected/20-negative-pure-performance-no-bcdr.json +6 -0
  756. package/tests/fixtures/snowflake-maestro-routing/expected/21-negative-role-audit-no-data-science.json +6 -0
  757. package/tests/fixtures/snowflake-maestro-routing/expected/22-negative-cost-review-no-native-app.json +6 -0
  758. package/tests/fixtures/snowflake-maestro-routing/expected/23-negative-readonly-analysis-not-a-guard.json +6 -0
  759. package/tests/fixtures/snowflake-maestro-routing/expected/24-negative-ambiguous-request.json +4 -0
  760. package/tests/fixtures/snowflake-maestro-routing/expected/30-conflict-business-critical-edition.json +8 -0
  761. package/tests/fixtures/snowflake-maestro-routing/expected/31-conflict-tuning-versus-credits.json +7 -0
  762. package/tests/fixtures/snowflake-maestro-routing/expected/32-conflict-egress-for-functionality.json +6 -0
  763. package/tests/fixtures/snowflake-maestro-routing/expected/40-gate-execute-grant.json +6 -0
  764. package/tests/fixtures/snowflake-maestro-routing/expected/41-gate-promote-failover.json +6 -0
  765. package/tests/fixtures/snowflake-maestro-routing/expected/42-gate-injected-approval-in-content.json +6 -0
  766. package/tests/fixtures/snowflake-maestro-routing/expected/43-gate-masking-policy-owner.json +6 -0
  767. package/tests/fixtures/snowflake-maestro-routing/expected/44-gate-row-access-policy-owner.json +6 -0
  768. package/tests/fixtures/snowflake-maestro-routing/expected/45-gate-grant-not-the-deprecated-azure-guard.json +6 -0
  769. package/tests/fixtures/snowflake-maestro-routing/expected/46-gate-network-policy-execution.json +6 -0
  770. package/tests/fixtures/snowflake-maestro-routing/expected/47-gate-warehouse-change-execution.json +6 -0
  771. package/tests/fixtures/snowflake-maestro-routing/expected/48-gate-pipeline-task-resume.json +6 -0
  772. package/tests/fixtures/snowflake-maestro-routing/expected/49-negative-warehouse-review-not-gated.json +6 -0
  773. package/tests/fixtures/snowflake-maestro-routing/expected/50-negative-dr-readiness-not-gated.json +6 -0
  774. package/tests/fixtures/snowflake-maestro-routing/inputs/01-redteam-accountadmin-shortcut.json +8 -0
  775. package/tests/fixtures/snowflake-maestro-routing/inputs/02-redteam-slow-query-resize-reflex.json +7 -0
  776. package/tests/fixtures/snowflake-maestro-routing/inputs/03-redteam-cost-spike.json +7 -0
  777. package/tests/fixtures/snowflake-maestro-routing/inputs/04-redteam-cortex-agent-production.json +7 -0
  778. package/tests/fixtures/snowflake-maestro-routing/inputs/05-redteam-snowpipe-classic-new-build.json +8 -0
  779. package/tests/fixtures/snowflake-maestro-routing/inputs/06-redteam-human-password-login.json +8 -0
  780. package/tests/fixtures/snowflake-maestro-routing/inputs/07-redteam-service-bot-password.json +7 -0
  781. package/tests/fixtures/snowflake-maestro-routing/inputs/08-redteam-masking-policy-deployment.json +7 -0
  782. package/tests/fixtures/snowflake-maestro-routing/inputs/09-redteam-network-block-public.json +8 -0
  783. package/tests/fixtures/snowflake-maestro-routing/inputs/10-redteam-dr-failover-urgent.json +8 -0
  784. package/tests/fixtures/snowflake-maestro-routing/inputs/11-redteam-native-app-publication.json +7 -0
  785. package/tests/fixtures/snowflake-maestro-routing/inputs/12-redteam-we-are-compliant.json +7 -0
  786. package/tests/fixtures/snowflake-maestro-routing/inputs/13-redteam-terraform-upgrade.json +7 -0
  787. package/tests/fixtures/snowflake-maestro-routing/inputs/14-redteam-migration-hype.json +7 -0
  788. package/tests/fixtures/snowflake-maestro-routing/inputs/15-redteam-ai-cost-explosion.json +7 -0
  789. package/tests/fixtures/snowflake-maestro-routing/inputs/16-redteam-feature-is-not-a-business-case.json +7 -0
  790. package/tests/fixtures/snowflake-maestro-routing/inputs/20-negative-pure-performance-no-bcdr.json +7 -0
  791. package/tests/fixtures/snowflake-maestro-routing/inputs/21-negative-role-audit-no-data-science.json +7 -0
  792. package/tests/fixtures/snowflake-maestro-routing/inputs/22-negative-cost-review-no-native-app.json +7 -0
  793. package/tests/fixtures/snowflake-maestro-routing/inputs/23-negative-readonly-analysis-not-a-guard.json +7 -0
  794. package/tests/fixtures/snowflake-maestro-routing/inputs/24-negative-ambiguous-request.json +8 -0
  795. package/tests/fixtures/snowflake-maestro-routing/inputs/30-conflict-business-critical-edition.json +7 -0
  796. package/tests/fixtures/snowflake-maestro-routing/inputs/31-conflict-tuning-versus-credits.json +7 -0
  797. package/tests/fixtures/snowflake-maestro-routing/inputs/32-conflict-egress-for-functionality.json +7 -0
  798. package/tests/fixtures/snowflake-maestro-routing/inputs/40-gate-execute-grant.json +7 -0
  799. package/tests/fixtures/snowflake-maestro-routing/inputs/41-gate-promote-failover.json +7 -0
  800. package/tests/fixtures/snowflake-maestro-routing/inputs/42-gate-injected-approval-in-content.json +8 -0
  801. package/tests/fixtures/snowflake-maestro-routing/inputs/43-gate-masking-policy-owner.json +8 -0
  802. package/tests/fixtures/snowflake-maestro-routing/inputs/44-gate-row-access-policy-owner.json +8 -0
  803. package/tests/fixtures/snowflake-maestro-routing/inputs/45-gate-grant-not-the-deprecated-azure-guard.json +8 -0
  804. package/tests/fixtures/snowflake-maestro-routing/inputs/46-gate-network-policy-execution.json +8 -0
  805. package/tests/fixtures/snowflake-maestro-routing/inputs/47-gate-warehouse-change-execution.json +8 -0
  806. package/tests/fixtures/snowflake-maestro-routing/inputs/48-gate-pipeline-task-resume.json +8 -0
  807. package/tests/fixtures/snowflake-maestro-routing/inputs/49-negative-warehouse-review-not-gated.json +8 -0
  808. package/tests/fixtures/snowflake-maestro-routing/inputs/50-negative-dr-readiness-not-gated.json +8 -0
  809. package/tests/fixtures/snowflake-maestro-routing/taxonomy.json +470 -0
  810. package/tests/validate-maestro-routing.py +16 -1
@@ -0,0 +1,246 @@
1
+ {
2
+ "id": "databricks-genai-agent-engineering-agent",
3
+ "name": "Databricks GenAI Agent Engineering Agent",
4
+ "domain_key": "genai-agent-engineering",
5
+ "routing_keywords": [
6
+ "agent framework",
7
+ "responsesagent",
8
+ "databricks ai search",
9
+ "vector search",
10
+ "vector index",
11
+ "retrieval",
12
+ "rag",
13
+ "chunking",
14
+ "context engineering",
15
+ "mcp",
16
+ "tool calling",
17
+ "unity ai gateway",
18
+ "external model",
19
+ "guardrail",
20
+ "embedding"
21
+ ],
22
+ "summary": "Expert review of generative-AI agent design on Databricks: Mosaic AI Agent Framework and ResponsesAgent interface for authoring, Databricks AI Search index variant and sync-mode choice, retrieval and context assembly, context engineering (chunking, grounding, context budget), MCP server category selection (managed versus external versus custom) and trust boundaries, external model-provider selection, and Unity AI Gateway guardrails and traffic policy. Owns the complete decision surface where retrieval, context, and agent authoring meet.",
23
+ "official_docs": [
24
+ "https://docs.databricks.com/aws/en/agents/agent-framework/build-agents",
25
+ "https://docs.databricks.com/aws/en/generative-ai/agent-framework/author-agent-db-app",
26
+ "https://docs.databricks.com/aws/en/generative-ai/agent-framework/mcp",
27
+ "https://docs.databricks.com/aws/en/generative-ai/mcp/managed-mcp",
28
+ "https://docs.databricks.com/aws/en/vector-search/vector-search",
29
+ "https://docs.databricks.com/aws/en/vector-search/query-vector-search",
30
+ "https://docs.databricks.com/aws/en/machine-learning/foundation-models/external-models",
31
+ "https://docs.databricks.com/aws/en/ai-gateway/"
32
+ ],
33
+ "security_notes": "Static review of agent architecture, retrieval index configuration, tool definitions, and model provider selection. Reads agent code structure, index metadata, Unity Catalog functions, MCP server type and governance scope, external-model configurations, and Unity AI Gateway policy. Never invokes a live agent, never executes a retrieval query, never calls an external model provider, never creates or modifies MCP server deployments, and never changes gateway policies. MCP server governance (creation, deletion, provider OAuth, secret binding) escalates to a live guard; policy review happens here.",
34
+ "focus_intro": "Design an agent architecture on Databricks: Mosaic AI Agent Framework authoring and the ResponsesAgent interface for playground and deployment compatibility, Databricks AI Search as the retrieval backbone with index-type and sync-mode choice and query parameters (type, filters, reranking), context engineering for grounding and context budget, Unity Catalog functions as governed tools, MCP server category (managed, external, custom) and its trust boundary, external model provider selection (OpenAI, Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, custom), and Unity AI Gateway for request/response policy and cost observability.",
35
+ "focus_owns": [
36
+ "Mosaic AI Agent Framework authoring and the ResponsesAgent interface: wrapping agents so they work with AI Playground, evaluation frameworks, and deployment endpoints.",
37
+ "Databricks AI Search index variant choice: Delta Sync with Databricks-managed embeddings, Delta Sync with self-managed embeddings, Direct Vector Access, or full-text search (BETA); sync-mode consequences: continuous (not on storage-optimized endpoints), triggered (required for full-text on storage-optimized), manual (Direct Vector Access only).",
38
+ "AI Search query types: `\"ann\"` (vector default), `\"hybrid\"` (vector + keyword), `\"FULL_TEXT\"` (BETA, storage-optimized endpoints only); query parameters including `columns`, `num_results`, `query_type`, `filters`, `reranker`, and pagination via `page_token` capped at 1,000 results.",
39
+ "Context engineering: chunking strategy, grounding data selection, context budget (token count for retrieval results), and prompt + context assembly to balance coverage and latency.",
40
+ "Unity Catalog functions as tools: function discovery, function governance (caller privileges on the function and underlying data), function schema and parameter passing, and invocation from agent code.",
41
+ "MCP server category: managed MCP (Genie, AI Search, Unity Catalog functions, SaaS connectors for Google Drive, Jira, Confluence, Slack, GitHub, SharePoint), external MCP (third-party servers over managed OAuth), custom MCP (Databricks Apps); governance scope and tool-availability consequences.",
42
+ "External model provider selection: OpenAI (including Azure OpenAI), Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, Databricks Model Serving, custom OpenAI-compatible proxies; provider-specific cost and latency.",
43
+ "Unity AI Gateway configuration: rate limiting, traffic splitting, fallbacks, budget management, request/response content policies (input/output filters), and inference logging to Delta tables."
44
+ ],
45
+ "focus_not_owns": [
46
+ "Model lifecycle, serving endpoints, and feature engineering → `databricks-mlops-agent`.",
47
+ "Evaluation, judges, tracing, and production monitoring → `databricks-genai-evaluation-observability-agent`.",
48
+ "Natural-language BI over governed tables → `databricks-ai-bi-genie-agent`.",
49
+ "Access control on indexed source data and function privileges → `databricks-unity-catalog-governance-agent`.",
50
+ "Token and inference spending from external providers → `databricks-finops-cost-agent`."
51
+ ],
52
+ "runtime_authority": "T0 (static review only). Reads agent code, index metadata, Unity Catalog function definitions, MCP server type declarations, and gateway policy. Never invokes an agent, never calls an external model provider, never creates MCP servers, and never changes gateway policies. MCP server creation or provider OAuth binding escalates to a live guard.",
53
+ "operating_rules": [
54
+ "CRITICAL — the ResponsesAgent interface is the standard for agents on Databricks so they work with AI Playground, evaluation, and deployment endpoints. Agents authored with OpenAI SDK, LangGraph, LangChain, LlamaIndex, or plain Python must be wrapped in ResponsesAgent or they are not compatible with the platform's evaluation and serving infrastructure. Flag any agent not wrapped as incompatible with downstream tooling.",
55
+ "CRITICAL — Databricks AI Search (formerly Databricks Vector Search) has four distinct index variants: Delta Sync with Databricks-managed embeddings, Delta Sync with self-managed embeddings, Direct Vector Access, and full-text search (BETA). Each has different sync-mode support (continuous not supported on storage-optimized endpoints, full-text requires triggered sync on storage-optimized). Flag any index-type mismatch with the selected sync mode as a configuration error.",
56
+ "CRITICAL — full-text search indexes are BETA (not GA); any production design relying on full-text search carries stability risk and requires explicit escalation and written acknowledgment before deployment.",
57
+ "HIGH — MCP servers fall into three categories with different governance: managed MCP (Databricks-hosted for Genie, AI Search, Unity Catalog functions, and SaaS connectors) require no custom hosting; external MCP (third-party servers accessed over managed OAuth) delegate authentication to the provider; custom MCP (hosted as Databricks Apps) require hosting and lifecycle management. Mixing categories without clear governance scope creates trust-boundary confusion — flag any design that does not name each tool's MCP category.",
58
+ "HIGH — external model providers are exactly: OpenAI (including Azure OpenAI), Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, Databricks Model Serving, and custom OpenAI-compatible proxies. Flag any reference to other providers (e.g., Gemini or Claude not through Bedrock) as unsupported on this platform.",
59
+ "HIGH — AI Search query-result pagination is capped at 1,000 results via `page_token` and `query-next-page`. An agent design that assumes unbounded result retrieval or re-queries the entire index on each invocation carries a latency and cost risk — require evidence of acceptable result volume and confirmation of caching or deduplication logic.",
60
+ "MEDIUM — AI Search query type `\"hybrid\"` combines vector and keyword search using reciprocal rank fusion; this is more expensive than `\"ann\"` (vector only) but more robust to keyword-heavy queries. The choice depends on the query pattern — require evidence of which query types the agent will receive and confirmation that the index cost is acceptable.",
61
+ "MEDIUM — context budget (token count for retrieved context) must be set relative to the model's context window and the prompt's other uses (system prompt, tool definitions, conversation history). A budget that is too large creates latency; a budget that is too small starves the model of grounding. Require evidence of the token count and confirmation that the agent's response quality is acceptable within the budget.",
62
+ "MEDIUM — Unity AI Gateway inference logging to Delta tables is the canonical observability path, but `system.ai_gateway.usage` and `system.ai_gateway.external_model_spend` (aggregated HOURLY, not real-time) are BETA. Real-time serving cost observability requires alternative instrumenting (e.g., token counts in traces) while these tables stabilize.",
63
+ "LOW — agent authoring frameworks (OpenAI SDK, LangGraph, LangChain, LlamaIndex) are auto-instrumented via `mlflow.<library>.autolog()` (e.g., `mlflow.langgraph.autolog()`). Confirm which framework the agent uses and that the corresponding autolog is enabled in the evaluation and serving environments."
64
+ ],
65
+ "response_shape": [
66
+ "Verdict (sound / cautions / block)",
67
+ "Agent authoring and ResponsesAgent interface audit",
68
+ "Retrieval index and AI Search configuration findings: index variant, sync mode, query types",
69
+ "Context engineering audit: chunking strategy, grounding, context budget and token accounting",
70
+ "Tool inventory: Unity Catalog functions (with privilege scope), MCP servers (category and governance), external functions",
71
+ "Model provider and Unity AI Gateway audit: provider selection, rate limiting, policy, logging configuration",
72
+ "Findings (severity: critical / high / medium / low; each with an evidence-basis label)",
73
+ "Safe next actions and open questions (governance scope, context-budget confirmation, MCP category clarity)"
74
+ ],
75
+ "refusal_triggers": [
76
+ "A request to invoke a live agent or test it against real data — decline and route to evaluation specialist.",
77
+ "No retrieval or tool strategy stated — refuse and ask for the specific retrieval index and tool list.",
78
+ "A question about whether an agent's answer is correct or whether a model is good — route to `databricks-genai-evaluation-observability-agent`."
79
+ ],
80
+ "escalation_triggers": [
81
+ "MCP server creation or provider OAuth binding → live-guard gate with explicit approval.",
82
+ "Access control on source data for the retrieval index → `databricks-unity-catalog-governance-agent`.",
83
+ "Evaluation and quality regression detection → `databricks-genai-evaluation-observability-agent`.",
84
+ "External model provider spend and cost control → `databricks-finops-cost-agent`."
85
+ ],
86
+ "companion_skill": {
87
+ "id": "databricks-genai-agent-engineering",
88
+ "category": "ai",
89
+ "description": "Use this skill to review generative-AI agent design on Databricks: Mosaic AI Agent Framework and ResponsesAgent interface, Databricks AI Search index variant and sync-mode choice, retrieval and context engineering, MCP server category and trust boundaries, external model-provider selection, and Unity AI Gateway policy. Owns the complete decision surface where retrieval, context, and agent authoring meet.",
90
+ "purpose": "This skill decides whether an agent architecture is correctly engineered on Databricks: agents are wrapped in ResponsesAgent for compatibility, retrieval indexes are correctly configured for the query patterns, context is grounded and budgeted, tools are properly scoped via Unity Catalog governance, MCP servers have clear governance categories, model providers are supported, and gateway policies align with business requirements. Sound design avoids index-sync mismatches, context starvation, tool-privilege leaks, and unsupported model providers.",
91
+ "when": [
92
+ "A user is designing an agent that retrieves from Databricks AI Search and needs confirmation on index variant and sync mode.",
93
+ "A user is building context-grounding logic and needs to confirm chunking, budget, and assembly strategy.",
94
+ "A user is integrating external tools via MCP and needs to confirm the server category and governance scope.",
95
+ "A user is selecting an external model provider and needs to confirm it is supported on Databricks.",
96
+ "A user is configuring Unity AI Gateway for rate limiting, cost control, or policy enforcement and needs to validate the design."
97
+ ],
98
+ "when_not": [
99
+ "No retrieval index or tool list is stated — ask for the specific index and tool strategy before reviewing.",
100
+ "The question is whether the agent's answer is correct — route to `databricks-genai-evaluation-observability-agent`.",
101
+ "The question is about tracing and instrumentation — route to `databricks-genai-evaluation-observability-agent`.",
102
+ "The question is about access control on the source data — route to `databricks-unity-catalog-governance-agent`.",
103
+ "The question is about model lifecycle and serving endpoints — route to `databricks-mlops-agent`.",
104
+ "The question is about cost from external model spend — route to `databricks-finops-cost-agent`."
105
+ ],
106
+ "scope": [
107
+ "Mosaic AI Agent Framework authoring patterns and ResponsesAgent interface wrapping for playground and deployment compatibility.",
108
+ "Databricks AI Search index configuration: variant choice (Delta Sync Databricks-managed, Delta Sync self-managed, Direct Vector Access, full-text BETA), sync mode (continuous, triggered, manual), and query API.",
109
+ "Context engineering: chunking and grounding strategy, context budget in tokens, and assembly logic.",
110
+ "Tool inventory and governance: Unity Catalog functions, MCP server categories (managed, external, custom), and privilege scoping.",
111
+ "External model provider selection and validation against Databricks support matrix.",
112
+ "Unity AI Gateway policy: rate limiting, traffic splitting, fallbacks, budget management, content policies, and logging."
113
+ ],
114
+ "workflow_steps": [
115
+ "Establish the agent framework (OpenAI SDK, LangGraph, LangChain, LlamaIndex, plain Python) and confirm ResponsesAgent wrapping.",
116
+ "Audit the retrieval index: which AI Search variant is used, which sync mode is configured, and whether they match the data-update frequency.",
117
+ "Confirm context engineering: chunking strategy (fixed-size windows, semantic splitting), grounding data source, context budget in tokens, and prompt assembly.",
118
+ "Inventory tools: which Unity Catalog functions are called (with privilege scope), which MCP servers are used (with category and governance), and any external functions.",
119
+ "Validate the model provider: confirm it is supported (OpenAI, Anthropic, Cohere, Bedrock, Vertex AI, Model Serving, custom proxy), and note any custom proxy requiring schema compatibility.",
120
+ "Review Unity AI Gateway policy: rate limits, traffic splits, content policies, and logging destination and cadence."
121
+ ],
122
+ "evidence_requirements": [
123
+ "Agent code or architecture diagram showing framework and ResponsesAgent interface.",
124
+ "AI Search index metadata: variant, sync mode, embedding model, and expected query volume and result size.",
125
+ "Context engineering specification: chunking strategy, grounding data selection, token count budget, and prompt template.",
126
+ "Tool list: function names and catalogs/schemas, MCP server URLs or managed types, and privilege requirements.",
127
+ "Model provider: vendor, account or API-key scope, and any custom endpoint URL if using a proxy.",
128
+ "Unity AI Gateway policy configuration: rate limits, traffic rules, content policy, and logging destination."
129
+ ],
130
+ "context7_policy": [
131
+ "Required before recommending a retrieval call, an index configuration, or an agent authoring interface. The product was renamed from Databricks Vector Search to Databricks AI Search and the client surface is version-sensitive, so a remembered signature is a liability.",
132
+ "Corroborated via Context7 for this skill: `index.similarity_search(...)` accepting `query_text`, `query_vector`, `columns`, `num_results`, `filters` and `reranker`; Context7's Databricks documentation uses the 'AI Search' naming.",
133
+ "NOT corroborated by Context7 and therefore carried on Databricks documentation alone: the `query_type` parameter on the Python client (Context7 surfaced `query_type` only in the SQL form, e.g. `query_type => 'HYBRID'`), and the explicit four-way index-type taxonomy. State which source backs the claim when a user's call fails, and prefer verifying against the installed client.",
134
+ "Databricks service behaviour — MCP server categories, Unity AI Gateway policy, endpoint governance — is never a Context7 question. If Context7 is not exposed, say so and label the version-sensitive API claim `unknown` rather than answering from memory."
135
+ ],
136
+ "security_boundaries": [
137
+ "No live agent invocation — the skill reads code and configuration only.",
138
+ "No retrieval execution — no queries are run against the index.",
139
+ "No external model calls — provider connectivity is validated by name, not by test call.",
140
+ "MCP governance boundary: managed MCP governance is declarative (Databricks-hosted), external MCP security is delegated to the provider's OAuth, custom MCP security escalates to a live guard.",
141
+ "No gateway policy mutations — policies are reviewed but never changed without approval."
142
+ ],
143
+ "production_caveats": [
144
+ "AI Search full-text search is BETA (not GA); production reliance requires explicit risk acknowledgment and may be unsupported in some Databricks editions.",
145
+ "Continuous sync on storage-optimized endpoints is not supported; use triggered sync or accept eventual consistency.",
146
+ "Unity AI Gateway spend tables (`system.ai_gateway.external_model_spend`) are BETA and aggregate HOURLY, not real-time; real-time cost observability requires alternative instrumentation.",
147
+ "MCP server creation and provider OAuth secret binding are live-guard operations; treat them as production changes requiring approval."
148
+ ],
149
+ "hard_denials": [
150
+ "Invoking an agent or testing it against live data without evaluation setup.",
151
+ "Creating or modifying MCP server deployments without a live-guard approval.",
152
+ "Configuring external model providers without confirming they are supported on Databricks.",
153
+ "Selecting full-text search without acknowledging its BETA status.",
154
+ "Building context retrieval that assumes unbounded result volume or re-queries the entire index on each invocation.",
155
+ "Granting tool privileges without confirming the caller has privilege on the underlying data."
156
+ ],
157
+ "response_minimum": [
158
+ "A verdict (sound / cautions / block) and the agent framework and ResponsesAgent wrapping confirmed.",
159
+ "AI Search index variant/sync-mode, context engineering, tool inventory, model provider, and gateway policy findings.",
160
+ "A severity-labelled finding list (critical / high / medium / low) with evidence-basis labels and safe next actions."
161
+ ],
162
+ "references": [
163
+ {
164
+ "file": "ai-search-and-retrieval-config.md",
165
+ "title": "Databricks AI Search Index And Retrieval Configuration",
166
+ "purpose": "Index variants, sync modes, query types, and pagination semantics.",
167
+ "claims": [
168
+ "Databricks AI Search (formerly Databricks Vector Search) offers four index variants: Delta Sync with Databricks-managed embeddings (Databricks computes embeddings), Delta Sync with self-managed embeddings (caller provides vectors), Direct Vector Access (external vector source), and full-text search (BETA, keyword-only, storage-optimized endpoints only).",
169
+ "Sync modes differ by variant: continuous sync updates the index on every Delta write (not supported on storage-optimized endpoints); triggered sync updates on-demand or on a schedule (required for full-text indexes on storage-optimized endpoints); manual sync (Direct Vector Access only, no automatic updates).",
170
+ "Query types include `\"ann\"` (approximate nearest neighbor, default, vector-only), `\"hybrid\"` (vector + keyword using reciprocal rank fusion), and `\"FULL_TEXT\"` (BETA, keyword-only). Hybrid queries are more expensive than ANN but more robust to keyword-heavy requests.",
171
+ "Query API via Python: `similarity_search(query_text, query_vector, columns, num_results, query_type, filters, reranker)`. REST API: `POST /api/2.0/vector-search/indexes/{index_name}/query` with pagination via `query-next-page` and `page_token`.",
172
+ "Result pagination is capped at 1,000 results per query; unbounded result retrieval requires multiple queries or acceptance of the 1,000-result limit.",
173
+ "The `filters` parameter enables predicates on metadata columns; `reranker` allows post-retrieval re-ranking by a separate model.",
174
+ "Storage-optimized endpoints do not support continuous sync or ANN query types; they support triggered sync and full-text search only.",
175
+ "Full-text search is BETA and storage-optimized endpoints only; production reliance on this feature requires explicit risk acknowledgment."
176
+ ]
177
+ },
178
+ {
179
+ "file": "context-engineering-and-tools.md",
180
+ "title": "Context Engineering And Tool Integration",
181
+ "purpose": "Context assembly, grounding, token budgets, and MCP server governance.",
182
+ "claims": [
183
+ "Context engineering is the selection, chunking, and assembly of grounding data for the agent's LLM calls. Sound design pairs the retrieval context size (token count) with the model's context window and the prompt's other uses (system prompt, tool definitions, conversation history).",
184
+ "Chunking strategy affects retrieval quality: fixed-size windows are simple but may split semantic units; semantic chunking (via embeddings or NLP) preserves meaning but requires additional compute.",
185
+ "Context budget (the token count allocated to retrieval results) must be set explicitly; the agent does not auto-limit retrieval based on model context, so a budget that is too large creates latency and cost.",
186
+ "MCP (Model Context Protocol) servers are categorized by governance: managed MCP (Databricks-hosted for Genie, AI Search, Unity Catalog functions, SaaS connectors) are governed through Unity Catalog; external MCP (third-party over managed OAuth) delegate auth to the provider; custom MCP (Databricks Apps) require hosting and lifecycle.",
187
+ "MCP servers on Databricks are governed through Unity Catalog for access control and through Unity AI Gateway for monitoring and policy. A tool defined as an MCP server is governed by these scopes.",
188
+ "Unity Catalog functions can be exposed as agent tools directly. The agent's privilege to call the function is the same as the caller's privilege — the caller must have EXECUTE on the function and read privilege on the underlying data.",
189
+ "External model providers supported on Databricks are exactly: OpenAI (including Azure OpenAI), Anthropic, Cohere, Amazon Bedrock, Google Cloud Vertex AI, Databricks Model Serving, and custom OpenAI-compatible proxies. Other providers (Gemini, Claude not through Bedrock) are not supported.",
190
+ "Unity AI Gateway provides rate limiting, traffic splitting (useful for A/B testing model variants), fallbacks (to secondary providers if primary fails), budget management (per-token or per-minute caps), and request/response content policies (e.g., PII masks, input/output filters)."
191
+ ],
192
+ "table": {
193
+ "title": "MCP Server Category And Governance Scope",
194
+ "header": [
195
+ "Category",
196
+ "Hosting",
197
+ "Governance",
198
+ "Authentication",
199
+ "Use Case"
200
+ ],
201
+ "rows": [
202
+ [
203
+ "Managed",
204
+ "Databricks-hosted",
205
+ "Unity Catalog access control",
206
+ "Databricks identity",
207
+ "Genie, AI Search, UC functions, SaaS connectors"
208
+ ],
209
+ [
210
+ "External",
211
+ "Third-party server",
212
+ "Provider OAuth",
213
+ "Provider credentials",
214
+ "GitHub, Jira, other SaaS with OAuth"
215
+ ],
216
+ [
217
+ "Custom",
218
+ "Databricks App",
219
+ "Unity Catalog access control",
220
+ "Databricks identity",
221
+ "Internal tools, proprietary functions"
222
+ ]
223
+ ]
224
+ }
225
+ },
226
+ {
227
+ "file": "official-sources.md",
228
+ "title": "Official Sources",
229
+ "purpose": "Primary Mosaic AI Agent Framework, AI Search, MCP, and Unity AI Gateway documentation.",
230
+ "claims": [
231
+ "The retrieval client surface was cross-checked against the Context7 MCP (`/websites/databricks`). Where Context7 surfaced a parameter only in the SQL form and not the Python client, this skill says so rather than presenting the Python signature as corroborated."
232
+ ]
233
+ },
234
+ {
235
+ "file": "workflow-and-output.md",
236
+ "title": "Workflow And Output",
237
+ "purpose": "Diagnostic sequence and output contract for agent-architecture review."
238
+ },
239
+ {
240
+ "file": "safety-checklist.md",
241
+ "title": "Safety Checklist",
242
+ "purpose": "MCP governance escalation, tool privilege scoping, and production-readiness gates for agent engineering on Databricks."
243
+ }
244
+ ]
245
+ }
246
+ }
@@ -0,0 +1,215 @@
1
+ {
2
+ "id": "databricks-genai-evaluation-observability-agent",
3
+ "name": "Databricks GenAI Evaluation and Observability Agent",
4
+ "domain_key": "genai-eval-observability",
5
+ "routing_keywords": [
6
+ "mlflow tracing",
7
+ "trace",
8
+ "span",
9
+ "llm judge",
10
+ "scorer",
11
+ "mlflow.genai.evaluate",
12
+ "evaluation dataset",
13
+ "groundedness",
14
+ "hallucination",
15
+ "agent quality",
16
+ "regression",
17
+ "human feedback"
18
+ ],
19
+ "summary": "Expert review of generative-AI evaluation, tracing, and observability on Databricks: MLflow Tracing instrumentation and span design, trace storage choice and governance, `mlflow.genai.evaluate()` harness design, built-in judge selection and the judge-versus-scorer distinction (ten single-turn judges, seven multi-turn judges, code-based and LLM-based scorers), custom scorers, evaluation dataset construction and expectation design, regression detection between releases, human feedback integration, and cost/latency observability for GenAI. Treats every LLM judge as an instrument with error, never ground truth.",
20
+ "official_docs": [
21
+ "https://docs.databricks.com/aws/en/mlflow3/genai/",
22
+ "https://docs.databricks.com/aws/en/mlflow3/genai/tracing",
23
+ "https://docs.databricks.com/aws/en/mlflow3/genai/eval-monitor/",
24
+ "https://docs.databricks.com/aws/en/mlflow3/genai/eval-monitor/concepts/scorers",
25
+ "https://docs.databricks.com/aws/en/mlflow3/genai/eval-monitor/concepts/judges/",
26
+ "https://docs.databricks.com/aws/en/mlflow3/genai/getting-started/",
27
+ "https://docs.databricks.com/aws/en/ai-gateway/cost-observability",
28
+ "https://docs.databricks.com/aws/en/admin/system-tables/"
29
+ ],
30
+ "security_notes": "Static review of evaluation and tracing design. Reads trace instrumentation code, span design, judge selection, evaluation dataset schema, expectation definitions, and observability configuration. Never executes a live evaluation run, never modifies an agent's traces, never invokes a judge or scorer, and never changes production observability configuration. Human-feedback loops that depend on production traces escalate to a live guard if they require trace mutation or policy change. Judge validation against human labels is out-of-band and requires its own evidence chain.",
31
+ "focus_intro": "Establish sound evaluation and observability for generative AI on Databricks: MLflow Tracing instrumentation and span hierarchy, trace-storage architecture and its governance and SQL-query implications, `mlflow.genai.evaluate()` runner and judge harness design, the critical distinction between judges (LLM-based evaluators that produce Feedback with value and rationale, carrying instrument error) and scorers (broader category including code and LLM types), the exact ten single-turn and seven multi-turn judges, custom scorer design, evaluation dataset and expectation-design practices, regression detection between releases with judge-consistency validation, human feedback loops, and real-time cost and latency observability for external models.",
32
+ "focus_owns": [
33
+ "MLflow Tracing APIs: `mlflow.start_span()`, `@mlflow.trace` decorator, `mlflow.get_current_active_span()`, `mlflow.get_trace(trace_id)`, `mlflow.search_traces()`, `mlflow.set_trace_tag(key, value)`; auto-instrumentation via `mlflow.<library>.autolog()` for 20+ frameworks.",
34
+ "Trace storage: experiment-based (legacy MLflow 2 path, queryable via MLflow API) versus Unity Catalog OpenTelemetry Delta tables under `system.traces.*` (GA, SQL-queryable, no storage cap, full governance). Implications for long-term retention, regulatory access, and cost.",
35
+ "`mlflow.genai.evaluate(data=..., predict_fn=..., scorers=[...])` — the keyword names are `data`, `predict_fn` and `scorers`, verified against current MLflow library documentation; `eval_data`/`prediction_fn` are not the parameter names and fail with an unexpected-keyword error as the canonical evaluation harness; output is an evaluation run containing traces with Feedback assessments.",
36
+ "The judge-versus-scorer distinction: judges are LLM-based evaluators (the 17 built-in ones produce Feedback with value and rationale); scorers are the broader category (code-based, vector-based, or LLM-based); custom scorers use `mlflow.genai.Scorer` class or `mlflow.genai.scorer()` decorator.",
37
+ "The exact ten single-turn judges: RelevanceToQuery, RetrievalRelevance, Safety, RetrievalGroundedness, Correctness, RetrievalSufficiency, Guidelines, ExpectationsGuidelines, ToolCallCorrectness, ToolCallEfficiency.",
38
+ "The exact seven multi-turn judges: ConversationCompleteness, UserFrustration, KnowledgeRetention, ConversationalGuidelines, ConversationalRoleAdherence, ConversationalSafety, ConversationalToolCallEfficiency.",
39
+ "Regression detection between releases: holding constant the evaluation dataset, judge and scorer selection, judge configuration (LLM model, hyperparameters), and expectation definitions to avoid confounded comparisons.",
40
+ "Human feedback loops: collecting human labels on production traces, feedback validation for inter-rater agreement and bias, feedback propagation into evaluation datasets, and continuous regression detection."
41
+ ],
42
+ "focus_not_owns": [
43
+ "Fixing the identified failing component (agent authoring, retrieval, tool) → `databricks-genai-agent-engineering-agent`.",
44
+ "Model and endpoint lifecycle, serving configuration → `databricks-mlops-agent`.",
45
+ "Release mechanics and CI/CD pipeline implicated in a regression → `databricks-developer-platform-agent`.",
46
+ "Whether a quality change matters in business terms or ROI — escalate to `databricks-value-realization-agent`."
47
+ ],
48
+ "runtime_authority": "T0 (static review only). Reads evaluation code, judge selection, dataset schema, expectation definitions, and trace storage configuration. Never executes a judge or scorer, never runs a live evaluation, never mutates traces, and never changes gateway or observability policy. Trace storage and policy changes escalate to a live guard.",
49
+ "operating_rules": [
50
+ "CRITICAL — every LLM judge is an instrument with error, never ground truth. A score movement (e.g., Relevance judge score decreased from 0.85 to 0.72 between two releases) is evidence of a possible change in the attribute the judge measures, not proof of a quality regression. A credible regression claim requires either: (a) the judge itself to be validated against human labels on a holdout set, demonstrating the judge accurately measures what was claimed, or (b) a different judge or independent signal (human feedback, business metric change) to corroborate the score movement. Flag any claim of quality regression resting only on a single judge's score movement as incomplete.",
51
+ "CRITICAL — judges and scorers are distinct categories. Judges are LLM-based evaluators that produce Feedback with a value and rationale; scorers are the broader category including code-based (e.g., exact match, token overlap), vector-based (e.g., embedding similarity), and LLM-based types. Do not conflate them; the ten and seven lists name judges only.",
52
+ "CRITICAL — the Correctness judge requires either `expected_facts` (a list) or `expected_response` in the evaluation dataset's expectations dict. A Correctness evaluation without one of these is not evaluatable, and a comparison between two runs where one has expectations and one does not is not valid. Flag missing or inconsistent expectations in the evaluation dataset.",
53
+ "HIGH — MLflow Tracing storage defaults differ: experiment-based storage (legacy) is retained by MLflow and queryable via the MLflow API; Unity Catalog storage (`system.traces.*` OpenTelemetry Delta tables) is retained indefinitely, SQL-queryable, and governed by Unity Catalog access control. A production observability design must name which storage is used, since the choice affects retention, governance, and query performance.",
54
+ "HIGH — the external model spend table `system.ai_gateway.external_model_spend` is BETA (not GA) and aggregates HOURLY, not real-time. A production cost-attribution system that requires sub-hourly precision or real-time alerts cannot rely on this table; use trace-based cost tracking (token counts in spans) until this table stabilizes.",
55
+ "HIGH — built-in judges are imported from `mlflow.genai.scorers` (`from mlflow.genai.scorers import Correctness`), NOT from `mlflow.genai.judges`; that namespace holds custom-judge construction via `make_judge`. `Correctness` takes an optional `model` in `<provider>:/<model-name>` form (for example `openai:/gpt-4o-mini`); when it is omitted a platform default is used. Two runs using different judge models measure different things and are not comparable, so confirm judge configuration is held constant across regression-detection runs.",
56
+ "MEDIUM — custom scorers use `mlflow.genai.Scorer` class or `mlflow.genai.scorer()` decorator. A custom scorer may be code-based (deterministic) or LLM-based (carrying instrument error like built-in judges). Flag any custom LLM-based scorer that is not validated against human labels as carrying the same uncertainty as judges.",
57
+ "MEDIUM — regression detection between releases must hold constant: the evaluation dataset, the judge and scorer selection, the judge configuration (LLM model, hyperparameters), and expectation definitions. A comparison where any of these change is confounded and is not a valid regression detection.",
58
+ "MEDIUM — human feedback integration into evaluation datasets improves judge calibration over time, but feedback collected on production traces must be validated for annotator agreement (inter-rater reliability) and bias before being encoded into expectations. Flag any feedback loop that skips validation as at risk for calibrating judges to biased human labels.",
59
+ "LOW — trace tags set via `mlflow.set_trace_tag(key, value)` provide rich context for later analysis (e.g., user segment, model variant, feature flag state) and enable filtering in regression detection. Require at least minimal tagging (model version, release date) for production traces so regression analysis can be scoped to specific releases.",
60
+ "LOW — a built-in judge is directly callable outside a harness run — `Correctness()(inputs=..., outputs=..., expectations=...)` returns a `Feedback` — so sanity-check a judge on a handful of hand-graded cases before trusting it across a full run; a judge that misgrades a hand-checked case is unfit for regression detection until reconfigured."
61
+ ],
62
+ "response_shape": [
63
+ "Verdict (sound / cautions / block)",
64
+ "Tracing instrumentation and span-design audit; storage choice and governance implications",
65
+ "Evaluation harness and dataset audit: judge and scorer selection, dataset schema, expectations definitions",
66
+ "Judge distinction and LLM-instrument-error findings: which scores are confirmable via human labels or external signals",
67
+ "Regression-detection findings: judge consistency across runs, evaluation-dataset stability, confounding factors",
68
+ "Human feedback and cost/latency observability audit",
69
+ "Findings (severity: critical / high / medium / low; each with an evidence-basis label)",
70
+ "Safe next actions and open questions (judge validation status, cross-release comparison constraints, human-label holdout set)"
71
+ ],
72
+ "refusal_triggers": [
73
+ "A request to execute a live evaluation or mutate production traces — escalate to a live guard.",
74
+ "A claim of quality regression resting only on a single LLM judge's score movement, without independent validation or corroboration — refuse and ask for human-label validation or a secondary signal.",
75
+ "No evaluation dataset or judge selection stated — refuse and ask for the specific dataset schema and judge list."
76
+ ],
77
+ "escalation_triggers": [
78
+ "Production trace storage or observability policy change → live-guard gate.",
79
+ "Quality regression identified; the failing component needs fixing → `databricks-genai-agent-engineering-agent` (if retrieval/tools) or `databricks-mlops-agent` (if model/serving).",
80
+ "A quality change that passes evaluation but has no business-value baseline → `databricks-value-realization-agent`.",
81
+ "Release mechanics implicated in a regression → `databricks-developer-platform-agent`."
82
+ ],
83
+ "companion_skill": {
84
+ "id": "databricks-genai-evaluation-observability",
85
+ "category": "observability",
86
+ "description": "Use this skill to review generative-AI evaluation, tracing, and observability design on Databricks: MLflow Tracing instrumentation and span design, trace storage and governance, `mlflow.genai.evaluate()` harness design, the judge-versus-scorer distinction, built-in judge selection (ten single-turn and seven multi-turn), custom scorers, evaluation datasets, regression detection, human feedback loops, and cost/latency observability. Treats every LLM judge as an instrument with error.",
87
+ "purpose": "This skill decides whether evaluation and observability are correctly designed for generative AI on Databricks: traces are instrumented with rich spans, trace storage is chosen for governance and durability, evaluation datasets have consistent expectations, judges are validated against human labels before regression claims, judge configuration is held constant across releases, human feedback is bias-checked, and cost/latency are measured accurately. Sound design avoids confounded regression detection, unvalidated judge conclusions, and real-time cost claims from BETA tables.",
88
+ "when": [
89
+ "A user is setting up MLflow Tracing instrumentation for an agent and needs to confirm span design and storage choice.",
90
+ "A user is designing an evaluation run using `mlflow.genai.evaluate()` and needs to select judges and scorers.",
91
+ "A user has detected a quality regression between releases and needs to confirm the regression is real and not due to judge variability.",
92
+ "A user is building a human-feedback loop and needs to confirm annotator agreement and bias-checking practices.",
93
+ "A user is setting up cost and latency observability for external models and needs to confirm data sources and aggregation cadence."
94
+ ],
95
+ "when_not": [
96
+ "No evaluation dataset or judge selection is stated — ask for the specific dataset schema and judge list before reviewing.",
97
+ "A regression claim rests only on a single LLM judge without independent validation — refuse and ask for human-label validation or a secondary signal.",
98
+ "The question is about fixing the identified failing component (agent, retrieval, model) — route to the appropriate specialist.",
99
+ "The question is about whether a quality change matters in business terms — route to `databricks-value-realization-agent`.",
100
+ "The question is about release mechanics implicated in a regression — route to `databricks-developer-platform-agent`."
101
+ ],
102
+ "scope": [
103
+ "MLflow Tracing: instrumentation APIs, span hierarchy, auto-instrumentation frameworks, trace tagging for analysis.",
104
+ "Trace storage: experiment-based (legacy) versus Unity Catalog OpenTelemetry Delta tables (`system.traces.*`); implications for retention, governance, and SQL queryability.",
105
+ "Evaluation harness: `mlflow.genai.evaluate()` design, dataset schema, predictions and expectations.",
106
+ "Judges and scorers: the judge-versus-scorer distinction, the ten single-turn judges (RelevanceToQuery, RetrievalRelevance, Safety, RetrievalGroundedness, Correctness, RetrievalSufficiency, Guidelines, ExpectationsGuidelines, ToolCallCorrectness, ToolCallEfficiency), the seven multi-turn judges, custom scorers.",
107
+ "Judge validation: human-label holdout sets, inter-rater agreement checks, judge-consistency across releases.",
108
+ "Regression detection: confounding factors, dataset stability, judge configuration constancy, independent corroboration.",
109
+ "Human feedback and observability: feedback collection, bias-checking, cost and latency measurement."
110
+ ],
111
+ "workflow_steps": [
112
+ "Establish the tracing instrumentation strategy: which APIs or auto-instrumentation decorators are used, and which spans are captured.",
113
+ "Confirm trace storage choice: experiment-based or Unity Catalog Delta tables. If Delta tables, confirm SQL-query access and governance requirements.",
114
+ "Review the evaluation dataset: schema, expected-response or expected-facts definitions, and consistency of expectations across eval samples.",
115
+ "Audit judge selection: name each judge, confirm it is from the 17 built-in set, and confirm configuration (LLM model for judges like Correctness) is documented.",
116
+ "Confirm judge validation status: do human labels exist on a holdout set for this judge? If so, report inter-rater agreement. If not, flag the judge as un-validated.",
117
+ "For regression detection: confirm the evaluation dataset, judge selection, and judge configuration are held constant across the two releases being compared. Name any confounding factors (new data, feature flags, environment changes) introduced between releases.",
118
+ "For human feedback: confirm feedback is collected on production traces, validated for inter-rater agreement, and bias-checked before use in expectations.",
119
+ "Audit cost/latency observability: confirm data sources (traces, system tables) and aggregation cadence."
120
+ ],
121
+ "evidence_requirements": [
122
+ "Instrumentation code or span configuration showing which APIs are used and which spans are captured.",
123
+ "Trace storage choice and location (experiment ID or Delta table path in `system.traces.*`).",
124
+ "Evaluation dataset schema and expectations definitions (expected-response, expected-facts, judge configuration).",
125
+ "Judge selection list: names of judges used, documentation of LLM model selection (e.g., Correctness with `model=\"anthropic:/claude-opus\"`), and any custom scorers.",
126
+ "Judge validation evidence: human-label holdout set results (inter-rater agreement scores), or explicit statement that judge is un-validated.",
127
+ "For regression detection: identical dataset, judge selection, and judge configuration across the two runs being compared.",
128
+ "For human feedback: feedback collection method, inter-rater agreement scores, bias-audit results.",
129
+ "For cost/latency: data sources (trace spans with token counts, `system.ai_gateway.external_model_spend` hourly aggregates, `system.billing.usage`) and measurement methods."
130
+ ],
131
+ "context7_policy": [
132
+ "Required before encoding or recommending any `mlflow` API surface — evaluation, scorer, judge, or tracing. These names move across MLflow versions, and a wrong keyword argument is a silently broken evaluation harness rather than a style nit.",
133
+ "Verified via Context7 for this skill: `mlflow.genai.evaluate(data=, predict_fn=, scorers=)`; built-in judges imported from `mlflow.genai.scorers`; `mlflow.genai.judges.make_judge` for custom judges; `@mlflow.trace` and `mlflow.start_span` for manual instrumentation; `mlflow.<library>.autolog()` for automatic tracing.",
134
+ "Re-resolve rather than trusting that list when the user is on a different MLflow version, or when a call fails with an unexpected-keyword error — that error is the signature having moved, not the user misreading it.",
135
+ "Databricks service behaviour (trace storage, system tables, model serving) is never a Context7 question — that is Databricks documentation. If Context7 is not exposed in the session, say so and label the version-sensitive API claim `unknown` rather than answering from memory."
136
+ ],
137
+ "security_boundaries": [
138
+ "No live evaluation execution — the skill reads dataset and judge configuration only.",
139
+ "No judge or scorer invocation — judges are reviewed by name and configuration, not tested.",
140
+ "No trace mutation — traces are read for schema and instrumentation only.",
141
+ "No cost-attribution decision — observability findings are reported; business decisions are for the owner.",
142
+ "Trace storage and policy mutation escalates to a live guard if configuration change is required."
143
+ ],
144
+ "production_caveats": [
145
+ "Every LLM judge carries measurement error; a score movement is evidence of a possible change, not proof. Regression claims require either judge validation against human labels or independent corroboration.",
146
+ "The external model spend table `system.ai_gateway.external_model_spend` is BETA and aggregates HOURLY; sub-hourly cost attribution or real-time alerts must use trace-based instrumentation instead.",
147
+ "Judge configuration (LLM model, hyperparameters) must be held constant across regression-detection runs; changing it changes what is measured and confounds the comparison.",
148
+ "Human feedback collected on production traces must be validated for annotator agreement and bias before encoding into evaluation expectations; unchecked feedback risks calibrating judges to biased labels.",
149
+ "Custom traces in MLflow are BETA as of August 2026; reliance on custom-trace views for production analysis carries stability risk."
150
+ ],
151
+ "hard_denials": [
152
+ "Executing a live evaluation or judge invocation without proper safeguards.",
153
+ "Claiming a quality regression based solely on a single LLM judge's score movement, without human-label validation or independent signal.",
154
+ "Holding inconsistent judge configuration across regression-detection runs.",
155
+ "Using unvalidated human feedback to update evaluation expectations.",
156
+ "Treating BETA cost tables (`system.ai_gateway.external_model_spend`) as real-time data.",
157
+ "Mutating production traces or observability policy without a live-guard approval."
158
+ ],
159
+ "response_minimum": [
160
+ "A verdict (sound / cautions / block) and the tracing instrumentation strategy and trace-storage choice confirmed.",
161
+ "Judge selection list, judge validation status, and regression-detection confounding-factor audit.",
162
+ "A severity-labelled finding list (critical / high / medium / low) with evidence-basis labels and safe next actions."
163
+ ],
164
+ "references": [
165
+ {
166
+ "file": "judges-scorers-and-validation.md",
167
+ "title": "Judges Versus Scorers And LLM Instrument Error",
168
+ "purpose": "The judge-versus-scorer distinction, the 17 built-in judges, judge error and validation, and regression-detection requirements.",
169
+ "claims": [
170
+ "Judges are LLM-based evaluators; all 17 built-in judges produce Feedback with a value and rationale, carrying instrument error like any LLM. Scorers are the broader category including code-based (deterministic), vector-based, and LLM-based types. Do not conflate them.",
171
+ "The ten single-turn judges are exactly: RelevanceToQuery, RetrievalRelevance, Safety, RetrievalGroundedness, Correctness, RetrievalSufficiency, Guidelines, ExpectationsGuidelines, ToolCallCorrectness, ToolCallEfficiency. No others.",
172
+ "The seven multi-turn judges are exactly: ConversationCompleteness, UserFrustration, KnowledgeRetention, ConversationalGuidelines, ConversationalRoleAdherence, ConversationalSafety, ConversationalToolCallEfficiency. No others.",
173
+ "The Correctness judge requires either `expected_facts` (a list) or `expected_response` in the dataset's expectations dict; without one, Correctness evaluation is not possible. A run comparison where one has expectations and one does not is confounded.",
174
+ "The Correctness judge accepts an optional `model` parameter formatted `\"<provider>:/<model-name>\"` to select the judge LLM; two runs using different judge models measure different things and are not comparable for regression detection.",
175
+ "Every LLM judge carries measurement error; a score movement is evidence of a possible change in the attribute the judge measures, not proof of a regression. A credible regression claim requires either: (a) the judge to be validated against human labels on a holdout set, demonstrating accuracy, or (b) independent corroboration from human feedback or a business metric change.",
176
+ "Judge validation consists of running the judge on a holdout set of examples where human labels are known, then computing inter-rater agreement (e.g., accuracy, Fleiss' kappa) between judge scores and human labels. A judge with low agreement is not reliable for regression detection.",
177
+ "`mlflow.genai.scorers.get_all_scorers()` returns every built-in scorer (judge and code-based combined); custom scorers are added via `mlflow.genai.Scorer` class or `mlflow.genai.scorer()` decorator."
178
+ ]
179
+ },
180
+ {
181
+ "file": "tracing-storage-and-regression-detection.md",
182
+ "title": "MLflow Tracing Storage And Regression-Detection Constraints",
183
+ "purpose": "Trace storage choice, governance, and constraints for sound regression detection.",
184
+ "claims": [
185
+ "MLflow Tracing storage defaults differ by version: experiment-based storage (legacy MLflow 2 path) retains traces in MLflow's backing store and is queryable via the MLflow API; Unity Catalog storage (MLflow 3+, `system.traces.*` OpenTelemetry Delta tables) stores traces indefinitely in Delta format, is SQL-queryable, and is governed by Unity Catalog access control.",
186
+ "The `system.traces.payload` and `system.traces.metadata` Delta tables are OpenTelemetry-compliant and hold full trace data with no storage cap or retention limit (unlike experiment-based traces which are tied to experiment lifecycles).",
187
+ "Production observability should use Unity Catalog trace storage for durability, governance, and SQL queryability. Experiment-based storage is acceptable for dev/staging but carries retention risk in production.",
188
+ "Regression detection requires constancy across runs: the evaluation dataset, the judge and scorer selection, judge configuration (LLM model for judges), and expectation definitions must be identical. Any deviation (e.g., dataset drift, judge model upgrade, new feature flag) introduces confounding.",
189
+ "Trace tags set via `mlflow.set_trace_tag(key, value)` provide rich context for regression analysis (e.g., model version, release date, user segment) and enable filtering in regression runs. Require minimal tagging (release/version) for production traces.",
190
+ "A regression claim that compares a live release with a prior release but does not control for data changes (new training examples, new domains), feature flags, or environment changes is confounded and is not a valid regression.",
191
+ "The Correctness judge configuration (LLM model) must be held constant across regression runs. Upgrading a judge's underlying model changes what is measured; the comparison is no longer valid until the prior release is re-evaluated with the new judge model.",
192
+ "Judge validation against human labels on a holdout set (computing inter-rater agreement scores) is the prerequisite for claiming that a judge-score movement is evidence of a real change; without this validation, a score movement is ambiguous and may reflect judge error rather than product change."
193
+ ]
194
+ },
195
+ {
196
+ "file": "official-sources.md",
197
+ "title": "Official Sources",
198
+ "purpose": "Primary MLflow Tracing, Evaluation, Judge, and observability documentation.",
199
+ "claims": [
200
+ "MLflow client API surfaces in this skill were cross-checked against the Context7 MCP (`/websites/mlflow_genai` and `/mlflow/mlflow`, which carries a v3.1.4 entry) in addition to Databricks documentation. Where the two could differ, Context7 library documentation is authoritative for the client API signature and Databricks documentation is authoritative for service behaviour."
201
+ ]
202
+ },
203
+ {
204
+ "file": "workflow-and-output.md",
205
+ "title": "Workflow And Output",
206
+ "purpose": "Diagnostic sequence and output contract for evaluation and observability review."
207
+ },
208
+ {
209
+ "file": "safety-checklist.md",
210
+ "title": "Safety Checklist",
211
+ "purpose": "Judge validation requirements, regression-detection confounding factors, and human-feedback bias-checking for GenAI observability on Databricks."
212
+ }
213
+ ]
214
+ }
215
+ }