ace-llm 0.36.3 → 0.38.2
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- checksums.yaml +4 -4
- data/.ace-defaults/llm/providers/anthropic.yml +29 -14
- data/.ace-defaults/llm/providers/google.yml +28 -14
- data/.ace-defaults/llm/providers/groq.yml +60 -5
- data/.ace-defaults/llm/providers/lmstudio.yml +1 -1
- data/.ace-defaults/llm/providers/mistral.yml +45 -14
- data/.ace-defaults/llm/providers/openai.yml +46 -14
- data/.ace-defaults/llm/providers/openrouter.yml +229 -32
- data/.ace-defaults/llm/providers/togetherai.yml +43 -10
- data/.ace-defaults/llm/providers/xai.yml +21 -8
- data/.ace-defaults/llm/providers/zai.yml +17 -6
- data/CHANGELOG.md +43 -0
- data/README.md +2 -0
- data/exe/ace-llm +2 -1
- data/handbook/guides/llm-query-tool-reference.g.md +78 -626
- data/lib/ace/llm/atoms/provider_config_validator.rb +53 -0
- data/lib/ace/llm/cli/commands/query.rb +10 -2
- data/lib/ace/llm/molecules/model_limit_resolver.rb +88 -0
- data/lib/ace/llm/version.rb +1 -1
- data/lib/ace/llm.rb +1 -0
- metadata +5 -4
|
@@ -1,26 +1,59 @@
|
|
|
1
1
|
name: togetherai
|
|
2
|
-
last_synced:
|
|
2
|
+
last_synced: 2026-04-24
|
|
3
3
|
class: Ace::LLM::Organisms::TogetherAIClient
|
|
4
4
|
gem: ace-llm
|
|
5
|
-
context_limit: 128000 # Default context limit; varies by model
|
|
6
5
|
models:
|
|
7
|
-
|
|
8
|
-
|
|
9
|
-
|
|
10
|
-
|
|
6
|
+
- MiniMaxAI/MiniMax-M2.5
|
|
7
|
+
- MiniMaxAI/MiniMax-M2.7
|
|
8
|
+
- Qwen/Qwen3-Coder-480B-A35B-Instruct-FP8
|
|
9
|
+
- Qwen/Qwen3-Coder-Next-FP8
|
|
10
|
+
- Qwen/Qwen3.5-397B-A17B
|
|
11
|
+
- essentialai/Rnj-1-Instruct
|
|
12
|
+
- google/gemma-4-31B-it
|
|
13
|
+
- moonshotai/Kimi-K2.5
|
|
14
|
+
- moonshotai/Kimi-K2.6
|
|
15
|
+
- openai/gpt-oss-120b
|
|
16
|
+
- zai-org/GLM-5.1
|
|
17
|
+
limits:
|
|
18
|
+
default:
|
|
19
|
+
context: 262144
|
|
20
|
+
output: 262144
|
|
21
|
+
models:
|
|
22
|
+
MiniMaxAI/MiniMax-M2.5:
|
|
23
|
+
context: 204800
|
|
24
|
+
output: 131072
|
|
25
|
+
MiniMaxAI/MiniMax-M2.7:
|
|
26
|
+
context: 202752
|
|
27
|
+
output: 131072
|
|
28
|
+
Qwen/Qwen3.5-397B-A17B:
|
|
29
|
+
output: 130000
|
|
30
|
+
essentialai/Rnj-1-Instruct:
|
|
31
|
+
context: 32768
|
|
32
|
+
output: 32768
|
|
33
|
+
google/gemma-4-31B-it:
|
|
34
|
+
output: 131072
|
|
35
|
+
moonshotai/Kimi-K2.6:
|
|
36
|
+
output: 131000
|
|
37
|
+
openai/gpt-oss-120b:
|
|
38
|
+
context: 131072
|
|
39
|
+
output: 131072
|
|
40
|
+
zai-org/GLM-5.1:
|
|
41
|
+
context: 202752
|
|
42
|
+
output: 131072
|
|
11
43
|
aliases:
|
|
12
44
|
model:
|
|
13
|
-
qwen: Qwen/Qwen3-Coder-480B-A35B-Instruct-FP8
|
|
14
|
-
deepseek: deepseek-ai/DeepSeek-V3
|
|
45
|
+
qwen-coder: Qwen/Qwen3-Coder-480B-A35B-Instruct-FP8
|
|
46
|
+
deepseek: deepseek-ai/DeepSeek-V3.1
|
|
15
47
|
kimi: moonshotai/Kimi-K2-Instruct
|
|
48
|
+
kimi-T: moonshotai/Kimi-K2-Thinking
|
|
16
49
|
oss: openai/gpt-oss-120b
|
|
17
50
|
api_key:
|
|
18
51
|
env: TOGETHER_API_KEY
|
|
19
52
|
required: true
|
|
20
53
|
description: Together AI API key
|
|
21
54
|
capabilities:
|
|
22
|
-
|
|
23
|
-
|
|
55
|
+
- text_generation
|
|
56
|
+
- streaming
|
|
24
57
|
default_options:
|
|
25
58
|
temperature: 0.7
|
|
26
59
|
max_tokens: 16384
|
|
@@ -1,14 +1,27 @@
|
|
|
1
1
|
name: xai
|
|
2
|
-
last_synced:
|
|
2
|
+
last_synced: 2026-04-24
|
|
3
3
|
class: Ace::LLM::Organisms::XAIClient
|
|
4
4
|
gem: ace-llm
|
|
5
|
-
context_limit: 131072 # Grok models have 131K context window
|
|
6
5
|
models:
|
|
7
|
-
|
|
8
|
-
|
|
9
|
-
|
|
10
|
-
|
|
11
|
-
|
|
6
|
+
- grok-4
|
|
7
|
+
- grok-4-1-fast
|
|
8
|
+
- grok-4-1-fast-non-reasoning
|
|
9
|
+
- grok-4-fast-non-reasoning
|
|
10
|
+
- grok-4.20-0309-non-reasoning
|
|
11
|
+
- grok-4.20-0309-reasoning
|
|
12
|
+
- grok-4.20-multi-agent-0309
|
|
13
|
+
- grok-code-fast-1
|
|
14
|
+
limits:
|
|
15
|
+
default:
|
|
16
|
+
context: 2000000
|
|
17
|
+
output: 30000
|
|
18
|
+
models:
|
|
19
|
+
grok-4:
|
|
20
|
+
context: 256000
|
|
21
|
+
output: 64000
|
|
22
|
+
grok-code-fast-1:
|
|
23
|
+
context: 256000
|
|
24
|
+
output: 10000
|
|
12
25
|
aliases:
|
|
13
26
|
global:
|
|
14
27
|
grok: xai:grok-4
|
|
@@ -24,7 +37,7 @@ api_key:
|
|
|
24
37
|
required: true
|
|
25
38
|
description: x.ai API key
|
|
26
39
|
capabilities:
|
|
27
|
-
|
|
40
|
+
- text_generation
|
|
28
41
|
default_options:
|
|
29
42
|
temperature: 0.7
|
|
30
43
|
max_tokens: 16384
|
|
@@ -1,18 +1,29 @@
|
|
|
1
1
|
name: zai
|
|
2
|
-
last_synced: 2026-
|
|
2
|
+
last_synced: 2026-04-24
|
|
3
3
|
class: Ace::LLM::Organisms::ZaiClient
|
|
4
4
|
gem: ace-llm
|
|
5
|
-
context_limit: 128000
|
|
6
5
|
models:
|
|
7
|
-
|
|
8
|
-
|
|
9
|
-
|
|
6
|
+
- glm-4.7
|
|
7
|
+
- glm-4.7-flashx
|
|
8
|
+
- glm-5
|
|
9
|
+
- glm-5-turbo
|
|
10
|
+
- glm-5.1
|
|
11
|
+
- glm-5v-turbo
|
|
12
|
+
limits:
|
|
13
|
+
default:
|
|
14
|
+
context: 200000
|
|
15
|
+
output: 131072
|
|
16
|
+
models:
|
|
17
|
+
glm-4.7:
|
|
18
|
+
context: 204800
|
|
19
|
+
glm-5:
|
|
20
|
+
context: 204800
|
|
10
21
|
api_key:
|
|
11
22
|
env: ZAI_API_KEY
|
|
12
23
|
required: true
|
|
13
24
|
description: Z.AI API key
|
|
14
25
|
capabilities:
|
|
15
|
-
|
|
26
|
+
- text_generation
|
|
16
27
|
default_options:
|
|
17
28
|
temperature: 0.7
|
|
18
29
|
max_tokens: 16384
|
data/CHANGELOG.md
CHANGED
|
@@ -6,6 +6,49 @@ The format is based on [Keep a Changelog](https://keepachangelog.com/en/1.0.0/),
|
|
|
6
6
|
and this project adheres to [Semantic Versioning](https://semver.org/spec/v2.0.0.html).
|
|
7
7
|
|
|
8
8
|
## [Unreleased]
|
|
9
|
+
### Fixed
|
|
10
|
+
- Stubbed the Z.ai endpoint in query command CLI-routing tests so provider fallback no longer triggers blocked real HTTP calls during deterministic suite runs.
|
|
11
|
+
|
|
12
|
+
## [0.38.2] - 2026-04-24
|
|
13
|
+
|
|
14
|
+
### Technical
|
|
15
|
+
- Tightened `TS-LLM-002` verification so `ace-llm --list-providers` remains a discovery-only public-surface check instead of implying provider readiness.
|
|
16
|
+
|
|
17
|
+
## [0.38.1] - 2026-04-24
|
|
18
|
+
|
|
19
|
+
### Changed
|
|
20
|
+
- Clarified provider-listing documentation so `ace-llm --list-providers` is described as discovery plus setup hints, with setup readiness delegated to `ace-config doctor`.
|
|
21
|
+
|
|
22
|
+
## [0.38.0] - 2026-04-24
|
|
23
|
+
|
|
24
|
+
### Added
|
|
25
|
+
- Added shared resolved-model limit lookup so alias, role, and preset-expanded concrete provider/model targets can read merged context and output limits from provider config.
|
|
26
|
+
|
|
27
|
+
### Changed
|
|
28
|
+
- Updated provider config validation and downstream callers to use the new `limits` schema and the expanded concrete model target instead of provider-wide limit assumptions.
|
|
29
|
+
|
|
30
|
+
## [0.37.1] - 2026-04-23
|
|
31
|
+
|
|
32
|
+
|
|
33
|
+
### Changed
|
|
34
|
+
- Added concurrent provider ping support in query execution and refreshed fallback orchestration behavior used by provider-driven query flows.
|
|
35
|
+
- Updated provider ping coverage so live query failures no longer block local diagnostics and stale fallback assumptions are exercised more explicitly.
|
|
36
|
+
|
|
37
|
+
|
|
38
|
+
## [0.37.0] - 2026-04-23
|
|
39
|
+
|
|
40
|
+
### Added
|
|
41
|
+
- Added `ace-llm --no-fallback` so callers can run ordinary prompts such as `ping` against the exact requested provider/model without fallback routing.
|
|
42
|
+
|
|
43
|
+
## [0.36.5] - 2026-04-22
|
|
44
|
+
|
|
45
|
+
### Fixed
|
|
46
|
+
- Stabilized deterministic verification by stubbing CLI query routing away from provider fallback HTTP calls and removing retry sleep from the provider/model fallback test.
|
|
47
|
+
|
|
48
|
+
## [0.36.4] - 2026-04-22
|
|
49
|
+
|
|
50
|
+
### Technical
|
|
51
|
+
- Added regression coverage to keep fresh `commit` role defaults from reintroducing stale `codex:gpt-5-mini` mappings while preserving `codex:mini` fallback behavior.
|
|
9
52
|
|
|
10
53
|
## [0.36.3] - 2026-04-19
|
|
11
54
|
|
data/README.md
CHANGED
|
@@ -51,6 +51,8 @@ Deterministic coverage lives in `test/fast/` and `test/feat/`. Scenario assets s
|
|
|
51
51
|
|
|
52
52
|
**Build resilient prompt workflows** - configure fallback chains and retry behavior through the [config cascade](.ace-defaults/llm/config.yml) so transient provider issues do not block work.
|
|
53
53
|
|
|
54
|
+
**Check exact provider reachability** - use `ace-llm gemini:pro "ping" --no-fallback --timeout 15 --max-tokens 4` to verify the requested provider/model without fallback routing.
|
|
55
|
+
|
|
54
56
|
**Power LLM-enhanced flows in sibling packages** - serve as the execution backend for [ace-git-commit](../ace-git-commit), [ace-idea](../ace-idea), [ace-review](../ace-review), [ace-sim](../ace-sim), [ace-prompt-prep](../ace-prompt-prep), and more.
|
|
55
57
|
|
|
56
58
|
---
|
data/exe/ace-llm
CHANGED
|
@@ -18,7 +18,8 @@ args = ARGV.empty? ? ["--help"] : ARGV
|
|
|
18
18
|
|
|
19
19
|
# Start ace-support-cli single-command entrypoint with exception-based exit code handling (per ADR-023)
|
|
20
20
|
begin
|
|
21
|
-
Ace::Support::Cli::Runner.new(Ace::LLM::CLI::Commands::Query).call(args: args)
|
|
21
|
+
exit_code = Ace::Support::Cli::Runner.new(Ace::LLM::CLI::Commands::Query).call(args: args)
|
|
22
|
+
exit(exit_code) if exit_code.to_i.positive?
|
|
22
23
|
rescue Ace::Support::Cli::Error => e
|
|
23
24
|
warn e.message
|
|
24
25
|
exit(e.exit_code)
|