@yanlinglabs/winter-provider-catalog 0.0.20 → 0.0.22

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -0,0 +1,243 @@
1
+ {
2
+ "$comment": "HAND-AUTHORED provider/model price evidence. This is intentionally a narrow patch layer rather than a model overlay: replacing an extracted model row merely to add a price would discard its endpoint, modalities, context window, and effort evidence. All values are USD per million tokens. Where a first-party table is CNY-only, its sourceRef records the exact selected CNY tier and the ECB reference-rate conversion used; `ModelPricing` cannot represent currency, region, or multiple context tiers.",
3
+ "pricing": {
4
+ "baidu/ernie-4.5-turbo-128k": {
5
+ "value": { "inputPerMTokUsd": 0.119445, "outputPerMTokUsd": 0.47778, "cacheReadPerMTokUsd": 0.029861 },
6
+ "source": "official-doc",
7
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — ERNIE-4.5-Turbo-128K online list rate: CNY 0.0008 input, CNY 0.0032 output, and CNY 0.0002 cache-hit input per 1K tokens. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
8
+ "confidence": "inferred",
9
+ "observedAt": "2026-09-19T00:00:00Z"
10
+ },
11
+ "baidu/ernie-4.5-turbo-32k": {
12
+ "value": { "inputPerMTokUsd": 0.119445, "outputPerMTokUsd": 0.47778, "cacheReadPerMTokUsd": 0.029861 },
13
+ "source": "official-doc",
14
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — ERNIE-4.5-Turbo-32K online list rate: CNY 0.0008 input, CNY 0.0032 output, and CNY 0.0002 cache-hit input per 1K tokens. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
15
+ "confidence": "inferred",
16
+ "observedAt": "2026-09-19T00:00:00Z"
17
+ },
18
+ "baidu/ernie-4.5-turbo-vl": {
19
+ "value": { "inputPerMTokUsd": 0.447919, "outputPerMTokUsd": 1.343756 },
20
+ "source": "official-doc",
21
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — ERNIE-4.5-Turbo-VL online list rate: CNY 0.003 input and CNY 0.009 output per 1K tokens. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
22
+ "confidence": "inferred",
23
+ "observedAt": "2026-09-19T00:00:00Z"
24
+ },
25
+ "baidu/ernie-5.0": {
26
+ "value": { "inputPerMTokUsd": 1.493062, "outputPerMTokUsd": 5.972249 },
27
+ "source": "official-doc",
28
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — ERNIE-5.0 online list rate is tiered by input length: CNY 0.006/0.024 per 1K tokens at <=32K and CNY 0.01/0.04 above 32K through 128K. This records the higher CNY 0.01 input / CNY 0.04 output tier so estimates do not under-report long prompts. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
29
+ "confidence": "inferred",
30
+ "observedAt": "2026-09-19T00:00:00Z"
31
+ },
32
+ "baidu/ernie-5.1": {
33
+ "value": { "inputPerMTokUsd": 0.895837, "outputPerMTokUsd": 3.284737 },
34
+ "source": "official-doc",
35
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — ERNIE-5.1 online list rate is tiered by input length: CNY 0.004/0.018 per 1K tokens at <=32K and CNY 0.006/0.022 above 32K through 128K. This records the higher CNY 0.006 input / CNY 0.022 output tier so estimates do not under-report long prompts. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
36
+ "confidence": "inferred",
37
+ "observedAt": "2026-09-19T00:00:00Z"
38
+ },
39
+ "qianfan/ernie-5.0-thinking-latest": {
40
+ "value": { "inputPerMTokUsd": 1.493062, "outputPerMTokUsd": 5.972249 },
41
+ "source": "official-doc",
42
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — the ERNIE-5.0 row explicitly includes ERNIE-5.0-Thinking-Latest. Its list rate is CNY 0.006/0.024 per 1K tokens at <=32K and CNY 0.01/0.04 above 32K through 128K. This records the higher CNY 0.01 input / CNY 0.04 output tier. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
43
+ "confidence": "inferred",
44
+ "observedAt": "2026-09-19T00:00:00Z"
45
+ },
46
+ "qianfan/ernie-5.1": {
47
+ "value": { "inputPerMTokUsd": 0.895837, "outputPerMTokUsd": 3.284737 },
48
+ "source": "official-doc",
49
+ "sourceRef": "https://cloud.baidu.com/doc/qianfan-docs/s/Jm8r1826a — ERNIE-5.1 online list rate is CNY 0.004/0.018 per 1K tokens at <=32K and CNY 0.006/0.022 above 32K through 128K. This records the higher CNY 0.006 input / CNY 0.022 output tier. Converted to USD/MTok using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
50
+ "confidence": "inferred",
51
+ "observedAt": "2026-09-19T00:00:00Z"
52
+ },
53
+ "qwen-cloud/qwen3.5-122b-a10b": {
54
+ "value": { "inputPerMTokUsd": 0.4, "outputPerMTokUsd": 3.2 },
55
+ "source": "official-doc",
56
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.5-122b-a10b — QwenCloud lists $0.40 input and $3.20 output per 1M tokens for this exact API model (retrieved 2026-09-19).",
57
+ "confidence": "declared",
58
+ "observedAt": "2026-09-19T00:00:00Z"
59
+ },
60
+ "qwen-cloud/qwen3.5-397b-a17b": {
61
+ "value": { "inputPerMTokUsd": 0.6, "outputPerMTokUsd": 3.6 },
62
+ "source": "official-doc",
63
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.5-397b-a17b — QwenCloud lists $0.60 input and $3.60 output per 1M tokens for this exact API model (retrieved 2026-09-19).",
64
+ "confidence": "declared",
65
+ "observedAt": "2026-09-19T00:00:00Z"
66
+ },
67
+ "qwen-cloud/qwen3.5-plus-2026-04-20": {
68
+ "value": { "inputPerMTokUsd": 0.4, "outputPerMTokUsd": 2.4, "cacheReadPerMTokUsd": 0.04, "cacheWritePerMTokUsd": 0.5 },
69
+ "source": "official-doc",
70
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.5-plus-2026-04-20 — QwenCloud's <=256K tier lists $0.40 input, $2.40 output, $0.50 explicit cache creation, and $0.04 explicit cache read per 1M tokens. The page also advertises a 1M context, but publishes no rate above 256K; ModelPricing therefore records the published <=256K rate only (retrieved 2026-09-19).",
71
+ "confidence": "declared",
72
+ "observedAt": "2026-09-19T00:00:00Z"
73
+ },
74
+ "qwen-cloud/qwen3.6-27b": {
75
+ "value": { "inputPerMTokUsd": 0.6, "outputPerMTokUsd": 3.6 },
76
+ "source": "official-doc",
77
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.6-27b — QwenCloud lists $0.60 input and $3.60 output per 1M tokens for this exact API model (retrieved 2026-09-19).",
78
+ "confidence": "declared",
79
+ "observedAt": "2026-09-19T00:00:00Z"
80
+ },
81
+ "qwen-cloud/qwen3.6-35b-a3b": {
82
+ "value": { "inputPerMTokUsd": 0.375, "outputPerMTokUsd": 2.25 },
83
+ "source": "official-doc",
84
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.6-35b-a3b — QwenCloud lists $0.375 input and $2.25 output per 1M tokens for this exact API model (retrieved 2026-09-19).",
85
+ "confidence": "declared",
86
+ "observedAt": "2026-09-19T00:00:00Z"
87
+ },
88
+ "qwen-cloud/qwen3.6-plus": {
89
+ "value": { "inputPerMTokUsd": 0.5, "outputPerMTokUsd": 3, "cacheReadPerMTokUsd": 0.05, "cacheWritePerMTokUsd": 0.625 },
90
+ "source": "official-doc",
91
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.6-plus — QwenCloud's <=256K tier lists $0.50 input, $3.00 output, $0.625 explicit cache creation, and $0.05 explicit cache read per 1M tokens. The page also advertises a 1M context, but publishes no rate above 256K; ModelPricing therefore records the published <=256K rate only (retrieved 2026-09-19).",
92
+ "confidence": "declared",
93
+ "observedAt": "2026-09-19T00:00:00Z"
94
+ },
95
+ "qwen-cloud/qwen3.7-max-2026-06-08": {
96
+ "value": { "inputPerMTokUsd": 2.5, "outputPerMTokUsd": 7.5, "cacheReadPerMTokUsd": 0.5, "cacheWritePerMTokUsd": 3.125 },
97
+ "source": "official-doc",
98
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.7-max-2026-06-08 — QwenCloud lists $2.50 input, $7.50 output, $0.50 implicit-cache input, and $3.125 explicit-cache creation per 1M tokens. It separately lists $0.25 explicit cache read; the one cacheRead field records the higher implicit-cache rate so estimates do not under-report either cache mode (retrieved 2026-09-19).",
99
+ "confidence": "declared",
100
+ "observedAt": "2026-09-19T00:00:00Z"
101
+ },
102
+ "qwen-cloud/qwen3.7-plus": {
103
+ "value": { "inputPerMTokUsd": 0.32, "outputPerMTokUsd": 1.28, "cacheReadPerMTokUsd": 0.064, "cacheWritePerMTokUsd": 0.4 },
104
+ "source": "official-doc",
105
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.7-plus — QwenCloud's current <=256K 20%-off price is $0.32 input, $1.28 output, $0.064 implicit-cache input, and $0.40 explicit-cache creation per 1M tokens. It separately lists $0.032 explicit cache read; the one cacheRead field records the higher implicit-cache rate. The undiscounted $0.40/$1.60 list price and the current promotion are both disclosed by the page (retrieved 2026-09-19).",
106
+ "confidence": "declared",
107
+ "observedAt": "2026-09-19T00:00:00Z"
108
+ },
109
+ "qwen-cloud/qwen3.8-max": {
110
+ "value": { "inputPerMTokUsd": 2, "outputPerMTokUsd": 6, "cacheReadPerMTokUsd": 0.25, "cacheWritePerMTokUsd": 2.5 },
111
+ "source": "official-doc",
112
+ "sourceRef": "https://www.qwencloud.com/models/qwen3.8-max — QwenCloud lists $2 input, $6 output, $0.25 implicit-cache input, and $2.50 explicit-cache creation per 1M tokens. It separately lists $0.17 explicit cache read; the one cacheRead field records the higher implicit-cache rate so estimates do not under-report either cache mode (retrieved 2026-09-19).",
113
+ "confidence": "declared",
114
+ "observedAt": "2026-09-19T00:00:00Z"
115
+ },
116
+ "alibaba/qwen3.5-122b-a10b": {
117
+ "value": { "inputPerMTokUsd": 0.438363, "outputPerMTokUsd": 3.506606 },
118
+ "source": "official-doc",
119
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.5-122b-a10b at CNY 2.936 input / CNY 23.486 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
120
+ "confidence": "inferred",
121
+ "observedAt": "2026-09-19T00:00:00Z"
122
+ },
123
+ "alibaba/qwen3.5-397b-a17b": {
124
+ "value": { "inputPerMTokUsd": 0.657545, "outputPerMTokUsd": 3.94482 },
125
+ "source": "official-doc",
126
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.5-397b-a17b at CNY 4.404 input / CNY 26.421 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
127
+ "confidence": "inferred",
128
+ "observedAt": "2026-09-19T00:00:00Z"
129
+ },
130
+ "alibaba/qwen3.5-plus": {
131
+ "value": { "inputPerMTokUsd": 0.547954, "outputPerMTokUsd": 3.287425 },
132
+ "source": "official-doc",
133
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.5-plus at CNY 2.936/17.614 per MTok through 256K and CNY 3.67/22.018 at 256K-1M. ModelPricing records the higher CNY 3.67 input / CNY 22.018 output tier so estimates do not under-report long prompts. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
134
+ "confidence": "inferred",
135
+ "observedAt": "2026-09-19T00:00:00Z"
136
+ },
137
+ "alibaba/qwen3.6-27b": {
138
+ "value": { "inputPerMTokUsd": 0.671358, "outputPerMTokUsd": 4.028151 },
139
+ "source": "official-doc",
140
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.6-27b at CNY 4.49652 input / CNY 26.97912 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
141
+ "confidence": "inferred",
142
+ "observedAt": "2026-09-19T00:00:00Z"
143
+ },
144
+ "alibaba/qwen3.6-35b-a3b": {
145
+ "value": { "inputPerMTokUsd": 0.419599, "outputPerMTokUsd": 2.517594 },
146
+ "source": "official-doc",
147
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.6-35b-a3b at CNY 2.810325 input / CNY 16.86195 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
148
+ "confidence": "inferred",
149
+ "observedAt": "2026-09-19T00:00:00Z"
150
+ },
151
+ "alibaba/qwen3.6-plus": {
152
+ "value": { "inputPerMTokUsd": 2.237862, "outputPerMTokUsd": 6.713555 },
153
+ "source": "official-doc",
154
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.6-plus at CNY 3.7471/22.4826 per MTok through 256K and CNY 14.9884/44.965 at 256K-1M. ModelPricing records the higher CNY 14.9884 input / CNY 44.965 output tier so estimates do not under-report long prompts. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
155
+ "confidence": "inferred",
156
+ "observedAt": "2026-09-19T00:00:00Z"
157
+ },
158
+ "alibaba/qwen3.7-max": {
159
+ "value": { "inputPerMTokUsd": 2.797402, "outputPerMTokUsd": 8.392056 },
160
+ "source": "official-doc",
161
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.7-max at CNY 18.736 input / CNY 56.207 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
162
+ "confidence": "inferred",
163
+ "observedAt": "2026-09-19T00:00:00Z"
164
+ },
165
+ "alibaba/qwen3.7-plus": {
166
+ "value": { "inputPerMTokUsd": 1.342711, "outputPerMTokUsd": 5.370844 },
167
+ "source": "official-doc",
168
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.7-plus at CNY 2.998/11.991 per MTok through 256K and CNY 8.993/35.972 at 256K-1M, marked as a limited-time 20%-off list price. ModelPricing records the higher CNY 8.993 input / CNY 35.972 output list tier so estimates do not under-report long prompts; it may over-report while the promotion applies. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
169
+ "confidence": "inferred",
170
+ "observedAt": "2026-09-19T00:00:00Z"
171
+ },
172
+ "alibaba/qwen3.8-max": {
173
+ "value": { "inputPerMTokUsd": 2.237802, "outputPerMTokUsd": 6.713555 },
174
+ "source": "official-doc",
175
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — Model Studio International (Singapore) lists qwen3.8-max at CNY 14.988 input / CNY 44.965 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
176
+ "confidence": "inferred",
177
+ "observedAt": "2026-09-19T00:00:00Z"
178
+ },
179
+ "alibaba-cn/qwen3.5-122b-a10b": {
180
+ "value": { "inputPerMTokUsd": 0.298612, "outputPerMTokUsd": 2.3889 },
181
+ "source": "official-doc",
182
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.5-122b-a10b at CNY 0.8/6.4 per MTok through 128K and CNY 2/16 at 128K-256K. ModelPricing records the higher CNY 2 input / CNY 16 output tier. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
183
+ "confidence": "inferred",
184
+ "observedAt": "2026-09-19T00:00:00Z"
185
+ },
186
+ "alibaba-cn/qwen3.5-397b-a17b": {
187
+ "value": { "inputPerMTokUsd": 0.447919, "outputPerMTokUsd": 2.687512 },
188
+ "source": "official-doc",
189
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.5-397b-a17b at CNY 1.2/7.2 per MTok through 128K and CNY 3/18 at 128K-256K. ModelPricing records the higher CNY 3 input / CNY 18 output tier. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
190
+ "confidence": "inferred",
191
+ "observedAt": "2026-09-19T00:00:00Z"
192
+ },
193
+ "alibaba-cn/qwen3.5-plus": {
194
+ "value": { "inputPerMTokUsd": 0.597225, "outputPerMTokUsd": 3.58335 },
195
+ "source": "official-doc",
196
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.5-plus at CNY 0.8/4.8 per MTok through 128K, CNY 2/12 through 256K, and CNY 4/24 at 256K-1M. ModelPricing records the higher CNY 4 input / CNY 24 output tier. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
197
+ "confidence": "inferred",
198
+ "observedAt": "2026-09-19T00:00:00Z"
199
+ },
200
+ "alibaba-cn/qwen3.6-27b": {
201
+ "value": { "inputPerMTokUsd": 0.447919, "outputPerMTokUsd": 2.687512 },
202
+ "source": "official-doc",
203
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.6-27b at CNY 3 input / CNY 18 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
204
+ "confidence": "inferred",
205
+ "observedAt": "2026-09-19T00:00:00Z"
206
+ },
207
+ "alibaba-cn/qwen3.6-35b-a3b": {
208
+ "value": { "inputPerMTokUsd": 0.268751, "outputPerMTokUsd": 1.612507 },
209
+ "source": "official-doc",
210
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.6-35b-a3b at CNY 1.8 input / CNY 10.8 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
211
+ "confidence": "inferred",
212
+ "observedAt": "2026-09-19T00:00:00Z"
213
+ },
214
+ "alibaba-cn/qwen3.6-plus": {
215
+ "value": { "inputPerMTokUsd": 1.19445, "outputPerMTokUsd": 7.166699 },
216
+ "source": "official-doc",
217
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.6-plus at CNY 2/12 per MTok through 256K and CNY 8/48 at 256K-1M. ModelPricing records the higher CNY 8 input / CNY 48 output tier. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
218
+ "confidence": "inferred",
219
+ "observedAt": "2026-09-19T00:00:00Z"
220
+ },
221
+ "alibaba-cn/qwen3.7-max": {
222
+ "value": { "inputPerMTokUsd": 1.791675, "outputPerMTokUsd": 5.375024 },
223
+ "source": "official-doc",
224
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.7-max at CNY 12 input / CNY 36 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
225
+ "confidence": "inferred",
226
+ "observedAt": "2026-09-19T00:00:00Z"
227
+ },
228
+ "alibaba-cn/qwen3.7-plus": {
229
+ "value": { "inputPerMTokUsd": 0.895837, "outputPerMTokUsd": 3.58335 },
230
+ "source": "official-doc",
231
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.7-plus at list CNY 2/8 per MTok through 256K and CNY 6/24 at 256K-1M, both marked as limited-time 20%-off. ModelPricing records the higher CNY 6 input / CNY 24 output list tier so estimates do not under-report long prompts; it may over-report while the promotion applies. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
232
+ "confidence": "inferred",
233
+ "observedAt": "2026-09-19T00:00:00Z"
234
+ },
235
+ "alibaba-cn/qwen3.8-max": {
236
+ "value": { "inputPerMTokUsd": 1.791675, "outputPerMTokUsd": 5.375024 },
237
+ "source": "official-doc",
238
+ "sourceRef": "https://help.aliyun.com/en/model-studio/model-pricing — China (Beijing) lists qwen3.8-max at CNY 12 input / CNY 36 output per MTok. Converted using the ECB 2026-09-18 reference rates (EUR 1 = USD 1.1460; CNY 7.6755): https://www.ecb.europa.eu/stats/shared/pdf/eurofxref.pdf?57eae2dc36bcb4af234d5e0b7e1a28c3= (retrieved 2026-09-19).",
239
+ "confidence": "inferred",
240
+ "observedAt": "2026-09-19T00:00:00Z"
241
+ }
242
+ }
243
+ }
@@ -884,6 +884,42 @@
884
884
  },
885
885
  "scope": "llm"
886
886
  },
887
+ {
888
+ "id": "kimi-coding-openai",
889
+ "displayName": "Kimi Code (OpenAI dialect)",
890
+ "$comment": "The Kimi Code membership endpoint supports both documented dialects. This separate subscription row is the OpenAI-compatible sibling of `kimi-coding`, not Moonshot's token-billed platform API.",
891
+ "protocols": [
892
+ "openai-chat-completions"
893
+ ],
894
+ "authKinds": [
895
+ "api-key"
896
+ ],
897
+ "defaultEndpoints": {
898
+ "api": "https://api.kimi.com/coding/v1"
899
+ },
900
+ "modelDiscovery": "none",
901
+ "liveCatalogAuthority": "unknown",
902
+ "adapterId": "winter.openai-chat-completions",
903
+ "family": "openai",
904
+ "upstream": {
905
+ "project": "winter",
906
+ "commit": "",
907
+ "sourcePaths": []
908
+ },
909
+ "risk": {
910
+ "class": "review-required",
911
+ "reasons": [
912
+ "Membership API-key access, plan-dependent model entitlements, and the OpenAI-compatible endpoint are documented, but no Winter live gate has exercised this subscription surface."
913
+ ]
914
+ },
915
+ "pricingBasis": "subscription",
916
+ "admission": {
917
+ "basis": "api-key",
918
+ "citation": "https://www.kimi.com/code/docs/en/kimi-code/models.html — Kimi Code's current model configuration documents the OpenAI-compatible Coding endpoint at https://api.kimi.com/coding/v1, membership API keys, and third-party tool use (retrieved 2026-09-19).",
919
+ "tier": "fetched-document"
920
+ },
921
+ "scope": "llm"
922
+ },
887
923
  {
888
924
  "id": "lemonade",
889
925
  "displayName": "Lemonade Server (local)",
@@ -2273,16 +2309,16 @@
2273
2309
  },
2274
2310
  {
2275
2311
  "id": "zai",
2276
- "displayName": "Z.AI GLM Coding (OpenAI dialect)",
2312
+ "displayName": "Z.AI API (OpenAI dialect)",
2277
2313
  "protocols": [
2278
2314
  "openai-chat-completions"
2279
2315
  ],
2280
- "$comment": "R6b-5 (WS-13b §2): a DIALECT SIBLING. The vendor documents two wire dialects at two different base URLs, and a provider row carries exactly one `adapterId`, so two dialects are two rows — each with its own `defaultEndpoints.api`, its own model keys, and the dialect in its `displayName`. `winter.anthropic-messages` is multi-provider in fact, measured rather than assumed: `runtime/src/provider/catalog-endpoint-shape.test.ts` drives two sibling rows against a loopback fake and each reaches ITS OWN `<root>/v1/messages`. A WINTER-OWNED row (`upstream.project: \"winter\"`, empty commit): it is not extracted, so nothing here came from the pinned tree's executable half. The endpoint and dialect are the VENDOR'S OWN documented values, fetched and read on 2026-09-06 — the `fetched-document` citation tier (PROVENANCE.md). The OPENAI half of z.ai's GLM Coding Plan. Its Anthropic sibling is the EXTRACTED `zai-anthropic` row — upstream's own `zai` entry is `format: \"claude\"` at `https://api.z.ai/api/anthropic/v1/messages`, which z.ai's Claude-client doc confirms verbatim, so the allowlist admits that id under the sibling id rather than under `zai` (the `gemini` -> `google` precedent). Upstream's `glm`/`glm-cn`/`glmt` ids reach this same endpoint through a bespoke `glm` executor Winter does not represent; they are in `blocked` pointing here.",
2316
+ "$comment": "Normal Z.AI token-billed API. This is deliberately distinct from `zai-coding`, whose Coding Plan key uses the different `/api/coding/paas/v4` endpoint and spends subscription credits rather than API balance. The provider documentation shows this ordinary OpenAI-compatible endpoint in its GLM-5.2 quick start; its pricing table is the basis for the token rates on this provider's models. The Anthropic-compatible route remains separately represented as `zai-anthropic` because adapters are dialect-specific.",
2281
2317
  "authKinds": [
2282
2318
  "api-key"
2283
2319
  ],
2284
2320
  "defaultEndpoints": {
2285
- "api": "https://api.z.ai/api/coding/paas/v4"
2321
+ "api": "https://api.z.ai/api/paas/v4"
2286
2322
  },
2287
2323
  "modelDiscovery": "openai-models",
2288
2324
  "liveCatalogAuthority": "unknown",
@@ -2296,13 +2332,87 @@
2296
2332
  "risk": {
2297
2333
  "class": "review-required",
2298
2334
  "reasons": [
2299
- "admission is fetched-document, but no live gate case has run against this endpoint; the model roster is the sibling's and is unverified on this dialect"
2335
+ "admission is fetched-document, but no live gate case has run against this endpoint"
2300
2336
  ]
2301
2337
  },
2302
2338
  "pricingBasis": "token",
2303
2339
  "admission": {
2304
2340
  "basis": "api-key",
2305
- "citation": "https://docs.z.ai/devpack/tool/others — z.ai's own GLM Coding Plan docs, retrieved 2026-09-06: the \"Coding Endpoint\" table lists `https://api.z.ai/api/coding/paas/v4` under \"OpenAI Chat Completions\", and the Cline configuration example repeats it.",
2341
+ "citation": "https://docs.z.ai/guides/llm/glm-5.2 — Z.AI's own normal API quick start, retrieved 2026-09-19: the OpenAI-compatible client uses an API key and base URL `https://api.z.ai/api/paas/v4/`; the separately documented Coding Plan endpoint is not used here.",
2342
+ "tier": "fetched-document"
2343
+ },
2344
+ "scope": "llm"
2345
+ },
2346
+ {
2347
+ "id": "zai-coding",
2348
+ "displayName": "Z.AI GLM Coding Plan (OpenAI dialect)",
2349
+ "protocols": [
2350
+ "openai-chat-completions"
2351
+ ],
2352
+ "$comment": "The GLM Coding Plan is a distinct subscription entitlement, not the normal Z.AI API balance. It uses a Coding Plan API key and `https://api.z.ai/api/coding/paas/v4`; model usage consumes plan credits subject to five-hour and weekly limits. `pricingBasis: subscription` intentionally keeps every model's token price undefined: ModelPricing cannot represent the plan's credits, rolling quotas, peak/off-peak discount, or plan tier. The provider's current plan documentation lists GLM-5.3 and GLM-5.3-Flash as the effective models; documented legacy ids are aliases on those two rows rather than separately priced models.",
2353
+ "authKinds": [
2354
+ "api-key"
2355
+ ],
2356
+ "defaultEndpoints": {
2357
+ "api": "https://api.z.ai/api/coding/paas/v4"
2358
+ },
2359
+ "modelDiscovery": "openai-models",
2360
+ "liveCatalogAuthority": "unknown",
2361
+ "adapterId": "winter.openai-chat-completions",
2362
+ "family": "openai",
2363
+ "upstream": {
2364
+ "project": "winter",
2365
+ "commit": "",
2366
+ "sourcePaths": []
2367
+ },
2368
+ "risk": {
2369
+ "class": "review-required",
2370
+ "reasons": [
2371
+ "The vendor limits subscription benefits to its listed tools and environments; although its current list names Codex, this application has not run a live plan-key gate against the endpoint.",
2372
+ "Plan entitlement, context, and rate limits vary by current vendor plan and are not discoverable from the normal API's `/models` response."
2373
+ ]
2374
+ },
2375
+ "pricingBasis": "subscription",
2376
+ "admission": {
2377
+ "basis": "api-key",
2378
+ "citation": "https://docs.z.ai/devpack/overview and https://docs.z.ai/devpack/tool/others — Z.AI's current GLM Coding Plan documentation, retrieved 2026-09-19: it is a subscription package; the OpenAI Chat Completions Coding Endpoint is `https://api.z.ai/api/coding/paas/v4`; supported tools include Codex; and plan usage is governed by five-hour and weekly credit limits.",
2379
+ "tier": "fetched-document"
2380
+ },
2381
+ "scope": "llm"
2382
+ },
2383
+ {
2384
+ "id": "zai-coding-anthropic",
2385
+ "displayName": "Z.AI GLM Coding Plan (Anthropic dialect)",
2386
+ "protocols": [
2387
+ "anthropic-messages"
2388
+ ],
2389
+ "$comment": "Dialect sibling of `zai-coding`. Z.AI's current GLM Coding Plan documentation explicitly supports Anthropic Messages at this root and gives a Claude Code configuration using the same plan API key. It is still subscription-priced: the plan's credits and limits are neither per-token list prices nor representable by ModelPricing.",
2390
+ "authKinds": [
2391
+ "api-key"
2392
+ ],
2393
+ "defaultEndpoints": {
2394
+ "api": "https://api.z.ai/api/anthropic"
2395
+ },
2396
+ "modelDiscovery": "none",
2397
+ "liveCatalogAuthority": "unknown",
2398
+ "adapterId": "winter.anthropic-messages",
2399
+ "family": "anthropic",
2400
+ "upstream": {
2401
+ "project": "winter",
2402
+ "commit": "",
2403
+ "sourcePaths": []
2404
+ },
2405
+ "risk": {
2406
+ "class": "review-required",
2407
+ "reasons": [
2408
+ "The vendor documents the Claude Code configuration but this application has not run a live plan-key gate against the Anthropic-compatible endpoint.",
2409
+ "The published OpenAI Chat reasoning replay contract must not be assumed for Anthropic Messages without a dialect-specific document or probe."
2410
+ ]
2411
+ },
2412
+ "pricingBasis": "subscription",
2413
+ "admission": {
2414
+ "basis": "api-key",
2415
+ "citation": "https://docs.z.ai/devpack/tool/claude and https://docs.z.ai/devpack/tool/others — Z.AI's current GLM Coding Plan documentation, retrieved 2026-09-19: the Claude Code configuration uses `ANTHROPIC_BASE_URL=https://api.z.ai/api/anthropic`; the Coding Plan supports the Anthropic protocol; and the plan is subscription credit based.",
2306
2416
  "tier": "fetched-document"
2307
2417
  },
2308
2418
  "scope": "llm"
@@ -2415,6 +2525,78 @@
2415
2525
  "tier": "audit"
2416
2526
  },
2417
2527
  "scope": "llm"
2528
+ },
2529
+ {
2530
+ "id": "minimax-token-plan",
2531
+ "displayName": "MiniMax Token Plan (OpenAI dialect)",
2532
+ "$comment": "Subscription-only sibling of `minimax`: MiniMax documents a Token Plan Subscription Key that is distinct from a pay-as-you-go API key. The shared OpenAI-compatible endpoint does not make the two billing boundaries interchangeable.",
2533
+ "protocols": [
2534
+ "openai-chat-completions"
2535
+ ],
2536
+ "authKinds": [
2537
+ "api-key"
2538
+ ],
2539
+ "defaultEndpoints": {
2540
+ "api": "https://api.minimax.io/v1"
2541
+ },
2542
+ "modelDiscovery": "none",
2543
+ "liveCatalogAuthority": "unknown",
2544
+ "adapterId": "winter.openai-chat-completions",
2545
+ "family": "openai",
2546
+ "upstream": {
2547
+ "project": "winter",
2548
+ "commit": "",
2549
+ "sourcePaths": []
2550
+ },
2551
+ "risk": {
2552
+ "class": "review-required",
2553
+ "reasons": [
2554
+ "MiniMax documents plan-key use in OpenAI-compatible tools, but Winter has not run a subscription-key live gate or verified entitlement differences among plan tiers."
2555
+ ]
2556
+ },
2557
+ "pricingBasis": "subscription",
2558
+ "admission": {
2559
+ "basis": "api-key",
2560
+ "citation": "https://platform.minimax.io/subscribe/coding-plan and https://platform.minimax.io/protocol/paid-agreement — MiniMax's Token Plan page documents use with MiniMax Code and OpenAI-compatible tools; its terms state that a Token Plan Subscription Key and normal pay-as-you-go API Key are separate (retrieved 2026-09-19).",
2561
+ "tier": "fetched-document"
2562
+ },
2563
+ "scope": "llm"
2564
+ },
2565
+ {
2566
+ "id": "novita-coding",
2567
+ "displayName": "Novita Coding Plan",
2568
+ "$comment": "Subscription-only sibling of `novita`: the Coding Plan publishes a fixed nine-model allowance on Novita's normal OpenAI-compatible API surface. It has no per-token price fields because plan credits, not ordinary API metering, govern it.",
2569
+ "protocols": [
2570
+ "openai-chat-completions"
2571
+ ],
2572
+ "authKinds": [
2573
+ "api-key"
2574
+ ],
2575
+ "defaultEndpoints": {
2576
+ "api": "https://api.novita.ai/openai/v1"
2577
+ },
2578
+ "modelDiscovery": "none",
2579
+ "liveCatalogAuthority": "unknown",
2580
+ "adapterId": "winter.openai-chat-completions",
2581
+ "family": "openai",
2582
+ "upstream": {
2583
+ "project": "winter",
2584
+ "commit": "",
2585
+ "sourcePaths": []
2586
+ },
2587
+ "risk": {
2588
+ "class": "review-required",
2589
+ "reasons": [
2590
+ "Novita documents the Coding Plan model set and its API-key OpenAI-compatible surface, but Winter has not run a plan-key live gate or established tier-specific rate limits."
2591
+ ]
2592
+ },
2593
+ "pricingBasis": "subscription",
2594
+ "admission": {
2595
+ "basis": "api-key",
2596
+ "citation": "https://novita.ai/coding-plan and https://docs.novita.ai/guides/introduction — Novita's Coding Plan lists its nine included models; Novita's current docs declare Bearer API-key authentication and the OpenAI-compatible base https://api.novita.ai/openai (retrieved 2026-09-19).",
2597
+ "tier": "fetched-document"
2598
+ },
2599
+ "scope": "llm"
2418
2600
  }
2419
2601
  ]
2420
2602
  }
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@yanlinglabs/winter-provider-catalog",
3
- "version": "0.0.20",
3
+ "version": "0.0.22",
4
4
  "license": "MIT",
5
5
  "type": "module",
6
6
  "engines": {