modelparams 0.0.13__tar.gz → 0.0.15__tar.gz
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- {modelparams-0.0.13 → modelparams-0.0.15}/PKG-INFO +1 -1
- {modelparams-0.0.13 → modelparams-0.0.15}/pyproject.toml +1 -1
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/catalog.json +195 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/model_ids.py +6 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/registry.py +3 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/alibaba.py +43 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/.gitignore +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/LICENSE +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/README.md +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/__init__.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/__init__.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/catalog.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/models.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/py.typed +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/__init__.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/anthropic.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/cerebras.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/cohere.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/deepseek.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/google.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/groq.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/meta.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/minimax.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/mistral.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/moonshot.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/nvidia.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/openai.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/perplexity.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/thinking_machines.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/xai.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/xiaomi.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/z_ai.py +0 -0
- {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/validation.py +0 -0
|
@@ -1,6 +1,6 @@
|
|
|
1
1
|
Metadata-Version: 2.5
|
|
2
2
|
Name: modelparams
|
|
3
|
-
Version: 0.0.
|
|
3
|
+
Version: 0.0.15
|
|
4
4
|
Summary: Typed model parameters for Python, generated from the modelparams.dev catalog.
|
|
5
5
|
Project-URL: Homepage, https://modelparams.dev
|
|
6
6
|
Project-URL: Repository, https://github.com/mnfst/modelparameters.dev
|
|
@@ -59,6 +59,66 @@
|
|
|
59
59
|
}
|
|
60
60
|
]
|
|
61
61
|
},
|
|
62
|
+
{
|
|
63
|
+
"provider": "alibaba",
|
|
64
|
+
"authType": "api_key",
|
|
65
|
+
"model": "qwen-max",
|
|
66
|
+
"params": [
|
|
67
|
+
{
|
|
68
|
+
"path": "max_tokens",
|
|
69
|
+
"label": "Max tokens",
|
|
70
|
+
"description": "Maximum number of output tokens the model may generate.",
|
|
71
|
+
"group": "generation_length",
|
|
72
|
+
"type": "integer",
|
|
73
|
+
"range": {
|
|
74
|
+
"min": 1
|
|
75
|
+
}
|
|
76
|
+
},
|
|
77
|
+
{
|
|
78
|
+
"path": "temperature",
|
|
79
|
+
"label": "Temperature",
|
|
80
|
+
"description": "Controls randomness. Lower values make outputs more focused; higher values make them more varied.",
|
|
81
|
+
"group": "sampling",
|
|
82
|
+
"type": "number",
|
|
83
|
+
"range": {
|
|
84
|
+
"min": 0,
|
|
85
|
+
"max": 1.9,
|
|
86
|
+
"step": 0.1
|
|
87
|
+
}
|
|
88
|
+
},
|
|
89
|
+
{
|
|
90
|
+
"path": "top_p",
|
|
91
|
+
"label": "Top P",
|
|
92
|
+
"description": "Controls nucleus sampling by limiting generation to tokens within the selected cumulative probability.",
|
|
93
|
+
"group": "sampling",
|
|
94
|
+
"type": "number",
|
|
95
|
+
"range": {
|
|
96
|
+
"min": 0,
|
|
97
|
+
"max": 1,
|
|
98
|
+
"step": 0.01
|
|
99
|
+
}
|
|
100
|
+
},
|
|
101
|
+
{
|
|
102
|
+
"path": "extra_body.top_k",
|
|
103
|
+
"label": "Top K",
|
|
104
|
+
"description": "Limits generation to the selected number of highest-probability tokens.",
|
|
105
|
+
"group": "sampling",
|
|
106
|
+
"type": "integer",
|
|
107
|
+
"default": 20,
|
|
108
|
+
"range": {
|
|
109
|
+
"min": 1
|
|
110
|
+
}
|
|
111
|
+
},
|
|
112
|
+
{
|
|
113
|
+
"path": "extra_body.chat_template_kwargs.enable_thinking",
|
|
114
|
+
"label": "Enable thinking",
|
|
115
|
+
"description": "Controls Qwen3 thinking mode when using OpenAI-compatible clients that pass provider-specific extra body fields.",
|
|
116
|
+
"group": "reasoning",
|
|
117
|
+
"type": "boolean",
|
|
118
|
+
"default": true
|
|
119
|
+
}
|
|
120
|
+
]
|
|
121
|
+
},
|
|
62
122
|
{
|
|
63
123
|
"provider": "alibaba",
|
|
64
124
|
"authType": "api_key",
|
|
@@ -119,6 +179,66 @@
|
|
|
119
179
|
}
|
|
120
180
|
]
|
|
121
181
|
},
|
|
182
|
+
{
|
|
183
|
+
"provider": "alibaba",
|
|
184
|
+
"authType": "api_key",
|
|
185
|
+
"model": "qwen-turbo",
|
|
186
|
+
"params": [
|
|
187
|
+
{
|
|
188
|
+
"path": "max_tokens",
|
|
189
|
+
"label": "Max tokens",
|
|
190
|
+
"description": "Maximum number of output tokens the model may generate.",
|
|
191
|
+
"group": "generation_length",
|
|
192
|
+
"type": "integer",
|
|
193
|
+
"range": {
|
|
194
|
+
"min": 1
|
|
195
|
+
}
|
|
196
|
+
},
|
|
197
|
+
{
|
|
198
|
+
"path": "temperature",
|
|
199
|
+
"label": "Temperature",
|
|
200
|
+
"description": "Controls randomness. Lower values make outputs more focused; higher values make them more varied.",
|
|
201
|
+
"group": "sampling",
|
|
202
|
+
"type": "number",
|
|
203
|
+
"range": {
|
|
204
|
+
"min": 0,
|
|
205
|
+
"max": 1.9,
|
|
206
|
+
"step": 0.1
|
|
207
|
+
}
|
|
208
|
+
},
|
|
209
|
+
{
|
|
210
|
+
"path": "top_p",
|
|
211
|
+
"label": "Top P",
|
|
212
|
+
"description": "Controls nucleus sampling by limiting generation to tokens within the selected cumulative probability.",
|
|
213
|
+
"group": "sampling",
|
|
214
|
+
"type": "number",
|
|
215
|
+
"range": {
|
|
216
|
+
"min": 0,
|
|
217
|
+
"max": 1,
|
|
218
|
+
"step": 0.01
|
|
219
|
+
}
|
|
220
|
+
},
|
|
221
|
+
{
|
|
222
|
+
"path": "extra_body.top_k",
|
|
223
|
+
"label": "Top K",
|
|
224
|
+
"description": "Limits generation to the selected number of highest-probability tokens.",
|
|
225
|
+
"group": "sampling",
|
|
226
|
+
"type": "integer",
|
|
227
|
+
"default": 20,
|
|
228
|
+
"range": {
|
|
229
|
+
"min": 1
|
|
230
|
+
}
|
|
231
|
+
},
|
|
232
|
+
{
|
|
233
|
+
"path": "extra_body.chat_template_kwargs.enable_thinking",
|
|
234
|
+
"label": "Enable thinking",
|
|
235
|
+
"description": "Controls Qwen3 thinking mode when using OpenAI-compatible clients that pass provider-specific extra body fields.",
|
|
236
|
+
"group": "reasoning",
|
|
237
|
+
"type": "boolean",
|
|
238
|
+
"default": true
|
|
239
|
+
}
|
|
240
|
+
]
|
|
241
|
+
},
|
|
122
242
|
{
|
|
123
243
|
"provider": "alibaba",
|
|
124
244
|
"authType": "api_key",
|
|
@@ -628,6 +748,81 @@
|
|
|
628
748
|
}
|
|
629
749
|
]
|
|
630
750
|
},
|
|
751
|
+
{
|
|
752
|
+
"provider": "alibaba",
|
|
753
|
+
"authType": "api_key",
|
|
754
|
+
"model": "qwen3.8-2.4t-a95b",
|
|
755
|
+
"params": [
|
|
756
|
+
{
|
|
757
|
+
"path": "max_completion_tokens",
|
|
758
|
+
"label": "Max tokens",
|
|
759
|
+
"description": "Maximum number of tokens to generate, including both reasoning and the final answer.",
|
|
760
|
+
"group": "generation_length",
|
|
761
|
+
"type": "integer",
|
|
762
|
+
"range": {
|
|
763
|
+
"min": 1
|
|
764
|
+
}
|
|
765
|
+
},
|
|
766
|
+
{
|
|
767
|
+
"path": "temperature",
|
|
768
|
+
"label": "Temperature",
|
|
769
|
+
"description": "Controls randomness. Lower values make outputs more focused; higher values make them more varied.",
|
|
770
|
+
"group": "sampling",
|
|
771
|
+
"type": "number",
|
|
772
|
+
"range": {
|
|
773
|
+
"min": 0,
|
|
774
|
+
"max": 2,
|
|
775
|
+
"step": 0.1
|
|
776
|
+
}
|
|
777
|
+
},
|
|
778
|
+
{
|
|
779
|
+
"path": "top_p",
|
|
780
|
+
"label": "Top P",
|
|
781
|
+
"description": "Controls nucleus sampling by limiting generation to tokens within the selected cumulative probability.",
|
|
782
|
+
"group": "sampling",
|
|
783
|
+
"type": "number",
|
|
784
|
+
"range": {
|
|
785
|
+
"min": 0,
|
|
786
|
+
"max": 1,
|
|
787
|
+
"step": 0.01
|
|
788
|
+
}
|
|
789
|
+
},
|
|
790
|
+
{
|
|
791
|
+
"path": "extra_body.top_k",
|
|
792
|
+
"label": "Top K",
|
|
793
|
+
"description": "Limits generation to the selected number of highest-probability tokens. Values above 100 disable top-k sampling.",
|
|
794
|
+
"group": "sampling",
|
|
795
|
+
"type": "integer",
|
|
796
|
+
"default": 20,
|
|
797
|
+
"range": {
|
|
798
|
+
"min": 0
|
|
799
|
+
}
|
|
800
|
+
},
|
|
801
|
+
{
|
|
802
|
+
"path": "extra_body.enable_thinking",
|
|
803
|
+
"label": "Enable thinking",
|
|
804
|
+
"description": "Toggles the model's hybrid thinking mode, sent as a provider-specific extra body field on OpenAI-compatible clients.",
|
|
805
|
+
"group": "reasoning",
|
|
806
|
+
"type": "boolean",
|
|
807
|
+
"default": true
|
|
808
|
+
},
|
|
809
|
+
{
|
|
810
|
+
"path": "extra_body.thinking_budget",
|
|
811
|
+
"label": "Thinking budget",
|
|
812
|
+
"description": "Maximum number of tokens the model may spend on reasoning before it starts the final answer; defaults to the model's maximum reasoning length.",
|
|
813
|
+
"group": "reasoning",
|
|
814
|
+
"applicability": {
|
|
815
|
+
"only": {
|
|
816
|
+
"extra_body.enable_thinking": true
|
|
817
|
+
}
|
|
818
|
+
},
|
|
819
|
+
"type": "integer",
|
|
820
|
+
"range": {
|
|
821
|
+
"min": 1
|
|
822
|
+
}
|
|
823
|
+
}
|
|
824
|
+
]
|
|
825
|
+
},
|
|
631
826
|
{
|
|
632
827
|
"provider": "alibaba",
|
|
633
828
|
"authType": "api_key",
|
|
@@ -5,7 +5,9 @@ from typing import Literal
|
|
|
5
5
|
|
|
6
6
|
ModelId = Literal[
|
|
7
7
|
"alibaba/qwen-flash",
|
|
8
|
+
"alibaba/qwen-max",
|
|
8
9
|
"alibaba/qwen-plus",
|
|
10
|
+
"alibaba/qwen-turbo",
|
|
9
11
|
"alibaba/qwen3-coder-flash",
|
|
10
12
|
"alibaba/qwen3-coder-plus",
|
|
11
13
|
"alibaba/qwen3-max",
|
|
@@ -14,6 +16,7 @@ ModelId = Literal[
|
|
|
14
16
|
"alibaba/qwen3.6-flash",
|
|
15
17
|
"alibaba/qwen3.7-max",
|
|
16
18
|
"alibaba/qwen3.7-plus",
|
|
19
|
+
"alibaba/qwen3.8-2.4t-a95b",
|
|
17
20
|
"alibaba/qwen3.8-max",
|
|
18
21
|
"alibaba/qwq-plus",
|
|
19
22
|
"anthropic/claude-3-5-haiku-20241022",
|
|
@@ -263,7 +266,9 @@ ModelId = Literal[
|
|
|
263
266
|
|
|
264
267
|
MODEL_IDS: tuple[ModelId, ...] = (
|
|
265
268
|
"alibaba/qwen-flash",
|
|
269
|
+
"alibaba/qwen-max",
|
|
266
270
|
"alibaba/qwen-plus",
|
|
271
|
+
"alibaba/qwen-turbo",
|
|
267
272
|
"alibaba/qwen3-coder-flash",
|
|
268
273
|
"alibaba/qwen3-coder-plus",
|
|
269
274
|
"alibaba/qwen3-max",
|
|
@@ -272,6 +277,7 @@ MODEL_IDS: tuple[ModelId, ...] = (
|
|
|
272
277
|
"alibaba/qwen3.6-flash",
|
|
273
278
|
"alibaba/qwen3.7-max",
|
|
274
279
|
"alibaba/qwen3.7-plus",
|
|
280
|
+
"alibaba/qwen3.8-2.4t-a95b",
|
|
275
281
|
"alibaba/qwen3.8-max",
|
|
276
282
|
"alibaba/qwq-plus",
|
|
277
283
|
"anthropic/claude-3-5-haiku-20241022",
|
|
@@ -26,7 +26,9 @@ from .model_ids import ModelId
|
|
|
26
26
|
|
|
27
27
|
PARAM_TYPES: dict[ModelId, Any] = {
|
|
28
28
|
"alibaba/qwen-flash": alibaba.Qwen_FlashParams,
|
|
29
|
+
"alibaba/qwen-max": alibaba.Qwen_MaxParams,
|
|
29
30
|
"alibaba/qwen-plus": alibaba.Qwen_PlusParams,
|
|
31
|
+
"alibaba/qwen-turbo": alibaba.Qwen_TurboParams,
|
|
30
32
|
"alibaba/qwen3-coder-flash": alibaba.Qwen3_Coder_FlashParams,
|
|
31
33
|
"alibaba/qwen3-coder-plus": alibaba.Qwen3_Coder_PlusParams,
|
|
32
34
|
"alibaba/qwen3-max": alibaba.Qwen3_MaxParams,
|
|
@@ -35,6 +37,7 @@ PARAM_TYPES: dict[ModelId, Any] = {
|
|
|
35
37
|
"alibaba/qwen3.6-flash": alibaba.Qwen3_6_FlashParams,
|
|
36
38
|
"alibaba/qwen3.7-max": alibaba.Qwen3_7_MaxParams,
|
|
37
39
|
"alibaba/qwen3.7-plus": alibaba.Qwen3_7_PlusParams,
|
|
40
|
+
"alibaba/qwen3.8-2.4t-a95b": alibaba.Qwen3_8_2_4t_A95bParams,
|
|
38
41
|
"alibaba/qwen3.8-max": alibaba.Qwen3_8_MaxParams,
|
|
39
42
|
"alibaba/qwq-plus": alibaba.Qwq_PlusParams,
|
|
40
43
|
"anthropic/claude-3-5-haiku-20241022": anthropic.Claude_3_5_Haiku_20241022Params,
|
|
@@ -23,6 +23,19 @@ Qwen_FlashParams = TypedDict(
|
|
|
23
23
|
)
|
|
24
24
|
setattr(Qwen_FlashParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
25
25
|
|
|
26
|
+
Qwen_MaxParams = TypedDict(
|
|
27
|
+
"Qwen_MaxParams",
|
|
28
|
+
{
|
|
29
|
+
"max_tokens": Annotated[int, Field(ge=1)],
|
|
30
|
+
"temperature": Annotated[float, Field(ge=0, le=1.9)],
|
|
31
|
+
"top_p": Annotated[float, Field(ge=0, le=1)],
|
|
32
|
+
"extra_body.top_k": Annotated[int, Field(ge=1)],
|
|
33
|
+
"extra_body.chat_template_kwargs.enable_thinking": bool,
|
|
34
|
+
},
|
|
35
|
+
total=False,
|
|
36
|
+
)
|
|
37
|
+
setattr(Qwen_MaxParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
38
|
+
|
|
26
39
|
Qwen_PlusParams = TypedDict(
|
|
27
40
|
"Qwen_PlusParams",
|
|
28
41
|
{
|
|
@@ -36,6 +49,19 @@ Qwen_PlusParams = TypedDict(
|
|
|
36
49
|
)
|
|
37
50
|
setattr(Qwen_PlusParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
38
51
|
|
|
52
|
+
Qwen_TurboParams = TypedDict(
|
|
53
|
+
"Qwen_TurboParams",
|
|
54
|
+
{
|
|
55
|
+
"max_tokens": Annotated[int, Field(ge=1)],
|
|
56
|
+
"temperature": Annotated[float, Field(ge=0, le=1.9)],
|
|
57
|
+
"top_p": Annotated[float, Field(ge=0, le=1)],
|
|
58
|
+
"extra_body.top_k": Annotated[int, Field(ge=1)],
|
|
59
|
+
"extra_body.chat_template_kwargs.enable_thinking": bool,
|
|
60
|
+
},
|
|
61
|
+
total=False,
|
|
62
|
+
)
|
|
63
|
+
setattr(Qwen_TurboParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
64
|
+
|
|
39
65
|
Qwen3_Coder_FlashParams = TypedDict(
|
|
40
66
|
"Qwen3_Coder_FlashParams",
|
|
41
67
|
{
|
|
@@ -141,6 +167,20 @@ Qwen3_7_PlusParams = TypedDict(
|
|
|
141
167
|
)
|
|
142
168
|
setattr(Qwen3_7_PlusParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
143
169
|
|
|
170
|
+
Qwen3_8_2_4t_A95bParams = TypedDict(
|
|
171
|
+
"Qwen3_8_2_4t_A95bParams",
|
|
172
|
+
{
|
|
173
|
+
"max_completion_tokens": Annotated[int, Field(ge=1)],
|
|
174
|
+
"temperature": Annotated[float, Field(ge=0, le=2)],
|
|
175
|
+
"top_p": Annotated[float, Field(ge=0, le=1)],
|
|
176
|
+
"extra_body.top_k": Annotated[int, Field(ge=0)],
|
|
177
|
+
"extra_body.enable_thinking": bool,
|
|
178
|
+
"extra_body.thinking_budget": Annotated[int, Field(ge=1)],
|
|
179
|
+
},
|
|
180
|
+
total=False,
|
|
181
|
+
)
|
|
182
|
+
setattr(Qwen3_8_2_4t_A95bParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
183
|
+
|
|
144
184
|
Qwen3_8_MaxParams = TypedDict(
|
|
145
185
|
"Qwen3_8_MaxParams",
|
|
146
186
|
{
|
|
@@ -169,7 +209,9 @@ setattr(Qwq_PlusParams, "__pydantic_config__", _PARAMS_CONFIG)
|
|
|
169
209
|
|
|
170
210
|
__all__ = [
|
|
171
211
|
"Qwen_FlashParams",
|
|
212
|
+
"Qwen_MaxParams",
|
|
172
213
|
"Qwen_PlusParams",
|
|
214
|
+
"Qwen_TurboParams",
|
|
173
215
|
"Qwen3_Coder_FlashParams",
|
|
174
216
|
"Qwen3_Coder_PlusParams",
|
|
175
217
|
"Qwen3_MaxParams",
|
|
@@ -178,6 +220,7 @@ __all__ = [
|
|
|
178
220
|
"Qwen3_6_FlashParams",
|
|
179
221
|
"Qwen3_7_MaxParams",
|
|
180
222
|
"Qwen3_7_PlusParams",
|
|
223
|
+
"Qwen3_8_2_4t_A95bParams",
|
|
181
224
|
"Qwen3_8_MaxParams",
|
|
182
225
|
"Qwq_PlusParams",
|
|
183
226
|
]
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|