mindstudio-probe 1.0.3__py3-none-any.whl → 1.0.4__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (262) hide show
  1. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/LICENSE +201 -201
  2. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/METADATA +36 -34
  3. mindstudio_probe-1.0.4.dist-info/RECORD +276 -0
  4. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/WHEEL +1 -1
  5. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/entry_points.txt +1 -0
  6. msprobe/README.md +101 -237
  7. msprobe/{config/config.json → config.json} +49 -49
  8. msprobe/core/advisor/advisor.py +124 -124
  9. msprobe/core/advisor/advisor_const.py +59 -59
  10. msprobe/core/advisor/advisor_result.py +58 -58
  11. msprobe/core/common/const.py +341 -318
  12. msprobe/core/common/exceptions.py +99 -99
  13. msprobe/core/common/{file_check.py → file_utils.py} +478 -283
  14. msprobe/core/common/log.py +76 -69
  15. msprobe/core/common/utils.py +385 -616
  16. msprobe/core/common_config.py +85 -71
  17. msprobe/core/compare/acc_compare.py +299 -298
  18. msprobe/core/compare/check.py +95 -95
  19. msprobe/core/compare/compare_cli.py +49 -49
  20. msprobe/core/compare/highlight.py +223 -222
  21. msprobe/core/compare/multiprocessing_compute.py +149 -149
  22. msprobe/core/compare/npy_compare.py +295 -295
  23. msprobe/core/compare/utils.py +430 -429
  24. msprobe/core/data_dump/data_collector.py +154 -144
  25. msprobe/core/data_dump/data_processor/base.py +314 -293
  26. msprobe/core/data_dump/data_processor/factory.py +59 -59
  27. msprobe/core/data_dump/data_processor/mindspore_processor.py +186 -198
  28. msprobe/core/data_dump/data_processor/pytorch_processor.py +366 -389
  29. msprobe/core/data_dump/json_writer.py +96 -116
  30. msprobe/core/data_dump/scope.py +178 -178
  31. msprobe/core/grad_probe/constant.py +70 -70
  32. msprobe/core/grad_probe/grad_compare.py +171 -175
  33. msprobe/core/grad_probe/utils.py +64 -52
  34. msprobe/docs/01.installation.md +89 -0
  35. msprobe/docs/02.config_introduction.md +165 -0
  36. msprobe/docs/03.config_examples.md +247 -0
  37. msprobe/docs/04.acl_config_examples.md +76 -0
  38. msprobe/docs/05.data_dump_PyTorch.md +198 -0
  39. msprobe/docs/06.data_dump_MindSpore.md +243 -0
  40. msprobe/docs/07.accuracy_checker_PyTorch.md +274 -0
  41. msprobe/docs/08.accuracy_checker_online_PyTorch.md +198 -0
  42. msprobe/docs/09.accuracy_checker_MindSpore.md +68 -0
  43. msprobe/docs/10.accuracy_compare_PyTorch.md +245 -0
  44. msprobe/docs/11.accuracy_compare_MindSpore.md +202 -0
  45. msprobe/docs/12.overflow_check_PyTorch.md +79 -0
  46. msprobe/docs/13.overflow_check_MindSpore.md +31 -0
  47. msprobe/{pytorch/doc/parse_tool.md → docs/14.data_parse_PyTorch.md} +283 -286
  48. msprobe/docs/15.free_benchmarking_PyTorch.md +164 -0
  49. msprobe/{doc/grad_probe/grad_probe.md → docs/17.grad_probe.md} +207 -207
  50. msprobe/docs/FAQ_PyTorch.md +177 -0
  51. msprobe/docs/S02.report_free_benchmarking_validation_performance_baseline.md +146 -0
  52. msprobe/docs/img/free_benchmark_framework.png +0 -0
  53. msprobe/mindspore/__init__.py +1 -1
  54. msprobe/mindspore/api_accuracy_checker/api_accuracy_checker.py +254 -245
  55. msprobe/mindspore/api_accuracy_checker/api_info.py +69 -69
  56. msprobe/mindspore/api_accuracy_checker/api_runner.py +155 -151
  57. msprobe/mindspore/api_accuracy_checker/base_compare_algorithm.py +196 -196
  58. msprobe/mindspore/api_accuracy_checker/cmd_parser.py +6 -0
  59. msprobe/mindspore/api_accuracy_checker/compute_element.py +238 -223
  60. msprobe/mindspore/api_accuracy_checker/main.py +8 -15
  61. msprobe/mindspore/api_accuracy_checker/type_mapping.py +113 -113
  62. msprobe/mindspore/api_accuracy_checker/utils.py +79 -62
  63. msprobe/mindspore/cell_processor.py +34 -34
  64. msprobe/mindspore/common/const.py +106 -87
  65. msprobe/mindspore/common/log.py +37 -37
  66. msprobe/mindspore/common/utils.py +81 -57
  67. msprobe/mindspore/compare/distributed_compare.py +75 -75
  68. msprobe/mindspore/compare/ms_compare.py +219 -117
  69. msprobe/mindspore/compare/ms_graph_compare.py +348 -317
  70. msprobe/mindspore/compare/ms_to_pt_api.yaml +399 -399
  71. msprobe/mindspore/debugger/debugger_config.py +66 -74
  72. msprobe/mindspore/debugger/precision_debugger.py +126 -107
  73. msprobe/mindspore/dump/dump_tool_factory.py +35 -35
  74. msprobe/mindspore/dump/hook_cell/api_registry.py +118 -104
  75. msprobe/mindspore/dump/hook_cell/hook_cell.py +55 -53
  76. msprobe/mindspore/dump/hook_cell/support_wrap_ops.yaml +922 -925
  77. msprobe/mindspore/dump/hook_cell/wrap_api.py +113 -0
  78. msprobe/mindspore/dump/jit_dump.py +72 -56
  79. msprobe/mindspore/dump/kernel_graph_dump.py +59 -60
  80. msprobe/mindspore/dump/kernel_kbyk_dump.py +64 -65
  81. msprobe/mindspore/free_benchmark/api_pynative_self_check.py +116 -116
  82. msprobe/mindspore/free_benchmark/common/config.py +12 -12
  83. msprobe/mindspore/free_benchmark/common/handler_params.py +17 -17
  84. msprobe/mindspore/free_benchmark/common/utils.py +71 -71
  85. msprobe/mindspore/free_benchmark/data/support_wrap_ops.yaml +842 -842
  86. msprobe/mindspore/free_benchmark/decorator/dec_forward.py +43 -42
  87. msprobe/mindspore/free_benchmark/decorator/decorator_factory.py +107 -107
  88. msprobe/mindspore/free_benchmark/handler/base_handler.py +90 -90
  89. msprobe/mindspore/free_benchmark/handler/check_handler.py +41 -41
  90. msprobe/mindspore/free_benchmark/handler/fix_handler.py +36 -36
  91. msprobe/mindspore/free_benchmark/handler/handler_factory.py +21 -21
  92. msprobe/mindspore/free_benchmark/perturbation/add_noise.py +67 -67
  93. msprobe/mindspore/free_benchmark/perturbation/base_perturbation.py +21 -21
  94. msprobe/mindspore/free_benchmark/perturbation/bit_noise.py +63 -63
  95. msprobe/mindspore/free_benchmark/perturbation/exchange_value.py +51 -0
  96. msprobe/mindspore/free_benchmark/perturbation/improve_precision.py +35 -34
  97. msprobe/mindspore/free_benchmark/perturbation/no_change.py +12 -12
  98. msprobe/mindspore/free_benchmark/perturbation/perturbation_factory.py +29 -27
  99. msprobe/mindspore/free_benchmark/self_check_tool_factory.py +33 -33
  100. msprobe/mindspore/grad_probe/global_context.py +90 -91
  101. msprobe/mindspore/grad_probe/grad_analyzer.py +231 -231
  102. msprobe/mindspore/grad_probe/grad_monitor.py +27 -27
  103. msprobe/mindspore/grad_probe/grad_stat_csv.py +131 -131
  104. msprobe/mindspore/grad_probe/hook.py +94 -92
  105. msprobe/mindspore/grad_probe/utils.py +29 -28
  106. msprobe/mindspore/ms_config.py +128 -126
  107. msprobe/mindspore/overflow_check/kernel_graph_overflow_check.py +44 -45
  108. msprobe/mindspore/overflow_check/overflow_check_tool_factory.py +34 -34
  109. msprobe/mindspore/runtime.py +4 -4
  110. msprobe/mindspore/service.py +378 -354
  111. msprobe/mindspore/task_handler_factory.py +24 -24
  112. msprobe/msprobe.py +105 -107
  113. msprobe/pytorch/__init__.py +3 -3
  114. msprobe/pytorch/api_accuracy_checker/common/config.py +53 -55
  115. msprobe/pytorch/api_accuracy_checker/common/utils.py +214 -165
  116. msprobe/pytorch/api_accuracy_checker/compare/algorithm.py +213 -213
  117. msprobe/pytorch/api_accuracy_checker/compare/api_precision_compare.py +606 -581
  118. msprobe/pytorch/api_accuracy_checker/compare/api_precision_standard.yaml +132 -132
  119. msprobe/pytorch/api_accuracy_checker/compare/api_precision_threshold.yaml +390 -390
  120. msprobe/pytorch/api_accuracy_checker/compare/compare.py +386 -381
  121. msprobe/pytorch/api_accuracy_checker/compare/compare_column.py +73 -73
  122. msprobe/pytorch/api_accuracy_checker/compare/compare_utils.py +245 -244
  123. msprobe/pytorch/api_accuracy_checker/config.yaml +10 -10
  124. msprobe/pytorch/api_accuracy_checker/run_ut/data_generate.py +335 -332
  125. msprobe/pytorch/api_accuracy_checker/run_ut/multi_run_ut.py +200 -199
  126. msprobe/pytorch/api_accuracy_checker/run_ut/run_overflow_check.py +133 -134
  127. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut.py +592 -581
  128. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut_utils.py +70 -74
  129. msprobe/pytorch/api_accuracy_checker/run_ut/torch_ut_setting.json +7 -4
  130. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/attl.py +197 -202
  131. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/client.py +325 -324
  132. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/device_dispatch.py +204 -204
  133. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/server.py +219 -218
  134. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/ssl_config.py +10 -10
  135. msprobe/pytorch/bench_functions/__init__.py +15 -15
  136. msprobe/pytorch/bench_functions/apply_adam_w.py +28 -28
  137. msprobe/pytorch/bench_functions/confusion_transpose.py +19 -19
  138. msprobe/pytorch/bench_functions/fast_gelu.py +55 -55
  139. msprobe/pytorch/bench_functions/layer_norm_eval.py +6 -6
  140. msprobe/pytorch/bench_functions/linear.py +12 -12
  141. msprobe/pytorch/bench_functions/matmul_backward.py +48 -48
  142. msprobe/pytorch/bench_functions/npu_fusion_attention.py +509 -421
  143. msprobe/pytorch/bench_functions/rms_norm.py +15 -15
  144. msprobe/pytorch/bench_functions/rotary_mul.py +52 -52
  145. msprobe/pytorch/bench_functions/scaled_mask_softmax.py +26 -26
  146. msprobe/pytorch/bench_functions/swiglu.py +55 -55
  147. msprobe/pytorch/common/__init__.py +2 -2
  148. msprobe/pytorch/common/compare_script.template +14 -14
  149. msprobe/pytorch/common/log.py +20 -31
  150. msprobe/pytorch/common/parse_json.py +39 -39
  151. msprobe/pytorch/common/utils.py +305 -300
  152. msprobe/pytorch/compare/distributed_compare.py +66 -66
  153. msprobe/pytorch/compare/mapping.yaml +607 -607
  154. msprobe/pytorch/compare/match.py +34 -33
  155. msprobe/pytorch/compare/pt_compare.py +50 -40
  156. msprobe/pytorch/debugger/debugger_config.py +95 -95
  157. msprobe/pytorch/debugger/precision_debugger.py +125 -125
  158. msprobe/pytorch/free_benchmark/__init__.py +8 -8
  159. msprobe/pytorch/free_benchmark/common/constant.py +70 -70
  160. msprobe/pytorch/free_benchmark/common/counter.py +71 -71
  161. msprobe/pytorch/free_benchmark/common/enums.py +37 -37
  162. msprobe/pytorch/free_benchmark/common/params.py +129 -129
  163. msprobe/pytorch/free_benchmark/common/utils.py +102 -102
  164. msprobe/pytorch/free_benchmark/compare/grad_saver.py +179 -179
  165. msprobe/pytorch/free_benchmark/compare/single_benchmark.py +104 -104
  166. msprobe/pytorch/free_benchmark/main.py +105 -105
  167. msprobe/pytorch/free_benchmark/perturbed_layers/base_layer.py +13 -13
  168. msprobe/pytorch/free_benchmark/perturbed_layers/layer_factory.py +41 -41
  169. msprobe/pytorch/free_benchmark/perturbed_layers/npu/add_noise.py +90 -90
  170. msprobe/pytorch/free_benchmark/perturbed_layers/npu/bit_noise.py +104 -104
  171. msprobe/pytorch/free_benchmark/perturbed_layers/npu/change_value.py +63 -63
  172. msprobe/pytorch/free_benchmark/perturbed_layers/npu/improve_precision.py +68 -68
  173. msprobe/pytorch/free_benchmark/perturbed_layers/npu/no_change.py +28 -28
  174. msprobe/pytorch/free_benchmark/perturbed_layers/npu/npu_base_layser.py +45 -45
  175. msprobe/pytorch/free_benchmark/perturbed_layers/run_cpu.py +19 -19
  176. msprobe/pytorch/free_benchmark/result_handlers/base_handler.py +217 -217
  177. msprobe/pytorch/free_benchmark/result_handlers/check_handler.py +39 -39
  178. msprobe/pytorch/free_benchmark/result_handlers/fix_handler.py +23 -23
  179. msprobe/pytorch/free_benchmark/result_handlers/handler_factory.py +30 -30
  180. msprobe/pytorch/free_benchmark/result_handlers/preheat_handler.py +170 -170
  181. msprobe/pytorch/function_factory.py +76 -75
  182. msprobe/pytorch/functional/dump_module.py +39 -39
  183. msprobe/pytorch/grad_probe/grad_monitor.py +91 -90
  184. msprobe/pytorch/grad_probe/grad_stat_csv.py +128 -128
  185. msprobe/pytorch/hook_module/api_registry.py +161 -161
  186. msprobe/pytorch/hook_module/hook_module.py +120 -120
  187. msprobe/pytorch/hook_module/support_wrap_ops.yaml +1879 -1877
  188. msprobe/pytorch/hook_module/utils.py +30 -29
  189. msprobe/pytorch/hook_module/wrap_aten.py +110 -110
  190. msprobe/pytorch/hook_module/wrap_distributed.py +78 -78
  191. msprobe/pytorch/hook_module/wrap_functional.py +105 -105
  192. msprobe/pytorch/hook_module/wrap_npu_custom.py +93 -84
  193. msprobe/pytorch/hook_module/wrap_tensor.py +71 -71
  194. msprobe/pytorch/hook_module/wrap_torch.py +86 -86
  195. msprobe/pytorch/hook_module/wrap_vf.py +62 -62
  196. msprobe/pytorch/module_processer.py +138 -138
  197. msprobe/pytorch/online_dispatch/__init__.py +20 -20
  198. msprobe/pytorch/online_dispatch/compare.py +236 -236
  199. msprobe/pytorch/online_dispatch/dispatch.py +271 -271
  200. msprobe/pytorch/online_dispatch/dump_compare.py +155 -156
  201. msprobe/pytorch/online_dispatch/single_compare.py +391 -391
  202. msprobe/pytorch/online_dispatch/torch_ops_config.yaml +49 -49
  203. msprobe/pytorch/online_dispatch/utils.py +130 -146
  204. msprobe/pytorch/parse.py +4 -4
  205. msprobe/pytorch/parse_tool/cli.py +32 -32
  206. msprobe/pytorch/parse_tool/lib/compare.py +260 -271
  207. msprobe/pytorch/parse_tool/lib/config.py +52 -52
  208. msprobe/pytorch/parse_tool/lib/file_desc.py +31 -31
  209. msprobe/pytorch/parse_tool/lib/interactive_cli.py +102 -102
  210. msprobe/pytorch/parse_tool/lib/parse_exception.py +54 -54
  211. msprobe/pytorch/parse_tool/lib/parse_tool.py +158 -158
  212. msprobe/pytorch/parse_tool/lib/utils.py +316 -321
  213. msprobe/pytorch/parse_tool/lib/visualization.py +85 -91
  214. msprobe/pytorch/pt_config.py +188 -187
  215. msprobe/pytorch/service.py +246 -252
  216. mindstudio_probe-1.0.3.dist-info/RECORD +0 -272
  217. msprobe/config/README.md +0 -539
  218. msprobe/mindspore/doc/compare.md +0 -58
  219. msprobe/mindspore/doc/dump.md +0 -217
  220. msprobe/mindspore/dump/hook_cell/wrap_functional.py +0 -91
  221. msprobe/mindspore/dump/hook_cell/wrap_tensor.py +0 -63
  222. msprobe/pytorch/doc/FAQ.md +0 -193
  223. msprobe/pytorch/doc/api_accuracy_checker.md +0 -313
  224. msprobe/pytorch/doc/api_accuracy_checker_online.md +0 -187
  225. msprobe/pytorch/doc/dump.md +0 -260
  226. msprobe/pytorch/doc/msprobe/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/206/320/245/342/226/221/321/206/320/235/320/276dump/321/206/320/260/320/227/321/205/320/227/320/226/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -182
  227. msprobe/pytorch/doc/ptdbg_ascend_compare.md +0 -240
  228. msprobe/pytorch/doc/ptdbg_ascend_overview.md +0 -68
  229. msprobe/pytorch/doc/ptdbg_ascend_quickstart.md +0 -381
  230. msprobe/pytorch/doc/run_overflow_check.md +0 -25
  231. msprobe/pytorch/doc//321/205/320/254/320/270/321/207/342/225/221/342/224/220/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/206/320/277/320/244/321/205/320/277/342/225/243.md +0 -90
  232. msprobe/pytorch/doc//321/206/320/247/320/260/321/206/320/260/320/227/321/206/320/255/320/226/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/205/320/254/342/225/221/321/206/320/251/320/277/321/211/320/272/320/234/321/210/320/277/320/221/321/205/320/242/320/234/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -151
  233. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/top_level.txt +0 -0
  234. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_1.png +0 -0
  235. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_2.png +0 -0
  236. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_3.png +0 -0
  237. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_4.png +0 -0
  238. /msprobe/{pytorch/doc → docs}/img/GPT-3_1.png +0 -0
  239. /msprobe/{pytorch/doc → docs}/img/GPT-3_2.png +0 -0
  240. /msprobe/{pytorch/doc → docs}/img/GPT-3_3.png +0 -0
  241. /msprobe/{pytorch/doc → docs}/img/GPT-3_4.png +0 -0
  242. /msprobe/{pytorch/doc → docs}/img/GPT-3_5.png +0 -0
  243. /msprobe/{pytorch/doc → docs}/img/GPT-3_6.png +0 -0
  244. /msprobe/{pytorch/doc → docs}/img/GPT-3_7.png +0 -0
  245. /msprobe/{pytorch/doc → docs}/img/GPT-3_8.png +0 -0
  246. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_1.png +0 -0
  247. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_2.png +0 -0
  248. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_details.png +0 -0
  249. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_result.png +0 -0
  250. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_details.png +0 -0
  251. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_result.png +0 -0
  252. /msprobe/{pytorch/doc → docs}/img/auto_analyze_log.png +0 -0
  253. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl.png +0 -0
  254. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl_md5.png.png +0 -0
  255. /msprobe/{pytorch/doc → docs}/img/cpu_info.png +0 -0
  256. /msprobe/{config → docs}/img/free_benchmark.png +0 -0
  257. /msprobe/{doc/grad_probe/img/image-1.png → docs/img/grad_probe_image-1.png} +0 -0
  258. /msprobe/{doc/grad_probe/img/image-2.png → docs/img/grad_probe_image-2.png} +0 -0
  259. /msprobe/{doc/grad_probe/img/image-3.png → docs/img/grad_probe_image-3.png} +0 -0
  260. /msprobe/{doc/grad_probe/img/image-4.png → docs/img/grad_probe_image-4.png} +0 -0
  261. /msprobe/{doc/grad_probe/img/image.png → docs/img/grad_probe_image.png} +0 -0
  262. /msprobe/{pytorch/doc → docs}/img/module_compare.png +0 -0
@@ -1,70 +1,70 @@
1
- from typing import Dict
2
-
3
- import numpy as np
4
- import torch
5
- from msprobe.pytorch.free_benchmark.common.enums import FuzzThreshold
6
- from msprobe.pytorch.free_benchmark.common.params import BenchmarkThd
7
-
8
-
9
- class CommonField:
10
- DEVICE = "device"
11
- META = "meta"
12
- FUZZ_TENSOR = "fuzz_tensor"
13
- REQUIRES_GRAD = "requires_grad"
14
- HOLD_PLACE = "hold_place"
15
- DISTRIBUTED_OP = "torch.distributed"
16
- GRADSAVER = "grad_saver"
17
-
18
-
19
- class ThresholdConfig:
20
- PERTURBATION_VALUE_DICT: Dict = {
21
- torch.bfloat16: FuzzThreshold.BF16_THD,
22
- torch.float16: FuzzThreshold.F16_THD,
23
- torch.float32: FuzzThreshold.F32_THD,
24
- torch.float64: FuzzThreshold.F64_THD,
25
- }
26
-
27
- ABS_TOL_VALUE_DICT: Dict = {
28
- torch.bfloat16: FuzzThreshold.BF16_THD,
29
- torch.float16: FuzzThreshold.F16_THD,
30
- torch.float32: FuzzThreshold.F32_THD,
31
- torch.float64: FuzzThreshold.F64_THD,
32
- }
33
-
34
- # bit翻转需要匹配到等长或更长的整型
35
- PERTURBATION_BIT_DICT = {
36
- torch.bfloat16: torch.int16,
37
- torch.float16: torch.int16,
38
- torch.float32: torch.int32,
39
- torch.float64: torch.int64,
40
- }
41
-
42
- # 输入噪声下界
43
- NOISE_INPUT_LOWER_BOUND = 1e-8
44
- COMP_CONSISTENT = 1.0
45
- COMP_NAN = np.nan
46
- SYMBOL_FLIPPING = "symbol_flipping"
47
- BACKWARD_OUTPUT_LOWER_BOUND = 1e-3
48
- SMALL_VALUE = 1.0
49
- # 预热初始阈值
50
- PREHEAT_INITIAL_THD = 2.05
51
- API_THD_STEP = 2.0
52
-
53
- DTYPE_PER_THD = {
54
- torch.float16: 1.002,
55
- torch.bfloat16: 1.004,
56
- torch.float32: 1.0002,
57
- }
58
- BENCHMARK_THD_DICT = {
59
- torch.float32: BenchmarkThd(2**-14, 1.0, 2**-14, 1e-4),
60
- torch.float16: BenchmarkThd(2**-11, 1.0, 2**-11, 1e-4),
61
- torch.bfloat16: BenchmarkThd(2**-8, 1.0, 2**-8, 1e-4),
62
- }
63
-
64
- TENSOR_SPLIT_MAX_CHUNK = 128
65
-
66
-
67
- class PreheatConfig:
68
- IF_PREHEAT = "if_preheat"
69
- PREHEAT_STEP = "preheat_step"
70
- MAX_SAMPLE = "max_sample"
1
+ from typing import Dict
2
+
3
+ import numpy as np
4
+ import torch
5
+ from msprobe.pytorch.free_benchmark.common.enums import FuzzThreshold
6
+ from msprobe.pytorch.free_benchmark.common.params import BenchmarkThd
7
+
8
+
9
+ class CommonField:
10
+ DEVICE = "device"
11
+ META = "meta"
12
+ FUZZ_TENSOR = "fuzz_tensor"
13
+ REQUIRES_GRAD = "requires_grad"
14
+ HOLD_PLACE = "hold_place"
15
+ DISTRIBUTED_OP = "torch.distributed"
16
+ GRADSAVER = "grad_saver"
17
+
18
+
19
+ class ThresholdConfig:
20
+ PERTURBATION_VALUE_DICT: Dict = {
21
+ torch.bfloat16: FuzzThreshold.BF16_THD,
22
+ torch.float16: FuzzThreshold.F16_THD,
23
+ torch.float32: FuzzThreshold.F32_THD,
24
+ torch.float64: FuzzThreshold.F64_THD,
25
+ }
26
+
27
+ ABS_TOL_VALUE_DICT: Dict = {
28
+ torch.bfloat16: FuzzThreshold.BF16_THD,
29
+ torch.float16: FuzzThreshold.F16_THD,
30
+ torch.float32: FuzzThreshold.F32_THD,
31
+ torch.float64: FuzzThreshold.F64_THD,
32
+ }
33
+
34
+ # bit翻转需要匹配到等长或更长的整型
35
+ PERTURBATION_BIT_DICT = {
36
+ torch.bfloat16: torch.int16,
37
+ torch.float16: torch.int16,
38
+ torch.float32: torch.int32,
39
+ torch.float64: torch.int64,
40
+ }
41
+
42
+ # 输入噪声下界
43
+ NOISE_INPUT_LOWER_BOUND = 1e-8
44
+ COMP_CONSISTENT = 1.0
45
+ COMP_NAN = np.nan
46
+ SYMBOL_FLIPPING = "symbol_flipping"
47
+ BACKWARD_OUTPUT_LOWER_BOUND = 1e-3
48
+ SMALL_VALUE = 1.0
49
+ # 预热初始阈值
50
+ PREHEAT_INITIAL_THD = 2.05
51
+ API_THD_STEP = 2.0
52
+
53
+ DTYPE_PER_THD = {
54
+ torch.float16: 1.002,
55
+ torch.bfloat16: 1.004,
56
+ torch.float32: 1.0002,
57
+ }
58
+ BENCHMARK_THD_DICT = {
59
+ torch.float32: BenchmarkThd(2**-14, 1.0, 2**-14, 1e-4),
60
+ torch.float16: BenchmarkThd(2**-11, 1.0, 2**-11, 1e-4),
61
+ torch.bfloat16: BenchmarkThd(2**-8, 1.0, 2**-8, 1e-4),
62
+ }
63
+
64
+ TENSOR_SPLIT_MAX_CHUNK = 128
65
+
66
+
67
+ class PreheatConfig:
68
+ IF_PREHEAT = "if_preheat"
69
+ PREHEAT_STEP = "preheat_step"
70
+ MAX_SAMPLE = "max_sample"
@@ -1,72 +1,72 @@
1
- from collections import defaultdict
2
- from msprobe.pytorch.free_benchmark.common.constant import ThresholdConfig
3
-
4
-
5
- class PreheatCounter:
6
- def __init__(self) -> None:
7
- self.api_called_time: dict = defaultdict(int)
8
- self.api_sample_time: dict = defaultdict(int)
9
- self.one_step_used_api: dict = defaultdict(int)
10
- self.api_thd: dict = defaultdict(dict)
11
- self.preheat_record: dict = defaultdict(dict)
12
- self.dtype_map: dict = {}
13
- self.if_preheat: dict = defaultdict(dict)
14
- self.step = 0
15
-
16
- def clear_step(self):
17
- self.preheat_record.clear()
18
- self.api_called_time.clear()
19
- self.api_sample_time.clear()
20
-
21
- def check_step(self, current_step):
22
- if current_step != self.step:
23
- self.clear_step()
24
- self.step = current_step
25
-
26
- def add_api_called_time(self, api_name: str):
27
- self.api_called_time[api_name] += 1
28
-
29
- def get_api_called_time(self, api_name: str) -> int:
30
- return self.api_called_time[api_name]
31
-
32
- def add_api_sample_time(self, api_name: str):
33
- self.api_sample_time[api_name] += 1
34
-
35
- def get_api_sample_time(self, api_name: str) -> int:
36
- return self.api_sample_time[api_name]
37
-
38
- def add_one_step_used_api(self, api_name: str):
39
- self.one_step_used_api[api_name] += 1
40
-
41
- def get_one_step_used_api(self, api_name: str):
42
- return self.one_step_used_api[api_name]
43
-
44
- def update_preheat_record(self, api_name, dtype, cmp_result):
45
- # 记录预热阶段CPU标杆比对的结果
46
- if str(dtype) not in self.preheat_record[api_name].keys():
47
- self.preheat_record[api_name][str(dtype)] = list()
48
- self.preheat_record[api_name][str(dtype)].append(cmp_result)
49
- self.dtype_map[str(dtype)] = dtype
50
-
51
- def update_api_thd(self, api_name, dtype, threshold, dthreshold):
52
- self.api_thd[api_name][str(dtype)] = (
53
- threshold if threshold > dthreshold else dthreshold
54
- )
55
-
56
- def get_api_thd(self, api_name, dtype):
57
- if not str(dtype) in self.api_thd[api_name]:
58
- self.api_thd[api_name][str(dtype)] = ThresholdConfig.PREHEAT_INITIAL_THD
59
- self.dtype_map[str(dtype)] = dtype
60
- return self.api_thd[api_name][str(dtype)]
61
-
62
- def set_api_preheat(self, api_name, dtype_str, is_preheat=True):
63
- # 标记cpu不一致的dtype 不再进行预热
64
- self.if_preheat[api_name][dtype_str] = is_preheat
65
-
66
- def get_api_preheat(self, api_name, dtype):
67
- # 标记cpu不一致的dtype 不再进行预热
68
- if str(dtype) not in self.if_preheat[api_name]:
69
- return True
70
- return self.if_preheat[api_name][str(dtype)]
71
-
1
+ from collections import defaultdict
2
+ from msprobe.pytorch.free_benchmark.common.constant import ThresholdConfig
3
+
4
+
5
+ class PreheatCounter:
6
+ def __init__(self) -> None:
7
+ self.api_called_time: dict = defaultdict(int)
8
+ self.api_sample_time: dict = defaultdict(int)
9
+ self.one_step_used_api: dict = defaultdict(int)
10
+ self.api_thd: dict = defaultdict(dict)
11
+ self.preheat_record: dict = defaultdict(dict)
12
+ self.dtype_map: dict = {}
13
+ self.if_preheat: dict = defaultdict(dict)
14
+ self.step = 0
15
+
16
+ def clear_step(self):
17
+ self.preheat_record.clear()
18
+ self.api_called_time.clear()
19
+ self.api_sample_time.clear()
20
+
21
+ def check_step(self, current_step):
22
+ if current_step != self.step:
23
+ self.clear_step()
24
+ self.step = current_step
25
+
26
+ def add_api_called_time(self, api_name: str):
27
+ self.api_called_time[api_name] += 1
28
+
29
+ def get_api_called_time(self, api_name: str) -> int:
30
+ return self.api_called_time[api_name]
31
+
32
+ def add_api_sample_time(self, api_name: str):
33
+ self.api_sample_time[api_name] += 1
34
+
35
+ def get_api_sample_time(self, api_name: str) -> int:
36
+ return self.api_sample_time[api_name]
37
+
38
+ def add_one_step_used_api(self, api_name: str):
39
+ self.one_step_used_api[api_name] += 1
40
+
41
+ def get_one_step_used_api(self, api_name: str):
42
+ return self.one_step_used_api[api_name]
43
+
44
+ def update_preheat_record(self, api_name, dtype, cmp_result):
45
+ # 记录预热阶段CPU标杆比对的结果
46
+ if str(dtype) not in self.preheat_record[api_name].keys():
47
+ self.preheat_record[api_name][str(dtype)] = list()
48
+ self.preheat_record[api_name][str(dtype)].append(cmp_result)
49
+ self.dtype_map[str(dtype)] = dtype
50
+
51
+ def update_api_thd(self, api_name, dtype, threshold, dthreshold):
52
+ self.api_thd[api_name][str(dtype)] = (
53
+ threshold if threshold > dthreshold else dthreshold
54
+ )
55
+
56
+ def get_api_thd(self, api_name, dtype):
57
+ if not str(dtype) in self.api_thd[api_name]:
58
+ self.api_thd[api_name][str(dtype)] = ThresholdConfig.PREHEAT_INITIAL_THD
59
+ self.dtype_map[str(dtype)] = dtype
60
+ return self.api_thd[api_name][str(dtype)]
61
+
62
+ def set_api_preheat(self, api_name, dtype_str, is_preheat=True):
63
+ # 标记cpu不一致的dtype 不再进行预热
64
+ self.if_preheat[api_name][dtype_str] = is_preheat
65
+
66
+ def get_api_preheat(self, api_name, dtype):
67
+ # 标记cpu不一致的dtype 不再进行预热
68
+ if str(dtype) not in self.if_preheat[api_name]:
69
+ return True
70
+ return self.if_preheat[api_name][str(dtype)]
71
+
72
72
  preheat_counter = PreheatCounter()
@@ -1,37 +1,37 @@
1
- class PerturbationMode:
2
- ADD_NOISE = "add_noise"
3
- CHANGE_VALUE = "change_value"
4
- IMPROVE_PRECISION = "improve_precision"
5
- NO_CHANGE = "no_change"
6
- BIT_NOISE = "bit_noise"
7
- TO_CPU = "to_cpu"
8
-
9
-
10
- class DeviceType:
11
- NPU = "npu"
12
- CPU = "cpu"
13
-
14
-
15
- class FuzzThreshold:
16
- BF16_THD = 1e-4
17
- F16_THD = 1e-6
18
- F32_THD = 1e-8
19
- F64_THD = 1e-16
20
-
21
-
22
- class NormType:
23
- ONE_NORM = (1, "one_norm")
24
- TWO_NORM = (2, "two_norm")
25
- ENDLESS_NORM = (3, "endless_norm")
26
-
27
-
28
- class HandlerType:
29
- CHECK = "check"
30
- PREHEAT = "preheat"
31
- FIX = "fix"
32
-
33
-
34
- class FuzzLevel:
35
- BASE_LEVEL = "L1"
36
- ADV_LEVEL = "L2"
37
- REAL_LEVEL = "L3"
1
+ class PerturbationMode:
2
+ ADD_NOISE = "add_noise"
3
+ CHANGE_VALUE = "change_value"
4
+ IMPROVE_PRECISION = "improve_precision"
5
+ NO_CHANGE = "no_change"
6
+ BIT_NOISE = "bit_noise"
7
+ TO_CPU = "to_cpu"
8
+
9
+
10
+ class DeviceType:
11
+ NPU = "npu"
12
+ CPU = "cpu"
13
+
14
+
15
+ class FuzzThreshold:
16
+ BF16_THD = 1e-4
17
+ F16_THD = 1e-6
18
+ F32_THD = 1e-8
19
+ F64_THD = 1e-16
20
+
21
+
22
+ class NormType:
23
+ ONE_NORM = (1, "one_norm")
24
+ TWO_NORM = (2, "two_norm")
25
+ ENDLESS_NORM = (3, "endless_norm")
26
+
27
+
28
+ class HandlerType:
29
+ CHECK = "check"
30
+ PREHEAT = "preheat"
31
+ FIX = "fix"
32
+
33
+
34
+ class FuzzLevel:
35
+ BASE_LEVEL = "L1"
36
+ ADV_LEVEL = "L2"
37
+ REAL_LEVEL = "L3"
@@ -1,129 +1,129 @@
1
- from dataclasses import dataclass
2
- from typing import Any, Callable, Dict, List, Optional, Tuple
3
-
4
- import torch
5
- from msprobe.pytorch.free_benchmark import logger
6
- from msprobe.pytorch.free_benchmark.common.enums import (
7
- DeviceType,
8
- FuzzLevel,
9
- PerturbationMode,
10
- )
11
- from msprobe.pytorch.free_benchmark.common.utils import Tools
12
-
13
-
14
- @dataclass
15
- class DataParams:
16
- args: Optional[Tuple] = None
17
- kwargs: Optional[Dict] = None
18
- valid_input_index: Optional[int] = None
19
- original_result: Optional[Any] = None
20
- perturbed_result: Optional[Any] = None
21
- is_consistent: Optional[bool] = True
22
- perturbed_value: Optional[Any] = None
23
- origin_func: Optional[Callable] = None
24
- api_type: Optional[str] = None
25
- fuzz_stage: Optional[str] = None
26
- grad_unequal_flag: Optional[bool] = True
27
-
28
-
29
- @dataclass
30
- class HandlerParams:
31
- handler_type: Optional[str] = None
32
- api_name: Optional[str] = None
33
- pert_mode: Optional[PerturbationMode] = None
34
- step: Optional[int] = None
35
- fuzz_stage: Optional[str] = None
36
- fuzz_device: Optional[DeviceType] = None
37
- preheat_config: Optional[Dict] = None
38
- fuzz_level: Optional[str] = None
39
-
40
-
41
- @dataclass
42
- class UnequalRow:
43
- rank: Optional[int] = None
44
- pert_mode: Optional[PerturbationMode] = None
45
- stage: Optional[str] = None
46
- step: Optional[int] = None
47
- api_name: Optional[str] = None
48
- max_rel: Optional[float] = None
49
- dtype: Optional[str] = None
50
- shape: Optional[str] = None
51
- output_index: Optional[int] = None
52
-
53
-
54
- @dataclass
55
- class BenchmarkThd:
56
- rtol: Optional[float] = None # 相对误差阈值
57
- small_value: Optional[float] = None # 小值域
58
- small_value_atol: Optional[float] = None # 小值域绝对阈值
59
- err_balance: Optional[float] = None # 误差均衡性
60
-
61
-
62
- def check_args_type(args: Tuple) -> int:
63
- for i, arg in enumerate(args):
64
- if torch.is_tensor(arg):
65
- if arg.is_meta:
66
- continue
67
- if not torch.is_floating_point(arg):
68
- continue
69
- return i
70
- if isinstance(arg, (List, Tuple, Dict)):
71
- return i
72
- return -1
73
-
74
-
75
- def data_pre_deal(name, func, args, kwargs):
76
- data_params = DataParams(args=args, kwargs=kwargs, origin_func=func)
77
- index = check_args_type(args)
78
- data_params.valid_input_index = index
79
- if index == -1:
80
- logger.warning_on_rank_0(
81
- f"[msprobe] Free benchmark: 无标杆工具不支持当前算子的输入类型 {name}."
82
- )
83
- return data_params
84
-
85
-
86
- def make_handler_params(name, config, step):
87
- handler_params = HandlerParams()
88
- handler_params.api_name = name
89
- handler_params.step = step
90
- handler_params.handler_type = config.handler_type
91
- handler_params.fuzz_stage = config.fuzz_stage
92
- handler_params.fuzz_device = config.fuzz_device
93
- handler_params.preheat_config = config.preheat_config
94
- handler_params.fuzz_level = config.fuzz_level
95
- handler_params.pert_mode = config.pert_mode
96
- return handler_params
97
-
98
-
99
- def make_unequal_row(
100
- data_params: DataParams,
101
- handle_params: HandlerParams,
102
- ratio: float = None,
103
- index: int = None,
104
- ):
105
- row = UnequalRow(
106
- api_name=handle_params.api_name,
107
- pert_mode=handle_params.pert_mode,
108
- output_index=index,
109
- stage=handle_params.fuzz_stage,
110
- step=handle_params.step,
111
- )
112
- if isinstance(ratio, float):
113
- row.max_rel = ratio - 1
114
- origin_tensor = data_params.original_result
115
- perturbed_tensor = data_params.perturbed_result
116
- if index:
117
- origin_tensor = origin_tensor[index]
118
- perturbed_tensor = perturbed_tensor[index]
119
- row.output_index = index
120
- if isinstance(origin_tensor, torch.Tensor):
121
- row.dtype = origin_tensor.dtype
122
- row.shape = origin_tensor.shape
123
- row.rank = Tools.get_dist_rank()
124
- # 以下暂不支持
125
- if handle_params.fuzz_level == FuzzLevel.ADV_LEVEL:
126
- pass
127
- if handle_params.fuzz_level == FuzzLevel.REAL_LEVEL:
128
- pass
129
- return row
1
+ from dataclasses import dataclass
2
+ from typing import Any, Callable, Dict, List, Optional, Tuple
3
+
4
+ import torch
5
+ from msprobe.pytorch.free_benchmark import logger
6
+ from msprobe.pytorch.free_benchmark.common.enums import (
7
+ DeviceType,
8
+ FuzzLevel,
9
+ PerturbationMode,
10
+ )
11
+ from msprobe.pytorch.free_benchmark.common.utils import Tools
12
+
13
+
14
+ @dataclass
15
+ class DataParams:
16
+ args: Optional[Tuple] = None
17
+ kwargs: Optional[Dict] = None
18
+ valid_input_index: Optional[int] = None
19
+ original_result: Optional[Any] = None
20
+ perturbed_result: Optional[Any] = None
21
+ is_consistent: Optional[bool] = True
22
+ perturbed_value: Optional[Any] = None
23
+ origin_func: Optional[Callable] = None
24
+ api_type: Optional[str] = None
25
+ fuzz_stage: Optional[str] = None
26
+ grad_unequal_flag: Optional[bool] = True
27
+
28
+
29
+ @dataclass
30
+ class HandlerParams:
31
+ handler_type: Optional[str] = None
32
+ api_name: Optional[str] = None
33
+ pert_mode: Optional[PerturbationMode] = None
34
+ step: Optional[int] = None
35
+ fuzz_stage: Optional[str] = None
36
+ fuzz_device: Optional[DeviceType] = None
37
+ preheat_config: Optional[Dict] = None
38
+ fuzz_level: Optional[str] = None
39
+
40
+
41
+ @dataclass
42
+ class UnequalRow:
43
+ rank: Optional[int] = None
44
+ pert_mode: Optional[PerturbationMode] = None
45
+ stage: Optional[str] = None
46
+ step: Optional[int] = None
47
+ api_name: Optional[str] = None
48
+ max_rel: Optional[float] = None
49
+ dtype: Optional[str] = None
50
+ shape: Optional[str] = None
51
+ output_index: Optional[int] = None
52
+
53
+
54
+ @dataclass
55
+ class BenchmarkThd:
56
+ rtol: Optional[float] = None # 相对误差阈值
57
+ small_value: Optional[float] = None # 小值域
58
+ small_value_atol: Optional[float] = None # 小值域绝对阈值
59
+ err_balance: Optional[float] = None # 误差均衡性
60
+
61
+
62
+ def check_args_type(args: Tuple) -> int:
63
+ for i, arg in enumerate(args):
64
+ if torch.is_tensor(arg):
65
+ if arg.is_meta:
66
+ continue
67
+ if not torch.is_floating_point(arg):
68
+ continue
69
+ return i
70
+ if isinstance(arg, (List, Tuple, Dict)):
71
+ return i
72
+ return -1
73
+
74
+
75
+ def data_pre_deal(name, func, args, kwargs):
76
+ data_params = DataParams(args=args, kwargs=kwargs, origin_func=func)
77
+ index = check_args_type(args)
78
+ data_params.valid_input_index = index
79
+ if index == -1:
80
+ logger.warning_on_rank_0(
81
+ f"[msprobe] Free benchmark: 无标杆工具不支持当前算子的输入类型 {name}."
82
+ )
83
+ return data_params
84
+
85
+
86
+ def make_handler_params(name, config, step):
87
+ handler_params = HandlerParams()
88
+ handler_params.api_name = name
89
+ handler_params.step = step
90
+ handler_params.handler_type = config.handler_type
91
+ handler_params.fuzz_stage = config.fuzz_stage
92
+ handler_params.fuzz_device = config.fuzz_device
93
+ handler_params.preheat_config = config.preheat_config
94
+ handler_params.fuzz_level = config.fuzz_level
95
+ handler_params.pert_mode = config.pert_mode
96
+ return handler_params
97
+
98
+
99
+ def make_unequal_row(
100
+ data_params: DataParams,
101
+ handle_params: HandlerParams,
102
+ ratio: float = None,
103
+ index: int = None,
104
+ ):
105
+ row = UnequalRow(
106
+ api_name=handle_params.api_name,
107
+ pert_mode=handle_params.pert_mode,
108
+ output_index=index,
109
+ stage=handle_params.fuzz_stage,
110
+ step=handle_params.step,
111
+ )
112
+ if isinstance(ratio, float):
113
+ row.max_rel = ratio - 1
114
+ origin_tensor = data_params.original_result
115
+ perturbed_tensor = data_params.perturbed_result
116
+ if index:
117
+ origin_tensor = origin_tensor[index]
118
+ perturbed_tensor = perturbed_tensor[index]
119
+ row.output_index = index
120
+ if isinstance(origin_tensor, torch.Tensor):
121
+ row.dtype = origin_tensor.dtype
122
+ row.shape = origin_tensor.shape
123
+ row.rank = Tools.get_dist_rank()
124
+ # 以下暂不支持
125
+ if handle_params.fuzz_level == FuzzLevel.ADV_LEVEL:
126
+ pass
127
+ if handle_params.fuzz_level == FuzzLevel.REAL_LEVEL:
128
+ pass
129
+ return row