mindstudio-probe 1.0.3__py3-none-any.whl → 1.0.4__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (262) hide show
  1. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/LICENSE +201 -201
  2. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/METADATA +36 -34
  3. mindstudio_probe-1.0.4.dist-info/RECORD +276 -0
  4. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/WHEEL +1 -1
  5. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/entry_points.txt +1 -0
  6. msprobe/README.md +101 -237
  7. msprobe/{config/config.json → config.json} +49 -49
  8. msprobe/core/advisor/advisor.py +124 -124
  9. msprobe/core/advisor/advisor_const.py +59 -59
  10. msprobe/core/advisor/advisor_result.py +58 -58
  11. msprobe/core/common/const.py +341 -318
  12. msprobe/core/common/exceptions.py +99 -99
  13. msprobe/core/common/{file_check.py → file_utils.py} +478 -283
  14. msprobe/core/common/log.py +76 -69
  15. msprobe/core/common/utils.py +385 -616
  16. msprobe/core/common_config.py +85 -71
  17. msprobe/core/compare/acc_compare.py +299 -298
  18. msprobe/core/compare/check.py +95 -95
  19. msprobe/core/compare/compare_cli.py +49 -49
  20. msprobe/core/compare/highlight.py +223 -222
  21. msprobe/core/compare/multiprocessing_compute.py +149 -149
  22. msprobe/core/compare/npy_compare.py +295 -295
  23. msprobe/core/compare/utils.py +430 -429
  24. msprobe/core/data_dump/data_collector.py +154 -144
  25. msprobe/core/data_dump/data_processor/base.py +314 -293
  26. msprobe/core/data_dump/data_processor/factory.py +59 -59
  27. msprobe/core/data_dump/data_processor/mindspore_processor.py +186 -198
  28. msprobe/core/data_dump/data_processor/pytorch_processor.py +366 -389
  29. msprobe/core/data_dump/json_writer.py +96 -116
  30. msprobe/core/data_dump/scope.py +178 -178
  31. msprobe/core/grad_probe/constant.py +70 -70
  32. msprobe/core/grad_probe/grad_compare.py +171 -175
  33. msprobe/core/grad_probe/utils.py +64 -52
  34. msprobe/docs/01.installation.md +89 -0
  35. msprobe/docs/02.config_introduction.md +165 -0
  36. msprobe/docs/03.config_examples.md +247 -0
  37. msprobe/docs/04.acl_config_examples.md +76 -0
  38. msprobe/docs/05.data_dump_PyTorch.md +198 -0
  39. msprobe/docs/06.data_dump_MindSpore.md +243 -0
  40. msprobe/docs/07.accuracy_checker_PyTorch.md +274 -0
  41. msprobe/docs/08.accuracy_checker_online_PyTorch.md +198 -0
  42. msprobe/docs/09.accuracy_checker_MindSpore.md +68 -0
  43. msprobe/docs/10.accuracy_compare_PyTorch.md +245 -0
  44. msprobe/docs/11.accuracy_compare_MindSpore.md +202 -0
  45. msprobe/docs/12.overflow_check_PyTorch.md +79 -0
  46. msprobe/docs/13.overflow_check_MindSpore.md +31 -0
  47. msprobe/{pytorch/doc/parse_tool.md → docs/14.data_parse_PyTorch.md} +283 -286
  48. msprobe/docs/15.free_benchmarking_PyTorch.md +164 -0
  49. msprobe/{doc/grad_probe/grad_probe.md → docs/17.grad_probe.md} +207 -207
  50. msprobe/docs/FAQ_PyTorch.md +177 -0
  51. msprobe/docs/S02.report_free_benchmarking_validation_performance_baseline.md +146 -0
  52. msprobe/docs/img/free_benchmark_framework.png +0 -0
  53. msprobe/mindspore/__init__.py +1 -1
  54. msprobe/mindspore/api_accuracy_checker/api_accuracy_checker.py +254 -245
  55. msprobe/mindspore/api_accuracy_checker/api_info.py +69 -69
  56. msprobe/mindspore/api_accuracy_checker/api_runner.py +155 -151
  57. msprobe/mindspore/api_accuracy_checker/base_compare_algorithm.py +196 -196
  58. msprobe/mindspore/api_accuracy_checker/cmd_parser.py +6 -0
  59. msprobe/mindspore/api_accuracy_checker/compute_element.py +238 -223
  60. msprobe/mindspore/api_accuracy_checker/main.py +8 -15
  61. msprobe/mindspore/api_accuracy_checker/type_mapping.py +113 -113
  62. msprobe/mindspore/api_accuracy_checker/utils.py +79 -62
  63. msprobe/mindspore/cell_processor.py +34 -34
  64. msprobe/mindspore/common/const.py +106 -87
  65. msprobe/mindspore/common/log.py +37 -37
  66. msprobe/mindspore/common/utils.py +81 -57
  67. msprobe/mindspore/compare/distributed_compare.py +75 -75
  68. msprobe/mindspore/compare/ms_compare.py +219 -117
  69. msprobe/mindspore/compare/ms_graph_compare.py +348 -317
  70. msprobe/mindspore/compare/ms_to_pt_api.yaml +399 -399
  71. msprobe/mindspore/debugger/debugger_config.py +66 -74
  72. msprobe/mindspore/debugger/precision_debugger.py +126 -107
  73. msprobe/mindspore/dump/dump_tool_factory.py +35 -35
  74. msprobe/mindspore/dump/hook_cell/api_registry.py +118 -104
  75. msprobe/mindspore/dump/hook_cell/hook_cell.py +55 -53
  76. msprobe/mindspore/dump/hook_cell/support_wrap_ops.yaml +922 -925
  77. msprobe/mindspore/dump/hook_cell/wrap_api.py +113 -0
  78. msprobe/mindspore/dump/jit_dump.py +72 -56
  79. msprobe/mindspore/dump/kernel_graph_dump.py +59 -60
  80. msprobe/mindspore/dump/kernel_kbyk_dump.py +64 -65
  81. msprobe/mindspore/free_benchmark/api_pynative_self_check.py +116 -116
  82. msprobe/mindspore/free_benchmark/common/config.py +12 -12
  83. msprobe/mindspore/free_benchmark/common/handler_params.py +17 -17
  84. msprobe/mindspore/free_benchmark/common/utils.py +71 -71
  85. msprobe/mindspore/free_benchmark/data/support_wrap_ops.yaml +842 -842
  86. msprobe/mindspore/free_benchmark/decorator/dec_forward.py +43 -42
  87. msprobe/mindspore/free_benchmark/decorator/decorator_factory.py +107 -107
  88. msprobe/mindspore/free_benchmark/handler/base_handler.py +90 -90
  89. msprobe/mindspore/free_benchmark/handler/check_handler.py +41 -41
  90. msprobe/mindspore/free_benchmark/handler/fix_handler.py +36 -36
  91. msprobe/mindspore/free_benchmark/handler/handler_factory.py +21 -21
  92. msprobe/mindspore/free_benchmark/perturbation/add_noise.py +67 -67
  93. msprobe/mindspore/free_benchmark/perturbation/base_perturbation.py +21 -21
  94. msprobe/mindspore/free_benchmark/perturbation/bit_noise.py +63 -63
  95. msprobe/mindspore/free_benchmark/perturbation/exchange_value.py +51 -0
  96. msprobe/mindspore/free_benchmark/perturbation/improve_precision.py +35 -34
  97. msprobe/mindspore/free_benchmark/perturbation/no_change.py +12 -12
  98. msprobe/mindspore/free_benchmark/perturbation/perturbation_factory.py +29 -27
  99. msprobe/mindspore/free_benchmark/self_check_tool_factory.py +33 -33
  100. msprobe/mindspore/grad_probe/global_context.py +90 -91
  101. msprobe/mindspore/grad_probe/grad_analyzer.py +231 -231
  102. msprobe/mindspore/grad_probe/grad_monitor.py +27 -27
  103. msprobe/mindspore/grad_probe/grad_stat_csv.py +131 -131
  104. msprobe/mindspore/grad_probe/hook.py +94 -92
  105. msprobe/mindspore/grad_probe/utils.py +29 -28
  106. msprobe/mindspore/ms_config.py +128 -126
  107. msprobe/mindspore/overflow_check/kernel_graph_overflow_check.py +44 -45
  108. msprobe/mindspore/overflow_check/overflow_check_tool_factory.py +34 -34
  109. msprobe/mindspore/runtime.py +4 -4
  110. msprobe/mindspore/service.py +378 -354
  111. msprobe/mindspore/task_handler_factory.py +24 -24
  112. msprobe/msprobe.py +105 -107
  113. msprobe/pytorch/__init__.py +3 -3
  114. msprobe/pytorch/api_accuracy_checker/common/config.py +53 -55
  115. msprobe/pytorch/api_accuracy_checker/common/utils.py +214 -165
  116. msprobe/pytorch/api_accuracy_checker/compare/algorithm.py +213 -213
  117. msprobe/pytorch/api_accuracy_checker/compare/api_precision_compare.py +606 -581
  118. msprobe/pytorch/api_accuracy_checker/compare/api_precision_standard.yaml +132 -132
  119. msprobe/pytorch/api_accuracy_checker/compare/api_precision_threshold.yaml +390 -390
  120. msprobe/pytorch/api_accuracy_checker/compare/compare.py +386 -381
  121. msprobe/pytorch/api_accuracy_checker/compare/compare_column.py +73 -73
  122. msprobe/pytorch/api_accuracy_checker/compare/compare_utils.py +245 -244
  123. msprobe/pytorch/api_accuracy_checker/config.yaml +10 -10
  124. msprobe/pytorch/api_accuracy_checker/run_ut/data_generate.py +335 -332
  125. msprobe/pytorch/api_accuracy_checker/run_ut/multi_run_ut.py +200 -199
  126. msprobe/pytorch/api_accuracy_checker/run_ut/run_overflow_check.py +133 -134
  127. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut.py +592 -581
  128. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut_utils.py +70 -74
  129. msprobe/pytorch/api_accuracy_checker/run_ut/torch_ut_setting.json +7 -4
  130. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/attl.py +197 -202
  131. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/client.py +325 -324
  132. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/device_dispatch.py +204 -204
  133. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/server.py +219 -218
  134. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/ssl_config.py +10 -10
  135. msprobe/pytorch/bench_functions/__init__.py +15 -15
  136. msprobe/pytorch/bench_functions/apply_adam_w.py +28 -28
  137. msprobe/pytorch/bench_functions/confusion_transpose.py +19 -19
  138. msprobe/pytorch/bench_functions/fast_gelu.py +55 -55
  139. msprobe/pytorch/bench_functions/layer_norm_eval.py +6 -6
  140. msprobe/pytorch/bench_functions/linear.py +12 -12
  141. msprobe/pytorch/bench_functions/matmul_backward.py +48 -48
  142. msprobe/pytorch/bench_functions/npu_fusion_attention.py +509 -421
  143. msprobe/pytorch/bench_functions/rms_norm.py +15 -15
  144. msprobe/pytorch/bench_functions/rotary_mul.py +52 -52
  145. msprobe/pytorch/bench_functions/scaled_mask_softmax.py +26 -26
  146. msprobe/pytorch/bench_functions/swiglu.py +55 -55
  147. msprobe/pytorch/common/__init__.py +2 -2
  148. msprobe/pytorch/common/compare_script.template +14 -14
  149. msprobe/pytorch/common/log.py +20 -31
  150. msprobe/pytorch/common/parse_json.py +39 -39
  151. msprobe/pytorch/common/utils.py +305 -300
  152. msprobe/pytorch/compare/distributed_compare.py +66 -66
  153. msprobe/pytorch/compare/mapping.yaml +607 -607
  154. msprobe/pytorch/compare/match.py +34 -33
  155. msprobe/pytorch/compare/pt_compare.py +50 -40
  156. msprobe/pytorch/debugger/debugger_config.py +95 -95
  157. msprobe/pytorch/debugger/precision_debugger.py +125 -125
  158. msprobe/pytorch/free_benchmark/__init__.py +8 -8
  159. msprobe/pytorch/free_benchmark/common/constant.py +70 -70
  160. msprobe/pytorch/free_benchmark/common/counter.py +71 -71
  161. msprobe/pytorch/free_benchmark/common/enums.py +37 -37
  162. msprobe/pytorch/free_benchmark/common/params.py +129 -129
  163. msprobe/pytorch/free_benchmark/common/utils.py +102 -102
  164. msprobe/pytorch/free_benchmark/compare/grad_saver.py +179 -179
  165. msprobe/pytorch/free_benchmark/compare/single_benchmark.py +104 -104
  166. msprobe/pytorch/free_benchmark/main.py +105 -105
  167. msprobe/pytorch/free_benchmark/perturbed_layers/base_layer.py +13 -13
  168. msprobe/pytorch/free_benchmark/perturbed_layers/layer_factory.py +41 -41
  169. msprobe/pytorch/free_benchmark/perturbed_layers/npu/add_noise.py +90 -90
  170. msprobe/pytorch/free_benchmark/perturbed_layers/npu/bit_noise.py +104 -104
  171. msprobe/pytorch/free_benchmark/perturbed_layers/npu/change_value.py +63 -63
  172. msprobe/pytorch/free_benchmark/perturbed_layers/npu/improve_precision.py +68 -68
  173. msprobe/pytorch/free_benchmark/perturbed_layers/npu/no_change.py +28 -28
  174. msprobe/pytorch/free_benchmark/perturbed_layers/npu/npu_base_layser.py +45 -45
  175. msprobe/pytorch/free_benchmark/perturbed_layers/run_cpu.py +19 -19
  176. msprobe/pytorch/free_benchmark/result_handlers/base_handler.py +217 -217
  177. msprobe/pytorch/free_benchmark/result_handlers/check_handler.py +39 -39
  178. msprobe/pytorch/free_benchmark/result_handlers/fix_handler.py +23 -23
  179. msprobe/pytorch/free_benchmark/result_handlers/handler_factory.py +30 -30
  180. msprobe/pytorch/free_benchmark/result_handlers/preheat_handler.py +170 -170
  181. msprobe/pytorch/function_factory.py +76 -75
  182. msprobe/pytorch/functional/dump_module.py +39 -39
  183. msprobe/pytorch/grad_probe/grad_monitor.py +91 -90
  184. msprobe/pytorch/grad_probe/grad_stat_csv.py +128 -128
  185. msprobe/pytorch/hook_module/api_registry.py +161 -161
  186. msprobe/pytorch/hook_module/hook_module.py +120 -120
  187. msprobe/pytorch/hook_module/support_wrap_ops.yaml +1879 -1877
  188. msprobe/pytorch/hook_module/utils.py +30 -29
  189. msprobe/pytorch/hook_module/wrap_aten.py +110 -110
  190. msprobe/pytorch/hook_module/wrap_distributed.py +78 -78
  191. msprobe/pytorch/hook_module/wrap_functional.py +105 -105
  192. msprobe/pytorch/hook_module/wrap_npu_custom.py +93 -84
  193. msprobe/pytorch/hook_module/wrap_tensor.py +71 -71
  194. msprobe/pytorch/hook_module/wrap_torch.py +86 -86
  195. msprobe/pytorch/hook_module/wrap_vf.py +62 -62
  196. msprobe/pytorch/module_processer.py +138 -138
  197. msprobe/pytorch/online_dispatch/__init__.py +20 -20
  198. msprobe/pytorch/online_dispatch/compare.py +236 -236
  199. msprobe/pytorch/online_dispatch/dispatch.py +271 -271
  200. msprobe/pytorch/online_dispatch/dump_compare.py +155 -156
  201. msprobe/pytorch/online_dispatch/single_compare.py +391 -391
  202. msprobe/pytorch/online_dispatch/torch_ops_config.yaml +49 -49
  203. msprobe/pytorch/online_dispatch/utils.py +130 -146
  204. msprobe/pytorch/parse.py +4 -4
  205. msprobe/pytorch/parse_tool/cli.py +32 -32
  206. msprobe/pytorch/parse_tool/lib/compare.py +260 -271
  207. msprobe/pytorch/parse_tool/lib/config.py +52 -52
  208. msprobe/pytorch/parse_tool/lib/file_desc.py +31 -31
  209. msprobe/pytorch/parse_tool/lib/interactive_cli.py +102 -102
  210. msprobe/pytorch/parse_tool/lib/parse_exception.py +54 -54
  211. msprobe/pytorch/parse_tool/lib/parse_tool.py +158 -158
  212. msprobe/pytorch/parse_tool/lib/utils.py +316 -321
  213. msprobe/pytorch/parse_tool/lib/visualization.py +85 -91
  214. msprobe/pytorch/pt_config.py +188 -187
  215. msprobe/pytorch/service.py +246 -252
  216. mindstudio_probe-1.0.3.dist-info/RECORD +0 -272
  217. msprobe/config/README.md +0 -539
  218. msprobe/mindspore/doc/compare.md +0 -58
  219. msprobe/mindspore/doc/dump.md +0 -217
  220. msprobe/mindspore/dump/hook_cell/wrap_functional.py +0 -91
  221. msprobe/mindspore/dump/hook_cell/wrap_tensor.py +0 -63
  222. msprobe/pytorch/doc/FAQ.md +0 -193
  223. msprobe/pytorch/doc/api_accuracy_checker.md +0 -313
  224. msprobe/pytorch/doc/api_accuracy_checker_online.md +0 -187
  225. msprobe/pytorch/doc/dump.md +0 -260
  226. msprobe/pytorch/doc/msprobe/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/206/320/245/342/226/221/321/206/320/235/320/276dump/321/206/320/260/320/227/321/205/320/227/320/226/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -182
  227. msprobe/pytorch/doc/ptdbg_ascend_compare.md +0 -240
  228. msprobe/pytorch/doc/ptdbg_ascend_overview.md +0 -68
  229. msprobe/pytorch/doc/ptdbg_ascend_quickstart.md +0 -381
  230. msprobe/pytorch/doc/run_overflow_check.md +0 -25
  231. msprobe/pytorch/doc//321/205/320/254/320/270/321/207/342/225/221/342/224/220/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/206/320/277/320/244/321/205/320/277/342/225/243.md +0 -90
  232. msprobe/pytorch/doc//321/206/320/247/320/260/321/206/320/260/320/227/321/206/320/255/320/226/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/205/320/254/342/225/221/321/206/320/251/320/277/321/211/320/272/320/234/321/210/320/277/320/221/321/205/320/242/320/234/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -151
  233. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.0.4.dist-info}/top_level.txt +0 -0
  234. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_1.png +0 -0
  235. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_2.png +0 -0
  236. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_3.png +0 -0
  237. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_4.png +0 -0
  238. /msprobe/{pytorch/doc → docs}/img/GPT-3_1.png +0 -0
  239. /msprobe/{pytorch/doc → docs}/img/GPT-3_2.png +0 -0
  240. /msprobe/{pytorch/doc → docs}/img/GPT-3_3.png +0 -0
  241. /msprobe/{pytorch/doc → docs}/img/GPT-3_4.png +0 -0
  242. /msprobe/{pytorch/doc → docs}/img/GPT-3_5.png +0 -0
  243. /msprobe/{pytorch/doc → docs}/img/GPT-3_6.png +0 -0
  244. /msprobe/{pytorch/doc → docs}/img/GPT-3_7.png +0 -0
  245. /msprobe/{pytorch/doc → docs}/img/GPT-3_8.png +0 -0
  246. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_1.png +0 -0
  247. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_2.png +0 -0
  248. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_details.png +0 -0
  249. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_result.png +0 -0
  250. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_details.png +0 -0
  251. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_result.png +0 -0
  252. /msprobe/{pytorch/doc → docs}/img/auto_analyze_log.png +0 -0
  253. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl.png +0 -0
  254. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl_md5.png.png +0 -0
  255. /msprobe/{pytorch/doc → docs}/img/cpu_info.png +0 -0
  256. /msprobe/{config → docs}/img/free_benchmark.png +0 -0
  257. /msprobe/{doc/grad_probe/img/image-1.png → docs/img/grad_probe_image-1.png} +0 -0
  258. /msprobe/{doc/grad_probe/img/image-2.png → docs/img/grad_probe_image-2.png} +0 -0
  259. /msprobe/{doc/grad_probe/img/image-3.png → docs/img/grad_probe_image-3.png} +0 -0
  260. /msprobe/{doc/grad_probe/img/image-4.png → docs/img/grad_probe_image-4.png} +0 -0
  261. /msprobe/{doc/grad_probe/img/image.png → docs/img/grad_probe_image.png} +0 -0
  262. /msprobe/{pytorch/doc → docs}/img/module_compare.png +0 -0
@@ -1,59 +1,59 @@
1
- from msprobe.core.common.const import Const
2
-
3
-
4
- class DataProcessorFactory:
5
- _data_processor = {}
6
- _module_processor = {}
7
-
8
- @classmethod
9
- def register_processor(cls, framework, task, processor_class):
10
- key = (framework, task)
11
- cls._data_processor[key] = processor_class
12
-
13
- @classmethod
14
- def register_module_processor(cls, framework, processor_class):
15
- cls._module_processor[framework] = processor_class
16
-
17
- @classmethod
18
- def get_module_processor(cls, framework):
19
- processor_class = cls._module_processor.get(framework)
20
- if not processor_class:
21
- raise ValueError(f"ModuleProcesser not found for framework: {framework}")
22
- return processor_class
23
-
24
- @classmethod
25
- def create_processor(cls, config, data_writer):
26
- cls.register_processors(config.framework)
27
- task = Const.KERNEL_DUMP if config.level == "L2" else config.task
28
- key = (config.framework, task)
29
- processor_class = cls._data_processor.get(key)
30
- if not processor_class:
31
- raise ValueError(f"Processor not found for framework: {config.framework}, task: {config.task}")
32
- return processor_class(config, data_writer)
33
-
34
- @classmethod
35
- def register_processors(cls, framework):
36
- if framework == Const.PT_FRAMEWORK:
37
- from .pytorch_processor import (
38
- StatisticsDataProcessor as PytorchStatisticsDataProcessor,
39
- TensorDataProcessor as PytorchTensorDataProcessor,
40
- OverflowCheckDataProcessor as PytorchOverflowCheckDataProcessor,
41
- FreeBenchmarkDataProcessor as PytorchFreeBenchmarkDataProcessor,
42
- KernelDumpDataProcessor as PytorchKernelDumpDataProcessor
43
- )
44
- from ....pytorch.module_processer import ModuleProcesser
45
- cls.register_processor(Const.PT_FRAMEWORK, Const.STATISTICS, PytorchStatisticsDataProcessor)
46
- cls.register_processor(Const.PT_FRAMEWORK, Const.TENSOR, PytorchTensorDataProcessor)
47
- cls.register_processor(Const.PT_FRAMEWORK, Const.OVERFLOW_CHECK, PytorchOverflowCheckDataProcessor)
48
- cls.register_processor(Const.PT_FRAMEWORK, Const.FREE_BENCHMARK, PytorchFreeBenchmarkDataProcessor)
49
- cls.register_processor(Const.PT_FRAMEWORK, Const.KERNEL_DUMP, PytorchKernelDumpDataProcessor)
50
- cls.register_module_processor(Const.PT_FRAMEWORK, ModuleProcesser)
51
- elif framework == Const.MS_FRAMEWORK:
52
- from .mindspore_processor import (
53
- StatisticsDataProcessor as MindsporeStatisticsDataProcessor,
54
- TensorDataProcessor as MindsporeTensorDataProcessor,
55
- OverflowCheckDataProcessor as MindsporeOverflowCheckDataProcessor
56
- )
57
- cls.register_processor(Const.MS_FRAMEWORK, Const.STATISTICS, MindsporeStatisticsDataProcessor)
58
- cls.register_processor(Const.MS_FRAMEWORK, Const.TENSOR, MindsporeTensorDataProcessor)
59
- cls.register_processor(Const.MS_FRAMEWORK, Const.OVERFLOW_CHECK, MindsporeOverflowCheckDataProcessor)
1
+ from msprobe.core.common.const import Const
2
+
3
+
4
+ class DataProcessorFactory:
5
+ _data_processor = {}
6
+ _module_processor = {}
7
+
8
+ @classmethod
9
+ def register_processor(cls, framework, task, processor_class):
10
+ key = (framework, task)
11
+ cls._data_processor[key] = processor_class
12
+
13
+ @classmethod
14
+ def register_module_processor(cls, framework, processor_class):
15
+ cls._module_processor[framework] = processor_class
16
+
17
+ @classmethod
18
+ def get_module_processor(cls, framework):
19
+ processor_class = cls._module_processor.get(framework)
20
+ if not processor_class:
21
+ raise ValueError(f"ModuleProcesser not found for framework: {framework}")
22
+ return processor_class
23
+
24
+ @classmethod
25
+ def create_processor(cls, config, data_writer):
26
+ cls.register_processors(config.framework)
27
+ task = Const.KERNEL_DUMP if config.level == "L2" else config.task
28
+ key = (config.framework, task)
29
+ processor_class = cls._data_processor.get(key)
30
+ if not processor_class:
31
+ raise ValueError(f"Processor not found for framework: {config.framework}, task: {config.task}")
32
+ return processor_class(config, data_writer)
33
+
34
+ @classmethod
35
+ def register_processors(cls, framework):
36
+ if framework == Const.PT_FRAMEWORK:
37
+ from .pytorch_processor import (
38
+ StatisticsDataProcessor as PytorchStatisticsDataProcessor,
39
+ TensorDataProcessor as PytorchTensorDataProcessor,
40
+ OverflowCheckDataProcessor as PytorchOverflowCheckDataProcessor,
41
+ FreeBenchmarkDataProcessor as PytorchFreeBenchmarkDataProcessor,
42
+ KernelDumpDataProcessor as PytorchKernelDumpDataProcessor
43
+ )
44
+ from ....pytorch.module_processer import ModuleProcesser
45
+ cls.register_processor(Const.PT_FRAMEWORK, Const.STATISTICS, PytorchStatisticsDataProcessor)
46
+ cls.register_processor(Const.PT_FRAMEWORK, Const.TENSOR, PytorchTensorDataProcessor)
47
+ cls.register_processor(Const.PT_FRAMEWORK, Const.OVERFLOW_CHECK, PytorchOverflowCheckDataProcessor)
48
+ cls.register_processor(Const.PT_FRAMEWORK, Const.FREE_BENCHMARK, PytorchFreeBenchmarkDataProcessor)
49
+ cls.register_processor(Const.PT_FRAMEWORK, Const.KERNEL_DUMP, PytorchKernelDumpDataProcessor)
50
+ cls.register_module_processor(Const.PT_FRAMEWORK, ModuleProcesser)
51
+ elif framework == Const.MS_FRAMEWORK:
52
+ from .mindspore_processor import (
53
+ StatisticsDataProcessor as MindsporeStatisticsDataProcessor,
54
+ TensorDataProcessor as MindsporeTensorDataProcessor,
55
+ OverflowCheckDataProcessor as MindsporeOverflowCheckDataProcessor
56
+ )
57
+ cls.register_processor(Const.MS_FRAMEWORK, Const.STATISTICS, MindsporeStatisticsDataProcessor)
58
+ cls.register_processor(Const.MS_FRAMEWORK, Const.TENSOR, MindsporeTensorDataProcessor)
59
+ cls.register_processor(Const.MS_FRAMEWORK, Const.OVERFLOW_CHECK, MindsporeOverflowCheckDataProcessor)
@@ -1,198 +1,186 @@
1
- # Copyright 2024 Huawei Technologies Co., Ltd
2
- #
3
- # Licensed under the Apache License, Version 2.0 (the "License");
4
- # you may not use this file except in compliance with the License.
5
- # You may obtain a copy of the License at
6
- #
7
- # http://www.apache.org/licenses/LICENSE-2.0
8
- #
9
- # Unless required by applicable law or agreed to in writing, software
10
- # distributed under the License is distributed on an "AS IS" BASIS,
11
- # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
12
- # See the License for the specific language governing permissions and
13
- # limitations under the License.
14
- # ============================================================================
15
-
16
- import zlib
17
-
18
- import mindspore as ms
19
- from mindspore import ops
20
- import numpy as np
21
-
22
- from msprobe.core.common.const import Const
23
- from msprobe.core.data_dump.data_processor.base import (BaseDataProcessor, TensorStatInfo,
24
- ModuleForwardInputsOutputs, ModuleBackwardInputsOutputs)
25
- from msprobe.core.common.file_check import path_len_exceeds_limit
26
- from msprobe.mindspore.dump.hook_cell.wrap_functional import load_ops_functions
27
- from msprobe.mindspore.common.utils import convert_bf16_to_fp32, save_tensor_as_npy
28
- from msprobe.mindspore.common.log import logger
29
- from msprobe.mindspore.dump.hook_cell.api_registry import api_register
30
-
31
-
32
- class MindsporeDataProcessor(BaseDataProcessor):
33
- mindspore_special_type = tuple([ms.Tensor])
34
- ops_func, mint_ops_func, _ = load_ops_functions()
35
-
36
- def __init__(self, config, data_writer):
37
- super().__init__(config, data_writer)
38
- self.mindspore_object_key = {
39
- "dtype": self.analyze_dtype_in_kwargs
40
- }
41
-
42
- @staticmethod
43
- def get_md5_for_tensor(x):
44
- x = convert_bf16_to_fp32(x)
45
- tensor_bytes = x.asnumpy().tobytes()
46
- crc32_hash = zlib.crc32(tensor_bytes)
47
- return f"{crc32_hash:08x}"
48
-
49
- @staticmethod
50
- def analyze_dtype_in_kwargs(element):
51
- return {"type": "mindspore.dtype", "value": str(element)}
52
-
53
- @staticmethod
54
- def _analyze_builtin(arg):
55
- single_arg = {}
56
- if isinstance(arg, slice):
57
- single_arg.update({"type": "slice"})
58
- # slice参数中可能存在tensor类型,json序列化,需要转换为python数值类型
59
- values = [
60
- value if not isinstance(value, ms.Tensor) else value.item()
61
- for value in [arg.start, arg.stop, arg.step]
62
- ]
63
- single_arg.update({"value": values})
64
- else:
65
- single_arg.update({"type": type(arg).__name__})
66
- single_arg.update({"value": arg})
67
- return single_arg
68
-
69
- @classmethod
70
- def get_special_types(cls):
71
- return super().get_special_types() + cls.mindspore_special_type
72
-
73
- def get_stat_info(self, data):
74
- tensor_stat = TensorStatInfo()
75
- if data.numel() == 0:
76
- return tensor_stat
77
- elif data.dtype == ms.bool_:
78
- data_np = data.asnumpy()
79
- tensor_stat.max = np.max(data_np).item()
80
- tensor_stat.min = np.min(data_np).item()
81
- elif not data.shape:
82
- tensor_stat.max = tensor_stat.min = tensor_stat.mean = tensor_stat.norm = data.item()
83
- elif data.dtype == ms.complex64 or data.dtype == ms.complex128:
84
- data_abs = np.abs(data.asnumpy())
85
- tensor_stat.max = np.max(data_abs).item()
86
- tensor_stat.min = np.min(data_abs).item()
87
- tensor_stat.mean = np.mean(data_abs).item()
88
- tensor_stat.norm = np.linalg.norm(data_abs).item()
89
- else:
90
- if data.dtype == ms.bfloat16 or not ops.is_floating_point(data):
91
- data = data.to(ms.float32)
92
- api_register.norm_inner_op_set_ori_func()
93
- tensor_stat.max = self.mint_ops_func["max"](data).item()
94
- tensor_stat.min = self.mint_ops_func["min"](data).item()
95
- tensor_stat.mean = self.mint_ops_func["mean"](data).item()
96
- tensor_stat.norm = self.ops_func["norm"](data).item()
97
- api_register.norm_inner_op_set_hook_func()
98
- return tensor_stat
99
-
100
- def analyze_single_element(self, element, suffix_stack):
101
- if suffix_stack and suffix_stack[-1] in self.mindspore_object_key:
102
- return self.mindspore_object_key[suffix_stack[-1]](element)
103
-
104
- converted_numpy, numpy_type = self._convert_numpy_to_builtin(element)
105
- if converted_numpy is not element:
106
- return self._analyze_numpy(converted_numpy, numpy_type)
107
- if isinstance(element, ms.Tensor):
108
- return self._analyze_tensor(element, Const.SEP.join(suffix_stack))
109
-
110
- if isinstance(element, (bool, int, float, str, slice)):
111
- return self._analyze_builtin(element)
112
- return {}
113
-
114
- def _analyze_tensor(self, tensor, suffix):
115
- tensor_stat = self.get_stat_info(tensor)
116
- tensor_json = {
117
- 'type': 'mindspore.Tensor',
118
- 'dtype': str(tensor.dtype),
119
- 'shape': tensor.shape,
120
- 'Max': self.transfer_type(tensor_stat.max),
121
- 'Min': self.transfer_type(tensor_stat.min),
122
- 'Mean': self.transfer_type(tensor_stat.mean),
123
- 'Norm': self.transfer_type(tensor_stat.norm),
124
- }
125
- if self.config.summary_mode == Const.MD5:
126
- tensor_md5 = self.get_md5_for_tensor(tensor)
127
- tensor_json.update({Const.MD5: tensor_md5})
128
- return tensor_json
129
-
130
-
131
- class StatisticsDataProcessor(MindsporeDataProcessor):
132
- pass
133
-
134
-
135
- class TensorDataProcessor(MindsporeDataProcessor):
136
- def _analyze_tensor(self, tensor, suffix):
137
- dump_data_name, file_path = self.get_save_file_path(suffix)
138
- single_arg = super()._analyze_tensor(tensor, suffix)
139
- single_arg.update({"data_name": dump_data_name})
140
- save_tensor_as_npy(tensor, file_path)
141
- return single_arg
142
-
143
-
144
- class OverflowCheckDataProcessor(MindsporeDataProcessor):
145
- __slots__ = ["cached_tensors_and_file_paths"]
146
-
147
- def __init__(self, config, data_writer):
148
- super().__init__(config, data_writer)
149
- self.cached_tensors_and_file_paths = {}
150
- self.real_overflow_nums = 0
151
- self.overflow_nums = config.overflow_nums
152
-
153
- @property
154
- def is_terminated(self):
155
- if self.overflow_nums == -1:
156
- return False
157
- if self.real_overflow_nums >= self.overflow_nums:
158
- logger.info(f"[msprobe] 超过预设溢出次数 当前溢出次数: {self.real_overflow_nums}")
159
- return True
160
- return False
161
-
162
- def analyze_forward(self, name, module, module_input_output: ModuleForwardInputsOutputs):
163
- self.has_overflow = False
164
- api_info_struct = super().analyze_forward(name, module, module_input_output)
165
- self.maybe_save_overflow_data()
166
- return api_info_struct if self.has_overflow else None
167
-
168
- def analyze_backward(self, name, module, module_input_output: ModuleBackwardInputsOutputs):
169
- self.has_overflow = False
170
- api_info_struct = super().analyze_backward(name, module, module_input_output)
171
- self.maybe_save_overflow_data()
172
- return api_info_struct if self.has_overflow else None
173
-
174
- def maybe_save_overflow_data(self):
175
- if self.has_overflow:
176
- for file_path, tensor in self.cached_tensors_and_file_paths.items():
177
- save_tensor_as_npy(tensor, file_path)
178
- self.real_overflow_nums += 1
179
- self.cached_tensors_and_file_paths = {}
180
-
181
- def _analyze_maybe_overflow_tensor(self, tensor_json):
182
- if tensor_json['Max'] is None:
183
- return
184
- if np.isinf(tensor_json['Max']) or np.isnan(tensor_json['Max']):
185
- self.has_overflow = True
186
- if np.isinf(tensor_json['Min']) or np.isnan(tensor_json['Min']):
187
- self.has_overflow = True
188
-
189
- def _analyze_tensor(self, tensor, suffix):
190
- dump_data_name, file_path = self.get_save_file_path(suffix)
191
- if not path_len_exceeds_limit(file_path):
192
- self.cached_tensors_and_file_paths.update({file_path: tensor})
193
- else:
194
- logger.warning(f'The file path {file_path} length exceeds limit.')
195
- single_arg = super()._analyze_tensor(tensor, suffix)
196
- self._analyze_maybe_overflow_tensor(single_arg)
197
- single_arg.update({"data_name": dump_data_name})
198
- return single_arg
1
+ # Copyright 2024 Huawei Technologies Co., Ltd
2
+ #
3
+ # Licensed under the Apache License, Version 2.0 (the "License");
4
+ # you may not use this file except in compliance with the License.
5
+ # You may obtain a copy of the License at
6
+ #
7
+ # http://www.apache.org/licenses/LICENSE-2.0
8
+ #
9
+ # Unless required by applicable law or agreed to in writing, software
10
+ # distributed under the License is distributed on an "AS IS" BASIS,
11
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
12
+ # See the License for the specific language governing permissions and
13
+ # limitations under the License.
14
+ # ============================================================================
15
+
16
+ import zlib
17
+
18
+ import mindspore as ms
19
+ from mindspore import mint, ops
20
+ import numpy as np
21
+
22
+ from msprobe.core.common.const import Const
23
+ from msprobe.core.data_dump.data_processor.base import (BaseDataProcessor, TensorStatInfo,
24
+ ModuleForwardInputsOutputs, ModuleBackwardInputsOutputs)
25
+ from msprobe.core.common.file_utils import path_len_exceeds_limit
26
+ from msprobe.mindspore.common.utils import convert_bf16_to_fp32, save_tensor_as_npy
27
+ from msprobe.mindspore.common.log import logger
28
+ from msprobe.mindspore.dump.hook_cell.api_registry import api_register
29
+
30
+
31
+ class MindsporeDataProcessor(BaseDataProcessor):
32
+ mindspore_special_type = tuple([ms.Tensor])
33
+
34
+ def __init__(self, config, data_writer):
35
+ super().__init__(config, data_writer)
36
+ self.mindspore_object_key = {
37
+ "dtype": self.analyze_dtype_in_kwargs
38
+ }
39
+
40
+ @staticmethod
41
+ def get_md5_for_tensor(x):
42
+ x = convert_bf16_to_fp32(x)
43
+ tensor_bytes = x.asnumpy().tobytes()
44
+ crc32_hash = zlib.crc32(tensor_bytes)
45
+ return f"{crc32_hash:08x}"
46
+
47
+ @staticmethod
48
+ def analyze_dtype_in_kwargs(element):
49
+ return {"type": "mindspore.dtype", "value": str(element)}
50
+
51
+ @classmethod
52
+ def get_special_types(cls):
53
+ return super().get_special_types() + cls.mindspore_special_type
54
+
55
+ def get_stat_info(self, data):
56
+ tensor_stat = TensorStatInfo()
57
+ if data.numel() == 0:
58
+ return tensor_stat
59
+ elif data.dtype == ms.bool_:
60
+ data_np = data.asnumpy()
61
+ tensor_stat.max = np.max(data_np).item()
62
+ tensor_stat.min = np.min(data_np).item()
63
+ elif not data.shape:
64
+ tensor_stat.max = tensor_stat.min = tensor_stat.mean = tensor_stat.norm = data.item()
65
+ elif data.dtype == ms.complex64 or data.dtype == ms.complex128:
66
+ data_abs = np.abs(data.asnumpy())
67
+ tensor_stat.max = np.max(data_abs).item()
68
+ tensor_stat.min = np.min(data_abs).item()
69
+ tensor_stat.mean = np.mean(data_abs).item()
70
+ tensor_stat.norm = np.linalg.norm(data_abs).item()
71
+ else:
72
+ if data.dtype == ms.bfloat16 or not ops.is_floating_point(data):
73
+ data = data.to(ms.float32)
74
+ api_register.norm_inner_op_set_ori_func()
75
+ get_max_value = api_register.mint_ops_ori_attr.get("max", mint.max)
76
+ get_min_value = api_register.mint_ops_ori_attr.get("min", mint.min)
77
+ get_mean_value = api_register.mint_ops_ori_attr.get("mean", mint.mean)
78
+ get_norm_value = api_register.functional_ori_attr.get("norm", ops.norm)
79
+ tensor_stat.max = get_max_value(data).item()
80
+ tensor_stat.min = get_min_value(data).item()
81
+ tensor_stat.mean = get_mean_value(data).item()
82
+ tensor_stat.norm = get_norm_value(data).item()
83
+ api_register.norm_inner_op_set_hook_func()
84
+ return tensor_stat
85
+
86
+ def analyze_single_element(self, element, suffix_stack):
87
+ if suffix_stack and suffix_stack[-1] in self.mindspore_object_key:
88
+ return self.mindspore_object_key[suffix_stack[-1]](element)
89
+
90
+ converted_numpy, numpy_type = self._convert_numpy_to_builtin(element)
91
+ if converted_numpy is not element:
92
+ return self._analyze_numpy(converted_numpy, numpy_type)
93
+ if isinstance(element, ms.Tensor):
94
+ return self._analyze_tensor(element, Const.SEP.join(suffix_stack))
95
+
96
+ if isinstance(element, (bool, int, float, str, slice, type(Ellipsis))):
97
+ return self._analyze_builtin(element)
98
+ return {}
99
+
100
+ def _analyze_tensor(self, tensor, suffix):
101
+ tensor_stat = self.get_stat_info(tensor)
102
+ tensor_json = {
103
+ 'type': 'mindspore.Tensor',
104
+ 'dtype': str(tensor.dtype),
105
+ 'shape': tensor.shape,
106
+ 'Max': self.transfer_type(tensor_stat.max),
107
+ 'Min': self.transfer_type(tensor_stat.min),
108
+ 'Mean': self.transfer_type(tensor_stat.mean),
109
+ 'Norm': self.transfer_type(tensor_stat.norm),
110
+ }
111
+ if self.config.summary_mode == Const.MD5:
112
+ tensor_md5 = self.get_md5_for_tensor(tensor)
113
+ tensor_json.update({Const.MD5: tensor_md5})
114
+ return tensor_json
115
+
116
+
117
+ class StatisticsDataProcessor(MindsporeDataProcessor):
118
+ pass
119
+
120
+
121
+ class TensorDataProcessor(MindsporeDataProcessor):
122
+ def _analyze_tensor(self, tensor, suffix):
123
+ dump_data_name, file_path = self.get_save_file_path(suffix)
124
+ single_arg = super()._analyze_tensor(tensor, suffix)
125
+ single_arg.update({"data_name": dump_data_name})
126
+ save_tensor_as_npy(tensor, file_path)
127
+ return single_arg
128
+
129
+
130
+ class OverflowCheckDataProcessor(MindsporeDataProcessor):
131
+ __slots__ = ["cached_tensors_and_file_paths"]
132
+
133
+ def __init__(self, config, data_writer):
134
+ super().__init__(config, data_writer)
135
+ self.has_overflow = False
136
+ self.cached_tensors_and_file_paths = {}
137
+ self.real_overflow_nums = 0
138
+ self.overflow_nums = config.overflow_nums
139
+
140
+ @property
141
+ def is_terminated(self):
142
+ if self.overflow_nums == -1:
143
+ return False
144
+ if self.real_overflow_nums >= self.overflow_nums:
145
+ return True
146
+ return False
147
+
148
+ def analyze_forward(self, name, module, module_input_output: ModuleForwardInputsOutputs):
149
+ self.has_overflow = False
150
+ api_info_struct = super().analyze_forward(name, module, module_input_output)
151
+ self.maybe_save_overflow_data()
152
+ return api_info_struct if self.has_overflow else None
153
+
154
+ def analyze_backward(self, name, module, module_input_output: ModuleBackwardInputsOutputs):
155
+ self.has_overflow = False
156
+ api_info_struct = super().analyze_backward(name, module, module_input_output)
157
+ self.maybe_save_overflow_data()
158
+ return api_info_struct if self.has_overflow else None
159
+
160
+ def maybe_save_overflow_data(self):
161
+ if self.has_overflow:
162
+ for file_path, tensor in self.cached_tensors_and_file_paths.items():
163
+ save_tensor_as_npy(tensor, file_path)
164
+ self.real_overflow_nums += 1
165
+ if self.overflow_nums != -1 and self.real_overflow_nums >= self.overflow_nums:
166
+ logger.info(f"[{Const.TOOL_NAME}] 超过预设溢出次数 当前溢出次数: {self.real_overflow_nums}")
167
+ self.cached_tensors_and_file_paths = {}
168
+
169
+ def _analyze_maybe_overflow_tensor(self, tensor_json):
170
+ if tensor_json['Max'] is None:
171
+ return
172
+ if np.isinf(tensor_json['Max']) or np.isnan(tensor_json['Max']):
173
+ self.has_overflow = True
174
+ if np.isinf(tensor_json['Min']) or np.isnan(tensor_json['Min']):
175
+ self.has_overflow = True
176
+
177
+ def _analyze_tensor(self, tensor, suffix):
178
+ dump_data_name, file_path = self.get_save_file_path(suffix)
179
+ if not path_len_exceeds_limit(file_path):
180
+ self.cached_tensors_and_file_paths.update({file_path: tensor})
181
+ else:
182
+ logger.warning(f'The file path {file_path} length exceeds limit.')
183
+ single_arg = super()._analyze_tensor(tensor, suffix)
184
+ self._analyze_maybe_overflow_tensor(single_arg)
185
+ single_arg.update({"data_name": dump_data_name})
186
+ return single_arg