mindstudio-probe 1.0.3__py3-none-any.whl → 1.1.0__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (278) hide show
  1. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/LICENSE +201 -201
  2. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/METADATA +36 -34
  3. mindstudio_probe-1.1.0.dist-info/RECORD +287 -0
  4. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/WHEEL +1 -1
  5. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/entry_points.txt +1 -0
  6. msprobe/README.md +131 -237
  7. msprobe/__init__.py +16 -1
  8. msprobe/{config/config.json → config.json} +47 -49
  9. msprobe/core/advisor/advisor.py +124 -124
  10. msprobe/core/advisor/advisor_const.py +58 -59
  11. msprobe/core/advisor/advisor_result.py +58 -58
  12. msprobe/core/common/const.py +402 -318
  13. msprobe/core/common/exceptions.py +99 -99
  14. msprobe/core/common/{file_check.py → file_utils.py} +523 -283
  15. msprobe/core/common/inplace_op_checker.py +38 -0
  16. msprobe/core/common/inplace_ops.yaml +251 -0
  17. msprobe/core/common/log.py +86 -69
  18. msprobe/core/common/utils.py +371 -616
  19. msprobe/core/common_config.py +78 -71
  20. msprobe/core/compare/acc_compare.py +472 -298
  21. msprobe/core/compare/check.py +180 -95
  22. msprobe/core/compare/compare_cli.py +69 -49
  23. msprobe/core/compare/highlight.py +259 -222
  24. msprobe/core/compare/multiprocessing_compute.py +174 -149
  25. msprobe/core/compare/npy_compare.py +310 -295
  26. msprobe/core/compare/utils.py +464 -429
  27. msprobe/core/data_dump/data_collector.py +153 -144
  28. msprobe/core/data_dump/data_processor/base.py +337 -293
  29. msprobe/core/data_dump/data_processor/factory.py +76 -59
  30. msprobe/core/data_dump/data_processor/mindspore_processor.py +192 -198
  31. msprobe/core/data_dump/data_processor/pytorch_processor.py +383 -389
  32. msprobe/core/data_dump/json_writer.py +117 -116
  33. msprobe/core/data_dump/scope.py +194 -178
  34. msprobe/core/grad_probe/constant.py +74 -70
  35. msprobe/core/grad_probe/grad_compare.py +170 -175
  36. msprobe/core/grad_probe/utils.py +77 -52
  37. msprobe/docs/01.installation.md +99 -0
  38. msprobe/docs/02.config_introduction.md +137 -0
  39. msprobe/docs/03.config_examples.md +237 -0
  40. msprobe/docs/04.acl_config_examples.md +78 -0
  41. msprobe/docs/05.data_dump_PyTorch.md +326 -0
  42. msprobe/docs/06.data_dump_MindSpore.md +285 -0
  43. msprobe/docs/07.accuracy_checker_PyTorch.md +297 -0
  44. msprobe/docs/08.accuracy_checker_online_PyTorch.md +238 -0
  45. msprobe/docs/09.accuracy_checker_MindSpore.md +68 -0
  46. msprobe/docs/10.accuracy_compare_PyTorch.md +327 -0
  47. msprobe/docs/11.accuracy_compare_MindSpore.md +333 -0
  48. msprobe/docs/12.overflow_check_PyTorch.md +79 -0
  49. msprobe/docs/13.overflow_check_MindSpore.md +31 -0
  50. msprobe/{pytorch/doc/parse_tool.md → docs/14.data_parse_PyTorch.md} +283 -286
  51. msprobe/docs/15.free_benchmarking_PyTorch.md +170 -0
  52. msprobe/docs/16.free_benchmarking_MindSpore.md +140 -0
  53. msprobe/{doc/grad_probe/grad_probe.md → docs/17.grad_probe.md} +205 -207
  54. msprobe/{pytorch/doc//321/205/320/254/320/270/321/207/342/225/221/342/224/220/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/206/320/277/320/244/321/205/320/277/342/225/243.md → docs/18.online_dispatch.md} +89 -90
  55. msprobe/docs/FAQ.md +189 -0
  56. msprobe/docs/S02.report_free_benchmarking_validation_performance_baseline.md +146 -0
  57. msprobe/docs/img/free_benchmark_framework.png +0 -0
  58. msprobe/docs/img/ms_dump.png +0 -0
  59. msprobe/docs/img/ms_layer.png +0 -0
  60. msprobe/docs/img/pt_dump.png +0 -0
  61. msprobe/mindspore/__init__.py +2 -1
  62. msprobe/mindspore/api_accuracy_checker/api_accuracy_checker.py +278 -245
  63. msprobe/mindspore/api_accuracy_checker/api_info.py +76 -69
  64. msprobe/mindspore/api_accuracy_checker/api_runner.py +155 -151
  65. msprobe/mindspore/api_accuracy_checker/base_compare_algorithm.py +196 -196
  66. msprobe/mindspore/api_accuracy_checker/cmd_parser.py +6 -0
  67. msprobe/mindspore/api_accuracy_checker/compute_element.py +238 -223
  68. msprobe/mindspore/api_accuracy_checker/main.py +8 -15
  69. msprobe/mindspore/api_accuracy_checker/type_mapping.py +113 -113
  70. msprobe/mindspore/api_accuracy_checker/utils.py +79 -62
  71. msprobe/mindspore/cell_processor.py +58 -34
  72. msprobe/mindspore/common/const.py +108 -87
  73. msprobe/mindspore/common/log.py +37 -37
  74. msprobe/mindspore/common/utils.py +97 -57
  75. msprobe/mindspore/compare/distributed_compare.py +62 -75
  76. msprobe/mindspore/compare/layer_mapping.py +146 -0
  77. msprobe/mindspore/compare/modify_mapping.py +107 -0
  78. msprobe/mindspore/compare/ms_compare.py +357 -117
  79. msprobe/mindspore/compare/ms_graph_compare.py +364 -317
  80. msprobe/mindspore/compare/ms_to_pt_api.yaml +399 -399
  81. msprobe/mindspore/debugger/debugger_config.py +69 -74
  82. msprobe/mindspore/debugger/precision_debugger.py +150 -107
  83. msprobe/mindspore/dump/dump_tool_factory.py +50 -35
  84. msprobe/mindspore/dump/hook_cell/api_registry.py +128 -104
  85. msprobe/mindspore/dump/hook_cell/hook_cell.py +55 -53
  86. msprobe/mindspore/dump/hook_cell/primitive_hooks.py +206 -0
  87. msprobe/mindspore/dump/hook_cell/support_wrap_ops.yaml +994 -925
  88. msprobe/mindspore/dump/hook_cell/wrap_api.py +121 -0
  89. msprobe/mindspore/dump/jit_dump.py +96 -56
  90. msprobe/mindspore/dump/kernel_graph_dump.py +75 -60
  91. msprobe/mindspore/dump/kernel_kbyk_dump.py +79 -65
  92. msprobe/mindspore/free_benchmark/api_pynative_self_check.py +131 -116
  93. msprobe/mindspore/free_benchmark/common/config.py +27 -12
  94. msprobe/mindspore/free_benchmark/common/handler_params.py +32 -17
  95. msprobe/mindspore/free_benchmark/common/utils.py +85 -71
  96. msprobe/mindspore/free_benchmark/data/support_wrap_ops.yaml +842 -842
  97. msprobe/mindspore/free_benchmark/decorator/dec_forward.py +57 -42
  98. msprobe/mindspore/free_benchmark/decorator/decorator_factory.py +122 -107
  99. msprobe/mindspore/free_benchmark/handler/base_handler.py +105 -90
  100. msprobe/mindspore/free_benchmark/handler/check_handler.py +56 -41
  101. msprobe/mindspore/free_benchmark/handler/fix_handler.py +51 -36
  102. msprobe/mindspore/free_benchmark/handler/handler_factory.py +36 -21
  103. msprobe/mindspore/free_benchmark/perturbation/add_noise.py +82 -67
  104. msprobe/mindspore/free_benchmark/perturbation/base_perturbation.py +36 -21
  105. msprobe/mindspore/free_benchmark/perturbation/bit_noise.py +78 -63
  106. msprobe/mindspore/free_benchmark/perturbation/exchange_value.py +77 -0
  107. msprobe/mindspore/free_benchmark/perturbation/improve_precision.py +49 -34
  108. msprobe/mindspore/free_benchmark/perturbation/no_change.py +27 -12
  109. msprobe/mindspore/free_benchmark/perturbation/perturbation_factory.py +44 -27
  110. msprobe/mindspore/free_benchmark/self_check_tool_factory.py +48 -33
  111. msprobe/mindspore/grad_probe/global_context.py +100 -91
  112. msprobe/mindspore/grad_probe/grad_analyzer.py +231 -231
  113. msprobe/mindspore/grad_probe/grad_monitor.py +27 -27
  114. msprobe/mindspore/grad_probe/grad_stat_csv.py +131 -131
  115. msprobe/mindspore/grad_probe/hook.py +94 -92
  116. msprobe/mindspore/grad_probe/utils.py +29 -28
  117. msprobe/mindspore/ms_config.py +128 -126
  118. msprobe/mindspore/overflow_check/kernel_graph_overflow_check.py +60 -45
  119. msprobe/mindspore/overflow_check/overflow_check_tool_factory.py +49 -34
  120. msprobe/mindspore/runtime.py +4 -4
  121. msprobe/mindspore/service.py +297 -354
  122. msprobe/mindspore/task_handler_factory.py +24 -24
  123. msprobe/msprobe.py +105 -107
  124. msprobe/pytorch/__init__.py +23 -4
  125. msprobe/pytorch/api_accuracy_checker/common/config.py +70 -55
  126. msprobe/pytorch/api_accuracy_checker/common/utils.py +246 -165
  127. msprobe/pytorch/api_accuracy_checker/compare/algorithm.py +230 -213
  128. msprobe/pytorch/api_accuracy_checker/compare/api_precision_compare.py +632 -581
  129. msprobe/pytorch/api_accuracy_checker/compare/api_precision_standard.yaml +132 -132
  130. msprobe/pytorch/api_accuracy_checker/compare/api_precision_threshold.yaml +390 -390
  131. msprobe/pytorch/api_accuracy_checker/compare/compare.py +416 -381
  132. msprobe/pytorch/api_accuracy_checker/compare/compare_column.py +90 -73
  133. msprobe/pytorch/api_accuracy_checker/compare/compare_utils.py +265 -244
  134. msprobe/pytorch/api_accuracy_checker/config.yaml +10 -10
  135. msprobe/pytorch/api_accuracy_checker/run_ut/data_generate.py +370 -332
  136. msprobe/pytorch/api_accuracy_checker/run_ut/multi_run_ut.py +221 -199
  137. msprobe/pytorch/api_accuracy_checker/run_ut/run_overflow_check.py +150 -134
  138. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut.py +518 -581
  139. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut_utils.py +213 -74
  140. msprobe/pytorch/api_accuracy_checker/run_ut/torch_ut_setting.json +7 -4
  141. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/attl.py +218 -202
  142. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/client.py +370 -324
  143. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/device_dispatch.py +227 -204
  144. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/dump_dispatch.py +110 -0
  145. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/server.py +244 -218
  146. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/torch_ops_config.yaml +63 -0
  147. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/utils.py +44 -0
  148. msprobe/pytorch/bench_functions/__init__.py +30 -15
  149. msprobe/pytorch/bench_functions/apply_adam_w.py +43 -28
  150. msprobe/pytorch/bench_functions/confusion_transpose.py +34 -19
  151. msprobe/pytorch/bench_functions/fast_gelu.py +70 -55
  152. msprobe/pytorch/bench_functions/layer_norm_eval.py +21 -6
  153. msprobe/pytorch/bench_functions/linear.py +27 -12
  154. msprobe/pytorch/bench_functions/matmul_backward.py +63 -48
  155. msprobe/pytorch/bench_functions/npu_fusion_attention.py +538 -421
  156. msprobe/pytorch/bench_functions/rms_norm.py +30 -15
  157. msprobe/pytorch/bench_functions/rotary_mul.py +71 -52
  158. msprobe/pytorch/bench_functions/scaled_mask_softmax.py +41 -26
  159. msprobe/pytorch/bench_functions/swiglu.py +70 -55
  160. msprobe/pytorch/common/__init__.py +17 -2
  161. msprobe/pytorch/common/compare_script.template +14 -14
  162. msprobe/pytorch/common/log.py +33 -32
  163. msprobe/pytorch/common/parse_json.py +54 -39
  164. msprobe/pytorch/common/utils.py +310 -300
  165. msprobe/pytorch/compare/distributed_compare.py +66 -66
  166. msprobe/pytorch/compare/mapping.yaml +607 -607
  167. msprobe/pytorch/compare/match.py +49 -33
  168. msprobe/pytorch/compare/pt_compare.py +82 -40
  169. msprobe/pytorch/debugger/debugger_config.py +108 -95
  170. msprobe/pytorch/debugger/precision_debugger.py +173 -125
  171. msprobe/pytorch/free_benchmark/__init__.py +23 -8
  172. msprobe/pytorch/free_benchmark/common/constant.py +70 -70
  173. msprobe/pytorch/free_benchmark/common/counter.py +71 -71
  174. msprobe/pytorch/free_benchmark/common/enums.py +65 -37
  175. msprobe/pytorch/free_benchmark/common/params.py +144 -129
  176. msprobe/pytorch/free_benchmark/common/utils.py +118 -102
  177. msprobe/pytorch/free_benchmark/compare/grad_saver.py +200 -179
  178. msprobe/pytorch/free_benchmark/compare/single_benchmark.py +119 -104
  179. msprobe/pytorch/free_benchmark/main.py +120 -105
  180. msprobe/pytorch/free_benchmark/perturbed_layers/base_layer.py +28 -13
  181. msprobe/pytorch/free_benchmark/perturbed_layers/layer_factory.py +56 -41
  182. msprobe/pytorch/free_benchmark/perturbed_layers/npu/add_noise.py +105 -90
  183. msprobe/pytorch/free_benchmark/perturbed_layers/npu/bit_noise.py +119 -104
  184. msprobe/pytorch/free_benchmark/perturbed_layers/npu/change_value.py +87 -63
  185. msprobe/pytorch/free_benchmark/perturbed_layers/npu/improve_precision.py +83 -68
  186. msprobe/pytorch/free_benchmark/perturbed_layers/npu/no_change.py +43 -28
  187. msprobe/pytorch/free_benchmark/perturbed_layers/npu/npu_base_layser.py +60 -45
  188. msprobe/pytorch/free_benchmark/perturbed_layers/run_cpu.py +34 -19
  189. msprobe/pytorch/free_benchmark/result_handlers/base_handler.py +256 -217
  190. msprobe/pytorch/free_benchmark/result_handlers/check_handler.py +54 -39
  191. msprobe/pytorch/free_benchmark/result_handlers/fix_handler.py +38 -23
  192. msprobe/pytorch/free_benchmark/result_handlers/handler_factory.py +45 -30
  193. msprobe/pytorch/free_benchmark/result_handlers/preheat_handler.py +185 -170
  194. msprobe/pytorch/function_factory.py +91 -75
  195. msprobe/pytorch/functional/module_dump.py +84 -0
  196. msprobe/pytorch/grad_probe/grad_monitor.py +91 -90
  197. msprobe/pytorch/grad_probe/grad_stat_csv.py +128 -128
  198. msprobe/pytorch/hook_module/__init__.py +16 -1
  199. msprobe/pytorch/hook_module/api_registry.py +166 -161
  200. msprobe/pytorch/hook_module/hook_module.py +118 -120
  201. msprobe/pytorch/hook_module/support_wrap_ops.yaml +1879 -1877
  202. msprobe/pytorch/hook_module/utils.py +28 -29
  203. msprobe/pytorch/hook_module/wrap_aten.py +111 -110
  204. msprobe/pytorch/hook_module/wrap_distributed.py +77 -78
  205. msprobe/pytorch/hook_module/wrap_functional.py +104 -105
  206. msprobe/pytorch/hook_module/wrap_npu_custom.py +85 -84
  207. msprobe/pytorch/hook_module/wrap_tensor.py +69 -71
  208. msprobe/pytorch/hook_module/wrap_torch.py +84 -86
  209. msprobe/pytorch/hook_module/wrap_vf.py +60 -62
  210. msprobe/pytorch/module_processer.py +153 -138
  211. msprobe/pytorch/online_dispatch/__init__.py +20 -20
  212. msprobe/pytorch/online_dispatch/compare.py +235 -236
  213. msprobe/pytorch/online_dispatch/dispatch.py +271 -271
  214. msprobe/pytorch/online_dispatch/dump_compare.py +155 -156
  215. msprobe/pytorch/online_dispatch/single_compare.py +391 -391
  216. msprobe/pytorch/online_dispatch/torch_ops_config.yaml +57 -49
  217. msprobe/pytorch/online_dispatch/utils.py +127 -146
  218. msprobe/pytorch/parse.py +19 -4
  219. msprobe/pytorch/parse_tool/cli.py +31 -32
  220. msprobe/pytorch/parse_tool/lib/compare.py +259 -271
  221. msprobe/pytorch/parse_tool/lib/config.py +52 -52
  222. msprobe/pytorch/parse_tool/lib/file_desc.py +31 -31
  223. msprobe/pytorch/parse_tool/lib/interactive_cli.py +102 -102
  224. msprobe/pytorch/parse_tool/lib/parse_exception.py +54 -54
  225. msprobe/pytorch/parse_tool/lib/parse_tool.py +161 -158
  226. msprobe/pytorch/parse_tool/lib/utils.py +320 -321
  227. msprobe/pytorch/parse_tool/lib/visualization.py +85 -91
  228. msprobe/pytorch/pt_config.py +317 -187
  229. msprobe/pytorch/service.py +311 -252
  230. mindstudio_probe-1.0.3.dist-info/RECORD +0 -272
  231. msprobe/config/README.md +0 -539
  232. msprobe/mindspore/doc/compare.md +0 -58
  233. msprobe/mindspore/doc/dump.md +0 -217
  234. msprobe/mindspore/dump/hook_cell/wrap_functional.py +0 -91
  235. msprobe/mindspore/dump/hook_cell/wrap_tensor.py +0 -63
  236. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/ssl_config.py +0 -10
  237. msprobe/pytorch/doc/FAQ.md +0 -193
  238. msprobe/pytorch/doc/api_accuracy_checker.md +0 -313
  239. msprobe/pytorch/doc/api_accuracy_checker_online.md +0 -187
  240. msprobe/pytorch/doc/dump.md +0 -260
  241. msprobe/pytorch/doc/msprobe/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/206/320/245/342/226/221/321/206/320/235/320/276dump/321/206/320/260/320/227/321/205/320/227/320/226/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -182
  242. msprobe/pytorch/doc/ptdbg_ascend_compare.md +0 -240
  243. msprobe/pytorch/doc/ptdbg_ascend_overview.md +0 -68
  244. msprobe/pytorch/doc/ptdbg_ascend_quickstart.md +0 -381
  245. msprobe/pytorch/doc/run_overflow_check.md +0 -25
  246. msprobe/pytorch/doc//321/206/320/247/320/260/321/206/320/260/320/227/321/206/320/255/320/226/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/205/320/254/342/225/221/321/206/320/251/320/277/321/211/320/272/320/234/321/210/320/277/320/221/321/205/320/242/320/234/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -151
  247. msprobe/pytorch/functional/data_processor.py +0 -0
  248. msprobe/pytorch/functional/dump_module.py +0 -39
  249. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/top_level.txt +0 -0
  250. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_1.png +0 -0
  251. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_2.png +0 -0
  252. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_3.png +0 -0
  253. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_4.png +0 -0
  254. /msprobe/{pytorch/doc → docs}/img/GPT-3_1.png +0 -0
  255. /msprobe/{pytorch/doc → docs}/img/GPT-3_2.png +0 -0
  256. /msprobe/{pytorch/doc → docs}/img/GPT-3_3.png +0 -0
  257. /msprobe/{pytorch/doc → docs}/img/GPT-3_4.png +0 -0
  258. /msprobe/{pytorch/doc → docs}/img/GPT-3_5.png +0 -0
  259. /msprobe/{pytorch/doc → docs}/img/GPT-3_6.png +0 -0
  260. /msprobe/{pytorch/doc → docs}/img/GPT-3_7.png +0 -0
  261. /msprobe/{pytorch/doc → docs}/img/GPT-3_8.png +0 -0
  262. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_1.png +0 -0
  263. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_2.png +0 -0
  264. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_details.png +0 -0
  265. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_result.png +0 -0
  266. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_details.png +0 -0
  267. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_result.png +0 -0
  268. /msprobe/{pytorch/doc → docs}/img/auto_analyze_log.png +0 -0
  269. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl.png +0 -0
  270. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl_md5.png.png +0 -0
  271. /msprobe/{pytorch/doc → docs}/img/cpu_info.png +0 -0
  272. /msprobe/{config → docs}/img/free_benchmark.png +0 -0
  273. /msprobe/{doc/grad_probe/img/image-1.png → docs/img/grad_probe_image-1.png} +0 -0
  274. /msprobe/{doc/grad_probe/img/image-2.png → docs/img/grad_probe_image-2.png} +0 -0
  275. /msprobe/{doc/grad_probe/img/image-3.png → docs/img/grad_probe_image-3.png} +0 -0
  276. /msprobe/{doc/grad_probe/img/image-4.png → docs/img/grad_probe_image-4.png} +0 -0
  277. /msprobe/{doc/grad_probe/img/image.png → docs/img/grad_probe_image.png} +0 -0
  278. /msprobe/{pytorch/doc → docs}/img/module_compare.png +0 -0
@@ -1,149 +1,174 @@
1
-
2
- import multiprocessing
3
- from dataclasses import dataclass
4
- from functools import partial
5
- import numpy as np
6
- import pandas as pd
7
- from msprobe.core.common.log import logger
8
- from msprobe.core.common.utils import CompareException
9
- from msprobe.core.common.const import CompareConst
10
-
11
-
12
- def _handle_multi_process(func, input_parma, result_df, lock):
13
- process_num = int((multiprocessing.cpu_count() + 1) / 2)
14
- op_name_mapping_dict = read_dump_data(result_df)
15
-
16
- df_chunk_size = len(result_df) // process_num
17
- if df_chunk_size > 0:
18
- df_chunks = [result_df.iloc[i:i + df_chunk_size] for i in range(0, len(result_df), df_chunk_size)]
19
- else:
20
- df_chunks = [result_df]
21
-
22
- results = []
23
- pool = multiprocessing.Pool(process_num)
24
-
25
- def err_call(args):
26
- logger.error('multiprocess compare failed! Reason: {}'.format(args))
27
- try:
28
- pool.terminate()
29
- except OSError as e:
30
- logger.error("pool terminate failed")
31
-
32
- for process_idx, df_chunk in enumerate(df_chunks):
33
- idx = df_chunk_size * process_idx
34
- result = pool.apply_async(func,
35
- args=(idx, op_name_mapping_dict, df_chunk, lock, input_parma),
36
- error_callback=err_call)
37
- results.append(result)
38
- final_results = [r.get() for r in results]
39
- pool.close()
40
- pool.join()
41
- return pd.concat(final_results, ignore_index=True)
42
-
43
-
44
- def _ms_graph_handle_multi_process(func, result_df, mode):
45
- process_num = int((multiprocessing.cpu_count() + 1) // 2)
46
- df_chunk_size = len(result_df) // process_num
47
- if df_chunk_size > 0:
48
- df_chunks = [result_df.iloc[i:i + df_chunk_size] for i in range(0, len(result_df), df_chunk_size)]
49
- else:
50
- df_chunks = [result_df]
51
-
52
- results = []
53
- pool = multiprocessing.Pool(process_num)
54
-
55
- def err_call(args):
56
- logger.error('multiprocess compare failed! Reason: {}'.format(args))
57
- try:
58
- pool.terminate()
59
- except OSError as e:
60
- logger.error("pool terminate failed")
61
-
62
- for df_chunk in df_chunks:
63
- result = pool.apply_async(func, args=(df_chunk, mode), error_callback=err_call)
64
- results.append(result)
65
- final_results = [r.get() for r in results]
66
- pool.close()
67
- pool.join()
68
- return pd.concat(final_results, ignore_index=True)
69
-
70
-
71
- def read_dump_data(result_df):
72
- try:
73
- npu_dump_name_list = result_df.iloc[0:, 0].tolist()
74
- npu_dump_tensor_list = result_df.iloc[0:, -1].tolist()
75
- op_name_mapping_dict = {}
76
- for index, _ in enumerate(npu_dump_name_list):
77
- npu_dump_name = npu_dump_name_list[index]
78
- npu_dump_tensor = npu_dump_tensor_list[index]
79
- op_name_mapping_dict[npu_dump_name] = [npu_dump_tensor, npu_dump_tensor]
80
- return op_name_mapping_dict
81
- except ValueError as e:
82
- logger.error('result dataframe is not found.')
83
- raise CompareException(CompareException.INVALID_DATA_ERROR) from e
84
- except IndexError as e:
85
- logger.error('result dataframe elements can not be access.')
86
- raise CompareException(CompareException.INDEX_OUT_OF_BOUNDS_ERROR) from e
87
-
88
- @dataclass
89
- class ComparisonResult:
90
- cos_result: list
91
- max_err_result: list
92
- max_relative_err_result: list
93
- err_msgs: list
94
- one_thousand_err_ratio_result: list
95
- five_thousand_err_ratio_result: list
96
-
97
-
98
- def _save_cmp_result(offset, result: ComparisonResult, result_df, lock):
99
- """
100
- Save comparison results into the result DataFrame with thread safety.
101
- Args:
102
- offset: offset for index
103
- result: data struct of ComparisonResult
104
- result_df: result of DataFrame
105
- lock: thread lock
106
-
107
- Returns:
108
- comparison results in DataFrame
109
- """
110
-
111
- lock.acquire()
112
- try:
113
- for i, _ in enumerate(result.cos_result):
114
- process_index = i + offset
115
- result_df.loc[process_index, CompareConst.COSINE] = result.cos_result[i]
116
- result_df.loc[process_index, CompareConst.MAX_ABS_ERR] = result.max_err_result[i]
117
- result_df.loc[process_index, CompareConst.MAX_RELATIVE_ERR] = result.max_relative_err_result[i]
118
- result_df.loc[process_index, CompareConst.ERROR_MESSAGE] = result.err_msgs[i]
119
- result_df.loc[process_index, CompareConst.ACCURACY] = check_accuracy(result.cos_result[i], result.max_err_result[i])
120
- result_df.loc[process_index, CompareConst.ONE_THOUSANDTH_ERR_RATIO] = result.one_thousand_err_ratio_result[i]
121
- result_df.loc[process_index, CompareConst.FIVE_THOUSANDTHS_ERR_RATIO] = result.five_thousand_err_ratio_result[i]
122
- return result_df
123
- except ValueError as e:
124
- logger.error('result dataframe is not found.')
125
- raise CompareException(CompareException.INVALID_DATA_ERROR) from e
126
- except IndexError as e:
127
- logger.error('result dataframe elements can not be access.')
128
- raise CompareException(CompareException.INDEX_OUT_OF_BOUNDS_ERROR) from e
129
- finally:
130
- lock.release()
131
-
132
-
133
- def check_accuracy(cos, max_abs_err):
134
- if cos == CompareConst.SHAPE_UNMATCH:
135
- return CompareConst.ACCURACY_CHECK_UNMATCH
136
- if cos == CompareConst.NONE or max_abs_err == CompareConst.NONE:
137
- return CompareConst.NONE
138
- if cos == "N/A" or max_abs_err == "N/A":
139
- return CompareConst.ACCURACY_CHECK_NO
140
- try:
141
- cos, max_abs_err = float(cos), float(max_abs_err)
142
- except ValueError:
143
- logger.warning("Cosine or MaxAbsErr can not get float value.")
144
- return CompareConst.NONE
145
- if cos < CompareConst.COS_THRESHOLD and max_abs_err > CompareConst.MAX_ABS_ERR_THRESHOLD:
146
- return CompareConst.ACCURACY_CHECK_NO
147
- if cos < CompareConst.COS_MAX_THRESHOLD or max_abs_err > CompareConst.MAX_ABS_ERR_MAX_THRESHOLD:
148
- return CompareConst.ACCURACY_CHECK_NO
149
- return CompareConst.ACCURACY_CHECK_YES
1
+ # Copyright (c) 2024-2024, Huawei Technologies Co., Ltd.
2
+ # All rights reserved.
3
+ #
4
+ # Licensed under the Apache License, Version 2.0 (the "License");
5
+ # you may not use this file except in compliance with the License.
6
+ # You may obtain a copy of the License at
7
+ #
8
+ # http://www.apache.org/licenses/LICENSE-2.0
9
+ #
10
+ # Unless required by applicable law or agreed to in writing, software
11
+ # distributed under the License is distributed on an "AS IS" BASIS,
12
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
13
+ # See the License for the specific language governing permissions and
14
+ # limitations under the License.
15
+
16
+ import multiprocessing
17
+ from dataclasses import dataclass
18
+ import pandas as pd
19
+ from tqdm import tqdm
20
+ from msprobe.core.common.log import logger
21
+ from msprobe.core.common.utils import CompareException
22
+ from msprobe.core.common.const import CompareConst
23
+
24
+
25
+ def _handle_multi_process(func, input_parma, result_df, lock):
26
+ process_num = int((multiprocessing.cpu_count() + 1) / 2)
27
+ op_name_mapping_dict = read_dump_data(result_df)
28
+
29
+ df_chunk_size = len(result_df) // process_num
30
+ if df_chunk_size > 0:
31
+ df_chunks = [result_df.iloc[i:i + df_chunk_size] for i in range(0, len(result_df), df_chunk_size)]
32
+ else:
33
+ df_chunks = [result_df]
34
+
35
+ results = []
36
+ pool = multiprocessing.Pool(process_num)
37
+
38
+ def err_call(args):
39
+ logger.error('multiprocess compare failed! Reason: {}'.format(args))
40
+ try:
41
+ pool.terminate()
42
+ except OSError as e:
43
+ logger.error("pool terminate failed")
44
+
45
+ progress_bar = tqdm(total=len(result_df), desc="API/Module Item Compare Process", unit="row", ncols=100)
46
+
47
+ def update_progress(size, progress_lock):
48
+ with progress_lock:
49
+ progress_bar.update(size)
50
+
51
+ for process_idx, df_chunk in enumerate(df_chunks):
52
+ idx = df_chunk_size * process_idx
53
+ chunk_size = len(df_chunk)
54
+ result = pool.apply_async(func,
55
+ args=(idx, op_name_mapping_dict, df_chunk, lock, input_parma),
56
+ error_callback=err_call,
57
+ callback=update_progress(chunk_size, lock))
58
+ results.append(result)
59
+ final_results = [r.get() for r in results]
60
+ pool.close()
61
+ pool.join()
62
+ return pd.concat(final_results, ignore_index=True)
63
+
64
+
65
+ def _ms_graph_handle_multi_process(func, result_df, mode):
66
+ process_num = int((multiprocessing.cpu_count() + 1) // 4)
67
+ df_chunk_size = len(result_df) // process_num
68
+ if df_chunk_size > 0:
69
+ df_chunks = [result_df.iloc[i:i + df_chunk_size] for i in range(0, len(result_df), df_chunk_size)]
70
+ else:
71
+ df_chunks = [result_df]
72
+
73
+ results = []
74
+ pool = multiprocessing.Pool(process_num)
75
+
76
+ def err_call(args):
77
+ logger.error('multiprocess compare failed! Reason: {}'.format(args))
78
+ try:
79
+ pool.terminate()
80
+ except OSError as e:
81
+ logger.error("pool terminate failed")
82
+
83
+ for df_chunk in df_chunks:
84
+ result = pool.apply_async(func, args=(df_chunk, mode), error_callback=err_call)
85
+ results.append(result)
86
+ final_results = [r.get() for r in results]
87
+ pool.close()
88
+ pool.join()
89
+ return pd.concat(final_results, ignore_index=True)
90
+
91
+
92
+ def read_dump_data(result_df):
93
+ try:
94
+ npu_dump_name_list = result_df.iloc[0:, 0].tolist()
95
+ npu_dump_tensor_list = result_df.iloc[0:, -1].tolist()
96
+ op_name_mapping_dict = {}
97
+ for index, _ in enumerate(npu_dump_name_list):
98
+ npu_dump_name = npu_dump_name_list[index]
99
+ npu_dump_tensor = npu_dump_tensor_list[index]
100
+ op_name_mapping_dict[npu_dump_name] = [npu_dump_tensor, npu_dump_tensor]
101
+ return op_name_mapping_dict
102
+ except ValueError as e:
103
+ logger.error('result dataframe is not found.')
104
+ raise CompareException(CompareException.INVALID_DATA_ERROR) from e
105
+ except IndexError as e:
106
+ logger.error('result dataframe elements can not be access.')
107
+ raise CompareException(CompareException.INDEX_OUT_OF_BOUNDS_ERROR) from e
108
+
109
+
110
+ @dataclass
111
+ class ComparisonResult:
112
+ cos_result: list
113
+ max_err_result: list
114
+ max_relative_err_result: list
115
+ err_msgs: list
116
+ one_thousand_err_ratio_result: list
117
+ five_thousand_err_ratio_result: list
118
+
119
+
120
+ def _save_cmp_result(offset, result: ComparisonResult, result_df, lock):
121
+ """
122
+ Save comparison results into the result DataFrame with thread safety.
123
+ Args:
124
+ offset: offset for index
125
+ result: data struct of ComparisonResult
126
+ result_df: result of DataFrame
127
+ lock: thread lock
128
+
129
+ Returns:
130
+ comparison results in DataFrame
131
+ """
132
+
133
+ lock.acquire()
134
+ try:
135
+ for i, _ in enumerate(result.cos_result):
136
+ process_index = i + offset
137
+ result_df.loc[process_index, CompareConst.COSINE] = result.cos_result[i]
138
+ result_df.loc[process_index, CompareConst.MAX_ABS_ERR] = result.max_err_result[i]
139
+ result_df.loc[process_index, CompareConst.MAX_RELATIVE_ERR] = result.max_relative_err_result[i]
140
+ result_df.loc[process_index, CompareConst.ERROR_MESSAGE] = result.err_msgs[i]
141
+ result_df.loc[process_index, CompareConst.ACCURACY] = (
142
+ check_accuracy(result.cos_result[i], result.max_err_result[i]))
143
+ result_df.loc[process_index, CompareConst.ONE_THOUSANDTH_ERR_RATIO] = (
144
+ result.one_thousand_err_ratio_result)[i]
145
+ result_df.loc[process_index, CompareConst.FIVE_THOUSANDTHS_ERR_RATIO] = (
146
+ result.five_thousand_err_ratio_result)[i]
147
+ return result_df
148
+ except ValueError as e:
149
+ logger.error('result dataframe is not found.')
150
+ raise CompareException(CompareException.INVALID_DATA_ERROR) from e
151
+ except IndexError as e:
152
+ logger.error('result dataframe elements can not be access.')
153
+ raise CompareException(CompareException.INDEX_OUT_OF_BOUNDS_ERROR) from e
154
+ finally:
155
+ lock.release()
156
+
157
+
158
+ def check_accuracy(cos, max_abs_err):
159
+ if cos == CompareConst.SHAPE_UNMATCH:
160
+ return CompareConst.ACCURACY_CHECK_UNMATCH
161
+ if cos == CompareConst.NONE or max_abs_err == CompareConst.NONE:
162
+ return CompareConst.NONE
163
+ if cos == "N/A" or max_abs_err == "N/A":
164
+ return CompareConst.ACCURACY_CHECK_NO
165
+ try:
166
+ cos, max_abs_err = float(cos), float(max_abs_err)
167
+ except ValueError:
168
+ logger.warning("Cosine or MaxAbsErr can not get float value.")
169
+ return CompareConst.NONE
170
+ if cos < CompareConst.COS_THRESHOLD and max_abs_err > CompareConst.MAX_ABS_ERR_THRESHOLD:
171
+ return CompareConst.ACCURACY_CHECK_NO
172
+ if cos < CompareConst.COS_MAX_THRESHOLD or max_abs_err > CompareConst.MAX_ABS_ERR_MAX_THRESHOLD:
173
+ return CompareConst.ACCURACY_CHECK_NO
174
+ return CompareConst.ACCURACY_CHECK_YES