crawlo 1.1.8__py3-none-any.whl → 1.2.0__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.

Potentially problematic release.


This version of crawlo might be problematic. Click here for more details.

Files changed (191) hide show
  1. crawlo/__init__.py +61 -61
  2. crawlo/__version__.py +1 -1
  3. crawlo/cleaners/__init__.py +60 -60
  4. crawlo/cleaners/data_formatter.py +225 -225
  5. crawlo/cleaners/encoding_converter.py +125 -125
  6. crawlo/cleaners/text_cleaner.py +232 -232
  7. crawlo/cli.py +65 -65
  8. crawlo/commands/__init__.py +14 -14
  9. crawlo/commands/check.py +594 -594
  10. crawlo/commands/genspider.py +151 -151
  11. crawlo/commands/help.py +132 -132
  12. crawlo/commands/list.py +155 -155
  13. crawlo/commands/run.py +292 -292
  14. crawlo/commands/startproject.py +418 -418
  15. crawlo/commands/stats.py +188 -188
  16. crawlo/commands/utils.py +186 -186
  17. crawlo/config.py +312 -312
  18. crawlo/config_validator.py +252 -252
  19. crawlo/core/__init__.py +2 -2
  20. crawlo/core/engine.py +354 -345
  21. crawlo/core/processor.py +40 -40
  22. crawlo/core/scheduler.py +143 -136
  23. crawlo/crawler.py +1027 -1027
  24. crawlo/downloader/__init__.py +266 -266
  25. crawlo/downloader/aiohttp_downloader.py +220 -220
  26. crawlo/downloader/cffi_downloader.py +256 -256
  27. crawlo/downloader/httpx_downloader.py +259 -259
  28. crawlo/downloader/hybrid_downloader.py +213 -213
  29. crawlo/downloader/playwright_downloader.py +402 -402
  30. crawlo/downloader/selenium_downloader.py +472 -472
  31. crawlo/event.py +11 -11
  32. crawlo/exceptions.py +81 -81
  33. crawlo/extension/__init__.py +37 -37
  34. crawlo/extension/health_check.py +141 -141
  35. crawlo/extension/log_interval.py +57 -57
  36. crawlo/extension/log_stats.py +81 -81
  37. crawlo/extension/logging_extension.py +43 -43
  38. crawlo/extension/memory_monitor.py +104 -104
  39. crawlo/extension/performance_profiler.py +133 -133
  40. crawlo/extension/request_recorder.py +107 -107
  41. crawlo/filters/__init__.py +154 -154
  42. crawlo/filters/aioredis_filter.py +280 -280
  43. crawlo/filters/memory_filter.py +269 -269
  44. crawlo/items/__init__.py +23 -23
  45. crawlo/items/base.py +21 -21
  46. crawlo/items/fields.py +53 -53
  47. crawlo/items/items.py +104 -104
  48. crawlo/middleware/__init__.py +21 -21
  49. crawlo/middleware/default_header.py +32 -32
  50. crawlo/middleware/download_delay.py +28 -28
  51. crawlo/middleware/middleware_manager.py +135 -135
  52. crawlo/middleware/proxy.py +272 -272
  53. crawlo/middleware/request_ignore.py +30 -30
  54. crawlo/middleware/response_code.py +18 -18
  55. crawlo/middleware/response_filter.py +26 -26
  56. crawlo/middleware/retry.py +124 -124
  57. crawlo/mode_manager.py +211 -211
  58. crawlo/network/__init__.py +21 -21
  59. crawlo/network/request.py +338 -338
  60. crawlo/network/response.py +359 -359
  61. crawlo/pipelines/__init__.py +21 -21
  62. crawlo/pipelines/bloom_dedup_pipeline.py +156 -156
  63. crawlo/pipelines/console_pipeline.py +39 -39
  64. crawlo/pipelines/csv_pipeline.py +316 -316
  65. crawlo/pipelines/database_dedup_pipeline.py +224 -224
  66. crawlo/pipelines/json_pipeline.py +218 -218
  67. crawlo/pipelines/memory_dedup_pipeline.py +115 -115
  68. crawlo/pipelines/mongo_pipeline.py +131 -131
  69. crawlo/pipelines/mysql_pipeline.py +316 -316
  70. crawlo/pipelines/pipeline_manager.py +61 -61
  71. crawlo/pipelines/redis_dedup_pipeline.py +167 -167
  72. crawlo/project.py +187 -187
  73. crawlo/queue/pqueue.py +37 -37
  74. crawlo/queue/queue_manager.py +337 -334
  75. crawlo/queue/redis_priority_queue.py +298 -298
  76. crawlo/settings/__init__.py +7 -7
  77. crawlo/settings/default_settings.py +219 -219
  78. crawlo/settings/setting_manager.py +122 -122
  79. crawlo/spider/__init__.py +639 -639
  80. crawlo/stats_collector.py +59 -59
  81. crawlo/subscriber.py +130 -130
  82. crawlo/task_manager.py +30 -30
  83. crawlo/templates/crawlo.cfg.tmpl +10 -10
  84. crawlo/templates/project/__init__.py.tmpl +3 -3
  85. crawlo/templates/project/items.py.tmpl +17 -17
  86. crawlo/templates/project/middlewares.py.tmpl +109 -109
  87. crawlo/templates/project/pipelines.py.tmpl +96 -96
  88. crawlo/templates/project/run.py.tmpl +45 -45
  89. crawlo/templates/project/settings.py.tmpl +326 -326
  90. crawlo/templates/project/settings_distributed.py.tmpl +119 -119
  91. crawlo/templates/project/settings_gentle.py.tmpl +94 -94
  92. crawlo/templates/project/settings_high_performance.py.tmpl +151 -151
  93. crawlo/templates/project/settings_simple.py.tmpl +68 -68
  94. crawlo/templates/project/spiders/__init__.py.tmpl +5 -5
  95. crawlo/templates/spider/spider.py.tmpl +141 -141
  96. crawlo/tools/__init__.py +182 -182
  97. crawlo/tools/anti_crawler.py +268 -268
  98. crawlo/tools/authenticated_proxy.py +240 -240
  99. crawlo/tools/data_validator.py +180 -180
  100. crawlo/tools/date_tools.py +35 -35
  101. crawlo/tools/distributed_coordinator.py +386 -386
  102. crawlo/tools/retry_mechanism.py +220 -220
  103. crawlo/tools/scenario_adapter.py +262 -262
  104. crawlo/utils/__init__.py +35 -35
  105. crawlo/utils/batch_processor.py +260 -260
  106. crawlo/utils/controlled_spider_mixin.py +439 -439
  107. crawlo/utils/date_tools.py +290 -290
  108. crawlo/utils/db_helper.py +343 -343
  109. crawlo/utils/enhanced_error_handler.py +359 -359
  110. crawlo/utils/env_config.py +105 -105
  111. crawlo/utils/error_handler.py +125 -125
  112. crawlo/utils/func_tools.py +82 -82
  113. crawlo/utils/large_scale_config.py +286 -286
  114. crawlo/utils/large_scale_helper.py +343 -343
  115. crawlo/utils/log.py +128 -128
  116. crawlo/utils/performance_monitor.py +284 -284
  117. crawlo/utils/queue_helper.py +175 -175
  118. crawlo/utils/redis_connection_pool.py +334 -334
  119. crawlo/utils/redis_key_validator.py +199 -199
  120. crawlo/utils/request.py +267 -267
  121. crawlo/utils/request_serializer.py +219 -219
  122. crawlo/utils/spider_loader.py +62 -62
  123. crawlo/utils/system.py +11 -11
  124. crawlo/utils/tools.py +4 -4
  125. crawlo/utils/url.py +39 -39
  126. crawlo-1.2.0.dist-info/METADATA +697 -0
  127. crawlo-1.2.0.dist-info/RECORD +190 -0
  128. examples/__init__.py +7 -7
  129. tests/DOUBLE_CRAWLO_PREFIX_FIX_REPORT.md +81 -81
  130. tests/__init__.py +7 -7
  131. tests/advanced_tools_example.py +275 -275
  132. tests/authenticated_proxy_example.py +236 -236
  133. tests/cleaners_example.py +160 -160
  134. tests/config_validation_demo.py +102 -102
  135. tests/controlled_spider_example.py +205 -205
  136. tests/date_tools_example.py +180 -180
  137. tests/dynamic_loading_example.py +523 -523
  138. tests/dynamic_loading_test.py +104 -104
  139. tests/env_config_example.py +133 -133
  140. tests/error_handling_example.py +171 -171
  141. tests/redis_key_validation_demo.py +130 -130
  142. tests/response_improvements_example.py +144 -144
  143. tests/test_advanced_tools.py +148 -148
  144. tests/test_all_redis_key_configs.py +145 -145
  145. tests/test_authenticated_proxy.py +141 -141
  146. tests/test_cleaners.py +54 -54
  147. tests/test_comprehensive.py +146 -146
  148. tests/test_config_validator.py +193 -193
  149. tests/test_date_tools.py +123 -123
  150. tests/test_double_crawlo_fix.py +207 -207
  151. tests/test_double_crawlo_fix_simple.py +124 -124
  152. tests/test_dynamic_downloaders_proxy.py +124 -124
  153. tests/test_dynamic_proxy.py +92 -92
  154. tests/test_dynamic_proxy_config.py +146 -146
  155. tests/test_dynamic_proxy_real.py +109 -109
  156. tests/test_edge_cases.py +303 -303
  157. tests/test_enhanced_error_handler.py +270 -270
  158. tests/test_env_config.py +121 -121
  159. tests/test_error_handler_compatibility.py +112 -112
  160. tests/test_final_validation.py +153 -153
  161. tests/test_framework_env_usage.py +103 -103
  162. tests/test_integration.py +356 -356
  163. tests/test_item_dedup_redis_key.py +122 -122
  164. tests/test_parsel.py +29 -29
  165. tests/test_performance.py +327 -327
  166. tests/test_proxy_health_check.py +32 -32
  167. tests/test_proxy_middleware_integration.py +136 -136
  168. tests/test_proxy_providers.py +56 -56
  169. tests/test_proxy_stats.py +19 -19
  170. tests/test_proxy_strategies.py +59 -59
  171. tests/test_queue_manager_double_crawlo.py +174 -231
  172. tests/test_queue_manager_redis_key.py +176 -176
  173. tests/test_redis_config.py +28 -28
  174. tests/test_redis_connection_pool.py +294 -294
  175. tests/test_redis_key_naming.py +181 -181
  176. tests/test_redis_key_validator.py +123 -123
  177. tests/test_redis_queue.py +224 -224
  178. tests/test_request_serialization.py +70 -70
  179. tests/test_response_improvements.py +152 -152
  180. tests/test_scheduler.py +241 -241
  181. tests/test_simple_response.py +61 -61
  182. tests/test_telecom_spider_redis_key.py +205 -205
  183. tests/test_template_content.py +87 -87
  184. tests/test_template_redis_key.py +134 -134
  185. tests/test_tools.py +153 -153
  186. tests/tools_example.py +257 -257
  187. crawlo-1.1.8.dist-info/METADATA +0 -626
  188. crawlo-1.1.8.dist-info/RECORD +0 -190
  189. {crawlo-1.1.8.dist-info → crawlo-1.2.0.dist-info}/WHEEL +0 -0
  190. {crawlo-1.1.8.dist-info → crawlo-1.2.0.dist-info}/entry_points.txt +0 -0
  191. {crawlo-1.1.8.dist-info → crawlo-1.2.0.dist-info}/top_level.txt +0 -0
@@ -1,290 +1,290 @@
1
- #!/usr/bin/python
2
- # -*- coding: UTF-8 -*-
3
- """
4
- # @Time : 2025-05-17 10:20
5
- # @Author : crawl-coder
6
- # @Desc : 智能时间工具库(专为爬虫场景设计)
7
- """
8
- import dateparser
9
- from typing import Optional, Union, Literal
10
- from datetime import datetime, timedelta, timezone
11
- from dateutil.relativedelta import relativedelta
12
- import pytz
13
- from pytz import timezone as pytz_timezone
14
-
15
- # 支持的单位类型
16
- TimeUnit = Literal["seconds", "minutes", "hours", "days"]
17
- # 时间输入类型
18
- TimeType = Union[str, datetime]
19
- # 时区类型
20
- TimezoneType = Union[str, timezone, pytz_timezone]
21
-
22
- # 常见时间格式列表(作为 dateparser 的后备方案)
23
- COMMON_FORMATS = [
24
- "%Y-%m-%d %H:%M:%S",
25
- "%Y/%m/%d %H:%M:%S",
26
- "%d-%m-%Y %H:%M:%S",
27
- "%d/%m/%Y %H:%M:%S",
28
- "%Y-%m-%d",
29
- "%Y/%m/%d",
30
- "%d-%m-%Y",
31
- "%d/%m/%Y",
32
- "%b %d, %Y",
33
- "%B %d, %Y",
34
- "%Y年%m月%d日",
35
- "%Y年%m月%d日 %H时%M分%S秒",
36
- "%a %b %d %H:%M:%S %Y",
37
- "%a, %d %b %Y %H:%M:%S",
38
- "%Y-%m-%dT%H:%M:%S.%f",
39
- "%Y-%m-%dT%H:%M:%S",
40
- ]
41
-
42
-
43
- class TimeUtils:
44
- """
45
- 时间处理工具类,提供日期解析、格式化、计算等一站式服务。
46
- 特别适用于爬虫中处理多语言、多格式、相对时间的场景。
47
- """
48
-
49
- @staticmethod
50
- def _try_strptime(time_str: str) -> Optional[datetime]:
51
- """尝试使用预定义格式解析,作为 dateparser 的后备"""
52
- for fmt in COMMON_FORMATS:
53
- try:
54
- return datetime.strptime(time_str.strip(), fmt)
55
- except ValueError:
56
- continue
57
- return None
58
-
59
- @classmethod
60
- def parse(cls, time_input: TimeType, *, default: Optional[datetime] = None) -> Optional[datetime]:
61
- """
62
- 智能解析时间输入(字符串或 datetime)。
63
-
64
- :param time_input: 时间字符串(支持各种语言、格式、相对时间)或 datetime 对象
65
- :param default: 解析失败时返回的默认值
66
- :return: 解析成功返回 datetime,失败返回 default
67
- """
68
- if isinstance(time_input, datetime):
69
- return time_input
70
-
71
- if not isinstance(time_input, str) or not time_input.strip():
72
- return default
73
-
74
- # 1. 优先使用 dateparser(支持多语言和相对时间)
75
- try:
76
- parsed = dateparser.parse(time_input.strip())
77
- if parsed:
78
- return parsed
79
- except Exception:
80
- pass # 忽略异常,尝试后备方案
81
-
82
- # 2. 尝试使用常见格式解析
83
- try:
84
- parsed = cls._try_strptime(time_input)
85
- if parsed:
86
- return parsed
87
- except Exception:
88
- pass
89
-
90
- return default
91
-
92
- @classmethod
93
- def format(cls, dt: TimeType, fmt: str = "%Y-%m-%d %H:%M:%S") -> Optional[str]:
94
- """
95
- 格式化时间。
96
-
97
- :param dt: datetime 对象或可解析的字符串
98
- :param fmt: 输出格式
99
- :return: 格式化后的字符串,失败返回 None
100
- """
101
- if isinstance(dt, str):
102
- dt = cls.parse(dt)
103
- if dt is None:
104
- return None
105
- try:
106
- return dt.strftime(fmt)
107
- except Exception:
108
- return None
109
-
110
- @classmethod
111
- def to_timestamp(cls, dt: TimeType) -> Optional[float]:
112
- """转换为时间戳(秒级)"""
113
- if isinstance(dt, str):
114
- dt = cls.parse(dt)
115
- if dt is None:
116
- return None
117
- try:
118
- return dt.timestamp()
119
- except Exception:
120
- return None
121
-
122
- @classmethod
123
- def from_timestamp(cls, ts: float) -> Optional[datetime]:
124
- """从时间戳创建 datetime"""
125
- try:
126
- return datetime.fromtimestamp(ts)
127
- except Exception:
128
- return None
129
-
130
- @classmethod
131
- def diff(cls, start: TimeType, end: TimeType, unit: TimeUnit = "seconds") -> Optional[int]:
132
- """
133
- 计算两个时间的差值(自动解析字符串)。
134
-
135
- :param start: 起始时间
136
- :param end: 结束时间
137
- :param unit: 单位 ('seconds', 'minutes', 'hours', 'days')
138
- :return: 差值(绝对值),失败返回 None
139
- """
140
- start_dt = cls.parse(start)
141
- end_dt = cls.parse(end)
142
- if not start_dt or not end_dt:
143
- return None
144
-
145
- delta = abs((end_dt - start_dt).total_seconds())
146
-
147
- unit_map = {
148
- "seconds": 1,
149
- "minutes": 60,
150
- "hours": 3600,
151
- "days": 86400,
152
- }
153
- return int(delta // unit_map.get(unit, 1))
154
-
155
- @classmethod
156
- def add_days(cls, dt: TimeType, days: int) -> Optional[datetime]:
157
- """日期加减(天)"""
158
- dt = cls.parse(dt)
159
- if dt is None:
160
- return None
161
- return dt + timedelta(days=days)
162
-
163
- @classmethod
164
- def add_months(cls, dt: TimeType, months: int) -> Optional[datetime]:
165
- """日期加减(月)"""
166
- dt = cls.parse(dt)
167
- if dt is None:
168
- return None
169
- return dt + relativedelta(months=months)
170
-
171
- @classmethod
172
- def days_between(cls, dt1: TimeType, dt2: TimeType) -> Optional[int]:
173
- """计算两个日期之间的天数差"""
174
- return cls.diff(dt1, dt2, "days")
175
-
176
- @classmethod
177
- def is_leap_year(cls, year: int) -> bool:
178
- """判断是否是闰年"""
179
- return (year % 4 == 0 and year % 100 != 0) or (year % 400 == 0)
180
-
181
- @classmethod
182
- def now(cls, fmt: Optional[str] = None) -> Union[datetime, str]:
183
- """
184
- 获取当前时间。
185
-
186
- :param fmt: 如果提供,则返回格式化字符串;否则返回 datetime 对象。
187
- :return: datetime 或 str
188
- """
189
- dt = datetime.now()
190
- if fmt is not None:
191
- return dt.strftime(fmt)
192
- return dt
193
-
194
- @classmethod
195
- def iso_format(cls, dt: TimeType) -> Optional[str]:
196
- """返回 ISO 8601 格式字符串"""
197
- dt = cls.parse(dt)
198
- if dt is None:
199
- return None
200
- return dt.isoformat()
201
-
202
- @classmethod
203
- def to_timezone(cls, dt: TimeType, tz: TimezoneType) -> Optional[datetime]:
204
- """将时间转换为指定时区"""
205
- dt = cls.parse(dt)
206
- if dt is None:
207
- return None
208
-
209
- try:
210
- if isinstance(tz, str):
211
- tz = pytz_timezone(tz)
212
- return dt.astimezone(tz)
213
- except Exception:
214
- return None
215
-
216
- @classmethod
217
- def to_utc(cls, dt: TimeType) -> Optional[datetime]:
218
- """将时间转换为 UTC 时区"""
219
- return cls.to_timezone(dt, pytz.UTC)
220
-
221
- @classmethod
222
- def to_local(cls, dt: TimeType) -> Optional[datetime]:
223
- """将时间转换为本地时区"""
224
- return cls.to_timezone(dt, pytz.timezone("Asia/Shanghai"))
225
-
226
- @classmethod
227
- def from_timestamp_with_tz(cls, ts: float, tz: TimezoneType = None) -> Optional[datetime]:
228
- """从时间戳创建 datetime,并可选择指定时区"""
229
- try:
230
- dt = datetime.fromtimestamp(ts)
231
- if tz:
232
- if isinstance(tz, str):
233
- tz = pytz_timezone(tz)
234
- dt = dt.replace(tzinfo=tz)
235
- return dt
236
- except Exception:
237
- return None
238
-
239
-
240
- # =======================对外接口=======================
241
-
242
- def parse_time(time_input: TimeType, default: Optional[datetime] = None) -> Optional[datetime]:
243
- """解析时间字符串或对象"""
244
- return TimeUtils.parse(time_input, default=default)
245
-
246
-
247
- def format_time(dt: TimeType, fmt: str = "%Y-%m-%d %H:%M:%S") -> Optional[str]:
248
- """格式化时间"""
249
- return TimeUtils.format(dt, fmt)
250
-
251
-
252
- def time_diff(start: TimeType, end: TimeType, unit: TimeUnit = "seconds") -> Optional[int]:
253
- """计算时间差"""
254
- return TimeUtils.diff(start, end, unit)
255
-
256
-
257
- def to_timestamp(dt: TimeType) -> Optional[float]:
258
- """转时间戳"""
259
- return TimeUtils.to_timestamp(dt)
260
-
261
-
262
- def to_datetime(ts: float) -> Optional[datetime]:
263
- """从时间戳转 datetime"""
264
- return TimeUtils.from_timestamp(ts)
265
-
266
-
267
- def now(fmt: Optional[str] = None) -> Union[datetime, str]:
268
- """获取当前时间"""
269
- return TimeUtils.now(fmt)
270
-
271
- def to_timezone(dt: TimeType, tz: TimezoneType) -> Optional[datetime]:
272
- """将时间转换为指定时区"""
273
- return TimeUtils.to_timezone(dt, tz)
274
-
275
- def to_utc(dt: TimeType) -> Optional[datetime]:
276
- """将时间转换为 UTC 时区"""
277
- return TimeUtils.to_utc(dt)
278
-
279
- def to_local(dt: TimeType) -> Optional[datetime]:
280
- """将时间转换为本地时区"""
281
- return TimeUtils.to_local(dt)
282
-
283
- def from_timestamp_with_tz(ts: float, tz: TimezoneType = None) -> Optional[datetime]:
284
- """从时间戳创建 datetime,并可选择指定时区"""
285
- return TimeUtils.from_timestamp_with_tz(ts, tz)
286
-
287
-
288
- if __name__ == '__main__':
289
- get_current_time = now(fmt="%Y-%m-%d %H:%M:%S")
290
- print(get_current_time)
1
+ #!/usr/bin/python
2
+ # -*- coding: UTF-8 -*-
3
+ """
4
+ # @Time : 2025-05-17 10:20
5
+ # @Author : crawl-coder
6
+ # @Desc : 智能时间工具库(专为爬虫场景设计)
7
+ """
8
+ import dateparser
9
+ from typing import Optional, Union, Literal
10
+ from datetime import datetime, timedelta, timezone
11
+ from dateutil.relativedelta import relativedelta
12
+ import pytz
13
+ from pytz import timezone as pytz_timezone
14
+
15
+ # 支持的单位类型
16
+ TimeUnit = Literal["seconds", "minutes", "hours", "days"]
17
+ # 时间输入类型
18
+ TimeType = Union[str, datetime]
19
+ # 时区类型
20
+ TimezoneType = Union[str, timezone, pytz_timezone]
21
+
22
+ # 常见时间格式列表(作为 dateparser 的后备方案)
23
+ COMMON_FORMATS = [
24
+ "%Y-%m-%d %H:%M:%S",
25
+ "%Y/%m/%d %H:%M:%S",
26
+ "%d-%m-%Y %H:%M:%S",
27
+ "%d/%m/%Y %H:%M:%S",
28
+ "%Y-%m-%d",
29
+ "%Y/%m/%d",
30
+ "%d-%m-%Y",
31
+ "%d/%m/%Y",
32
+ "%b %d, %Y",
33
+ "%B %d, %Y",
34
+ "%Y年%m月%d日",
35
+ "%Y年%m月%d日 %H时%M分%S秒",
36
+ "%a %b %d %H:%M:%S %Y",
37
+ "%a, %d %b %Y %H:%M:%S",
38
+ "%Y-%m-%dT%H:%M:%S.%f",
39
+ "%Y-%m-%dT%H:%M:%S",
40
+ ]
41
+
42
+
43
+ class TimeUtils:
44
+ """
45
+ 时间处理工具类,提供日期解析、格式化、计算等一站式服务。
46
+ 特别适用于爬虫中处理多语言、多格式、相对时间的场景。
47
+ """
48
+
49
+ @staticmethod
50
+ def _try_strptime(time_str: str) -> Optional[datetime]:
51
+ """尝试使用预定义格式解析,作为 dateparser 的后备"""
52
+ for fmt in COMMON_FORMATS:
53
+ try:
54
+ return datetime.strptime(time_str.strip(), fmt)
55
+ except ValueError:
56
+ continue
57
+ return None
58
+
59
+ @classmethod
60
+ def parse(cls, time_input: TimeType, *, default: Optional[datetime] = None) -> Optional[datetime]:
61
+ """
62
+ 智能解析时间输入(字符串或 datetime)。
63
+
64
+ :param time_input: 时间字符串(支持各种语言、格式、相对时间)或 datetime 对象
65
+ :param default: 解析失败时返回的默认值
66
+ :return: 解析成功返回 datetime,失败返回 default
67
+ """
68
+ if isinstance(time_input, datetime):
69
+ return time_input
70
+
71
+ if not isinstance(time_input, str) or not time_input.strip():
72
+ return default
73
+
74
+ # 1. 优先使用 dateparser(支持多语言和相对时间)
75
+ try:
76
+ parsed = dateparser.parse(time_input.strip())
77
+ if parsed:
78
+ return parsed
79
+ except Exception:
80
+ pass # 忽略异常,尝试后备方案
81
+
82
+ # 2. 尝试使用常见格式解析
83
+ try:
84
+ parsed = cls._try_strptime(time_input)
85
+ if parsed:
86
+ return parsed
87
+ except Exception:
88
+ pass
89
+
90
+ return default
91
+
92
+ @classmethod
93
+ def format(cls, dt: TimeType, fmt: str = "%Y-%m-%d %H:%M:%S") -> Optional[str]:
94
+ """
95
+ 格式化时间。
96
+
97
+ :param dt: datetime 对象或可解析的字符串
98
+ :param fmt: 输出格式
99
+ :return: 格式化后的字符串,失败返回 None
100
+ """
101
+ if isinstance(dt, str):
102
+ dt = cls.parse(dt)
103
+ if dt is None:
104
+ return None
105
+ try:
106
+ return dt.strftime(fmt)
107
+ except Exception:
108
+ return None
109
+
110
+ @classmethod
111
+ def to_timestamp(cls, dt: TimeType) -> Optional[float]:
112
+ """转换为时间戳(秒级)"""
113
+ if isinstance(dt, str):
114
+ dt = cls.parse(dt)
115
+ if dt is None:
116
+ return None
117
+ try:
118
+ return dt.timestamp()
119
+ except Exception:
120
+ return None
121
+
122
+ @classmethod
123
+ def from_timestamp(cls, ts: float) -> Optional[datetime]:
124
+ """从时间戳创建 datetime"""
125
+ try:
126
+ return datetime.fromtimestamp(ts)
127
+ except Exception:
128
+ return None
129
+
130
+ @classmethod
131
+ def diff(cls, start: TimeType, end: TimeType, unit: TimeUnit = "seconds") -> Optional[int]:
132
+ """
133
+ 计算两个时间的差值(自动解析字符串)。
134
+
135
+ :param start: 起始时间
136
+ :param end: 结束时间
137
+ :param unit: 单位 ('seconds', 'minutes', 'hours', 'days')
138
+ :return: 差值(绝对值),失败返回 None
139
+ """
140
+ start_dt = cls.parse(start)
141
+ end_dt = cls.parse(end)
142
+ if not start_dt or not end_dt:
143
+ return None
144
+
145
+ delta = abs((end_dt - start_dt).total_seconds())
146
+
147
+ unit_map = {
148
+ "seconds": 1,
149
+ "minutes": 60,
150
+ "hours": 3600,
151
+ "days": 86400,
152
+ }
153
+ return int(delta // unit_map.get(unit, 1))
154
+
155
+ @classmethod
156
+ def add_days(cls, dt: TimeType, days: int) -> Optional[datetime]:
157
+ """日期加减(天)"""
158
+ dt = cls.parse(dt)
159
+ if dt is None:
160
+ return None
161
+ return dt + timedelta(days=days)
162
+
163
+ @classmethod
164
+ def add_months(cls, dt: TimeType, months: int) -> Optional[datetime]:
165
+ """日期加减(月)"""
166
+ dt = cls.parse(dt)
167
+ if dt is None:
168
+ return None
169
+ return dt + relativedelta(months=months)
170
+
171
+ @classmethod
172
+ def days_between(cls, dt1: TimeType, dt2: TimeType) -> Optional[int]:
173
+ """计算两个日期之间的天数差"""
174
+ return cls.diff(dt1, dt2, "days")
175
+
176
+ @classmethod
177
+ def is_leap_year(cls, year: int) -> bool:
178
+ """判断是否是闰年"""
179
+ return (year % 4 == 0 and year % 100 != 0) or (year % 400 == 0)
180
+
181
+ @classmethod
182
+ def now(cls, fmt: Optional[str] = None) -> Union[datetime, str]:
183
+ """
184
+ 获取当前时间。
185
+
186
+ :param fmt: 如果提供,则返回格式化字符串;否则返回 datetime 对象。
187
+ :return: datetime 或 str
188
+ """
189
+ dt = datetime.now()
190
+ if fmt is not None:
191
+ return dt.strftime(fmt)
192
+ return dt
193
+
194
+ @classmethod
195
+ def iso_format(cls, dt: TimeType) -> Optional[str]:
196
+ """返回 ISO 8601 格式字符串"""
197
+ dt = cls.parse(dt)
198
+ if dt is None:
199
+ return None
200
+ return dt.isoformat()
201
+
202
+ @classmethod
203
+ def to_timezone(cls, dt: TimeType, tz: TimezoneType) -> Optional[datetime]:
204
+ """将时间转换为指定时区"""
205
+ dt = cls.parse(dt)
206
+ if dt is None:
207
+ return None
208
+
209
+ try:
210
+ if isinstance(tz, str):
211
+ tz = pytz_timezone(tz)
212
+ return dt.astimezone(tz)
213
+ except Exception:
214
+ return None
215
+
216
+ @classmethod
217
+ def to_utc(cls, dt: TimeType) -> Optional[datetime]:
218
+ """将时间转换为 UTC 时区"""
219
+ return cls.to_timezone(dt, pytz.UTC)
220
+
221
+ @classmethod
222
+ def to_local(cls, dt: TimeType) -> Optional[datetime]:
223
+ """将时间转换为本地时区"""
224
+ return cls.to_timezone(dt, pytz.timezone("Asia/Shanghai"))
225
+
226
+ @classmethod
227
+ def from_timestamp_with_tz(cls, ts: float, tz: TimezoneType = None) -> Optional[datetime]:
228
+ """从时间戳创建 datetime,并可选择指定时区"""
229
+ try:
230
+ dt = datetime.fromtimestamp(ts)
231
+ if tz:
232
+ if isinstance(tz, str):
233
+ tz = pytz_timezone(tz)
234
+ dt = dt.replace(tzinfo=tz)
235
+ return dt
236
+ except Exception:
237
+ return None
238
+
239
+
240
+ # =======================对外接口=======================
241
+
242
+ def parse_time(time_input: TimeType, default: Optional[datetime] = None) -> Optional[datetime]:
243
+ """解析时间字符串或对象"""
244
+ return TimeUtils.parse(time_input, default=default)
245
+
246
+
247
+ def format_time(dt: TimeType, fmt: str = "%Y-%m-%d %H:%M:%S") -> Optional[str]:
248
+ """格式化时间"""
249
+ return TimeUtils.format(dt, fmt)
250
+
251
+
252
+ def time_diff(start: TimeType, end: TimeType, unit: TimeUnit = "seconds") -> Optional[int]:
253
+ """计算时间差"""
254
+ return TimeUtils.diff(start, end, unit)
255
+
256
+
257
+ def to_timestamp(dt: TimeType) -> Optional[float]:
258
+ """转时间戳"""
259
+ return TimeUtils.to_timestamp(dt)
260
+
261
+
262
+ def to_datetime(ts: float) -> Optional[datetime]:
263
+ """从时间戳转 datetime"""
264
+ return TimeUtils.from_timestamp(ts)
265
+
266
+
267
+ def now(fmt: Optional[str] = None) -> Union[datetime, str]:
268
+ """获取当前时间"""
269
+ return TimeUtils.now(fmt)
270
+
271
+ def to_timezone(dt: TimeType, tz: TimezoneType) -> Optional[datetime]:
272
+ """将时间转换为指定时区"""
273
+ return TimeUtils.to_timezone(dt, tz)
274
+
275
+ def to_utc(dt: TimeType) -> Optional[datetime]:
276
+ """将时间转换为 UTC 时区"""
277
+ return TimeUtils.to_utc(dt)
278
+
279
+ def to_local(dt: TimeType) -> Optional[datetime]:
280
+ """将时间转换为本地时区"""
281
+ return TimeUtils.to_local(dt)
282
+
283
+ def from_timestamp_with_tz(ts: float, tz: TimezoneType = None) -> Optional[datetime]:
284
+ """从时间戳创建 datetime,并可选择指定时区"""
285
+ return TimeUtils.from_timestamp_with_tz(ts, tz)
286
+
287
+
288
+ if __name__ == '__main__':
289
+ get_current_time = now(fmt="%Y-%m-%d %H:%M:%S")
290
+ print(get_current_time)