duckdb 0.7.2-dev12.0 → 0.7.2-dev1244.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (631) hide show
  1. package/binding.gyp +12 -7
  2. package/lib/duckdb.d.ts +55 -2
  3. package/lib/duckdb.js +20 -1
  4. package/package.json +1 -1
  5. package/src/connection.cpp +1 -2
  6. package/src/database.cpp +1 -1
  7. package/src/duckdb/extension/icu/icu-extension.cpp +4 -0
  8. package/src/duckdb/extension/icu/icu-list-range.cpp +207 -0
  9. package/src/duckdb/extension/icu/icu-table-range.cpp +194 -0
  10. package/src/duckdb/extension/icu/include/icu-list-range.hpp +17 -0
  11. package/src/duckdb/extension/icu/include/icu-table-range.hpp +17 -0
  12. package/src/duckdb/extension/icu/third_party/icu/stubdata/stubdata.cpp +1 -1
  13. package/src/duckdb/extension/json/include/json_common.hpp +1 -0
  14. package/src/duckdb/extension/json/include/json_functions.hpp +2 -0
  15. package/src/duckdb/extension/json/include/json_serializer.hpp +77 -0
  16. package/src/duckdb/extension/json/json_functions/json_serialize_sql.cpp +147 -0
  17. package/src/duckdb/extension/json/json_functions/read_json.cpp +6 -5
  18. package/src/duckdb/extension/json/json_functions.cpp +12 -4
  19. package/src/duckdb/extension/json/json_scan.cpp +2 -2
  20. package/src/duckdb/extension/json/json_serializer.cpp +217 -0
  21. package/src/duckdb/extension/parquet/column_reader.cpp +94 -15
  22. package/src/duckdb/extension/parquet/column_writer.cpp +0 -1
  23. package/src/duckdb/extension/parquet/include/column_reader.hpp +1 -2
  24. package/src/duckdb/extension/parquet/include/decode_utils.hpp +5 -4
  25. package/src/duckdb/extension/parquet/include/generated_column_reader.hpp +1 -11
  26. package/src/duckdb/extension/parquet/include/parquet_timestamp.hpp +2 -1
  27. package/src/duckdb/extension/parquet/parquet-extension.cpp +14 -3
  28. package/src/duckdb/extension/parquet/parquet_reader.cpp +6 -1
  29. package/src/duckdb/extension/parquet/parquet_statistics.cpp +49 -36
  30. package/src/duckdb/extension/parquet/parquet_timestamp.cpp +16 -6
  31. package/src/duckdb/src/catalog/catalog.cpp +34 -5
  32. package/src/duckdb/src/catalog/catalog_entry/duck_schema_entry.cpp +4 -0
  33. package/src/duckdb/src/catalog/catalog_entry/duck_table_entry.cpp +2 -21
  34. package/src/duckdb/src/catalog/catalog_entry/scalar_function_catalog_entry.cpp +7 -6
  35. package/src/duckdb/src/catalog/catalog_entry/table_catalog_entry.cpp +3 -3
  36. package/src/duckdb/src/catalog/catalog_entry/table_function_catalog_entry.cpp +20 -1
  37. package/src/duckdb/src/catalog/catalog_entry/type_catalog_entry.cpp +8 -2
  38. package/src/duckdb/src/catalog/catalog_set.cpp +1 -0
  39. package/src/duckdb/src/catalog/default/default_functions.cpp +3 -0
  40. package/src/duckdb/src/catalog/dependency_list.cpp +12 -0
  41. package/src/duckdb/src/catalog/duck_catalog.cpp +34 -7
  42. package/src/duckdb/src/common/arrow/arrow_appender.cpp +48 -4
  43. package/src/duckdb/src/common/arrow/arrow_converter.cpp +1 -1
  44. package/src/duckdb/src/common/box_renderer.cpp +109 -23
  45. package/src/duckdb/src/common/enums/expression_type.cpp +8 -222
  46. package/src/duckdb/src/common/enums/join_type.cpp +3 -22
  47. package/src/duckdb/src/common/enums/logical_operator_type.cpp +2 -0
  48. package/src/duckdb/src/common/enums/statement_type.cpp +2 -0
  49. package/src/duckdb/src/common/exception.cpp +15 -1
  50. package/src/duckdb/src/common/field_writer.cpp +1 -0
  51. package/src/duckdb/src/common/hive_partitioning.cpp +3 -1
  52. package/src/duckdb/src/common/local_file_system.cpp +64 -7
  53. package/src/duckdb/src/common/operator/cast_operators.cpp +1 -1
  54. package/src/duckdb/src/common/preserved_error.cpp +7 -5
  55. package/src/duckdb/src/common/progress_bar/progress_bar.cpp +7 -0
  56. package/src/duckdb/src/common/serializer/buffered_deserializer.cpp +4 -0
  57. package/src/duckdb/src/common/serializer/buffered_file_reader.cpp +15 -2
  58. package/src/duckdb/src/common/serializer/enum_serializer.cpp +1176 -0
  59. package/src/duckdb/src/common/sort/comparators.cpp +14 -5
  60. package/src/duckdb/src/common/sort/sort_state.cpp +5 -7
  61. package/src/duckdb/src/common/sort/sorted_block.cpp +0 -1
  62. package/src/duckdb/src/common/string_util.cpp +18 -1
  63. package/src/duckdb/src/common/types/bit.cpp +166 -87
  64. package/src/duckdb/src/common/types/blob.cpp +1 -1
  65. package/src/duckdb/src/common/types/chunk_collection.cpp +2 -2
  66. package/src/duckdb/src/common/types/column_data_collection.cpp +39 -2
  67. package/src/duckdb/src/common/types/column_data_collection_segment.cpp +12 -10
  68. package/src/duckdb/src/common/types/data_chunk.cpp +1 -1
  69. package/src/duckdb/src/common/types/interval.cpp +0 -41
  70. package/src/duckdb/src/common/types/list_segment.cpp +658 -0
  71. package/src/duckdb/src/common/types/string_heap.cpp +1 -1
  72. package/src/duckdb/src/common/types/string_type.cpp +1 -1
  73. package/src/duckdb/src/common/types/time.cpp +13 -0
  74. package/src/duckdb/src/common/types/validity_mask.cpp +24 -7
  75. package/src/duckdb/src/common/types/value.cpp +320 -154
  76. package/src/duckdb/src/common/types/vector.cpp +158 -134
  77. package/src/duckdb/src/common/types.cpp +313 -153
  78. package/src/duckdb/src/common/value_operations/comparison_operations.cpp +14 -22
  79. package/src/duckdb/src/common/vector_operations/comparison_operators.cpp +10 -10
  80. package/src/duckdb/src/common/vector_operations/is_distinct_from.cpp +11 -10
  81. package/src/duckdb/src/common/vector_operations/vector_cast.cpp +2 -1
  82. package/src/duckdb/src/execution/aggregate_hashtable.cpp +98 -74
  83. package/src/duckdb/src/execution/column_binding_resolver.cpp +21 -5
  84. package/src/duckdb/src/execution/expression_executor/execute_cast.cpp +2 -1
  85. package/src/duckdb/src/execution/expression_executor/execute_comparison.cpp +2 -2
  86. package/src/duckdb/src/execution/index/art/art.cpp +19 -5
  87. package/src/duckdb/src/execution/join_hashtable.cpp +3 -1
  88. package/src/duckdb/src/execution/operator/aggregate/physical_hash_aggregate.cpp +1 -1
  89. package/src/duckdb/src/execution/operator/aggregate/physical_perfecthash_aggregate.cpp +4 -5
  90. package/src/duckdb/src/execution/operator/aggregate/physical_window.cpp +117 -26
  91. package/src/duckdb/src/execution/operator/helper/physical_limit.cpp +3 -0
  92. package/src/duckdb/src/execution/operator/helper/physical_vacuum.cpp +5 -3
  93. package/src/duckdb/src/execution/operator/join/physical_blockwise_nl_join.cpp +64 -17
  94. package/src/duckdb/src/execution/operator/join/physical_hash_join.cpp +2 -0
  95. package/src/duckdb/src/execution/operator/join/physical_iejoin.cpp +2 -2
  96. package/src/duckdb/src/execution/operator/join/physical_index_join.cpp +13 -4
  97. package/src/duckdb/src/execution/operator/join/physical_join.cpp +0 -3
  98. package/src/duckdb/src/execution/operator/join/physical_piecewise_merge_join.cpp +6 -11
  99. package/src/duckdb/src/execution/operator/join/physical_range_join.cpp +3 -1
  100. package/src/duckdb/src/execution/operator/persistent/base_csv_reader.cpp +11 -4
  101. package/src/duckdb/src/execution/operator/persistent/buffered_csv_reader.cpp +24 -19
  102. package/src/duckdb/src/execution/operator/persistent/csv_reader_options.cpp +3 -0
  103. package/src/duckdb/src/execution/operator/persistent/physical_batch_insert.cpp +2 -1
  104. package/src/duckdb/src/execution/operator/persistent/physical_copy_to_file.cpp +2 -2
  105. package/src/duckdb/src/execution/operator/persistent/physical_delete.cpp +1 -3
  106. package/src/duckdb/src/execution/operator/persistent/physical_insert.cpp +1 -0
  107. package/src/duckdb/src/execution/operator/projection/physical_projection.cpp +34 -0
  108. package/src/duckdb/src/execution/operator/scan/physical_positional_scan.cpp +20 -5
  109. package/src/duckdb/src/execution/operator/schema/physical_create_type.cpp +20 -40
  110. package/src/duckdb/src/execution/operator/set/physical_recursive_cte.cpp +2 -5
  111. package/src/duckdb/src/execution/partitionable_hashtable.cpp +20 -5
  112. package/src/duckdb/src/execution/physical_plan/plan_aggregate.cpp +22 -16
  113. package/src/duckdb/src/execution/physical_plan/plan_asof_join.cpp +97 -0
  114. package/src/duckdb/src/execution/physical_plan/plan_comparison_join.cpp +95 -47
  115. package/src/duckdb/src/execution/physical_plan/plan_create_index.cpp +2 -1
  116. package/src/duckdb/src/execution/physical_plan/plan_distinct.cpp +5 -8
  117. package/src/duckdb/src/execution/physical_plan/plan_positional_join.cpp +14 -5
  118. package/src/duckdb/src/execution/physical_plan_generator.cpp +3 -0
  119. package/src/duckdb/src/execution/radix_partitioned_hashtable.cpp +23 -15
  120. package/src/duckdb/src/execution/window_segment_tree.cpp +173 -1
  121. package/src/duckdb/src/function/aggregate/algebraic/avg.cpp +0 -6
  122. package/src/duckdb/src/function/aggregate/distributive/bitagg.cpp +99 -95
  123. package/src/duckdb/src/function/aggregate/distributive/bitstring_agg.cpp +269 -0
  124. package/src/duckdb/src/function/aggregate/distributive/bool.cpp +2 -0
  125. package/src/duckdb/src/function/aggregate/distributive/count.cpp +3 -4
  126. package/src/duckdb/src/function/aggregate/distributive/first.cpp +1 -0
  127. package/src/duckdb/src/function/aggregate/distributive/minmax.cpp +2 -0
  128. package/src/duckdb/src/function/aggregate/distributive/sum.cpp +19 -16
  129. package/src/duckdb/src/function/aggregate/distributive_functions.cpp +1 -0
  130. package/src/duckdb/src/function/aggregate/holistic/approximate_quantile.cpp +5 -2
  131. package/src/duckdb/src/function/aggregate/holistic/mode.cpp +1 -1
  132. package/src/duckdb/src/function/aggregate/holistic/quantile.cpp +16 -1
  133. package/src/duckdb/src/function/aggregate/nested/list.cpp +6 -712
  134. package/src/duckdb/src/function/aggregate/sorted_aggregate_function.cpp +138 -45
  135. package/src/duckdb/src/function/cast/bit_cast.cpp +0 -2
  136. package/src/duckdb/src/function/cast/blob_cast.cpp +0 -1
  137. package/src/duckdb/src/function/cast/cast_function_set.cpp +1 -1
  138. package/src/duckdb/src/function/cast/enum_casts.cpp +25 -3
  139. package/src/duckdb/src/function/cast/list_casts.cpp +17 -4
  140. package/src/duckdb/src/function/cast/map_cast.cpp +5 -2
  141. package/src/duckdb/src/function/cast/string_cast.cpp +36 -10
  142. package/src/duckdb/src/function/cast/struct_cast.cpp +24 -4
  143. package/src/duckdb/src/function/cast/time_casts.cpp +2 -2
  144. package/src/duckdb/src/function/cast/union_casts.cpp +33 -7
  145. package/src/duckdb/src/function/cast_rules.cpp +9 -4
  146. package/src/duckdb/src/function/function_binder.cpp +1 -8
  147. package/src/duckdb/src/function/pragma/pragma_queries.cpp +24 -1
  148. package/src/duckdb/src/function/scalar/bit/bitstring.cpp +100 -0
  149. package/src/duckdb/src/function/scalar/date/current.cpp +0 -2
  150. package/src/duckdb/src/function/scalar/date/date_diff.cpp +0 -1
  151. package/src/duckdb/src/function/scalar/date/date_part.cpp +18 -26
  152. package/src/duckdb/src/function/scalar/date/date_sub.cpp +0 -1
  153. package/src/duckdb/src/function/scalar/date/date_trunc.cpp +10 -14
  154. package/src/duckdb/src/function/scalar/generic/stats.cpp +2 -4
  155. package/src/duckdb/src/function/scalar/list/contains_or_position.cpp +4 -146
  156. package/src/duckdb/src/function/scalar/list/flatten.cpp +5 -12
  157. package/src/duckdb/src/function/scalar/list/list_aggregates.cpp +1 -1
  158. package/src/duckdb/src/function/scalar/list/list_concat.cpp +8 -12
  159. package/src/duckdb/src/function/scalar/list/list_extract.cpp +5 -12
  160. package/src/duckdb/src/function/scalar/list/list_lambdas.cpp +7 -3
  161. package/src/duckdb/src/function/scalar/list/list_sort.cpp +25 -18
  162. package/src/duckdb/src/function/scalar/list/list_value.cpp +6 -10
  163. package/src/duckdb/src/function/scalar/map/map.cpp +47 -1
  164. package/src/duckdb/src/function/scalar/map/map_entries.cpp +61 -0
  165. package/src/duckdb/src/function/scalar/map/map_extract.cpp +68 -26
  166. package/src/duckdb/src/function/scalar/map/map_keys_values.cpp +97 -0
  167. package/src/duckdb/src/function/scalar/math/numeric.cpp +101 -17
  168. package/src/duckdb/src/function/scalar/math_functions.cpp +3 -0
  169. package/src/duckdb/src/function/scalar/nested_functions.cpp +3 -0
  170. package/src/duckdb/src/function/scalar/operators/add.cpp +0 -9
  171. package/src/duckdb/src/function/scalar/operators/arithmetic.cpp +29 -48
  172. package/src/duckdb/src/function/scalar/operators/bitwise.cpp +0 -63
  173. package/src/duckdb/src/function/scalar/operators/multiply.cpp +5 -6
  174. package/src/duckdb/src/function/scalar/operators/subtract.cpp +0 -6
  175. package/src/duckdb/src/function/scalar/string/caseconvert.cpp +2 -6
  176. package/src/duckdb/src/function/scalar/string/hex.cpp +201 -0
  177. package/src/duckdb/src/function/scalar/string/instr.cpp +2 -6
  178. package/src/duckdb/src/function/scalar/string/length.cpp +2 -6
  179. package/src/duckdb/src/function/scalar/string/like.cpp +2 -6
  180. package/src/duckdb/src/function/scalar/string/regexp/regexp_extract_all.cpp +243 -0
  181. package/src/duckdb/src/function/scalar/string/regexp/regexp_util.cpp +79 -0
  182. package/src/duckdb/src/function/scalar/string/regexp.cpp +21 -80
  183. package/src/duckdb/src/function/scalar/string/substring.cpp +2 -6
  184. package/src/duckdb/src/function/scalar/string_functions.cpp +2 -0
  185. package/src/duckdb/src/function/scalar/struct/struct_extract.cpp +5 -10
  186. package/src/duckdb/src/function/scalar/struct/struct_insert.cpp +11 -14
  187. package/src/duckdb/src/function/scalar/struct/struct_pack.cpp +6 -7
  188. package/src/duckdb/src/function/table/arrow.cpp +5 -2
  189. package/src/duckdb/src/function/table/arrow_conversion.cpp +25 -1
  190. package/src/duckdb/src/function/table/checkpoint.cpp +5 -1
  191. package/src/duckdb/src/function/table/read_csv.cpp +60 -0
  192. package/src/duckdb/src/function/table/system/duckdb_constraints.cpp +2 -2
  193. package/src/duckdb/src/function/table/system/test_all_types.cpp +2 -2
  194. package/src/duckdb/src/function/table/table_scan.cpp +9 -12
  195. package/src/duckdb/src/function/table/version/pragma_version.cpp +2 -2
  196. package/src/duckdb/src/function/table_function.cpp +30 -11
  197. package/src/duckdb/src/include/duckdb/catalog/catalog.hpp +6 -0
  198. package/src/duckdb/src/include/duckdb/catalog/catalog_entry/duck_table_entry.hpp +1 -1
  199. package/src/duckdb/src/include/duckdb/catalog/catalog_entry/table_function_catalog_entry.hpp +6 -8
  200. package/src/duckdb/src/include/duckdb/catalog/dependency_list.hpp +3 -0
  201. package/src/duckdb/src/include/duckdb/catalog/duck_catalog.hpp +2 -1
  202. package/src/duckdb/src/include/duckdb/common/box_renderer.hpp +8 -2
  203. package/src/duckdb/src/include/duckdb/common/constants.hpp +0 -19
  204. package/src/duckdb/src/include/duckdb/common/enums/aggregate_handling.hpp +2 -0
  205. package/src/duckdb/src/include/duckdb/common/enums/expression_type.hpp +2 -3
  206. package/src/duckdb/src/include/duckdb/common/enums/joinref_type.hpp +7 -4
  207. package/src/duckdb/src/include/duckdb/common/enums/logical_operator_type.hpp +1 -0
  208. package/src/duckdb/src/include/duckdb/common/enums/order_type.hpp +2 -0
  209. package/src/duckdb/src/include/duckdb/common/enums/set_operation_type.hpp +2 -1
  210. package/src/duckdb/src/include/duckdb/common/enums/statement_type.hpp +2 -1
  211. package/src/duckdb/src/include/duckdb/common/enums/tableref_type.hpp +2 -1
  212. package/src/duckdb/src/include/duckdb/common/exception.hpp +69 -2
  213. package/src/duckdb/src/include/duckdb/common/field_writer.hpp +12 -4
  214. package/src/duckdb/src/include/duckdb/common/helper.hpp +1 -1
  215. package/src/duckdb/src/include/duckdb/common/{http_stats.hpp → http_state.hpp} +18 -4
  216. package/src/duckdb/src/include/duckdb/common/operator/comparison_operators.hpp +45 -149
  217. package/src/duckdb/src/include/duckdb/common/operator/multiply.hpp +2 -0
  218. package/src/duckdb/src/include/duckdb/common/optional_ptr.hpp +45 -0
  219. package/src/duckdb/src/include/duckdb/common/preserved_error.hpp +6 -1
  220. package/src/duckdb/src/include/duckdb/common/progress_bar/progress_bar.hpp +2 -0
  221. package/src/duckdb/src/include/duckdb/common/serializer/buffered_deserializer.hpp +4 -2
  222. package/src/duckdb/src/include/duckdb/common/serializer/buffered_file_reader.hpp +8 -2
  223. package/src/duckdb/src/include/duckdb/common/serializer/enum_serializer.hpp +113 -0
  224. package/src/duckdb/src/include/duckdb/common/serializer/format_deserializer.hpp +336 -0
  225. package/src/duckdb/src/include/duckdb/common/serializer/format_serializer.hpp +268 -0
  226. package/src/duckdb/src/include/duckdb/common/serializer/serialization_traits.hpp +126 -0
  227. package/src/duckdb/src/include/duckdb/common/serializer.hpp +13 -0
  228. package/src/duckdb/src/include/duckdb/common/string_util.hpp +27 -0
  229. package/src/duckdb/src/include/duckdb/common/types/bit.hpp +12 -7
  230. package/src/duckdb/src/include/duckdb/common/types/interval.hpp +39 -3
  231. package/src/duckdb/src/include/duckdb/common/types/list_segment.hpp +70 -0
  232. package/src/duckdb/src/include/duckdb/common/types/string_type.hpp +73 -3
  233. package/src/duckdb/src/include/duckdb/common/types/time.hpp +3 -0
  234. package/src/duckdb/src/include/duckdb/common/types/validity_mask.hpp +4 -1
  235. package/src/duckdb/src/include/duckdb/common/types/value.hpp +17 -48
  236. package/src/duckdb/src/include/duckdb/common/types/value_map.hpp +1 -1
  237. package/src/duckdb/src/include/duckdb/common/types/vector.hpp +3 -1
  238. package/src/duckdb/src/include/duckdb/common/types.hpp +45 -8
  239. package/src/duckdb/src/include/duckdb/common/vector_operations/unary_executor.hpp +2 -2
  240. package/src/duckdb/src/include/duckdb/execution/aggregate_hashtable.hpp +35 -20
  241. package/src/duckdb/src/include/duckdb/execution/index/art/art.hpp +3 -14
  242. package/src/duckdb/src/include/duckdb/execution/operator/aggregate/physical_perfecthash_aggregate.hpp +1 -1
  243. package/src/duckdb/src/include/duckdb/execution/operator/join/physical_cross_product.hpp +2 -0
  244. package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_file_handle.hpp +1 -0
  245. package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_reader_options.hpp +10 -0
  246. package/src/duckdb/src/include/duckdb/execution/operator/projection/physical_projection.hpp +5 -0
  247. package/src/duckdb/src/include/duckdb/execution/partitionable_hashtable.hpp +5 -1
  248. package/src/duckdb/src/include/duckdb/execution/physical_plan_generator.hpp +1 -3
  249. package/src/duckdb/src/include/duckdb/execution/window_segment_tree.hpp +54 -0
  250. package/src/duckdb/src/include/duckdb/function/aggregate/distributive_functions.hpp +5 -0
  251. package/src/duckdb/src/include/duckdb/function/aggregate_function.hpp +18 -6
  252. package/src/duckdb/src/include/duckdb/function/cast/bound_cast_data.hpp +84 -0
  253. package/src/duckdb/src/include/duckdb/function/cast/cast_function_set.hpp +2 -2
  254. package/src/duckdb/src/include/duckdb/function/cast/default_casts.hpp +28 -64
  255. package/src/duckdb/src/include/duckdb/function/function_binder.hpp +3 -6
  256. package/src/duckdb/src/include/duckdb/function/scalar/bit_functions.hpp +4 -0
  257. package/src/duckdb/src/include/duckdb/function/scalar/list/contains_or_position.hpp +138 -0
  258. package/src/duckdb/src/include/duckdb/function/scalar/math_functions.hpp +8 -0
  259. package/src/duckdb/src/include/duckdb/function/scalar/nested_functions.hpp +59 -0
  260. package/src/duckdb/src/include/duckdb/function/scalar/regexp.hpp +81 -1
  261. package/src/duckdb/src/include/duckdb/function/scalar/string_functions.hpp +4 -0
  262. package/src/duckdb/src/include/duckdb/function/scalar_function.hpp +2 -2
  263. package/src/duckdb/src/include/duckdb/function/table/arrow.hpp +12 -1
  264. package/src/duckdb/src/include/duckdb/function/table_function.hpp +10 -0
  265. package/src/duckdb/src/include/duckdb/main/capi/capi_internal.hpp +2 -0
  266. package/src/duckdb/src/include/duckdb/main/client_config.hpp +2 -0
  267. package/src/duckdb/src/include/duckdb/main/client_data.hpp +3 -3
  268. package/src/duckdb/src/include/duckdb/main/config.hpp +3 -0
  269. package/src/duckdb/src/include/duckdb/main/connection_manager.hpp +2 -0
  270. package/src/duckdb/src/include/duckdb/main/database.hpp +1 -0
  271. package/src/duckdb/src/include/duckdb/main/extension_entries.hpp +2 -0
  272. package/src/duckdb/src/include/duckdb/main/prepared_statement.hpp +2 -0
  273. package/src/duckdb/src/include/duckdb/main/relation/explain_relation.hpp +2 -1
  274. package/src/duckdb/src/include/duckdb/main/relation.hpp +2 -1
  275. package/src/duckdb/src/include/duckdb/optimizer/filter_pushdown.hpp +2 -0
  276. package/src/duckdb/src/include/duckdb/optimizer/join_order/cardinality_estimator.hpp +2 -2
  277. package/src/duckdb/src/include/duckdb/optimizer/rule/list.hpp +1 -0
  278. package/src/duckdb/src/include/duckdb/optimizer/rule/ordered_aggregate_optimizer.hpp +24 -0
  279. package/src/duckdb/src/include/duckdb/parser/common_table_expression_info.hpp +4 -0
  280. package/src/duckdb/src/include/duckdb/parser/expression/between_expression.hpp +3 -0
  281. package/src/duckdb/src/include/duckdb/parser/expression/bound_expression.hpp +2 -0
  282. package/src/duckdb/src/include/duckdb/parser/expression/case_expression.hpp +5 -0
  283. package/src/duckdb/src/include/duckdb/parser/expression/cast_expression.hpp +2 -0
  284. package/src/duckdb/src/include/duckdb/parser/expression/collate_expression.hpp +2 -0
  285. package/src/duckdb/src/include/duckdb/parser/expression/columnref_expression.hpp +2 -0
  286. package/src/duckdb/src/include/duckdb/parser/expression/comparison_expression.hpp +2 -0
  287. package/src/duckdb/src/include/duckdb/parser/expression/conjunction_expression.hpp +2 -0
  288. package/src/duckdb/src/include/duckdb/parser/expression/constant_expression.hpp +3 -0
  289. package/src/duckdb/src/include/duckdb/parser/expression/default_expression.hpp +1 -0
  290. package/src/duckdb/src/include/duckdb/parser/expression/function_expression.hpp +4 -2
  291. package/src/duckdb/src/include/duckdb/parser/expression/lambda_expression.hpp +2 -0
  292. package/src/duckdb/src/include/duckdb/parser/expression/operator_expression.hpp +2 -0
  293. package/src/duckdb/src/include/duckdb/parser/expression/parameter_expression.hpp +2 -0
  294. package/src/duckdb/src/include/duckdb/parser/expression/positional_reference_expression.hpp +2 -0
  295. package/src/duckdb/src/include/duckdb/parser/expression/star_expression.hpp +4 -2
  296. package/src/duckdb/src/include/duckdb/parser/expression/subquery_expression.hpp +2 -0
  297. package/src/duckdb/src/include/duckdb/parser/expression/window_expression.hpp +5 -0
  298. package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_info.hpp +5 -1
  299. package/src/duckdb/src/include/duckdb/parser/parsed_data/{alter_function_info.hpp → alter_scalar_function_info.hpp} +13 -13
  300. package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_function_info.hpp +47 -0
  301. package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_info.hpp +6 -0
  302. package/src/duckdb/src/include/duckdb/parser/parsed_data/create_table_function_info.hpp +2 -1
  303. package/src/duckdb/src/include/duckdb/parser/parsed_data/sample_options.hpp +2 -0
  304. package/src/duckdb/src/include/duckdb/parser/parsed_expression.hpp +5 -0
  305. package/src/duckdb/src/include/duckdb/parser/query_node/recursive_cte_node.hpp +3 -0
  306. package/src/duckdb/src/include/duckdb/parser/query_node/select_node.hpp +5 -0
  307. package/src/duckdb/src/include/duckdb/parser/query_node/set_operation_node.hpp +3 -0
  308. package/src/duckdb/src/include/duckdb/parser/query_node.hpp +13 -2
  309. package/src/duckdb/src/include/duckdb/parser/result_modifier.hpp +24 -1
  310. package/src/duckdb/src/include/duckdb/parser/sql_statement.hpp +2 -1
  311. package/src/duckdb/src/include/duckdb/parser/statement/multi_statement.hpp +28 -0
  312. package/src/duckdb/src/include/duckdb/parser/statement/select_statement.hpp +6 -1
  313. package/src/duckdb/src/include/duckdb/parser/tableref/basetableref.hpp +4 -0
  314. package/src/duckdb/src/include/duckdb/parser/tableref/emptytableref.hpp +2 -0
  315. package/src/duckdb/src/include/duckdb/parser/tableref/expressionlistref.hpp +3 -0
  316. package/src/duckdb/src/include/duckdb/parser/tableref/joinref.hpp +3 -0
  317. package/src/duckdb/src/include/duckdb/parser/tableref/list.hpp +1 -0
  318. package/src/duckdb/src/include/duckdb/parser/tableref/pivotref.hpp +87 -0
  319. package/src/duckdb/src/include/duckdb/parser/tableref/subqueryref.hpp +3 -0
  320. package/src/duckdb/src/include/duckdb/parser/tableref/table_function_ref.hpp +3 -0
  321. package/src/duckdb/src/include/duckdb/parser/tableref.hpp +3 -1
  322. package/src/duckdb/src/include/duckdb/parser/tokens.hpp +2 -0
  323. package/src/duckdb/src/include/duckdb/parser/transformer.hpp +33 -0
  324. package/src/duckdb/src/include/duckdb/planner/bind_context.hpp +2 -0
  325. package/src/duckdb/src/include/duckdb/planner/binder.hpp +15 -4
  326. package/src/duckdb/src/include/duckdb/planner/bound_result_modifier.hpp +3 -0
  327. package/src/duckdb/src/include/duckdb/planner/expression/bound_aggregate_expression.hpp +3 -0
  328. package/src/duckdb/src/include/duckdb/planner/expression_binder/base_select_binder.hpp +64 -0
  329. package/src/duckdb/src/include/duckdb/planner/expression_binder/having_binder.hpp +2 -2
  330. package/src/duckdb/src/include/duckdb/planner/expression_binder/order_binder.hpp +4 -1
  331. package/src/duckdb/src/include/duckdb/planner/expression_binder/qualify_binder.hpp +2 -2
  332. package/src/duckdb/src/include/duckdb/planner/expression_binder/select_binder.hpp +9 -38
  333. package/src/duckdb/src/include/duckdb/planner/expression_binder.hpp +1 -1
  334. package/src/duckdb/src/include/duckdb/planner/logical_tokens.hpp +1 -0
  335. package/src/duckdb/src/include/duckdb/planner/operator/list.hpp +1 -0
  336. package/src/duckdb/src/include/duckdb/planner/operator/logical_asof_join.hpp +22 -0
  337. package/src/duckdb/src/include/duckdb/planner/operator/logical_comparison_join.hpp +5 -2
  338. package/src/duckdb/src/include/duckdb/planner/operator/logical_distinct.hpp +3 -0
  339. package/src/duckdb/src/include/duckdb/planner/query_node/bound_select_node.hpp +8 -2
  340. package/src/duckdb/src/include/duckdb/storage/buffer/block_handle.hpp +2 -0
  341. package/src/duckdb/src/include/duckdb/storage/buffer_manager.hpp +76 -44
  342. package/src/duckdb/src/include/duckdb/storage/checkpoint/table_data_writer.hpp +3 -2
  343. package/src/duckdb/src/include/duckdb/storage/checkpoint_manager.hpp +1 -1
  344. package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_compress.hpp +2 -2
  345. package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_fetch.hpp +1 -1
  346. package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_scan.hpp +2 -1
  347. package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_compress.hpp +2 -2
  348. package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_fetch.hpp +1 -1
  349. package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_scan.hpp +2 -1
  350. package/src/duckdb/src/include/duckdb/storage/data_pointer.hpp +4 -3
  351. package/src/duckdb/src/include/duckdb/storage/data_table.hpp +4 -3
  352. package/src/duckdb/src/include/duckdb/storage/index.hpp +5 -4
  353. package/src/duckdb/src/include/duckdb/storage/meta_block_reader.hpp +7 -0
  354. package/src/duckdb/src/include/duckdb/storage/statistics/base_statistics.hpp +93 -29
  355. package/src/duckdb/src/include/duckdb/storage/statistics/column_statistics.hpp +22 -3
  356. package/src/duckdb/src/include/duckdb/storage/statistics/distinct_statistics.hpp +8 -6
  357. package/src/duckdb/src/include/duckdb/storage/statistics/list_stats.hpp +41 -0
  358. package/src/duckdb/src/include/duckdb/storage/statistics/node_statistics.hpp +26 -0
  359. package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats.hpp +114 -0
  360. package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats_union.hpp +62 -0
  361. package/src/duckdb/src/include/duckdb/storage/statistics/segment_statistics.hpp +2 -7
  362. package/src/duckdb/src/include/duckdb/storage/statistics/string_stats.hpp +74 -0
  363. package/src/duckdb/src/include/duckdb/storage/statistics/struct_stats.hpp +42 -0
  364. package/src/duckdb/src/include/duckdb/storage/string_uncompressed.hpp +2 -3
  365. package/src/duckdb/src/include/duckdb/storage/table/column_checkpoint_state.hpp +2 -1
  366. package/src/duckdb/src/include/duckdb/storage/table/column_data.hpp +21 -7
  367. package/src/duckdb/src/include/duckdb/storage/table/column_data_checkpointer.hpp +3 -2
  368. package/src/duckdb/src/include/duckdb/storage/table/column_segment.hpp +5 -6
  369. package/src/duckdb/src/include/duckdb/storage/table/column_segment_tree.hpp +18 -0
  370. package/src/duckdb/src/include/duckdb/storage/table/list_column_data.hpp +1 -1
  371. package/src/duckdb/src/include/duckdb/storage/table/persistent_table_data.hpp +6 -3
  372. package/src/duckdb/src/include/duckdb/storage/table/row_group.hpp +41 -45
  373. package/src/duckdb/src/include/duckdb/storage/table/row_group_collection.hpp +23 -7
  374. package/src/duckdb/src/include/duckdb/storage/table/row_group_segment_tree.hpp +35 -0
  375. package/src/duckdb/src/include/duckdb/storage/table/scan_state.hpp +21 -29
  376. package/src/duckdb/src/include/duckdb/storage/table/segment_base.hpp +6 -6
  377. package/src/duckdb/src/include/duckdb/storage/table/segment_tree.hpp +281 -26
  378. package/src/duckdb/src/include/duckdb/storage/table/standard_column_data.hpp +0 -4
  379. package/src/duckdb/src/include/duckdb/storage/table/table_statistics.hpp +5 -0
  380. package/src/duckdb/src/include/duckdb/storage/table/update_segment.hpp +0 -1
  381. package/src/duckdb/src/include/duckdb/storage/write_ahead_log.hpp +1 -1
  382. package/src/duckdb/src/include/duckdb/transaction/local_storage.hpp +6 -3
  383. package/src/duckdb/src/include/duckdb.h +71 -2
  384. package/src/duckdb/src/include/duckdb.hpp +0 -1
  385. package/src/duckdb/src/main/capi/pending-c.cpp +16 -3
  386. package/src/duckdb/src/main/capi/result-c.cpp +27 -1
  387. package/src/duckdb/src/main/capi/stream-c.cpp +25 -0
  388. package/src/duckdb/src/main/capi/table_function-c.cpp +23 -0
  389. package/src/duckdb/src/main/client_context.cpp +38 -34
  390. package/src/duckdb/src/main/client_data.cpp +7 -6
  391. package/src/duckdb/src/main/config.cpp +70 -1
  392. package/src/duckdb/src/main/database.cpp +19 -2
  393. package/src/duckdb/src/main/extension/extension_install.cpp +7 -2
  394. package/src/duckdb/src/main/prepared_statement.cpp +4 -0
  395. package/src/duckdb/src/main/query_profiler.cpp +17 -15
  396. package/src/duckdb/src/main/relation/explain_relation.cpp +3 -3
  397. package/src/duckdb/src/main/relation.cpp +3 -2
  398. package/src/duckdb/src/main/settings/settings.cpp +20 -8
  399. package/src/duckdb/src/optimizer/column_lifetime_analyzer.cpp +1 -0
  400. package/src/duckdb/src/optimizer/deliminator.cpp +1 -1
  401. package/src/duckdb/src/optimizer/filter_combiner.cpp +3 -6
  402. package/src/duckdb/src/optimizer/filter_pullup.cpp +3 -1
  403. package/src/duckdb/src/optimizer/filter_pushdown.cpp +14 -8
  404. package/src/duckdb/src/optimizer/join_order/cardinality_estimator.cpp +107 -71
  405. package/src/duckdb/src/optimizer/join_order/join_order_optimizer.cpp +32 -12
  406. package/src/duckdb/src/optimizer/optimizer.cpp +1 -0
  407. package/src/duckdb/src/optimizer/pullup/pullup_from_left.cpp +2 -2
  408. package/src/duckdb/src/optimizer/pushdown/pushdown_aggregate.cpp +33 -5
  409. package/src/duckdb/src/optimizer/pushdown/pushdown_cross_product.cpp +1 -1
  410. package/src/duckdb/src/optimizer/pushdown/pushdown_inner_join.cpp +3 -0
  411. package/src/duckdb/src/optimizer/pushdown/pushdown_left_join.cpp +5 -12
  412. package/src/duckdb/src/optimizer/pushdown/pushdown_mark_join.cpp +2 -2
  413. package/src/duckdb/src/optimizer/pushdown/pushdown_single_join.cpp +1 -1
  414. package/src/duckdb/src/optimizer/remove_unused_columns.cpp +1 -0
  415. package/src/duckdb/src/optimizer/rule/move_constants.cpp +10 -4
  416. package/src/duckdb/src/optimizer/rule/ordered_aggregate_optimizer.cpp +30 -0
  417. package/src/duckdb/src/optimizer/rule/regex_optimizations.cpp +9 -2
  418. package/src/duckdb/src/optimizer/statistics/expression/propagate_aggregate.cpp +9 -3
  419. package/src/duckdb/src/optimizer/statistics/expression/propagate_and_compress.cpp +6 -7
  420. package/src/duckdb/src/optimizer/statistics/expression/propagate_cast.cpp +14 -11
  421. package/src/duckdb/src/optimizer/statistics/expression/propagate_columnref.cpp +1 -1
  422. package/src/duckdb/src/optimizer/statistics/expression/propagate_comparison.cpp +13 -15
  423. package/src/duckdb/src/optimizer/statistics/expression/propagate_conjunction.cpp +0 -1
  424. package/src/duckdb/src/optimizer/statistics/expression/propagate_constant.cpp +3 -75
  425. package/src/duckdb/src/optimizer/statistics/expression/propagate_function.cpp +7 -2
  426. package/src/duckdb/src/optimizer/statistics/expression/propagate_operator.cpp +10 -0
  427. package/src/duckdb/src/optimizer/statistics/operator/propagate_aggregate.cpp +2 -3
  428. package/src/duckdb/src/optimizer/statistics/operator/propagate_filter.cpp +29 -32
  429. package/src/duckdb/src/optimizer/statistics/operator/propagate_join.cpp +5 -5
  430. package/src/duckdb/src/optimizer/statistics/operator/propagate_set_operation.cpp +3 -3
  431. package/src/duckdb/src/optimizer/statistics_propagator.cpp +2 -1
  432. package/src/duckdb/src/optimizer/unnest_rewriter.cpp +2 -2
  433. package/src/duckdb/src/parallel/meta_pipeline.cpp +0 -7
  434. package/src/duckdb/src/parser/common_table_expression_info.cpp +19 -0
  435. package/src/duckdb/src/parser/expression/between_expression.cpp +17 -0
  436. package/src/duckdb/src/parser/expression/case_expression.cpp +28 -0
  437. package/src/duckdb/src/parser/expression/cast_expression.cpp +17 -0
  438. package/src/duckdb/src/parser/expression/collate_expression.cpp +16 -0
  439. package/src/duckdb/src/parser/expression/columnref_expression.cpp +15 -0
  440. package/src/duckdb/src/parser/expression/comparison_expression.cpp +16 -0
  441. package/src/duckdb/src/parser/expression/conjunction_expression.cpp +17 -0
  442. package/src/duckdb/src/parser/expression/constant_expression.cpp +14 -0
  443. package/src/duckdb/src/parser/expression/default_expression.cpp +7 -0
  444. package/src/duckdb/src/parser/expression/function_expression.cpp +35 -0
  445. package/src/duckdb/src/parser/expression/lambda_expression.cpp +16 -0
  446. package/src/duckdb/src/parser/expression/operator_expression.cpp +15 -0
  447. package/src/duckdb/src/parser/expression/parameter_expression.cpp +15 -0
  448. package/src/duckdb/src/parser/expression/positional_reference_expression.cpp +14 -0
  449. package/src/duckdb/src/parser/expression/star_expression.cpp +26 -6
  450. package/src/duckdb/src/parser/expression/subquery_expression.cpp +20 -0
  451. package/src/duckdb/src/parser/expression/window_expression.cpp +43 -0
  452. package/src/duckdb/src/parser/parsed_data/alter_info.cpp +7 -3
  453. package/src/duckdb/src/parser/parsed_data/alter_scalar_function_info.cpp +56 -0
  454. package/src/duckdb/src/parser/parsed_data/alter_table_function_info.cpp +51 -0
  455. package/src/duckdb/src/parser/parsed_data/create_scalar_function_info.cpp +3 -2
  456. package/src/duckdb/src/parser/parsed_data/create_table_function_info.cpp +6 -0
  457. package/src/duckdb/src/parser/parsed_data/sample_options.cpp +22 -10
  458. package/src/duckdb/src/parser/parsed_expression.cpp +72 -0
  459. package/src/duckdb/src/parser/parsed_expression_iterator.cpp +15 -1
  460. package/src/duckdb/src/parser/query_node/recursive_cte_node.cpp +21 -0
  461. package/src/duckdb/src/parser/query_node/select_node.cpp +31 -0
  462. package/src/duckdb/src/parser/query_node/set_operation_node.cpp +17 -0
  463. package/src/duckdb/src/parser/query_node.cpp +51 -1
  464. package/src/duckdb/src/parser/result_modifier.cpp +78 -0
  465. package/src/duckdb/src/parser/statement/multi_statement.cpp +18 -0
  466. package/src/duckdb/src/parser/statement/select_statement.cpp +12 -0
  467. package/src/duckdb/src/parser/tableref/basetableref.cpp +21 -0
  468. package/src/duckdb/src/parser/tableref/emptytableref.cpp +4 -0
  469. package/src/duckdb/src/parser/tableref/expressionlistref.cpp +17 -0
  470. package/src/duckdb/src/parser/tableref/joinref.cpp +29 -0
  471. package/src/duckdb/src/parser/tableref/pivotref.cpp +373 -0
  472. package/src/duckdb/src/parser/tableref/subqueryref.cpp +15 -0
  473. package/src/duckdb/src/parser/tableref/table_function.cpp +17 -0
  474. package/src/duckdb/src/parser/tableref.cpp +49 -0
  475. package/src/duckdb/src/parser/transform/expression/transform_array_access.cpp +11 -0
  476. package/src/duckdb/src/parser/transform/expression/transform_bool_expr.cpp +1 -1
  477. package/src/duckdb/src/parser/transform/expression/transform_columnref.cpp +17 -2
  478. package/src/duckdb/src/parser/transform/expression/transform_function.cpp +85 -42
  479. package/src/duckdb/src/parser/transform/expression/transform_operator.cpp +1 -1
  480. package/src/duckdb/src/parser/transform/expression/transform_subquery.cpp +1 -1
  481. package/src/duckdb/src/parser/transform/helpers/transform_alias.cpp +12 -6
  482. package/src/duckdb/src/parser/transform/helpers/transform_cte.cpp +24 -0
  483. package/src/duckdb/src/parser/transform/helpers/transform_groupby.cpp +7 -0
  484. package/src/duckdb/src/parser/transform/helpers/transform_orderby.cpp +0 -7
  485. package/src/duckdb/src/parser/transform/helpers/transform_typename.cpp +3 -2
  486. package/src/duckdb/src/parser/transform/statement/transform_create_function.cpp +4 -0
  487. package/src/duckdb/src/parser/transform/statement/transform_create_view.cpp +4 -0
  488. package/src/duckdb/src/parser/transform/statement/transform_pivot_stmt.cpp +179 -0
  489. package/src/duckdb/src/parser/transform/statement/transform_rename.cpp +3 -4
  490. package/src/duckdb/src/parser/transform/statement/transform_select.cpp +8 -0
  491. package/src/duckdb/src/parser/transform/statement/transform_select_node.cpp +2 -3
  492. package/src/duckdb/src/parser/transform/tableref/transform_join.cpp +12 -1
  493. package/src/duckdb/src/parser/transform/tableref/transform_pivot.cpp +121 -0
  494. package/src/duckdb/src/parser/transform/tableref/transform_tableref.cpp +2 -0
  495. package/src/duckdb/src/parser/transformer.cpp +15 -3
  496. package/src/duckdb/src/planner/bind_context.cpp +18 -25
  497. package/src/duckdb/src/planner/binder/expression/bind_aggregate_expression.cpp +9 -7
  498. package/src/duckdb/src/planner/binder/expression/bind_columnref_expression.cpp +4 -3
  499. package/src/duckdb/src/planner/binder/expression/bind_function_expression.cpp +23 -12
  500. package/src/duckdb/src/planner/binder/expression/bind_lambda.cpp +3 -2
  501. package/src/duckdb/src/planner/binder/expression/bind_star_expression.cpp +176 -0
  502. package/src/duckdb/src/planner/binder/expression/bind_subquery_expression.cpp +4 -0
  503. package/src/duckdb/src/planner/binder/expression/bind_unnest_expression.cpp +163 -24
  504. package/src/duckdb/src/planner/binder/expression/bind_window_expression.cpp +2 -2
  505. package/src/duckdb/src/planner/binder/query_node/bind_select_node.cpp +109 -94
  506. package/src/duckdb/src/planner/binder/query_node/plan_query_node.cpp +11 -0
  507. package/src/duckdb/src/planner/binder/query_node/plan_select_node.cpp +9 -4
  508. package/src/duckdb/src/planner/binder/statement/bind_copy.cpp +5 -3
  509. package/src/duckdb/src/planner/binder/statement/bind_create.cpp +3 -2
  510. package/src/duckdb/src/planner/binder/statement/bind_create_table.cpp +10 -1
  511. package/src/duckdb/src/planner/binder/statement/bind_delete.cpp +1 -1
  512. package/src/duckdb/src/planner/binder/statement/bind_insert.cpp +12 -8
  513. package/src/duckdb/src/planner/binder/statement/bind_logical_plan.cpp +17 -0
  514. package/src/duckdb/src/planner/binder/statement/bind_update.cpp +4 -2
  515. package/src/duckdb/src/planner/binder/tableref/bind_joinref.cpp +19 -3
  516. package/src/duckdb/src/planner/binder/tableref/bind_pivot.cpp +366 -0
  517. package/src/duckdb/src/planner/binder/tableref/bind_table_function.cpp +11 -1
  518. package/src/duckdb/src/planner/binder/tableref/plan_cteref.cpp +1 -0
  519. package/src/duckdb/src/planner/binder/tableref/plan_joinref.cpp +61 -13
  520. package/src/duckdb/src/planner/binder.cpp +19 -24
  521. package/src/duckdb/src/planner/bound_result_modifier.cpp +27 -1
  522. package/src/duckdb/src/planner/expression/bound_aggregate_expression.cpp +9 -2
  523. package/src/duckdb/src/planner/expression/bound_expression.cpp +4 -0
  524. package/src/duckdb/src/planner/expression/bound_window_expression.cpp +1 -1
  525. package/src/duckdb/src/planner/expression_binder/base_select_binder.cpp +146 -0
  526. package/src/duckdb/src/planner/expression_binder/having_binder.cpp +6 -3
  527. package/src/duckdb/src/planner/expression_binder/qualify_binder.cpp +3 -3
  528. package/src/duckdb/src/planner/expression_binder/select_binder.cpp +1 -132
  529. package/src/duckdb/src/planner/expression_binder.cpp +10 -3
  530. package/src/duckdb/src/planner/expression_iterator.cpp +17 -10
  531. package/src/duckdb/src/planner/filter/constant_filter.cpp +4 -6
  532. package/src/duckdb/src/planner/logical_operator.cpp +7 -2
  533. package/src/duckdb/src/planner/logical_operator_visitor.cpp +6 -0
  534. package/src/duckdb/src/planner/operator/logical_asof_join.cpp +8 -0
  535. package/src/duckdb/src/planner/operator/logical_distinct.cpp +3 -0
  536. package/src/duckdb/src/planner/planner.cpp +2 -1
  537. package/src/duckdb/src/planner/pragma_handler.cpp +10 -2
  538. package/src/duckdb/src/planner/subquery/flatten_dependent_join.cpp +3 -1
  539. package/src/duckdb/src/storage/buffer_manager.cpp +44 -46
  540. package/src/duckdb/src/storage/checkpoint/row_group_writer.cpp +1 -1
  541. package/src/duckdb/src/storage/checkpoint/table_data_reader.cpp +4 -15
  542. package/src/duckdb/src/storage/checkpoint/table_data_writer.cpp +10 -4
  543. package/src/duckdb/src/storage/checkpoint_manager.cpp +9 -3
  544. package/src/duckdb/src/storage/compression/bitpacking.cpp +29 -25
  545. package/src/duckdb/src/storage/compression/fixed_size_uncompressed.cpp +45 -46
  546. package/src/duckdb/src/storage/compression/numeric_constant.cpp +10 -11
  547. package/src/duckdb/src/storage/compression/patas.cpp +1 -1
  548. package/src/duckdb/src/storage/compression/rle.cpp +20 -15
  549. package/src/duckdb/src/storage/compression/validity_uncompressed.cpp +6 -6
  550. package/src/duckdb/src/storage/data_table.cpp +23 -23
  551. package/src/duckdb/src/storage/index.cpp +12 -1
  552. package/src/duckdb/src/storage/local_storage.cpp +27 -23
  553. package/src/duckdb/src/storage/meta_block_reader.cpp +22 -0
  554. package/src/duckdb/src/storage/statistics/base_statistics.cpp +373 -128
  555. package/src/duckdb/src/storage/statistics/column_statistics.cpp +57 -3
  556. package/src/duckdb/src/storage/statistics/distinct_statistics.cpp +8 -9
  557. package/src/duckdb/src/storage/statistics/list_stats.cpp +121 -0
  558. package/src/duckdb/src/storage/statistics/numeric_stats.cpp +591 -0
  559. package/src/duckdb/src/storage/statistics/numeric_stats_union.cpp +65 -0
  560. package/src/duckdb/src/storage/statistics/segment_statistics.cpp +2 -11
  561. package/src/duckdb/src/storage/statistics/string_stats.cpp +273 -0
  562. package/src/duckdb/src/storage/statistics/struct_stats.cpp +133 -0
  563. package/src/duckdb/src/storage/storage_info.cpp +2 -2
  564. package/src/duckdb/src/storage/table/column_checkpoint_state.cpp +4 -10
  565. package/src/duckdb/src/storage/table/column_data.cpp +118 -62
  566. package/src/duckdb/src/storage/table/column_data_checkpointer.cpp +10 -9
  567. package/src/duckdb/src/storage/table/column_segment.cpp +30 -45
  568. package/src/duckdb/src/storage/table/list_column_data.cpp +50 -71
  569. package/src/duckdb/src/storage/table/persistent_table_data.cpp +2 -1
  570. package/src/duckdb/src/storage/table/row_group.cpp +213 -143
  571. package/src/duckdb/src/storage/table/row_group_collection.cpp +151 -105
  572. package/src/duckdb/src/storage/table/scan_state.cpp +45 -33
  573. package/src/duckdb/src/storage/table/standard_column_data.cpp +11 -12
  574. package/src/duckdb/src/storage/table/struct_column_data.cpp +27 -34
  575. package/src/duckdb/src/storage/table/table_statistics.cpp +27 -7
  576. package/src/duckdb/src/storage/table/update_segment.cpp +23 -18
  577. package/src/duckdb/src/storage/wal_replay.cpp +8 -5
  578. package/src/duckdb/src/storage/write_ahead_log.cpp +2 -2
  579. package/src/duckdb/src/transaction/commit_state.cpp +11 -7
  580. package/src/duckdb/src/verification/deserialized_statement_verifier.cpp +0 -1
  581. package/src/duckdb/third_party/libpg_query/include/nodes/nodes.hpp +35 -0
  582. package/src/duckdb/third_party/libpg_query/include/nodes/parsenodes.hpp +36 -2
  583. package/src/duckdb/third_party/libpg_query/include/nodes/primnodes.hpp +3 -3
  584. package/src/duckdb/third_party/libpg_query/include/parser/gram.hpp +1022 -530
  585. package/src/duckdb/third_party/libpg_query/include/parser/kwlist.hpp +8 -0
  586. package/src/duckdb/third_party/libpg_query/src_backend_parser_gram.cpp +24462 -22828
  587. package/src/duckdb/third_party/re2/re2/re2.cc +9 -0
  588. package/src/duckdb/third_party/re2/re2/re2.h +2 -0
  589. package/src/duckdb/ub_extension_icu_third_party_icu_i18n.cpp +4 -4
  590. package/src/duckdb/ub_extension_json_json_functions.cpp +2 -0
  591. package/src/duckdb/ub_src_common_serializer.cpp +2 -0
  592. package/src/duckdb/ub_src_common_types.cpp +2 -0
  593. package/src/duckdb/ub_src_execution_physical_plan.cpp +2 -0
  594. package/src/duckdb/ub_src_function_aggregate_distributive.cpp +2 -0
  595. package/src/duckdb/ub_src_function_scalar_bit.cpp +2 -0
  596. package/src/duckdb/ub_src_function_scalar_map.cpp +4 -0
  597. package/src/duckdb/ub_src_function_scalar_string.cpp +2 -0
  598. package/src/duckdb/ub_src_function_scalar_string_regexp.cpp +4 -0
  599. package/src/duckdb/ub_src_main_capi.cpp +2 -0
  600. package/src/duckdb/ub_src_optimizer_rule.cpp +2 -0
  601. package/src/duckdb/ub_src_parser.cpp +2 -0
  602. package/src/duckdb/ub_src_parser_parsed_data.cpp +4 -2
  603. package/src/duckdb/ub_src_parser_statement.cpp +2 -0
  604. package/src/duckdb/ub_src_parser_tableref.cpp +2 -0
  605. package/src/duckdb/ub_src_parser_transform_statement.cpp +2 -0
  606. package/src/duckdb/ub_src_parser_transform_tableref.cpp +2 -0
  607. package/src/duckdb/ub_src_planner_binder_expression.cpp +2 -0
  608. package/src/duckdb/ub_src_planner_binder_tableref.cpp +2 -0
  609. package/src/duckdb/ub_src_planner_expression_binder.cpp +2 -0
  610. package/src/duckdb/ub_src_planner_operator.cpp +2 -0
  611. package/src/duckdb/ub_src_storage_statistics.cpp +6 -6
  612. package/src/duckdb/ub_src_storage_table.cpp +0 -2
  613. package/src/duckdb_node.hpp +2 -1
  614. package/src/statement.cpp +5 -5
  615. package/src/utils.cpp +27 -2
  616. package/test/extension.test.ts +44 -26
  617. package/test/syntax_error.test.ts +3 -1
  618. package/filelist.cache +0 -0
  619. package/src/duckdb/src/include/duckdb/main/loadable_extension.hpp +0 -59
  620. package/src/duckdb/src/include/duckdb/storage/statistics/list_statistics.hpp +0 -36
  621. package/src/duckdb/src/include/duckdb/storage/statistics/numeric_statistics.hpp +0 -75
  622. package/src/duckdb/src/include/duckdb/storage/statistics/string_statistics.hpp +0 -49
  623. package/src/duckdb/src/include/duckdb/storage/statistics/struct_statistics.hpp +0 -36
  624. package/src/duckdb/src/include/duckdb/storage/statistics/validity_statistics.hpp +0 -45
  625. package/src/duckdb/src/parser/parsed_data/alter_function_info.cpp +0 -55
  626. package/src/duckdb/src/storage/statistics/list_statistics.cpp +0 -94
  627. package/src/duckdb/src/storage/statistics/numeric_statistics.cpp +0 -307
  628. package/src/duckdb/src/storage/statistics/string_statistics.cpp +0 -220
  629. package/src/duckdb/src/storage/statistics/struct_statistics.cpp +0 -108
  630. package/src/duckdb/src/storage/statistics/validity_statistics.cpp +0 -91
  631. package/src/duckdb/src/storage/table/segment_tree.cpp +0 -179
@@ -163,17 +163,26 @@ int Comparators::CompareStringAndAdvance(data_ptr_t &left_ptr, data_ptr_t &right
163
163
  if (!valid) {
164
164
  return 0;
165
165
  }
166
- // Construct the string_t
167
166
  uint32_t left_string_size = Load<uint32_t>(left_ptr);
168
167
  uint32_t right_string_size = Load<uint32_t>(right_ptr);
169
168
  left_ptr += sizeof(uint32_t);
170
169
  right_ptr += sizeof(uint32_t);
171
- string_t left_val((const char *)left_ptr, left_string_size);
172
- string_t right_val((const char *)right_ptr, right_string_size);
170
+ auto memcmp_res = memcmp((const char *)left_ptr, (const char *)right_ptr,
171
+ std::min<uint32_t>(left_string_size, right_string_size));
172
+
173
173
  left_ptr += left_string_size;
174
174
  right_ptr += right_string_size;
175
- // Compare
176
- return TemplatedCompareVal<string_t>((data_ptr_t)&left_val, (data_ptr_t)&right_val);
175
+
176
+ if (memcmp_res != 0) {
177
+ return memcmp_res;
178
+ }
179
+ if (left_string_size == right_string_size) {
180
+ return 0;
181
+ }
182
+ if (left_string_size < right_string_size) {
183
+ return -1;
184
+ }
185
+ return 1;
177
186
  }
178
187
 
179
188
  int Comparators::CompareStructAndAdvance(data_ptr_t &left_ptr, data_ptr_t &right_ptr,
@@ -3,7 +3,6 @@
3
3
  #include "duckdb/common/row_operations/row_operations.hpp"
4
4
  #include "duckdb/common/sort/sort.hpp"
5
5
  #include "duckdb/common/sort/sorted_block.hpp"
6
- #include "duckdb/storage/statistics/string_statistics.hpp"
7
6
 
8
7
  #include <algorithm>
9
8
  #include <numeric>
@@ -66,9 +65,8 @@ SortLayout::SortLayout(const vector<BoundOrderByNode> &orders)
66
65
  prefix_lengths.back() = GetNestedSortingColSize(col_size, expr.return_type);
67
66
  } else if (physical_type == PhysicalType::VARCHAR) {
68
67
  idx_t size_before = col_size;
69
- if (stats.back()) {
70
- auto &str_stats = (StringStatistics &)*stats.back();
71
- col_size += str_stats.max_string_length;
68
+ if (stats.back() && StringStats::HasMaxStringLength(*stats.back())) {
69
+ col_size += StringStats::MaxStringLength(*stats.back());
72
70
  if (col_size > 12) {
73
71
  col_size = 12;
74
72
  } else {
@@ -95,9 +93,9 @@ SortLayout::SortLayout(const vector<BoundOrderByNode> &orders)
95
93
  if (bytes_to_fill == 0) {
96
94
  break;
97
95
  }
98
- if (logical_types[col_idx].InternalType() == PhysicalType::VARCHAR && stats[col_idx]) {
99
- auto &str_stats = (StringStatistics &)*stats[col_idx];
100
- idx_t diff = str_stats.max_string_length - prefix_lengths[col_idx];
96
+ if (logical_types[col_idx].InternalType() == PhysicalType::VARCHAR && stats[col_idx] &&
97
+ StringStats::HasMaxStringLength(*stats[col_idx])) {
98
+ idx_t diff = StringStats::MaxStringLength(*stats[col_idx]) - prefix_lengths[col_idx];
101
99
  if (diff > 0) {
102
100
  // Increase all sizes accordingly
103
101
  idx_t increase = MinValue(bytes_to_fill, diff);
@@ -70,7 +70,6 @@ void SortedData::Unswizzle() {
70
70
  auto data_handle_p = buffer_manager.Pin(data_block->block);
71
71
  auto heap_handle_p = buffer_manager.Pin(heap_block->block);
72
72
  RowOperations::UnswizzlePointers(layout, data_handle_p.Ptr(), heap_handle_p.Ptr(), data_block->count);
73
- data_block->block->SetSwizzling("SortedData::Unswizzle");
74
73
  state.heap_blocks.push_back(std::move(heap_block));
75
74
  state.pinned_blocks.push_back(std::move(heap_handle_p));
76
75
  }
@@ -11,9 +11,23 @@
11
11
  #include <sstream>
12
12
  #include <stdarg.h>
13
13
  #include <string.h>
14
+ #include <random>
14
15
 
15
16
  namespace duckdb {
16
17
 
18
+ string StringUtil::GenerateRandomName(idx_t length) {
19
+ std::random_device rd;
20
+ std::mt19937 gen(rd());
21
+ std::uniform_int_distribution<> dis(0, 15);
22
+
23
+ std::stringstream ss;
24
+ ss << std::hex;
25
+ for (idx_t i = 0; i < length; i++) {
26
+ ss << dis(gen);
27
+ }
28
+ return ss.str();
29
+ }
30
+
17
31
  bool StringUtil::Contains(const string &haystack, const string &needle) {
18
32
  return (haystack.find(needle) != string::npos);
19
33
  }
@@ -191,11 +205,14 @@ vector<string> StringUtil::Split(const string &input, const string &split) {
191
205
 
192
206
  // Push the substring [last, next) on to splits
193
207
  string substr = input.substr(last, next - last);
194
- if (substr.empty() == false) {
208
+ if (!substr.empty()) {
195
209
  splits.push_back(substr);
196
210
  }
197
211
  last = next + split_len;
198
212
  }
213
+ if (splits.empty()) {
214
+ splits.push_back(input);
215
+ }
199
216
  return splits;
200
217
  }
201
218
 
@@ -4,71 +4,67 @@
4
4
 
5
5
  namespace duckdb {
6
6
 
7
- void Bit::SetEmptyBitString(string_t &target, string_t &input) {
8
- char *res_buf = target.GetDataWriteable();
9
- const char *buf = input.GetDataUnsafe();
10
- memset(res_buf, 0, input.GetSize());
11
- res_buf[0] = buf[0];
7
+ // **** helper functions ****
8
+ static char ComputePadding(idx_t len) {
9
+ return (8 - (len % 8)) % 8;
12
10
  }
13
11
 
14
- idx_t Bit::BitLength(string_t bits) {
15
- return ((bits.GetSize() - 1) * 8) - GetPadding(bits);
12
+ idx_t Bit::ComputeBitstringLen(idx_t len) {
13
+ idx_t result = len / 8;
14
+ if (len % 8 != 0) {
15
+ result++;
16
+ }
17
+ // additional first byte to store info on zero padding
18
+ result++;
19
+ return result;
16
20
  }
17
21
 
18
- idx_t Bit::OctetLength(string_t bits) {
19
- return bits.GetSize() - 1;
22
+ static inline idx_t GetBitPadding(const string_t &bit_string) {
23
+ auto data = (const_data_ptr_t)bit_string.GetDataUnsafe();
24
+ D_ASSERT(idx_t(data[0]) <= 8);
25
+ return data[0];
20
26
  }
21
27
 
22
- idx_t Bit::BitCount(string_t bits) {
23
- idx_t count = 0;
24
- const char *buf = bits.GetDataUnsafe();
25
- for (idx_t byte_idx = 1; byte_idx < OctetLength(bits) + 1; byte_idx++) {
26
- for (idx_t bit_idx = 0; bit_idx < 8; bit_idx++) {
27
- count += (buf[byte_idx] & (1 << bit_idx)) ? 1 : 0;
28
- }
28
+ static inline idx_t GetBitSize(const string_t &str) {
29
+ string error_message;
30
+ idx_t str_len;
31
+ if (!Bit::TryGetBitStringSize(str, str_len, &error_message)) {
32
+ throw ConversionException(error_message);
29
33
  }
30
- return count;
34
+ return str_len;
31
35
  }
32
36
 
33
- idx_t Bit::BitPosition(string_t substring, string_t bits) {
34
- const char *buf = bits.GetDataUnsafe();
35
- auto len = bits.GetSize();
36
- auto substr_len = BitLength(substring);
37
- idx_t substr_idx = 0;
38
-
39
- for (idx_t bit_idx = GetPadding(bits); bit_idx < 8; bit_idx++) {
40
- idx_t bit = buf[1] & (1 << (7 - bit_idx)) ? 1 : 0;
41
- if (bit == GetBit(substring, substr_idx)) {
42
- substr_idx++;
43
- if (substr_idx == substr_len) {
44
- return (bit_idx - GetPadding(bits)) - substr_len + 2;
45
- }
46
- } else {
47
- substr_idx = 0;
48
- }
37
+ void Bit::Finalize(string_t &str) {
38
+ // bit strings require all padding bits to be set to 1
39
+ // this method sets all padding bits to 1
40
+ auto padding = GetBitPadding(str);
41
+ for (idx_t i = 0; i < idx_t(padding); i++) {
42
+ Bit::SetBitInternal(str, i, 1);
49
43
  }
44
+ Bit::Verify(str);
45
+ }
50
46
 
51
- for (idx_t byte_idx = 2; byte_idx < len; byte_idx++) {
52
- for (idx_t bit_idx = 0; bit_idx < 8; bit_idx++) {
53
- idx_t bit = buf[byte_idx] & (1 << (7 - bit_idx)) ? 1 : 0;
54
- if (bit == GetBit(substring, substr_idx)) {
55
- substr_idx++;
56
- if (substr_idx == substr_len) {
57
- return (((byte_idx - 1) * 8) + bit_idx - GetPadding(bits)) - substr_len + 2;
58
- }
59
- } else {
60
- substr_idx = 0;
61
- }
62
- }
63
- }
64
- return 0;
47
+ void Bit::SetEmptyBitString(string_t &target, string_t &input) {
48
+ char *res_buf = target.GetDataWriteable();
49
+ const char *buf = input.GetDataUnsafe();
50
+ memset(res_buf, 0, input.GetSize());
51
+ res_buf[0] = buf[0];
52
+ Bit::Finalize(target);
65
53
  }
66
54
 
55
+ void Bit::SetEmptyBitString(string_t &target, idx_t len) {
56
+ char *res_buf = target.GetDataWriteable();
57
+ memset(res_buf, 0, target.GetSize());
58
+ res_buf[0] = ComputePadding(len);
59
+ Bit::Finalize(target);
60
+ }
61
+
62
+ // **** casting functions ****
67
63
  void Bit::ToString(string_t bits, char *output) {
68
64
  auto data = (const_data_ptr_t)bits.GetDataUnsafe();
69
65
  auto len = bits.GetSize();
70
66
 
71
- idx_t padding = GetPadding(bits);
67
+ idx_t padding = GetBitPadding(bits);
72
68
  idx_t output_idx = 0;
73
69
  for (idx_t bit_idx = padding; bit_idx < 8; bit_idx++) {
74
70
  output[output_idx++] = data[1] & (1 << (7 - bit_idx)) ? '1' : '0';
@@ -101,23 +97,19 @@ bool Bit::TryGetBitStringSize(string_t str, idx_t &str_len, string *error_messag
101
97
  return false;
102
98
  }
103
99
  }
104
- str_len = str_len % 8 ? (str_len / 8) + 1 : str_len / 8;
105
- str_len++; // additional first byte to store info on zero padding
106
- return true;
107
- }
108
-
109
- idx_t Bit::GetBitSize(string_t str) {
110
- string error_message;
111
- idx_t str_len;
112
- if (!Bit::TryGetBitStringSize(str, str_len, &error_message)) {
113
- throw ConversionException(error_message);
100
+ if (str_len == 0) {
101
+ string error = "Cannot cast empty string to BIT";
102
+ HandleCastError::AssignError(error, error_message);
103
+ return false;
114
104
  }
115
- return str_len;
105
+ str_len = ComputeBitstringLen(str_len);
106
+ return true;
116
107
  }
117
108
 
118
- void Bit::ToBit(string_t str, data_ptr_t output) {
109
+ void Bit::ToBit(string_t str, string_t &output_str) {
119
110
  auto data = (const_data_ptr_t)str.GetDataUnsafe();
120
111
  auto len = str.GetSize();
112
+ auto output = output_str.GetDataWriteable();
121
113
 
122
114
  char byte = 0;
123
115
  idx_t padded_byte = len % 8;
@@ -142,57 +134,124 @@ void Bit::ToBit(string_t str, data_ptr_t output) {
142
134
  }
143
135
  *(output++) = byte;
144
136
  }
137
+ Bit::Finalize(output_str);
138
+ Bit::Verify(output_str);
145
139
  }
146
140
 
147
141
  string Bit::ToBit(string_t str) {
148
142
  auto bit_len = GetBitSize(str);
149
143
  auto buffer = std::unique_ptr<char[]>(new char[bit_len]);
150
- Bit::ToBit(str, (data_ptr_t)buffer.get());
151
- return string(buffer.get(), bit_len);
144
+ string_t output_str(buffer.get(), bit_len);
145
+ Bit::ToBit(str, output_str);
146
+ return output_str.GetString();
152
147
  }
153
148
 
154
- idx_t Bit::GetBit(string_t bit_string, idx_t n) {
155
- const char *buf = bit_string.GetDataUnsafe();
156
- n += GetPadding(bit_string);
149
+ // **** scalar functions ****
150
+ void Bit::BitString(const string_t &input, const idx_t &bit_length, string_t &result) {
151
+ char *res_buf = result.GetDataWriteable();
152
+ const char *buf = input.GetDataUnsafe();
157
153
 
158
- char byte = buf[(n / 8) + 1] >> (7 - (n % 8));
159
- return (byte & 1 ? 1 : 0);
154
+ auto padding = ComputePadding(bit_length);
155
+ res_buf[0] = padding;
156
+ for (idx_t i = 0; i < bit_length; i++) {
157
+ if (i < bit_length - input.GetSize()) {
158
+ Bit::SetBit(result, i, 0);
159
+ } else {
160
+ idx_t bit = buf[i - (bit_length - input.GetSize())] == '1' ? 1 : 0;
161
+ Bit::SetBit(result, i, bit);
162
+ }
163
+ }
164
+ Bit::Finalize(result);
160
165
  }
161
166
 
162
- void Bit::SetBit(const string_t &bit_string, idx_t n, idx_t new_value, string_t &result) {
163
- char *result_buf = result.GetDataWriteable();
164
- const char *buf = bit_string.GetDataUnsafe();
165
- n += GetPadding(bit_string);
167
+ idx_t Bit::BitLength(string_t bits) {
168
+ return ((bits.GetSize() - 1) * 8) - GetBitPadding(bits);
169
+ }
166
170
 
167
- memcpy(result_buf, buf, bit_string.GetSize());
168
- char shift_byte = 1 << (7 - (n % 8));
169
- if (new_value == 0) {
170
- shift_byte = ~shift_byte;
171
- result_buf[(n / 8) + 1] = buf[(n / 8) + 1] & shift_byte;
172
- } else {
173
- result_buf[(n / 8) + 1] = buf[(n / 8) + 1] | shift_byte;
171
+ idx_t Bit::OctetLength(string_t bits) {
172
+ return bits.GetSize() - 1;
173
+ }
174
+
175
+ idx_t Bit::BitCount(string_t bits) {
176
+ idx_t count = 0;
177
+ const char *buf = bits.GetDataUnsafe();
178
+ for (idx_t byte_idx = 1; byte_idx < OctetLength(bits) + 1; byte_idx++) {
179
+ for (idx_t bit_idx = 0; bit_idx < 8; bit_idx++) {
180
+ count += (buf[byte_idx] & (1 << bit_idx)) ? 1 : 0;
181
+ }
174
182
  }
183
+ return count - GetBitPadding(bits);
184
+ }
185
+
186
+ idx_t Bit::BitPosition(string_t substring, string_t bits) {
187
+ const char *buf = bits.GetDataUnsafe();
188
+ auto len = bits.GetSize();
189
+ auto substr_len = BitLength(substring);
190
+ idx_t substr_idx = 0;
191
+
192
+ for (idx_t bit_idx = GetBitPadding(bits); bit_idx < 8; bit_idx++) {
193
+ idx_t bit = buf[1] & (1 << (7 - bit_idx)) ? 1 : 0;
194
+ if (bit == GetBit(substring, substr_idx)) {
195
+ substr_idx++;
196
+ if (substr_idx == substr_len) {
197
+ return (bit_idx - GetBitPadding(bits)) - substr_len + 2;
198
+ }
199
+ } else {
200
+ substr_idx = 0;
201
+ }
202
+ }
203
+
204
+ for (idx_t byte_idx = 2; byte_idx < len; byte_idx++) {
205
+ for (idx_t bit_idx = 0; bit_idx < 8; bit_idx++) {
206
+ idx_t bit = buf[byte_idx] & (1 << (7 - bit_idx)) ? 1 : 0;
207
+ if (bit == GetBit(substring, substr_idx)) {
208
+ substr_idx++;
209
+ if (substr_idx == substr_len) {
210
+ return (((byte_idx - 1) * 8) + bit_idx - GetBitPadding(bits)) - substr_len + 2;
211
+ }
212
+ } else {
213
+ substr_idx = 0;
214
+ }
215
+ }
216
+ }
217
+ return 0;
218
+ }
219
+
220
+ idx_t Bit::GetBit(string_t bit_string, idx_t n) {
221
+ return Bit::GetBitInternal(bit_string, n + GetBitPadding(bit_string));
222
+ }
223
+
224
+ idx_t Bit::GetBitIndex(idx_t n) {
225
+ return n / 8 + 1;
226
+ }
227
+
228
+ idx_t Bit::GetBitInternal(string_t bit_string, idx_t n) {
229
+ const char *buf = bit_string.GetDataUnsafe();
230
+ auto idx = Bit::GetBitIndex(n);
231
+ D_ASSERT(idx < bit_string.GetSize());
232
+ char byte = buf[idx] >> (7 - (n % 8));
233
+ return (byte & 1 ? 1 : 0);
175
234
  }
176
235
 
177
236
  void Bit::SetBit(string_t &bit_string, idx_t n, idx_t new_value) {
237
+ SetBitInternal(bit_string, n + GetBitPadding(bit_string), new_value);
238
+ }
239
+
240
+ void Bit::SetBitInternal(string_t &bit_string, idx_t n, idx_t new_value) {
178
241
  char *buf = bit_string.GetDataWriteable();
179
- n += GetPadding(bit_string);
180
242
 
243
+ auto idx = Bit::GetBitIndex(n);
244
+ D_ASSERT(idx < bit_string.GetSize());
181
245
  char shift_byte = 1 << (7 - (n % 8));
182
246
  if (new_value == 0) {
183
247
  shift_byte = ~shift_byte;
184
- buf[(n / 8) + 1] &= shift_byte;
248
+ buf[idx] &= shift_byte;
185
249
  } else {
186
- buf[(n / 8) + 1] |= shift_byte;
250
+ buf[idx] |= shift_byte;
187
251
  }
188
252
  }
189
253
 
190
- inline idx_t Bit::GetPadding(const string_t &bit_string) {
191
- auto data = (const_data_ptr_t)bit_string.GetDataUnsafe();
192
- return data[0];
193
- }
194
-
195
- // **** BITWISE OPERATORS ****
254
+ // **** BITWISE operators ****
196
255
  void Bit::RightShift(const string_t &bit_string, const idx_t &shift, string_t &result) {
197
256
  char *res_buf = result.GetDataWriteable();
198
257
  const char *buf = bit_string.GetDataUnsafe();
@@ -205,6 +264,7 @@ void Bit::RightShift(const string_t &bit_string, const idx_t &shift, string_t &r
205
264
  Bit::SetBit(result, i, bit);
206
265
  }
207
266
  }
267
+ Bit::Finalize(result);
208
268
  }
209
269
 
210
270
  void Bit::LeftShift(const string_t &bit_string, const idx_t &shift, string_t &result) {
@@ -219,6 +279,8 @@ void Bit::LeftShift(const string_t &bit_string, const idx_t &shift, string_t &re
219
279
  Bit::SetBit(result, i, 0);
220
280
  }
221
281
  }
282
+ Bit::Finalize(result);
283
+ Bit::Verify(result);
222
284
  }
223
285
 
224
286
  void Bit::BitwiseAnd(const string_t &rhs, const string_t &lhs, string_t &result) {
@@ -234,6 +296,8 @@ void Bit::BitwiseAnd(const string_t &rhs, const string_t &lhs, string_t &result)
234
296
  for (idx_t i = 1; i < lhs.GetSize(); i++) {
235
297
  buf[i] = l_buf[i] & r_buf[i];
236
298
  }
299
+ // and should preserve padding bits
300
+ Bit::Verify(result);
237
301
  }
238
302
 
239
303
  void Bit::BitwiseOr(const string_t &rhs, const string_t &lhs, string_t &result) {
@@ -249,6 +313,8 @@ void Bit::BitwiseOr(const string_t &rhs, const string_t &lhs, string_t &result)
249
313
  for (idx_t i = 1; i < lhs.GetSize(); i++) {
250
314
  buf[i] = l_buf[i] | r_buf[i];
251
315
  }
316
+ // or should preserve padding bits
317
+ Bit::Verify(result);
252
318
  }
253
319
 
254
320
  void Bit::BitwiseXor(const string_t &rhs, const string_t &lhs, string_t &result) {
@@ -264,6 +330,7 @@ void Bit::BitwiseXor(const string_t &rhs, const string_t &lhs, string_t &result)
264
330
  for (idx_t i = 1; i < lhs.GetSize(); i++) {
265
331
  buf[i] = l_buf[i] ^ r_buf[i];
266
332
  }
333
+ Bit::Finalize(result);
267
334
  }
268
335
 
269
336
  void Bit::BitwiseNot(const string_t &input, string_t &result) {
@@ -274,5 +341,17 @@ void Bit::BitwiseNot(const string_t &input, string_t &result) {
274
341
  for (idx_t i = 1; i < input.GetSize(); i++) {
275
342
  result_buf[i] = ~buf[i];
276
343
  }
344
+ Bit::Finalize(result);
277
345
  }
346
+
347
+ void Bit::Verify(const string_t &input) {
348
+ #ifdef DEBUG
349
+ // bit strings require all padding bits to be set to 1
350
+ auto padding = GetBitPadding(input);
351
+ for (idx_t i = 0; i < padding; i++) {
352
+ D_ASSERT(Bit::GetBitInternal(input, i));
353
+ }
354
+ #endif
355
+ }
356
+
278
357
  } // namespace duckdb
@@ -20,7 +20,7 @@ const int Blob::HEX_MAP[256] = {
20
20
  -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1, -1};
21
21
 
22
22
  bool IsRegularCharacter(data_t c) {
23
- return c >= 32 && c <= 127 && c != '\\' && c != '\'' && c != '"';
23
+ return c >= 32 && c <= 126 && c != '\\' && c != '\'' && c != '"';
24
24
  }
25
25
 
26
26
  idx_t Blob::GetStringSize(string_t blob) {
@@ -143,7 +143,7 @@ void ChunkCollection::Fuse(ChunkCollection &other) {
143
143
  auto &rhs = other.GetChunk(chunk_idx);
144
144
  lhs->data.reserve(rhs.data.size());
145
145
  for (auto &v : rhs.data) {
146
- lhs->data.emplace_back(Vector(v));
146
+ lhs->data.emplace_back(v);
147
147
  }
148
148
  lhs->SetCardinality(rhs.size());
149
149
  chunks.push_back(std::move(lhs));
@@ -156,7 +156,7 @@ void ChunkCollection::Fuse(ChunkCollection &other) {
156
156
  auto &rhs = other.GetChunk(chunk_idx);
157
157
  D_ASSERT(lhs.size() == rhs.size());
158
158
  for (auto &v : rhs.data) {
159
- lhs.data.emplace_back(Vector(v));
159
+ lhs.data.emplace_back(v);
160
160
  }
161
161
  }
162
162
  }
@@ -5,6 +5,7 @@
5
5
  #include "duckdb/common/types/column_data_collection_segment.hpp"
6
6
  #include "duckdb/common/vector_operations/vector_operations.hpp"
7
7
  #include "duckdb/storage/buffer_manager.hpp"
8
+ #include "duckdb/common/types/value_map.hpp"
8
9
 
9
10
  namespace duckdb {
10
11
 
@@ -944,6 +945,12 @@ void ColumnDataCollection::Reset() {
944
945
  allocator = make_shared<ColumnDataAllocator>(*allocator);
945
946
  }
946
947
 
948
+ struct ValueResultEquals {
949
+ bool operator()(const Value &a, const Value &b) const {
950
+ return Value::DefaultValuesAreEqual(a, b);
951
+ }
952
+ };
953
+
947
954
  bool ColumnDataCollection::ResultEquals(const ColumnDataCollection &left, const ColumnDataCollection &right,
948
955
  string &error_message) {
949
956
  if (left.ColumnCount() != right.ColumnCount()) {
@@ -959,13 +966,43 @@ bool ColumnDataCollection::ResultEquals(const ColumnDataCollection &left, const
959
966
  for (idx_t r = 0; r < left.Count(); r++) {
960
967
  for (idx_t c = 0; c < left.ColumnCount(); c++) {
961
968
  auto lvalue = left_rows.GetValue(c, r);
962
- auto rvalue = left_rows.GetValue(c, r);
969
+ auto rvalue = right_rows.GetValue(c, r);
963
970
  if (!Value::DefaultValuesAreEqual(lvalue, rvalue)) {
964
971
  error_message =
965
972
  StringUtil::Format("%s <> %s (row: %lld, col: %lld)\n", lvalue.ToString(), rvalue.ToString(), r, c);
966
- return false;
973
+ break;
967
974
  }
968
975
  }
976
+ if (!error_message.empty()) {
977
+ break;
978
+ }
979
+ }
980
+ if (!error_message.empty()) {
981
+ // do an unordered comparison
982
+ bool found_all = true;
983
+ for (idx_t c = 0; c < left.ColumnCount(); c++) {
984
+ std::unordered_multiset<Value, ValueHashFunction, ValueResultEquals> lvalues;
985
+ for (idx_t r = 0; r < left.Count(); r++) {
986
+ auto lvalue = left_rows.GetValue(c, r);
987
+ lvalues.insert(lvalue);
988
+ }
989
+ for (idx_t r = 0; r < right.Count(); r++) {
990
+ auto rvalue = right_rows.GetValue(c, r);
991
+ auto entry = lvalues.find(rvalue);
992
+ if (entry == lvalues.end()) {
993
+ found_all = false;
994
+ break;
995
+ }
996
+ lvalues.erase(entry);
997
+ }
998
+ if (!found_all) {
999
+ break;
1000
+ }
1001
+ }
1002
+ if (!found_all) {
1003
+ return false;
1004
+ }
1005
+ error_message = string();
969
1006
  }
970
1007
  return true;
971
1008
  }
@@ -1,4 +1,5 @@
1
1
  #include "duckdb/common/types/column_data_collection_segment.hpp"
2
+ #include "duckdb/common/vector_operations/vector_operations.hpp"
2
3
 
3
4
  namespace duckdb {
4
5
 
@@ -168,11 +169,8 @@ idx_t ColumnDataCollectionSegment::ReadVectorInternal(ChunkManagementState &stat
168
169
  if (type_size > 0) {
169
170
  memcpy(target_data + current_offset * type_size, base_ptr, current_vdata.count * type_size);
170
171
  }
171
- // FIXME: use bitwise operations here
172
172
  ValidityMask current_validity(validity_data);
173
- for (idx_t k = 0; k < current_vdata.count; k++) {
174
- target_validity.Set(current_offset + k, current_validity.RowIsValid(k));
175
- }
173
+ target_validity.SliceInPlace(current_validity, current_offset, 0, current_vdata.count);
176
174
  current_offset += current_vdata.count;
177
175
  next_index = current_vdata.next_data;
178
176
  }
@@ -202,12 +200,16 @@ idx_t ColumnDataCollectionSegment::ReadVector(ChunkManagementState &state, Vecto
202
200
  throw InternalException("Column Data Collection: mismatch in struct child sizes");
203
201
  }
204
202
  }
205
- } else if (internal_type == PhysicalType::VARCHAR &&
206
- allocator->GetType() == ColumnDataAllocatorType::BUFFER_MANAGER_ALLOCATOR) {
207
- for (auto &swizzle_segment : vdata.swizzle_data) {
208
- auto &string_heap_segment = GetVectorData(swizzle_segment.child_index);
209
- allocator->UnswizzlePointers(state, result, swizzle_segment.offset, swizzle_segment.count,
210
- string_heap_segment.block_id, string_heap_segment.offset);
203
+ } else if (internal_type == PhysicalType::VARCHAR) {
204
+ if (allocator->GetType() == ColumnDataAllocatorType::BUFFER_MANAGER_ALLOCATOR) {
205
+ for (auto &swizzle_segment : vdata.swizzle_data) {
206
+ auto &string_heap_segment = GetVectorData(swizzle_segment.child_index);
207
+ allocator->UnswizzlePointers(state, result, swizzle_segment.offset, swizzle_segment.count,
208
+ string_heap_segment.block_id, string_heap_segment.offset);
209
+ }
210
+ }
211
+ if (state.properties == ColumnDataScanProperties::DISALLOW_ZERO_COPY) {
212
+ VectorOperations::Copy(result, result, vdata.count, 0, 0);
211
213
  }
212
214
  }
213
215
  return vcount;
@@ -62,7 +62,7 @@ void DataChunk::InitializeEmpty(vector<LogicalType>::const_iterator begin, vecto
62
62
  D_ASSERT(data.empty()); // can only be initialized once
63
63
  D_ASSERT(std::distance(begin, end) != 0); // empty chunk not allowed
64
64
  for (; begin != end; begin++) {
65
- data.emplace_back(Vector(*begin, nullptr));
65
+ data.emplace_back(*begin, nullptr);
66
66
  }
67
67
  }
68
68
 
@@ -397,47 +397,6 @@ interval_t Interval::FromMicro(int64_t delta_us) {
397
397
  return result;
398
398
  }
399
399
 
400
- static void NormalizeIntervalEntries(interval_t input, int64_t &months, int64_t &days, int64_t &micros) {
401
- int64_t extra_months_d = input.days / Interval::DAYS_PER_MONTH;
402
- int64_t extra_months_micros = input.micros / Interval::MICROS_PER_MONTH;
403
- input.days -= extra_months_d * Interval::DAYS_PER_MONTH;
404
- input.micros -= extra_months_micros * Interval::MICROS_PER_MONTH;
405
-
406
- int64_t extra_days_micros = input.micros / Interval::MICROS_PER_DAY;
407
- input.micros -= extra_days_micros * Interval::MICROS_PER_DAY;
408
-
409
- months = input.months + extra_months_d + extra_months_micros;
410
- days = input.days + extra_days_micros;
411
- micros = input.micros;
412
- }
413
-
414
- bool Interval::Equals(interval_t left, interval_t right) {
415
- return left.months == right.months && left.days == right.days && left.micros == right.micros;
416
- }
417
-
418
- bool Interval::GreaterThan(interval_t left, interval_t right) {
419
- int64_t lmonths, ldays, lmicros;
420
- int64_t rmonths, rdays, rmicros;
421
- NormalizeIntervalEntries(left, lmonths, ldays, lmicros);
422
- NormalizeIntervalEntries(right, rmonths, rdays, rmicros);
423
-
424
- if (lmonths > rmonths) {
425
- return true;
426
- } else if (lmonths < rmonths) {
427
- return false;
428
- }
429
- if (ldays > rdays) {
430
- return true;
431
- } else if (ldays < rdays) {
432
- return false;
433
- }
434
- return lmicros > rmicros;
435
- }
436
-
437
- bool Interval::GreaterThanEquals(interval_t left, interval_t right) {
438
- return GreaterThan(left, right) || Equals(left, right);
439
- }
440
-
441
400
  interval_t Interval::Invert(interval_t interval) {
442
401
  interval.days = -interval.days;
443
402
  interval.micros = -interval.micros;