duckdb 0.7.2-dev12.0 → 0.7.2-dev1244.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (631) hide show
  1. package/binding.gyp +12 -7
  2. package/lib/duckdb.d.ts +55 -2
  3. package/lib/duckdb.js +20 -1
  4. package/package.json +1 -1
  5. package/src/connection.cpp +1 -2
  6. package/src/database.cpp +1 -1
  7. package/src/duckdb/extension/icu/icu-extension.cpp +4 -0
  8. package/src/duckdb/extension/icu/icu-list-range.cpp +207 -0
  9. package/src/duckdb/extension/icu/icu-table-range.cpp +194 -0
  10. package/src/duckdb/extension/icu/include/icu-list-range.hpp +17 -0
  11. package/src/duckdb/extension/icu/include/icu-table-range.hpp +17 -0
  12. package/src/duckdb/extension/icu/third_party/icu/stubdata/stubdata.cpp +1 -1
  13. package/src/duckdb/extension/json/include/json_common.hpp +1 -0
  14. package/src/duckdb/extension/json/include/json_functions.hpp +2 -0
  15. package/src/duckdb/extension/json/include/json_serializer.hpp +77 -0
  16. package/src/duckdb/extension/json/json_functions/json_serialize_sql.cpp +147 -0
  17. package/src/duckdb/extension/json/json_functions/read_json.cpp +6 -5
  18. package/src/duckdb/extension/json/json_functions.cpp +12 -4
  19. package/src/duckdb/extension/json/json_scan.cpp +2 -2
  20. package/src/duckdb/extension/json/json_serializer.cpp +217 -0
  21. package/src/duckdb/extension/parquet/column_reader.cpp +94 -15
  22. package/src/duckdb/extension/parquet/column_writer.cpp +0 -1
  23. package/src/duckdb/extension/parquet/include/column_reader.hpp +1 -2
  24. package/src/duckdb/extension/parquet/include/decode_utils.hpp +5 -4
  25. package/src/duckdb/extension/parquet/include/generated_column_reader.hpp +1 -11
  26. package/src/duckdb/extension/parquet/include/parquet_timestamp.hpp +2 -1
  27. package/src/duckdb/extension/parquet/parquet-extension.cpp +14 -3
  28. package/src/duckdb/extension/parquet/parquet_reader.cpp +6 -1
  29. package/src/duckdb/extension/parquet/parquet_statistics.cpp +49 -36
  30. package/src/duckdb/extension/parquet/parquet_timestamp.cpp +16 -6
  31. package/src/duckdb/src/catalog/catalog.cpp +34 -5
  32. package/src/duckdb/src/catalog/catalog_entry/duck_schema_entry.cpp +4 -0
  33. package/src/duckdb/src/catalog/catalog_entry/duck_table_entry.cpp +2 -21
  34. package/src/duckdb/src/catalog/catalog_entry/scalar_function_catalog_entry.cpp +7 -6
  35. package/src/duckdb/src/catalog/catalog_entry/table_catalog_entry.cpp +3 -3
  36. package/src/duckdb/src/catalog/catalog_entry/table_function_catalog_entry.cpp +20 -1
  37. package/src/duckdb/src/catalog/catalog_entry/type_catalog_entry.cpp +8 -2
  38. package/src/duckdb/src/catalog/catalog_set.cpp +1 -0
  39. package/src/duckdb/src/catalog/default/default_functions.cpp +3 -0
  40. package/src/duckdb/src/catalog/dependency_list.cpp +12 -0
  41. package/src/duckdb/src/catalog/duck_catalog.cpp +34 -7
  42. package/src/duckdb/src/common/arrow/arrow_appender.cpp +48 -4
  43. package/src/duckdb/src/common/arrow/arrow_converter.cpp +1 -1
  44. package/src/duckdb/src/common/box_renderer.cpp +109 -23
  45. package/src/duckdb/src/common/enums/expression_type.cpp +8 -222
  46. package/src/duckdb/src/common/enums/join_type.cpp +3 -22
  47. package/src/duckdb/src/common/enums/logical_operator_type.cpp +2 -0
  48. package/src/duckdb/src/common/enums/statement_type.cpp +2 -0
  49. package/src/duckdb/src/common/exception.cpp +15 -1
  50. package/src/duckdb/src/common/field_writer.cpp +1 -0
  51. package/src/duckdb/src/common/hive_partitioning.cpp +3 -1
  52. package/src/duckdb/src/common/local_file_system.cpp +64 -7
  53. package/src/duckdb/src/common/operator/cast_operators.cpp +1 -1
  54. package/src/duckdb/src/common/preserved_error.cpp +7 -5
  55. package/src/duckdb/src/common/progress_bar/progress_bar.cpp +7 -0
  56. package/src/duckdb/src/common/serializer/buffered_deserializer.cpp +4 -0
  57. package/src/duckdb/src/common/serializer/buffered_file_reader.cpp +15 -2
  58. package/src/duckdb/src/common/serializer/enum_serializer.cpp +1176 -0
  59. package/src/duckdb/src/common/sort/comparators.cpp +14 -5
  60. package/src/duckdb/src/common/sort/sort_state.cpp +5 -7
  61. package/src/duckdb/src/common/sort/sorted_block.cpp +0 -1
  62. package/src/duckdb/src/common/string_util.cpp +18 -1
  63. package/src/duckdb/src/common/types/bit.cpp +166 -87
  64. package/src/duckdb/src/common/types/blob.cpp +1 -1
  65. package/src/duckdb/src/common/types/chunk_collection.cpp +2 -2
  66. package/src/duckdb/src/common/types/column_data_collection.cpp +39 -2
  67. package/src/duckdb/src/common/types/column_data_collection_segment.cpp +12 -10
  68. package/src/duckdb/src/common/types/data_chunk.cpp +1 -1
  69. package/src/duckdb/src/common/types/interval.cpp +0 -41
  70. package/src/duckdb/src/common/types/list_segment.cpp +658 -0
  71. package/src/duckdb/src/common/types/string_heap.cpp +1 -1
  72. package/src/duckdb/src/common/types/string_type.cpp +1 -1
  73. package/src/duckdb/src/common/types/time.cpp +13 -0
  74. package/src/duckdb/src/common/types/validity_mask.cpp +24 -7
  75. package/src/duckdb/src/common/types/value.cpp +320 -154
  76. package/src/duckdb/src/common/types/vector.cpp +158 -134
  77. package/src/duckdb/src/common/types.cpp +313 -153
  78. package/src/duckdb/src/common/value_operations/comparison_operations.cpp +14 -22
  79. package/src/duckdb/src/common/vector_operations/comparison_operators.cpp +10 -10
  80. package/src/duckdb/src/common/vector_operations/is_distinct_from.cpp +11 -10
  81. package/src/duckdb/src/common/vector_operations/vector_cast.cpp +2 -1
  82. package/src/duckdb/src/execution/aggregate_hashtable.cpp +98 -74
  83. package/src/duckdb/src/execution/column_binding_resolver.cpp +21 -5
  84. package/src/duckdb/src/execution/expression_executor/execute_cast.cpp +2 -1
  85. package/src/duckdb/src/execution/expression_executor/execute_comparison.cpp +2 -2
  86. package/src/duckdb/src/execution/index/art/art.cpp +19 -5
  87. package/src/duckdb/src/execution/join_hashtable.cpp +3 -1
  88. package/src/duckdb/src/execution/operator/aggregate/physical_hash_aggregate.cpp +1 -1
  89. package/src/duckdb/src/execution/operator/aggregate/physical_perfecthash_aggregate.cpp +4 -5
  90. package/src/duckdb/src/execution/operator/aggregate/physical_window.cpp +117 -26
  91. package/src/duckdb/src/execution/operator/helper/physical_limit.cpp +3 -0
  92. package/src/duckdb/src/execution/operator/helper/physical_vacuum.cpp +5 -3
  93. package/src/duckdb/src/execution/operator/join/physical_blockwise_nl_join.cpp +64 -17
  94. package/src/duckdb/src/execution/operator/join/physical_hash_join.cpp +2 -0
  95. package/src/duckdb/src/execution/operator/join/physical_iejoin.cpp +2 -2
  96. package/src/duckdb/src/execution/operator/join/physical_index_join.cpp +13 -4
  97. package/src/duckdb/src/execution/operator/join/physical_join.cpp +0 -3
  98. package/src/duckdb/src/execution/operator/join/physical_piecewise_merge_join.cpp +6 -11
  99. package/src/duckdb/src/execution/operator/join/physical_range_join.cpp +3 -1
  100. package/src/duckdb/src/execution/operator/persistent/base_csv_reader.cpp +11 -4
  101. package/src/duckdb/src/execution/operator/persistent/buffered_csv_reader.cpp +24 -19
  102. package/src/duckdb/src/execution/operator/persistent/csv_reader_options.cpp +3 -0
  103. package/src/duckdb/src/execution/operator/persistent/physical_batch_insert.cpp +2 -1
  104. package/src/duckdb/src/execution/operator/persistent/physical_copy_to_file.cpp +2 -2
  105. package/src/duckdb/src/execution/operator/persistent/physical_delete.cpp +1 -3
  106. package/src/duckdb/src/execution/operator/persistent/physical_insert.cpp +1 -0
  107. package/src/duckdb/src/execution/operator/projection/physical_projection.cpp +34 -0
  108. package/src/duckdb/src/execution/operator/scan/physical_positional_scan.cpp +20 -5
  109. package/src/duckdb/src/execution/operator/schema/physical_create_type.cpp +20 -40
  110. package/src/duckdb/src/execution/operator/set/physical_recursive_cte.cpp +2 -5
  111. package/src/duckdb/src/execution/partitionable_hashtable.cpp +20 -5
  112. package/src/duckdb/src/execution/physical_plan/plan_aggregate.cpp +22 -16
  113. package/src/duckdb/src/execution/physical_plan/plan_asof_join.cpp +97 -0
  114. package/src/duckdb/src/execution/physical_plan/plan_comparison_join.cpp +95 -47
  115. package/src/duckdb/src/execution/physical_plan/plan_create_index.cpp +2 -1
  116. package/src/duckdb/src/execution/physical_plan/plan_distinct.cpp +5 -8
  117. package/src/duckdb/src/execution/physical_plan/plan_positional_join.cpp +14 -5
  118. package/src/duckdb/src/execution/physical_plan_generator.cpp +3 -0
  119. package/src/duckdb/src/execution/radix_partitioned_hashtable.cpp +23 -15
  120. package/src/duckdb/src/execution/window_segment_tree.cpp +173 -1
  121. package/src/duckdb/src/function/aggregate/algebraic/avg.cpp +0 -6
  122. package/src/duckdb/src/function/aggregate/distributive/bitagg.cpp +99 -95
  123. package/src/duckdb/src/function/aggregate/distributive/bitstring_agg.cpp +269 -0
  124. package/src/duckdb/src/function/aggregate/distributive/bool.cpp +2 -0
  125. package/src/duckdb/src/function/aggregate/distributive/count.cpp +3 -4
  126. package/src/duckdb/src/function/aggregate/distributive/first.cpp +1 -0
  127. package/src/duckdb/src/function/aggregate/distributive/minmax.cpp +2 -0
  128. package/src/duckdb/src/function/aggregate/distributive/sum.cpp +19 -16
  129. package/src/duckdb/src/function/aggregate/distributive_functions.cpp +1 -0
  130. package/src/duckdb/src/function/aggregate/holistic/approximate_quantile.cpp +5 -2
  131. package/src/duckdb/src/function/aggregate/holistic/mode.cpp +1 -1
  132. package/src/duckdb/src/function/aggregate/holistic/quantile.cpp +16 -1
  133. package/src/duckdb/src/function/aggregate/nested/list.cpp +6 -712
  134. package/src/duckdb/src/function/aggregate/sorted_aggregate_function.cpp +138 -45
  135. package/src/duckdb/src/function/cast/bit_cast.cpp +0 -2
  136. package/src/duckdb/src/function/cast/blob_cast.cpp +0 -1
  137. package/src/duckdb/src/function/cast/cast_function_set.cpp +1 -1
  138. package/src/duckdb/src/function/cast/enum_casts.cpp +25 -3
  139. package/src/duckdb/src/function/cast/list_casts.cpp +17 -4
  140. package/src/duckdb/src/function/cast/map_cast.cpp +5 -2
  141. package/src/duckdb/src/function/cast/string_cast.cpp +36 -10
  142. package/src/duckdb/src/function/cast/struct_cast.cpp +24 -4
  143. package/src/duckdb/src/function/cast/time_casts.cpp +2 -2
  144. package/src/duckdb/src/function/cast/union_casts.cpp +33 -7
  145. package/src/duckdb/src/function/cast_rules.cpp +9 -4
  146. package/src/duckdb/src/function/function_binder.cpp +1 -8
  147. package/src/duckdb/src/function/pragma/pragma_queries.cpp +24 -1
  148. package/src/duckdb/src/function/scalar/bit/bitstring.cpp +100 -0
  149. package/src/duckdb/src/function/scalar/date/current.cpp +0 -2
  150. package/src/duckdb/src/function/scalar/date/date_diff.cpp +0 -1
  151. package/src/duckdb/src/function/scalar/date/date_part.cpp +18 -26
  152. package/src/duckdb/src/function/scalar/date/date_sub.cpp +0 -1
  153. package/src/duckdb/src/function/scalar/date/date_trunc.cpp +10 -14
  154. package/src/duckdb/src/function/scalar/generic/stats.cpp +2 -4
  155. package/src/duckdb/src/function/scalar/list/contains_or_position.cpp +4 -146
  156. package/src/duckdb/src/function/scalar/list/flatten.cpp +5 -12
  157. package/src/duckdb/src/function/scalar/list/list_aggregates.cpp +1 -1
  158. package/src/duckdb/src/function/scalar/list/list_concat.cpp +8 -12
  159. package/src/duckdb/src/function/scalar/list/list_extract.cpp +5 -12
  160. package/src/duckdb/src/function/scalar/list/list_lambdas.cpp +7 -3
  161. package/src/duckdb/src/function/scalar/list/list_sort.cpp +25 -18
  162. package/src/duckdb/src/function/scalar/list/list_value.cpp +6 -10
  163. package/src/duckdb/src/function/scalar/map/map.cpp +47 -1
  164. package/src/duckdb/src/function/scalar/map/map_entries.cpp +61 -0
  165. package/src/duckdb/src/function/scalar/map/map_extract.cpp +68 -26
  166. package/src/duckdb/src/function/scalar/map/map_keys_values.cpp +97 -0
  167. package/src/duckdb/src/function/scalar/math/numeric.cpp +101 -17
  168. package/src/duckdb/src/function/scalar/math_functions.cpp +3 -0
  169. package/src/duckdb/src/function/scalar/nested_functions.cpp +3 -0
  170. package/src/duckdb/src/function/scalar/operators/add.cpp +0 -9
  171. package/src/duckdb/src/function/scalar/operators/arithmetic.cpp +29 -48
  172. package/src/duckdb/src/function/scalar/operators/bitwise.cpp +0 -63
  173. package/src/duckdb/src/function/scalar/operators/multiply.cpp +5 -6
  174. package/src/duckdb/src/function/scalar/operators/subtract.cpp +0 -6
  175. package/src/duckdb/src/function/scalar/string/caseconvert.cpp +2 -6
  176. package/src/duckdb/src/function/scalar/string/hex.cpp +201 -0
  177. package/src/duckdb/src/function/scalar/string/instr.cpp +2 -6
  178. package/src/duckdb/src/function/scalar/string/length.cpp +2 -6
  179. package/src/duckdb/src/function/scalar/string/like.cpp +2 -6
  180. package/src/duckdb/src/function/scalar/string/regexp/regexp_extract_all.cpp +243 -0
  181. package/src/duckdb/src/function/scalar/string/regexp/regexp_util.cpp +79 -0
  182. package/src/duckdb/src/function/scalar/string/regexp.cpp +21 -80
  183. package/src/duckdb/src/function/scalar/string/substring.cpp +2 -6
  184. package/src/duckdb/src/function/scalar/string_functions.cpp +2 -0
  185. package/src/duckdb/src/function/scalar/struct/struct_extract.cpp +5 -10
  186. package/src/duckdb/src/function/scalar/struct/struct_insert.cpp +11 -14
  187. package/src/duckdb/src/function/scalar/struct/struct_pack.cpp +6 -7
  188. package/src/duckdb/src/function/table/arrow.cpp +5 -2
  189. package/src/duckdb/src/function/table/arrow_conversion.cpp +25 -1
  190. package/src/duckdb/src/function/table/checkpoint.cpp +5 -1
  191. package/src/duckdb/src/function/table/read_csv.cpp +60 -0
  192. package/src/duckdb/src/function/table/system/duckdb_constraints.cpp +2 -2
  193. package/src/duckdb/src/function/table/system/test_all_types.cpp +2 -2
  194. package/src/duckdb/src/function/table/table_scan.cpp +9 -12
  195. package/src/duckdb/src/function/table/version/pragma_version.cpp +2 -2
  196. package/src/duckdb/src/function/table_function.cpp +30 -11
  197. package/src/duckdb/src/include/duckdb/catalog/catalog.hpp +6 -0
  198. package/src/duckdb/src/include/duckdb/catalog/catalog_entry/duck_table_entry.hpp +1 -1
  199. package/src/duckdb/src/include/duckdb/catalog/catalog_entry/table_function_catalog_entry.hpp +6 -8
  200. package/src/duckdb/src/include/duckdb/catalog/dependency_list.hpp +3 -0
  201. package/src/duckdb/src/include/duckdb/catalog/duck_catalog.hpp +2 -1
  202. package/src/duckdb/src/include/duckdb/common/box_renderer.hpp +8 -2
  203. package/src/duckdb/src/include/duckdb/common/constants.hpp +0 -19
  204. package/src/duckdb/src/include/duckdb/common/enums/aggregate_handling.hpp +2 -0
  205. package/src/duckdb/src/include/duckdb/common/enums/expression_type.hpp +2 -3
  206. package/src/duckdb/src/include/duckdb/common/enums/joinref_type.hpp +7 -4
  207. package/src/duckdb/src/include/duckdb/common/enums/logical_operator_type.hpp +1 -0
  208. package/src/duckdb/src/include/duckdb/common/enums/order_type.hpp +2 -0
  209. package/src/duckdb/src/include/duckdb/common/enums/set_operation_type.hpp +2 -1
  210. package/src/duckdb/src/include/duckdb/common/enums/statement_type.hpp +2 -1
  211. package/src/duckdb/src/include/duckdb/common/enums/tableref_type.hpp +2 -1
  212. package/src/duckdb/src/include/duckdb/common/exception.hpp +69 -2
  213. package/src/duckdb/src/include/duckdb/common/field_writer.hpp +12 -4
  214. package/src/duckdb/src/include/duckdb/common/helper.hpp +1 -1
  215. package/src/duckdb/src/include/duckdb/common/{http_stats.hpp → http_state.hpp} +18 -4
  216. package/src/duckdb/src/include/duckdb/common/operator/comparison_operators.hpp +45 -149
  217. package/src/duckdb/src/include/duckdb/common/operator/multiply.hpp +2 -0
  218. package/src/duckdb/src/include/duckdb/common/optional_ptr.hpp +45 -0
  219. package/src/duckdb/src/include/duckdb/common/preserved_error.hpp +6 -1
  220. package/src/duckdb/src/include/duckdb/common/progress_bar/progress_bar.hpp +2 -0
  221. package/src/duckdb/src/include/duckdb/common/serializer/buffered_deserializer.hpp +4 -2
  222. package/src/duckdb/src/include/duckdb/common/serializer/buffered_file_reader.hpp +8 -2
  223. package/src/duckdb/src/include/duckdb/common/serializer/enum_serializer.hpp +113 -0
  224. package/src/duckdb/src/include/duckdb/common/serializer/format_deserializer.hpp +336 -0
  225. package/src/duckdb/src/include/duckdb/common/serializer/format_serializer.hpp +268 -0
  226. package/src/duckdb/src/include/duckdb/common/serializer/serialization_traits.hpp +126 -0
  227. package/src/duckdb/src/include/duckdb/common/serializer.hpp +13 -0
  228. package/src/duckdb/src/include/duckdb/common/string_util.hpp +27 -0
  229. package/src/duckdb/src/include/duckdb/common/types/bit.hpp +12 -7
  230. package/src/duckdb/src/include/duckdb/common/types/interval.hpp +39 -3
  231. package/src/duckdb/src/include/duckdb/common/types/list_segment.hpp +70 -0
  232. package/src/duckdb/src/include/duckdb/common/types/string_type.hpp +73 -3
  233. package/src/duckdb/src/include/duckdb/common/types/time.hpp +3 -0
  234. package/src/duckdb/src/include/duckdb/common/types/validity_mask.hpp +4 -1
  235. package/src/duckdb/src/include/duckdb/common/types/value.hpp +17 -48
  236. package/src/duckdb/src/include/duckdb/common/types/value_map.hpp +1 -1
  237. package/src/duckdb/src/include/duckdb/common/types/vector.hpp +3 -1
  238. package/src/duckdb/src/include/duckdb/common/types.hpp +45 -8
  239. package/src/duckdb/src/include/duckdb/common/vector_operations/unary_executor.hpp +2 -2
  240. package/src/duckdb/src/include/duckdb/execution/aggregate_hashtable.hpp +35 -20
  241. package/src/duckdb/src/include/duckdb/execution/index/art/art.hpp +3 -14
  242. package/src/duckdb/src/include/duckdb/execution/operator/aggregate/physical_perfecthash_aggregate.hpp +1 -1
  243. package/src/duckdb/src/include/duckdb/execution/operator/join/physical_cross_product.hpp +2 -0
  244. package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_file_handle.hpp +1 -0
  245. package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_reader_options.hpp +10 -0
  246. package/src/duckdb/src/include/duckdb/execution/operator/projection/physical_projection.hpp +5 -0
  247. package/src/duckdb/src/include/duckdb/execution/partitionable_hashtable.hpp +5 -1
  248. package/src/duckdb/src/include/duckdb/execution/physical_plan_generator.hpp +1 -3
  249. package/src/duckdb/src/include/duckdb/execution/window_segment_tree.hpp +54 -0
  250. package/src/duckdb/src/include/duckdb/function/aggregate/distributive_functions.hpp +5 -0
  251. package/src/duckdb/src/include/duckdb/function/aggregate_function.hpp +18 -6
  252. package/src/duckdb/src/include/duckdb/function/cast/bound_cast_data.hpp +84 -0
  253. package/src/duckdb/src/include/duckdb/function/cast/cast_function_set.hpp +2 -2
  254. package/src/duckdb/src/include/duckdb/function/cast/default_casts.hpp +28 -64
  255. package/src/duckdb/src/include/duckdb/function/function_binder.hpp +3 -6
  256. package/src/duckdb/src/include/duckdb/function/scalar/bit_functions.hpp +4 -0
  257. package/src/duckdb/src/include/duckdb/function/scalar/list/contains_or_position.hpp +138 -0
  258. package/src/duckdb/src/include/duckdb/function/scalar/math_functions.hpp +8 -0
  259. package/src/duckdb/src/include/duckdb/function/scalar/nested_functions.hpp +59 -0
  260. package/src/duckdb/src/include/duckdb/function/scalar/regexp.hpp +81 -1
  261. package/src/duckdb/src/include/duckdb/function/scalar/string_functions.hpp +4 -0
  262. package/src/duckdb/src/include/duckdb/function/scalar_function.hpp +2 -2
  263. package/src/duckdb/src/include/duckdb/function/table/arrow.hpp +12 -1
  264. package/src/duckdb/src/include/duckdb/function/table_function.hpp +10 -0
  265. package/src/duckdb/src/include/duckdb/main/capi/capi_internal.hpp +2 -0
  266. package/src/duckdb/src/include/duckdb/main/client_config.hpp +2 -0
  267. package/src/duckdb/src/include/duckdb/main/client_data.hpp +3 -3
  268. package/src/duckdb/src/include/duckdb/main/config.hpp +3 -0
  269. package/src/duckdb/src/include/duckdb/main/connection_manager.hpp +2 -0
  270. package/src/duckdb/src/include/duckdb/main/database.hpp +1 -0
  271. package/src/duckdb/src/include/duckdb/main/extension_entries.hpp +2 -0
  272. package/src/duckdb/src/include/duckdb/main/prepared_statement.hpp +2 -0
  273. package/src/duckdb/src/include/duckdb/main/relation/explain_relation.hpp +2 -1
  274. package/src/duckdb/src/include/duckdb/main/relation.hpp +2 -1
  275. package/src/duckdb/src/include/duckdb/optimizer/filter_pushdown.hpp +2 -0
  276. package/src/duckdb/src/include/duckdb/optimizer/join_order/cardinality_estimator.hpp +2 -2
  277. package/src/duckdb/src/include/duckdb/optimizer/rule/list.hpp +1 -0
  278. package/src/duckdb/src/include/duckdb/optimizer/rule/ordered_aggregate_optimizer.hpp +24 -0
  279. package/src/duckdb/src/include/duckdb/parser/common_table_expression_info.hpp +4 -0
  280. package/src/duckdb/src/include/duckdb/parser/expression/between_expression.hpp +3 -0
  281. package/src/duckdb/src/include/duckdb/parser/expression/bound_expression.hpp +2 -0
  282. package/src/duckdb/src/include/duckdb/parser/expression/case_expression.hpp +5 -0
  283. package/src/duckdb/src/include/duckdb/parser/expression/cast_expression.hpp +2 -0
  284. package/src/duckdb/src/include/duckdb/parser/expression/collate_expression.hpp +2 -0
  285. package/src/duckdb/src/include/duckdb/parser/expression/columnref_expression.hpp +2 -0
  286. package/src/duckdb/src/include/duckdb/parser/expression/comparison_expression.hpp +2 -0
  287. package/src/duckdb/src/include/duckdb/parser/expression/conjunction_expression.hpp +2 -0
  288. package/src/duckdb/src/include/duckdb/parser/expression/constant_expression.hpp +3 -0
  289. package/src/duckdb/src/include/duckdb/parser/expression/default_expression.hpp +1 -0
  290. package/src/duckdb/src/include/duckdb/parser/expression/function_expression.hpp +4 -2
  291. package/src/duckdb/src/include/duckdb/parser/expression/lambda_expression.hpp +2 -0
  292. package/src/duckdb/src/include/duckdb/parser/expression/operator_expression.hpp +2 -0
  293. package/src/duckdb/src/include/duckdb/parser/expression/parameter_expression.hpp +2 -0
  294. package/src/duckdb/src/include/duckdb/parser/expression/positional_reference_expression.hpp +2 -0
  295. package/src/duckdb/src/include/duckdb/parser/expression/star_expression.hpp +4 -2
  296. package/src/duckdb/src/include/duckdb/parser/expression/subquery_expression.hpp +2 -0
  297. package/src/duckdb/src/include/duckdb/parser/expression/window_expression.hpp +5 -0
  298. package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_info.hpp +5 -1
  299. package/src/duckdb/src/include/duckdb/parser/parsed_data/{alter_function_info.hpp → alter_scalar_function_info.hpp} +13 -13
  300. package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_function_info.hpp +47 -0
  301. package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_info.hpp +6 -0
  302. package/src/duckdb/src/include/duckdb/parser/parsed_data/create_table_function_info.hpp +2 -1
  303. package/src/duckdb/src/include/duckdb/parser/parsed_data/sample_options.hpp +2 -0
  304. package/src/duckdb/src/include/duckdb/parser/parsed_expression.hpp +5 -0
  305. package/src/duckdb/src/include/duckdb/parser/query_node/recursive_cte_node.hpp +3 -0
  306. package/src/duckdb/src/include/duckdb/parser/query_node/select_node.hpp +5 -0
  307. package/src/duckdb/src/include/duckdb/parser/query_node/set_operation_node.hpp +3 -0
  308. package/src/duckdb/src/include/duckdb/parser/query_node.hpp +13 -2
  309. package/src/duckdb/src/include/duckdb/parser/result_modifier.hpp +24 -1
  310. package/src/duckdb/src/include/duckdb/parser/sql_statement.hpp +2 -1
  311. package/src/duckdb/src/include/duckdb/parser/statement/multi_statement.hpp +28 -0
  312. package/src/duckdb/src/include/duckdb/parser/statement/select_statement.hpp +6 -1
  313. package/src/duckdb/src/include/duckdb/parser/tableref/basetableref.hpp +4 -0
  314. package/src/duckdb/src/include/duckdb/parser/tableref/emptytableref.hpp +2 -0
  315. package/src/duckdb/src/include/duckdb/parser/tableref/expressionlistref.hpp +3 -0
  316. package/src/duckdb/src/include/duckdb/parser/tableref/joinref.hpp +3 -0
  317. package/src/duckdb/src/include/duckdb/parser/tableref/list.hpp +1 -0
  318. package/src/duckdb/src/include/duckdb/parser/tableref/pivotref.hpp +87 -0
  319. package/src/duckdb/src/include/duckdb/parser/tableref/subqueryref.hpp +3 -0
  320. package/src/duckdb/src/include/duckdb/parser/tableref/table_function_ref.hpp +3 -0
  321. package/src/duckdb/src/include/duckdb/parser/tableref.hpp +3 -1
  322. package/src/duckdb/src/include/duckdb/parser/tokens.hpp +2 -0
  323. package/src/duckdb/src/include/duckdb/parser/transformer.hpp +33 -0
  324. package/src/duckdb/src/include/duckdb/planner/bind_context.hpp +2 -0
  325. package/src/duckdb/src/include/duckdb/planner/binder.hpp +15 -4
  326. package/src/duckdb/src/include/duckdb/planner/bound_result_modifier.hpp +3 -0
  327. package/src/duckdb/src/include/duckdb/planner/expression/bound_aggregate_expression.hpp +3 -0
  328. package/src/duckdb/src/include/duckdb/planner/expression_binder/base_select_binder.hpp +64 -0
  329. package/src/duckdb/src/include/duckdb/planner/expression_binder/having_binder.hpp +2 -2
  330. package/src/duckdb/src/include/duckdb/planner/expression_binder/order_binder.hpp +4 -1
  331. package/src/duckdb/src/include/duckdb/planner/expression_binder/qualify_binder.hpp +2 -2
  332. package/src/duckdb/src/include/duckdb/planner/expression_binder/select_binder.hpp +9 -38
  333. package/src/duckdb/src/include/duckdb/planner/expression_binder.hpp +1 -1
  334. package/src/duckdb/src/include/duckdb/planner/logical_tokens.hpp +1 -0
  335. package/src/duckdb/src/include/duckdb/planner/operator/list.hpp +1 -0
  336. package/src/duckdb/src/include/duckdb/planner/operator/logical_asof_join.hpp +22 -0
  337. package/src/duckdb/src/include/duckdb/planner/operator/logical_comparison_join.hpp +5 -2
  338. package/src/duckdb/src/include/duckdb/planner/operator/logical_distinct.hpp +3 -0
  339. package/src/duckdb/src/include/duckdb/planner/query_node/bound_select_node.hpp +8 -2
  340. package/src/duckdb/src/include/duckdb/storage/buffer/block_handle.hpp +2 -0
  341. package/src/duckdb/src/include/duckdb/storage/buffer_manager.hpp +76 -44
  342. package/src/duckdb/src/include/duckdb/storage/checkpoint/table_data_writer.hpp +3 -2
  343. package/src/duckdb/src/include/duckdb/storage/checkpoint_manager.hpp +1 -1
  344. package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_compress.hpp +2 -2
  345. package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_fetch.hpp +1 -1
  346. package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_scan.hpp +2 -1
  347. package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_compress.hpp +2 -2
  348. package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_fetch.hpp +1 -1
  349. package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_scan.hpp +2 -1
  350. package/src/duckdb/src/include/duckdb/storage/data_pointer.hpp +4 -3
  351. package/src/duckdb/src/include/duckdb/storage/data_table.hpp +4 -3
  352. package/src/duckdb/src/include/duckdb/storage/index.hpp +5 -4
  353. package/src/duckdb/src/include/duckdb/storage/meta_block_reader.hpp +7 -0
  354. package/src/duckdb/src/include/duckdb/storage/statistics/base_statistics.hpp +93 -29
  355. package/src/duckdb/src/include/duckdb/storage/statistics/column_statistics.hpp +22 -3
  356. package/src/duckdb/src/include/duckdb/storage/statistics/distinct_statistics.hpp +8 -6
  357. package/src/duckdb/src/include/duckdb/storage/statistics/list_stats.hpp +41 -0
  358. package/src/duckdb/src/include/duckdb/storage/statistics/node_statistics.hpp +26 -0
  359. package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats.hpp +114 -0
  360. package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats_union.hpp +62 -0
  361. package/src/duckdb/src/include/duckdb/storage/statistics/segment_statistics.hpp +2 -7
  362. package/src/duckdb/src/include/duckdb/storage/statistics/string_stats.hpp +74 -0
  363. package/src/duckdb/src/include/duckdb/storage/statistics/struct_stats.hpp +42 -0
  364. package/src/duckdb/src/include/duckdb/storage/string_uncompressed.hpp +2 -3
  365. package/src/duckdb/src/include/duckdb/storage/table/column_checkpoint_state.hpp +2 -1
  366. package/src/duckdb/src/include/duckdb/storage/table/column_data.hpp +21 -7
  367. package/src/duckdb/src/include/duckdb/storage/table/column_data_checkpointer.hpp +3 -2
  368. package/src/duckdb/src/include/duckdb/storage/table/column_segment.hpp +5 -6
  369. package/src/duckdb/src/include/duckdb/storage/table/column_segment_tree.hpp +18 -0
  370. package/src/duckdb/src/include/duckdb/storage/table/list_column_data.hpp +1 -1
  371. package/src/duckdb/src/include/duckdb/storage/table/persistent_table_data.hpp +6 -3
  372. package/src/duckdb/src/include/duckdb/storage/table/row_group.hpp +41 -45
  373. package/src/duckdb/src/include/duckdb/storage/table/row_group_collection.hpp +23 -7
  374. package/src/duckdb/src/include/duckdb/storage/table/row_group_segment_tree.hpp +35 -0
  375. package/src/duckdb/src/include/duckdb/storage/table/scan_state.hpp +21 -29
  376. package/src/duckdb/src/include/duckdb/storage/table/segment_base.hpp +6 -6
  377. package/src/duckdb/src/include/duckdb/storage/table/segment_tree.hpp +281 -26
  378. package/src/duckdb/src/include/duckdb/storage/table/standard_column_data.hpp +0 -4
  379. package/src/duckdb/src/include/duckdb/storage/table/table_statistics.hpp +5 -0
  380. package/src/duckdb/src/include/duckdb/storage/table/update_segment.hpp +0 -1
  381. package/src/duckdb/src/include/duckdb/storage/write_ahead_log.hpp +1 -1
  382. package/src/duckdb/src/include/duckdb/transaction/local_storage.hpp +6 -3
  383. package/src/duckdb/src/include/duckdb.h +71 -2
  384. package/src/duckdb/src/include/duckdb.hpp +0 -1
  385. package/src/duckdb/src/main/capi/pending-c.cpp +16 -3
  386. package/src/duckdb/src/main/capi/result-c.cpp +27 -1
  387. package/src/duckdb/src/main/capi/stream-c.cpp +25 -0
  388. package/src/duckdb/src/main/capi/table_function-c.cpp +23 -0
  389. package/src/duckdb/src/main/client_context.cpp +38 -34
  390. package/src/duckdb/src/main/client_data.cpp +7 -6
  391. package/src/duckdb/src/main/config.cpp +70 -1
  392. package/src/duckdb/src/main/database.cpp +19 -2
  393. package/src/duckdb/src/main/extension/extension_install.cpp +7 -2
  394. package/src/duckdb/src/main/prepared_statement.cpp +4 -0
  395. package/src/duckdb/src/main/query_profiler.cpp +17 -15
  396. package/src/duckdb/src/main/relation/explain_relation.cpp +3 -3
  397. package/src/duckdb/src/main/relation.cpp +3 -2
  398. package/src/duckdb/src/main/settings/settings.cpp +20 -8
  399. package/src/duckdb/src/optimizer/column_lifetime_analyzer.cpp +1 -0
  400. package/src/duckdb/src/optimizer/deliminator.cpp +1 -1
  401. package/src/duckdb/src/optimizer/filter_combiner.cpp +3 -6
  402. package/src/duckdb/src/optimizer/filter_pullup.cpp +3 -1
  403. package/src/duckdb/src/optimizer/filter_pushdown.cpp +14 -8
  404. package/src/duckdb/src/optimizer/join_order/cardinality_estimator.cpp +107 -71
  405. package/src/duckdb/src/optimizer/join_order/join_order_optimizer.cpp +32 -12
  406. package/src/duckdb/src/optimizer/optimizer.cpp +1 -0
  407. package/src/duckdb/src/optimizer/pullup/pullup_from_left.cpp +2 -2
  408. package/src/duckdb/src/optimizer/pushdown/pushdown_aggregate.cpp +33 -5
  409. package/src/duckdb/src/optimizer/pushdown/pushdown_cross_product.cpp +1 -1
  410. package/src/duckdb/src/optimizer/pushdown/pushdown_inner_join.cpp +3 -0
  411. package/src/duckdb/src/optimizer/pushdown/pushdown_left_join.cpp +5 -12
  412. package/src/duckdb/src/optimizer/pushdown/pushdown_mark_join.cpp +2 -2
  413. package/src/duckdb/src/optimizer/pushdown/pushdown_single_join.cpp +1 -1
  414. package/src/duckdb/src/optimizer/remove_unused_columns.cpp +1 -0
  415. package/src/duckdb/src/optimizer/rule/move_constants.cpp +10 -4
  416. package/src/duckdb/src/optimizer/rule/ordered_aggregate_optimizer.cpp +30 -0
  417. package/src/duckdb/src/optimizer/rule/regex_optimizations.cpp +9 -2
  418. package/src/duckdb/src/optimizer/statistics/expression/propagate_aggregate.cpp +9 -3
  419. package/src/duckdb/src/optimizer/statistics/expression/propagate_and_compress.cpp +6 -7
  420. package/src/duckdb/src/optimizer/statistics/expression/propagate_cast.cpp +14 -11
  421. package/src/duckdb/src/optimizer/statistics/expression/propagate_columnref.cpp +1 -1
  422. package/src/duckdb/src/optimizer/statistics/expression/propagate_comparison.cpp +13 -15
  423. package/src/duckdb/src/optimizer/statistics/expression/propagate_conjunction.cpp +0 -1
  424. package/src/duckdb/src/optimizer/statistics/expression/propagate_constant.cpp +3 -75
  425. package/src/duckdb/src/optimizer/statistics/expression/propagate_function.cpp +7 -2
  426. package/src/duckdb/src/optimizer/statistics/expression/propagate_operator.cpp +10 -0
  427. package/src/duckdb/src/optimizer/statistics/operator/propagate_aggregate.cpp +2 -3
  428. package/src/duckdb/src/optimizer/statistics/operator/propagate_filter.cpp +29 -32
  429. package/src/duckdb/src/optimizer/statistics/operator/propagate_join.cpp +5 -5
  430. package/src/duckdb/src/optimizer/statistics/operator/propagate_set_operation.cpp +3 -3
  431. package/src/duckdb/src/optimizer/statistics_propagator.cpp +2 -1
  432. package/src/duckdb/src/optimizer/unnest_rewriter.cpp +2 -2
  433. package/src/duckdb/src/parallel/meta_pipeline.cpp +0 -7
  434. package/src/duckdb/src/parser/common_table_expression_info.cpp +19 -0
  435. package/src/duckdb/src/parser/expression/between_expression.cpp +17 -0
  436. package/src/duckdb/src/parser/expression/case_expression.cpp +28 -0
  437. package/src/duckdb/src/parser/expression/cast_expression.cpp +17 -0
  438. package/src/duckdb/src/parser/expression/collate_expression.cpp +16 -0
  439. package/src/duckdb/src/parser/expression/columnref_expression.cpp +15 -0
  440. package/src/duckdb/src/parser/expression/comparison_expression.cpp +16 -0
  441. package/src/duckdb/src/parser/expression/conjunction_expression.cpp +17 -0
  442. package/src/duckdb/src/parser/expression/constant_expression.cpp +14 -0
  443. package/src/duckdb/src/parser/expression/default_expression.cpp +7 -0
  444. package/src/duckdb/src/parser/expression/function_expression.cpp +35 -0
  445. package/src/duckdb/src/parser/expression/lambda_expression.cpp +16 -0
  446. package/src/duckdb/src/parser/expression/operator_expression.cpp +15 -0
  447. package/src/duckdb/src/parser/expression/parameter_expression.cpp +15 -0
  448. package/src/duckdb/src/parser/expression/positional_reference_expression.cpp +14 -0
  449. package/src/duckdb/src/parser/expression/star_expression.cpp +26 -6
  450. package/src/duckdb/src/parser/expression/subquery_expression.cpp +20 -0
  451. package/src/duckdb/src/parser/expression/window_expression.cpp +43 -0
  452. package/src/duckdb/src/parser/parsed_data/alter_info.cpp +7 -3
  453. package/src/duckdb/src/parser/parsed_data/alter_scalar_function_info.cpp +56 -0
  454. package/src/duckdb/src/parser/parsed_data/alter_table_function_info.cpp +51 -0
  455. package/src/duckdb/src/parser/parsed_data/create_scalar_function_info.cpp +3 -2
  456. package/src/duckdb/src/parser/parsed_data/create_table_function_info.cpp +6 -0
  457. package/src/duckdb/src/parser/parsed_data/sample_options.cpp +22 -10
  458. package/src/duckdb/src/parser/parsed_expression.cpp +72 -0
  459. package/src/duckdb/src/parser/parsed_expression_iterator.cpp +15 -1
  460. package/src/duckdb/src/parser/query_node/recursive_cte_node.cpp +21 -0
  461. package/src/duckdb/src/parser/query_node/select_node.cpp +31 -0
  462. package/src/duckdb/src/parser/query_node/set_operation_node.cpp +17 -0
  463. package/src/duckdb/src/parser/query_node.cpp +51 -1
  464. package/src/duckdb/src/parser/result_modifier.cpp +78 -0
  465. package/src/duckdb/src/parser/statement/multi_statement.cpp +18 -0
  466. package/src/duckdb/src/parser/statement/select_statement.cpp +12 -0
  467. package/src/duckdb/src/parser/tableref/basetableref.cpp +21 -0
  468. package/src/duckdb/src/parser/tableref/emptytableref.cpp +4 -0
  469. package/src/duckdb/src/parser/tableref/expressionlistref.cpp +17 -0
  470. package/src/duckdb/src/parser/tableref/joinref.cpp +29 -0
  471. package/src/duckdb/src/parser/tableref/pivotref.cpp +373 -0
  472. package/src/duckdb/src/parser/tableref/subqueryref.cpp +15 -0
  473. package/src/duckdb/src/parser/tableref/table_function.cpp +17 -0
  474. package/src/duckdb/src/parser/tableref.cpp +49 -0
  475. package/src/duckdb/src/parser/transform/expression/transform_array_access.cpp +11 -0
  476. package/src/duckdb/src/parser/transform/expression/transform_bool_expr.cpp +1 -1
  477. package/src/duckdb/src/parser/transform/expression/transform_columnref.cpp +17 -2
  478. package/src/duckdb/src/parser/transform/expression/transform_function.cpp +85 -42
  479. package/src/duckdb/src/parser/transform/expression/transform_operator.cpp +1 -1
  480. package/src/duckdb/src/parser/transform/expression/transform_subquery.cpp +1 -1
  481. package/src/duckdb/src/parser/transform/helpers/transform_alias.cpp +12 -6
  482. package/src/duckdb/src/parser/transform/helpers/transform_cte.cpp +24 -0
  483. package/src/duckdb/src/parser/transform/helpers/transform_groupby.cpp +7 -0
  484. package/src/duckdb/src/parser/transform/helpers/transform_orderby.cpp +0 -7
  485. package/src/duckdb/src/parser/transform/helpers/transform_typename.cpp +3 -2
  486. package/src/duckdb/src/parser/transform/statement/transform_create_function.cpp +4 -0
  487. package/src/duckdb/src/parser/transform/statement/transform_create_view.cpp +4 -0
  488. package/src/duckdb/src/parser/transform/statement/transform_pivot_stmt.cpp +179 -0
  489. package/src/duckdb/src/parser/transform/statement/transform_rename.cpp +3 -4
  490. package/src/duckdb/src/parser/transform/statement/transform_select.cpp +8 -0
  491. package/src/duckdb/src/parser/transform/statement/transform_select_node.cpp +2 -3
  492. package/src/duckdb/src/parser/transform/tableref/transform_join.cpp +12 -1
  493. package/src/duckdb/src/parser/transform/tableref/transform_pivot.cpp +121 -0
  494. package/src/duckdb/src/parser/transform/tableref/transform_tableref.cpp +2 -0
  495. package/src/duckdb/src/parser/transformer.cpp +15 -3
  496. package/src/duckdb/src/planner/bind_context.cpp +18 -25
  497. package/src/duckdb/src/planner/binder/expression/bind_aggregate_expression.cpp +9 -7
  498. package/src/duckdb/src/planner/binder/expression/bind_columnref_expression.cpp +4 -3
  499. package/src/duckdb/src/planner/binder/expression/bind_function_expression.cpp +23 -12
  500. package/src/duckdb/src/planner/binder/expression/bind_lambda.cpp +3 -2
  501. package/src/duckdb/src/planner/binder/expression/bind_star_expression.cpp +176 -0
  502. package/src/duckdb/src/planner/binder/expression/bind_subquery_expression.cpp +4 -0
  503. package/src/duckdb/src/planner/binder/expression/bind_unnest_expression.cpp +163 -24
  504. package/src/duckdb/src/planner/binder/expression/bind_window_expression.cpp +2 -2
  505. package/src/duckdb/src/planner/binder/query_node/bind_select_node.cpp +109 -94
  506. package/src/duckdb/src/planner/binder/query_node/plan_query_node.cpp +11 -0
  507. package/src/duckdb/src/planner/binder/query_node/plan_select_node.cpp +9 -4
  508. package/src/duckdb/src/planner/binder/statement/bind_copy.cpp +5 -3
  509. package/src/duckdb/src/planner/binder/statement/bind_create.cpp +3 -2
  510. package/src/duckdb/src/planner/binder/statement/bind_create_table.cpp +10 -1
  511. package/src/duckdb/src/planner/binder/statement/bind_delete.cpp +1 -1
  512. package/src/duckdb/src/planner/binder/statement/bind_insert.cpp +12 -8
  513. package/src/duckdb/src/planner/binder/statement/bind_logical_plan.cpp +17 -0
  514. package/src/duckdb/src/planner/binder/statement/bind_update.cpp +4 -2
  515. package/src/duckdb/src/planner/binder/tableref/bind_joinref.cpp +19 -3
  516. package/src/duckdb/src/planner/binder/tableref/bind_pivot.cpp +366 -0
  517. package/src/duckdb/src/planner/binder/tableref/bind_table_function.cpp +11 -1
  518. package/src/duckdb/src/planner/binder/tableref/plan_cteref.cpp +1 -0
  519. package/src/duckdb/src/planner/binder/tableref/plan_joinref.cpp +61 -13
  520. package/src/duckdb/src/planner/binder.cpp +19 -24
  521. package/src/duckdb/src/planner/bound_result_modifier.cpp +27 -1
  522. package/src/duckdb/src/planner/expression/bound_aggregate_expression.cpp +9 -2
  523. package/src/duckdb/src/planner/expression/bound_expression.cpp +4 -0
  524. package/src/duckdb/src/planner/expression/bound_window_expression.cpp +1 -1
  525. package/src/duckdb/src/planner/expression_binder/base_select_binder.cpp +146 -0
  526. package/src/duckdb/src/planner/expression_binder/having_binder.cpp +6 -3
  527. package/src/duckdb/src/planner/expression_binder/qualify_binder.cpp +3 -3
  528. package/src/duckdb/src/planner/expression_binder/select_binder.cpp +1 -132
  529. package/src/duckdb/src/planner/expression_binder.cpp +10 -3
  530. package/src/duckdb/src/planner/expression_iterator.cpp +17 -10
  531. package/src/duckdb/src/planner/filter/constant_filter.cpp +4 -6
  532. package/src/duckdb/src/planner/logical_operator.cpp +7 -2
  533. package/src/duckdb/src/planner/logical_operator_visitor.cpp +6 -0
  534. package/src/duckdb/src/planner/operator/logical_asof_join.cpp +8 -0
  535. package/src/duckdb/src/planner/operator/logical_distinct.cpp +3 -0
  536. package/src/duckdb/src/planner/planner.cpp +2 -1
  537. package/src/duckdb/src/planner/pragma_handler.cpp +10 -2
  538. package/src/duckdb/src/planner/subquery/flatten_dependent_join.cpp +3 -1
  539. package/src/duckdb/src/storage/buffer_manager.cpp +44 -46
  540. package/src/duckdb/src/storage/checkpoint/row_group_writer.cpp +1 -1
  541. package/src/duckdb/src/storage/checkpoint/table_data_reader.cpp +4 -15
  542. package/src/duckdb/src/storage/checkpoint/table_data_writer.cpp +10 -4
  543. package/src/duckdb/src/storage/checkpoint_manager.cpp +9 -3
  544. package/src/duckdb/src/storage/compression/bitpacking.cpp +29 -25
  545. package/src/duckdb/src/storage/compression/fixed_size_uncompressed.cpp +45 -46
  546. package/src/duckdb/src/storage/compression/numeric_constant.cpp +10 -11
  547. package/src/duckdb/src/storage/compression/patas.cpp +1 -1
  548. package/src/duckdb/src/storage/compression/rle.cpp +20 -15
  549. package/src/duckdb/src/storage/compression/validity_uncompressed.cpp +6 -6
  550. package/src/duckdb/src/storage/data_table.cpp +23 -23
  551. package/src/duckdb/src/storage/index.cpp +12 -1
  552. package/src/duckdb/src/storage/local_storage.cpp +27 -23
  553. package/src/duckdb/src/storage/meta_block_reader.cpp +22 -0
  554. package/src/duckdb/src/storage/statistics/base_statistics.cpp +373 -128
  555. package/src/duckdb/src/storage/statistics/column_statistics.cpp +57 -3
  556. package/src/duckdb/src/storage/statistics/distinct_statistics.cpp +8 -9
  557. package/src/duckdb/src/storage/statistics/list_stats.cpp +121 -0
  558. package/src/duckdb/src/storage/statistics/numeric_stats.cpp +591 -0
  559. package/src/duckdb/src/storage/statistics/numeric_stats_union.cpp +65 -0
  560. package/src/duckdb/src/storage/statistics/segment_statistics.cpp +2 -11
  561. package/src/duckdb/src/storage/statistics/string_stats.cpp +273 -0
  562. package/src/duckdb/src/storage/statistics/struct_stats.cpp +133 -0
  563. package/src/duckdb/src/storage/storage_info.cpp +2 -2
  564. package/src/duckdb/src/storage/table/column_checkpoint_state.cpp +4 -10
  565. package/src/duckdb/src/storage/table/column_data.cpp +118 -62
  566. package/src/duckdb/src/storage/table/column_data_checkpointer.cpp +10 -9
  567. package/src/duckdb/src/storage/table/column_segment.cpp +30 -45
  568. package/src/duckdb/src/storage/table/list_column_data.cpp +50 -71
  569. package/src/duckdb/src/storage/table/persistent_table_data.cpp +2 -1
  570. package/src/duckdb/src/storage/table/row_group.cpp +213 -143
  571. package/src/duckdb/src/storage/table/row_group_collection.cpp +151 -105
  572. package/src/duckdb/src/storage/table/scan_state.cpp +45 -33
  573. package/src/duckdb/src/storage/table/standard_column_data.cpp +11 -12
  574. package/src/duckdb/src/storage/table/struct_column_data.cpp +27 -34
  575. package/src/duckdb/src/storage/table/table_statistics.cpp +27 -7
  576. package/src/duckdb/src/storage/table/update_segment.cpp +23 -18
  577. package/src/duckdb/src/storage/wal_replay.cpp +8 -5
  578. package/src/duckdb/src/storage/write_ahead_log.cpp +2 -2
  579. package/src/duckdb/src/transaction/commit_state.cpp +11 -7
  580. package/src/duckdb/src/verification/deserialized_statement_verifier.cpp +0 -1
  581. package/src/duckdb/third_party/libpg_query/include/nodes/nodes.hpp +35 -0
  582. package/src/duckdb/third_party/libpg_query/include/nodes/parsenodes.hpp +36 -2
  583. package/src/duckdb/third_party/libpg_query/include/nodes/primnodes.hpp +3 -3
  584. package/src/duckdb/third_party/libpg_query/include/parser/gram.hpp +1022 -530
  585. package/src/duckdb/third_party/libpg_query/include/parser/kwlist.hpp +8 -0
  586. package/src/duckdb/third_party/libpg_query/src_backend_parser_gram.cpp +24462 -22828
  587. package/src/duckdb/third_party/re2/re2/re2.cc +9 -0
  588. package/src/duckdb/third_party/re2/re2/re2.h +2 -0
  589. package/src/duckdb/ub_extension_icu_third_party_icu_i18n.cpp +4 -4
  590. package/src/duckdb/ub_extension_json_json_functions.cpp +2 -0
  591. package/src/duckdb/ub_src_common_serializer.cpp +2 -0
  592. package/src/duckdb/ub_src_common_types.cpp +2 -0
  593. package/src/duckdb/ub_src_execution_physical_plan.cpp +2 -0
  594. package/src/duckdb/ub_src_function_aggregate_distributive.cpp +2 -0
  595. package/src/duckdb/ub_src_function_scalar_bit.cpp +2 -0
  596. package/src/duckdb/ub_src_function_scalar_map.cpp +4 -0
  597. package/src/duckdb/ub_src_function_scalar_string.cpp +2 -0
  598. package/src/duckdb/ub_src_function_scalar_string_regexp.cpp +4 -0
  599. package/src/duckdb/ub_src_main_capi.cpp +2 -0
  600. package/src/duckdb/ub_src_optimizer_rule.cpp +2 -0
  601. package/src/duckdb/ub_src_parser.cpp +2 -0
  602. package/src/duckdb/ub_src_parser_parsed_data.cpp +4 -2
  603. package/src/duckdb/ub_src_parser_statement.cpp +2 -0
  604. package/src/duckdb/ub_src_parser_tableref.cpp +2 -0
  605. package/src/duckdb/ub_src_parser_transform_statement.cpp +2 -0
  606. package/src/duckdb/ub_src_parser_transform_tableref.cpp +2 -0
  607. package/src/duckdb/ub_src_planner_binder_expression.cpp +2 -0
  608. package/src/duckdb/ub_src_planner_binder_tableref.cpp +2 -0
  609. package/src/duckdb/ub_src_planner_expression_binder.cpp +2 -0
  610. package/src/duckdb/ub_src_planner_operator.cpp +2 -0
  611. package/src/duckdb/ub_src_storage_statistics.cpp +6 -6
  612. package/src/duckdb/ub_src_storage_table.cpp +0 -2
  613. package/src/duckdb_node.hpp +2 -1
  614. package/src/statement.cpp +5 -5
  615. package/src/utils.cpp +27 -2
  616. package/test/extension.test.ts +44 -26
  617. package/test/syntax_error.test.ts +3 -1
  618. package/filelist.cache +0 -0
  619. package/src/duckdb/src/include/duckdb/main/loadable_extension.hpp +0 -59
  620. package/src/duckdb/src/include/duckdb/storage/statistics/list_statistics.hpp +0 -36
  621. package/src/duckdb/src/include/duckdb/storage/statistics/numeric_statistics.hpp +0 -75
  622. package/src/duckdb/src/include/duckdb/storage/statistics/string_statistics.hpp +0 -49
  623. package/src/duckdb/src/include/duckdb/storage/statistics/struct_statistics.hpp +0 -36
  624. package/src/duckdb/src/include/duckdb/storage/statistics/validity_statistics.hpp +0 -45
  625. package/src/duckdb/src/parser/parsed_data/alter_function_info.cpp +0 -55
  626. package/src/duckdb/src/storage/statistics/list_statistics.cpp +0 -94
  627. package/src/duckdb/src/storage/statistics/numeric_statistics.cpp +0 -307
  628. package/src/duckdb/src/storage/statistics/string_statistics.cpp +0 -220
  629. package/src/duckdb/src/storage/statistics/struct_statistics.cpp +0 -108
  630. package/src/duckdb/src/storage/statistics/validity_statistics.cpp +0 -91
  631. package/src/duckdb/src/storage/table/segment_tree.cpp +0 -179
@@ -1,569 +1,11 @@
1
1
  #include "duckdb/common/pair.hpp"
2
2
  #include "duckdb/common/types/chunk_collection.hpp"
3
+ #include "duckdb/common/types/list_segment.hpp"
3
4
  #include "duckdb/function/aggregate/nested_functions.hpp"
4
5
  #include "duckdb/planner/expression/bound_aggregate_expression.hpp"
5
6
 
6
7
  namespace duckdb {
7
8
 
8
- struct ListSegment {
9
- uint16_t count;
10
- uint16_t capacity;
11
- ListSegment *next;
12
- };
13
- struct LinkedList {
14
- LinkedList() {};
15
- LinkedList(idx_t total_capacity_p, ListSegment *first_segment_p, ListSegment *last_segment_p)
16
- : total_capacity(total_capacity_p), first_segment(first_segment_p), last_segment(last_segment_p) {
17
- }
18
-
19
- idx_t total_capacity = 0;
20
- ListSegment *first_segment = nullptr;
21
- ListSegment *last_segment = nullptr;
22
- };
23
-
24
- // forward declarations
25
- struct WriteDataToSegment;
26
- struct ReadDataFromSegment;
27
- struct CopyDataFromSegment;
28
- typedef ListSegment *(*create_segment_t)(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
29
- vector<AllocatedData> &owning_vector, const uint16_t &capacity);
30
- typedef void (*write_data_to_segment_t)(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
31
- vector<AllocatedData> &owning_vector, ListSegment *segment, Vector &input,
32
- idx_t &entry_idx, idx_t &count);
33
- typedef void (*read_data_from_segment_t)(ReadDataFromSegment &read_data_from_segment, const ListSegment *segment,
34
- Vector &result, idx_t &total_count);
35
- typedef ListSegment *(*copy_data_from_segment_t)(CopyDataFromSegment &copy_data_from_segment, const ListSegment *source,
36
- Allocator &allocator, vector<AllocatedData> &owning_vector);
37
-
38
- struct WriteDataToSegment {
39
- create_segment_t create_segment;
40
- write_data_to_segment_t segment_function;
41
- vector<WriteDataToSegment> child_functions;
42
- };
43
- struct ReadDataFromSegment {
44
- read_data_from_segment_t segment_function;
45
- vector<ReadDataFromSegment> child_functions;
46
- };
47
- struct CopyDataFromSegment {
48
- copy_data_from_segment_t segment_function;
49
- vector<CopyDataFromSegment> child_functions;
50
- };
51
-
52
- // forward declarations
53
- static void AppendRow(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
54
- vector<AllocatedData> &owning_vector, LinkedList *linked_list, Vector &input, idx_t &entry_idx,
55
- idx_t &count);
56
- static void BuildListVector(ReadDataFromSegment &read_data_from_segment, LinkedList *linked_list, Vector &result,
57
- idx_t &initial_total_count);
58
- static void CopyLinkedList(CopyDataFromSegment &copy_data_from_segment, const LinkedList *source_list,
59
- LinkedList &target_list, Allocator &allocator, vector<AllocatedData> &owning_vector);
60
-
61
- template <class T>
62
- static data_ptr_t AllocatePrimitiveData(Allocator &allocator, vector<AllocatedData> &owning_vector,
63
- const uint16_t &capacity) {
64
-
65
- owning_vector.emplace_back(allocator.Allocate(sizeof(ListSegment) + capacity * (sizeof(bool) + sizeof(T))));
66
- return owning_vector.back().get();
67
- }
68
-
69
- static data_ptr_t AllocateListData(Allocator &allocator, vector<AllocatedData> &owning_vector,
70
- const uint16_t &capacity) {
71
-
72
- owning_vector.emplace_back(
73
- allocator.Allocate(sizeof(ListSegment) + capacity * (sizeof(bool) + sizeof(uint64_t)) + sizeof(LinkedList)));
74
- return owning_vector.back().get();
75
- }
76
-
77
- static data_ptr_t AllocateStructData(Allocator &allocator, vector<AllocatedData> &owning_vector,
78
- const uint16_t &capacity, const idx_t &child_count) {
79
-
80
- owning_vector.emplace_back(
81
- allocator.Allocate(sizeof(ListSegment) + capacity * sizeof(bool) + child_count * sizeof(ListSegment *)));
82
- return owning_vector.back().get();
83
- }
84
-
85
- template <class T>
86
- static T *GetPrimitiveData(const ListSegment *segment) {
87
- return (T *)(((char *)segment) + sizeof(ListSegment) + segment->capacity * sizeof(bool));
88
- }
89
-
90
- static uint64_t *GetListLengthData(const ListSegment *segment) {
91
- return (uint64_t *)(((char *)segment) + sizeof(ListSegment) + segment->capacity * sizeof(bool));
92
- }
93
-
94
- static LinkedList *GetListChildData(const ListSegment *segment) {
95
- return (LinkedList *)(((char *)segment) + sizeof(ListSegment) +
96
- segment->capacity * (sizeof(bool) + sizeof(uint64_t)));
97
- }
98
-
99
- static ListSegment **GetStructData(const ListSegment *segment) {
100
- return (ListSegment **)(((char *)segment) + sizeof(ListSegment) + segment->capacity * sizeof(bool));
101
- }
102
-
103
- static bool *GetNullMask(const ListSegment *segment) {
104
- return (bool *)(((char *)segment) + sizeof(ListSegment));
105
- }
106
-
107
- static uint16_t GetCapacityForNewSegment(const LinkedList *linked_list) {
108
-
109
- // consecutive segments grow by the power of two
110
- uint16_t capacity = 4;
111
- if (linked_list->last_segment) {
112
- auto next_power_of_two = linked_list->last_segment->capacity * 2;
113
- capacity = next_power_of_two < 65536 ? next_power_of_two : linked_list->last_segment->capacity;
114
- }
115
- return capacity;
116
- }
117
-
118
- template <class T>
119
- static ListSegment *CreatePrimitiveSegment(WriteDataToSegment &, Allocator &allocator,
120
- vector<AllocatedData> &owning_vector, const uint16_t &capacity) {
121
-
122
- // allocate data and set the header
123
- auto segment = (ListSegment *)AllocatePrimitiveData<T>(allocator, owning_vector, capacity);
124
- segment->capacity = capacity;
125
- segment->count = 0;
126
- segment->next = nullptr;
127
- return segment;
128
- }
129
-
130
- static ListSegment *CreateListSegment(WriteDataToSegment &, Allocator &allocator, vector<AllocatedData> &owning_vector,
131
- const uint16_t &capacity) {
132
-
133
- // allocate data and set the header
134
- auto segment = (ListSegment *)AllocateListData(allocator, owning_vector, capacity);
135
- segment->capacity = capacity;
136
- segment->count = 0;
137
- segment->next = nullptr;
138
-
139
- // create an empty linked list for the child vector
140
- auto linked_child_list = GetListChildData(segment);
141
- LinkedList linked_list(0, nullptr, nullptr);
142
- Store<LinkedList>(linked_list, (data_ptr_t)linked_child_list);
143
-
144
- return segment;
145
- }
146
-
147
- static ListSegment *CreateStructSegment(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
148
- vector<AllocatedData> &owning_vector, const uint16_t &capacity) {
149
-
150
- // allocate data and set header
151
- auto segment = (ListSegment *)AllocateStructData(allocator, owning_vector, capacity,
152
- write_data_to_segment.child_functions.size());
153
- segment->capacity = capacity;
154
- segment->count = 0;
155
- segment->next = nullptr;
156
-
157
- // create a child ListSegment with exactly the same capacity for each child vector
158
- auto child_segments = GetStructData(segment);
159
- for (idx_t i = 0; i < write_data_to_segment.child_functions.size(); i++) {
160
- auto child_function = write_data_to_segment.child_functions[i];
161
- auto child_segment = child_function.create_segment(child_function, allocator, owning_vector, capacity);
162
- Store<ListSegment *>(child_segment, (data_ptr_t)(child_segments + i));
163
- }
164
-
165
- return segment;
166
- }
167
-
168
- static ListSegment *GetSegment(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
169
- vector<AllocatedData> &owning_vector, LinkedList *linked_list) {
170
-
171
- ListSegment *segment = nullptr;
172
-
173
- // determine segment
174
- if (!linked_list->last_segment) {
175
- // empty linked list, create the first (and last) segment
176
- auto capacity = GetCapacityForNewSegment(linked_list);
177
- segment = write_data_to_segment.create_segment(write_data_to_segment, allocator, owning_vector, capacity);
178
- linked_list->first_segment = segment;
179
- linked_list->last_segment = segment;
180
-
181
- } else if (linked_list->last_segment->capacity == linked_list->last_segment->count) {
182
- // the last segment of the linked list is full, create a new one and append it
183
- auto capacity = GetCapacityForNewSegment(linked_list);
184
- segment = write_data_to_segment.create_segment(write_data_to_segment, allocator, owning_vector, capacity);
185
- linked_list->last_segment->next = segment;
186
- linked_list->last_segment = segment;
187
-
188
- } else {
189
- // the last segment of the linked list is not full, append the data to it
190
- segment = linked_list->last_segment;
191
- }
192
-
193
- D_ASSERT(segment);
194
- return segment;
195
- }
196
-
197
- template <class T>
198
- static void WriteDataToPrimitiveSegment(WriteDataToSegment &, Allocator &allocator,
199
- vector<AllocatedData> &owning_vector, ListSegment *segment, Vector &input,
200
- idx_t &entry_idx, idx_t &count) {
201
-
202
- // get the vector data and the source index of the entry that we want to write
203
- auto input_data = FlatVector::GetData(input);
204
-
205
- // write null validity
206
- auto null_mask = GetNullMask(segment);
207
- auto is_null = FlatVector::IsNull(input, entry_idx);
208
- null_mask[segment->count] = is_null;
209
-
210
- // write value
211
- if (!is_null) {
212
- auto data = GetPrimitiveData<T>(segment);
213
- Store<T>(((T *)input_data)[entry_idx], (data_ptr_t)(data + segment->count));
214
- }
215
- }
216
-
217
- static void WriteDataToVarcharSegment(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
218
- vector<AllocatedData> &owning_vector, ListSegment *segment, Vector &input,
219
- idx_t &entry_idx, idx_t &count) {
220
-
221
- // get the vector data and the source index of the entry that we want to write
222
- auto input_data = FlatVector::GetData(input);
223
-
224
- // write null validity
225
- auto null_mask = GetNullMask(segment);
226
- auto is_null = FlatVector::IsNull(input, entry_idx);
227
- null_mask[segment->count] = is_null;
228
-
229
- // set the length of this string
230
- auto str_length_data = GetListLengthData(segment);
231
- uint64_t str_length = 0;
232
-
233
- // get the string
234
- string_t str_t;
235
- if (!is_null) {
236
- str_t = ((string_t *)input_data)[entry_idx];
237
- str_length = str_t.GetSize();
238
- }
239
-
240
- // we can reconstruct the offset from the length
241
- Store<uint64_t>(str_length, (data_ptr_t)(str_length_data + segment->count));
242
-
243
- if (is_null) {
244
- return;
245
- }
246
-
247
- // write the characters to the linked list of child segments
248
- auto child_segments = Load<LinkedList>((data_ptr_t)GetListChildData(segment));
249
- for (char &c : str_t.GetString()) {
250
- auto child_segment =
251
- GetSegment(write_data_to_segment.child_functions.back(), allocator, owning_vector, &child_segments);
252
- auto data = GetPrimitiveData<char>(child_segment);
253
- data[child_segment->count] = c;
254
- child_segment->count++;
255
- child_segments.total_capacity++;
256
- }
257
-
258
- // store the updated linked list
259
- Store<LinkedList>(child_segments, (data_ptr_t)GetListChildData(segment));
260
- }
261
-
262
- static void WriteDataToListSegment(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
263
- vector<AllocatedData> &owning_vector, ListSegment *segment, Vector &input,
264
- idx_t &entry_idx, idx_t &count) {
265
-
266
- // get the vector data and the source index of the entry that we want to write
267
- auto input_data = FlatVector::GetData(input);
268
-
269
- // write null validity
270
- auto null_mask = GetNullMask(segment);
271
- auto is_null = FlatVector::IsNull(input, entry_idx);
272
- null_mask[segment->count] = is_null;
273
-
274
- // set the length of this list
275
- auto list_length_data = GetListLengthData(segment);
276
- uint64_t list_length = 0;
277
-
278
- if (!is_null) {
279
- // get list entry information
280
- auto list_entries = (list_entry_t *)input_data;
281
- const auto &list_entry = list_entries[entry_idx];
282
- list_length = list_entry.length;
283
-
284
- // get the child vector and its data
285
- auto lists_size = ListVector::GetListSize(input);
286
- auto &child_vector = ListVector::GetEntry(input);
287
-
288
- // loop over the child vector entries and recurse on them
289
- auto child_segments = Load<LinkedList>((data_ptr_t)GetListChildData(segment));
290
- D_ASSERT(write_data_to_segment.child_functions.size() == 1);
291
- for (idx_t child_idx = 0; child_idx < list_entry.length; child_idx++) {
292
- auto source_idx_child = list_entry.offset + child_idx;
293
- AppendRow(write_data_to_segment.child_functions[0], allocator, owning_vector, &child_segments, child_vector,
294
- source_idx_child, lists_size);
295
- }
296
- // store the updated linked list
297
- Store<LinkedList>(child_segments, (data_ptr_t)GetListChildData(segment));
298
- }
299
-
300
- Store<uint64_t>(list_length, (data_ptr_t)(list_length_data + segment->count));
301
- }
302
-
303
- static void WriteDataToStructSegment(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
304
- vector<AllocatedData> &owning_vector, ListSegment *segment, Vector &input,
305
- idx_t &entry_idx, idx_t &count) {
306
-
307
- // write null validity
308
- auto null_mask = GetNullMask(segment);
309
- auto is_null = FlatVector::IsNull(input, entry_idx);
310
- null_mask[segment->count] = is_null;
311
-
312
- // write value
313
- auto &children = StructVector::GetEntries(input);
314
- D_ASSERT(children.size() == write_data_to_segment.child_functions.size());
315
- auto child_list = GetStructData(segment);
316
-
317
- // write the data of each of the children of the struct
318
- for (idx_t child_count = 0; child_count < children.size(); child_count++) {
319
- auto child_list_segment = Load<ListSegment *>((data_ptr_t)(child_list + child_count));
320
- auto &child_function = write_data_to_segment.child_functions[child_count];
321
- child_function.segment_function(child_function, allocator, owning_vector, child_list_segment,
322
- *children[child_count], entry_idx, count);
323
- child_list_segment->count++;
324
- }
325
- }
326
-
327
- static void AppendRow(WriteDataToSegment &write_data_to_segment, Allocator &allocator,
328
- vector<AllocatedData> &owning_vector, LinkedList *linked_list, Vector &input, idx_t &entry_idx,
329
- idx_t &count) {
330
-
331
- D_ASSERT(input.GetVectorType() == VectorType::FLAT_VECTOR);
332
-
333
- auto segment = GetSegment(write_data_to_segment, allocator, owning_vector, linked_list);
334
- write_data_to_segment.segment_function(write_data_to_segment, allocator, owning_vector, segment, input, entry_idx,
335
- count);
336
-
337
- linked_list->total_capacity++;
338
- segment->count++;
339
- }
340
-
341
- template <class T>
342
- static void ReadDataFromPrimitiveSegment(ReadDataFromSegment &, const ListSegment *segment, Vector &result,
343
- idx_t &total_count) {
344
-
345
- auto &aggr_vector_validity = FlatVector::Validity(result);
346
-
347
- // set NULLs
348
- auto null_mask = GetNullMask(segment);
349
- for (idx_t i = 0; i < segment->count; i++) {
350
- if (null_mask[i]) {
351
- aggr_vector_validity.SetInvalid(total_count + i);
352
- }
353
- }
354
-
355
- auto aggr_vector_data = FlatVector::GetData(result);
356
-
357
- // load values
358
- for (idx_t i = 0; i < segment->count; i++) {
359
- if (aggr_vector_validity.RowIsValid(total_count + i)) {
360
- auto data = GetPrimitiveData<T>(segment);
361
- ((T *)aggr_vector_data)[total_count + i] = Load<T>((data_ptr_t)(data + i));
362
- }
363
- }
364
- }
365
-
366
- static void ReadDataFromVarcharSegment(ReadDataFromSegment &, const ListSegment *segment, Vector &result,
367
- idx_t &total_count) {
368
-
369
- auto &aggr_vector_validity = FlatVector::Validity(result);
370
-
371
- // set NULLs
372
- auto null_mask = GetNullMask(segment);
373
- for (idx_t i = 0; i < segment->count; i++) {
374
- if (null_mask[i]) {
375
- aggr_vector_validity.SetInvalid(total_count + i);
376
- }
377
- }
378
-
379
- // append all the child chars to one string
380
- string str = "";
381
- auto linked_child_list = Load<LinkedList>((data_ptr_t)GetListChildData(segment));
382
- while (linked_child_list.first_segment) {
383
- auto child_segment = linked_child_list.first_segment;
384
- auto data = GetPrimitiveData<char>(child_segment);
385
- str.append(data, child_segment->count);
386
- linked_child_list.first_segment = child_segment->next;
387
- }
388
- linked_child_list.last_segment = nullptr;
389
-
390
- // use length and (reconstructed) offset to get the correct substrings
391
- auto aggr_vector_data = FlatVector::GetData(result);
392
- auto str_length_data = GetListLengthData(segment);
393
-
394
- // get the substrings and write them to the result vector
395
- idx_t offset = 0;
396
- for (idx_t i = 0; i < segment->count; i++) {
397
- if (!null_mask[i]) {
398
- auto str_length = Load<uint64_t>((data_ptr_t)(str_length_data + i));
399
- auto substr = str.substr(offset, str_length);
400
- auto str_t = StringVector::AddStringOrBlob(result, substr);
401
- ((string_t *)aggr_vector_data)[total_count + i] = str_t;
402
- offset += str_length;
403
- }
404
- }
405
- }
406
-
407
- static void ReadDataFromListSegment(ReadDataFromSegment &read_data_from_segment, const ListSegment *segment,
408
- Vector &result, idx_t &total_count) {
409
-
410
- auto &aggr_vector_validity = FlatVector::Validity(result);
411
-
412
- // set NULLs
413
- auto null_mask = GetNullMask(segment);
414
- for (idx_t i = 0; i < segment->count; i++) {
415
- if (null_mask[i]) {
416
- aggr_vector_validity.SetInvalid(total_count + i);
417
- }
418
- }
419
-
420
- auto list_vector_data = FlatVector::GetData<list_entry_t>(result);
421
-
422
- // get the starting offset
423
- idx_t offset = 0;
424
- if (total_count != 0) {
425
- offset = list_vector_data[total_count - 1].offset + list_vector_data[total_count - 1].length;
426
- }
427
- idx_t starting_offset = offset;
428
-
429
- // set length and offsets
430
- auto list_length_data = GetListLengthData(segment);
431
- for (idx_t i = 0; i < segment->count; i++) {
432
- auto list_length = Load<uint64_t>((data_ptr_t)(list_length_data + i));
433
- list_vector_data[total_count + i].length = list_length;
434
- list_vector_data[total_count + i].offset = offset;
435
- offset += list_length;
436
- }
437
-
438
- auto &child_vector = ListVector::GetEntry(result);
439
- auto linked_child_list = Load<LinkedList>((data_ptr_t)GetListChildData(segment));
440
- ListVector::Reserve(result, offset);
441
-
442
- // recurse into the linked list of child values
443
- D_ASSERT(read_data_from_segment.child_functions.size() == 1);
444
- BuildListVector(read_data_from_segment.child_functions[0], &linked_child_list, child_vector, starting_offset);
445
- }
446
-
447
- static void ReadDataFromStructSegment(ReadDataFromSegment &read_data_from_segment, const ListSegment *segment,
448
- Vector &result, idx_t &total_count) {
449
-
450
- auto &aggr_vector_validity = FlatVector::Validity(result);
451
-
452
- // set NULLs
453
- auto null_mask = GetNullMask(segment);
454
- for (idx_t i = 0; i < segment->count; i++) {
455
- if (null_mask[i]) {
456
- aggr_vector_validity.SetInvalid(total_count + i);
457
- }
458
- }
459
-
460
- auto &children = StructVector::GetEntries(result);
461
-
462
- // recurse into the child segments of each child of the struct
463
- D_ASSERT(children.size() == read_data_from_segment.child_functions.size());
464
- auto struct_children = GetStructData(segment);
465
- for (idx_t child_count = 0; child_count < children.size(); child_count++) {
466
- auto struct_children_segment = Load<ListSegment *>((data_ptr_t)(struct_children + child_count));
467
- auto &child_function = read_data_from_segment.child_functions[child_count];
468
- child_function.segment_function(child_function, struct_children_segment, *children[child_count], total_count);
469
- }
470
- }
471
-
472
- static void BuildListVector(ReadDataFromSegment &read_data_from_segment, LinkedList *linked_list, Vector &result,
473
- idx_t &initial_total_count) {
474
-
475
- idx_t total_count = initial_total_count;
476
- while (linked_list->first_segment) {
477
- auto segment = linked_list->first_segment;
478
- read_data_from_segment.segment_function(read_data_from_segment, segment, result, total_count);
479
-
480
- total_count += segment->count;
481
- linked_list->first_segment = segment->next;
482
- }
483
-
484
- linked_list->last_segment = nullptr;
485
- }
486
-
487
- template <class T>
488
- static ListSegment *CopyDataFromPrimitiveSegment(CopyDataFromSegment &, const ListSegment *source, Allocator &allocator,
489
- vector<AllocatedData> &owning_vector) {
490
-
491
- auto target = (ListSegment *)AllocatePrimitiveData<T>(allocator, owning_vector, source->capacity);
492
- memcpy(target, source, sizeof(ListSegment) + source->capacity * (sizeof(bool) + sizeof(T)));
493
- target->next = nullptr;
494
- return target;
495
- }
496
-
497
- static ListSegment *CopyDataFromListSegment(CopyDataFromSegment &copy_data_from_segment, const ListSegment *source,
498
- Allocator &allocator, vector<AllocatedData> &owning_vector) {
499
-
500
- // create an empty linked list for the child vector of target
501
- auto source_linked_child_list = Load<LinkedList>((data_ptr_t)GetListChildData(source));
502
-
503
- // create the segment
504
- auto target = (ListSegment *)AllocateListData(allocator, owning_vector, source->capacity);
505
- memcpy(target, source,
506
- sizeof(ListSegment) + source->capacity * (sizeof(bool) + sizeof(uint64_t)) + sizeof(LinkedList));
507
- target->next = nullptr;
508
-
509
- auto target_linked_list = GetListChildData(target);
510
- LinkedList linked_list(source_linked_child_list.total_capacity, nullptr, nullptr);
511
- Store<LinkedList>(linked_list, (data_ptr_t)target_linked_list);
512
-
513
- // recurse to copy the linked child list
514
- auto target_linked_child_list = Load<LinkedList>((data_ptr_t)GetListChildData(target));
515
- D_ASSERT(copy_data_from_segment.child_functions.size() == 1);
516
- CopyLinkedList(copy_data_from_segment.child_functions[0], &source_linked_child_list, target_linked_child_list,
517
- allocator, owning_vector);
518
-
519
- // store the updated linked list
520
- Store<LinkedList>(target_linked_child_list, (data_ptr_t)GetListChildData(target));
521
- return target;
522
- }
523
-
524
- static ListSegment *CopyDataFromStructSegment(CopyDataFromSegment &copy_data_from_segment, const ListSegment *source,
525
- Allocator &allocator, vector<AllocatedData> &owning_vector) {
526
-
527
- auto source_child_count = copy_data_from_segment.child_functions.size();
528
- auto target = (ListSegment *)AllocateStructData(allocator, owning_vector, source->capacity, source_child_count);
529
- memcpy(target, source,
530
- sizeof(ListSegment) + source->capacity * sizeof(bool) + source_child_count * sizeof(ListSegment *));
531
- target->next = nullptr;
532
-
533
- // recurse and copy the children
534
- auto source_child_segments = GetStructData(source);
535
- auto target_child_segments = GetStructData(target);
536
-
537
- for (idx_t i = 0; i < copy_data_from_segment.child_functions.size(); i++) {
538
- auto child_function = copy_data_from_segment.child_functions[i];
539
- auto source_child_segment = Load<ListSegment *>((data_ptr_t)(source_child_segments + i));
540
- auto target_child_segment =
541
- child_function.segment_function(child_function, source_child_segment, allocator, owning_vector);
542
- Store<ListSegment *>(target_child_segment, (data_ptr_t)(target_child_segments + i));
543
- }
544
- return target;
545
- }
546
-
547
- static void CopyLinkedList(CopyDataFromSegment &copy_data_from_segment, const LinkedList *source_list,
548
- LinkedList &target_list, Allocator &allocator, vector<AllocatedData> &owning_vector) {
549
-
550
- auto source_segment = source_list->first_segment;
551
-
552
- while (source_segment) {
553
- auto target_segment =
554
- copy_data_from_segment.segment_function(copy_data_from_segment, source_segment, allocator, owning_vector);
555
- source_segment = source_segment->next;
556
-
557
- if (!target_list.first_segment) {
558
- target_list.first_segment = target_segment;
559
- }
560
- if (target_list.last_segment) {
561
- target_list.last_segment->next = target_segment;
562
- }
563
- target_list.last_segment = target_segment;
564
- }
565
- }
566
-
567
9
  static void InitializeValidities(Vector &vector, idx_t &capacity) {
568
10
 
569
11
  auto &validity_mask = FlatVector::Validity(vector);
@@ -619,154 +61,6 @@ struct ListBindData : public FunctionData {
619
61
  }
620
62
  };
621
63
 
622
- static void GetSegmentDataFunctions(WriteDataToSegment &write_data_to_segment,
623
- ReadDataFromSegment &read_data_from_segment,
624
- CopyDataFromSegment &copy_data_from_segment, const LogicalType &type) {
625
-
626
- auto physical_type = type.InternalType();
627
- switch (physical_type) {
628
- case PhysicalType::BIT:
629
- case PhysicalType::BOOL: {
630
- write_data_to_segment.create_segment = CreatePrimitiveSegment<bool>;
631
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<bool>;
632
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<bool>;
633
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<bool>;
634
- break;
635
- }
636
- case PhysicalType::INT8: {
637
- write_data_to_segment.create_segment = CreatePrimitiveSegment<int8_t>;
638
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<int8_t>;
639
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<int8_t>;
640
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<int8_t>;
641
- break;
642
- }
643
- case PhysicalType::INT16: {
644
- write_data_to_segment.create_segment = CreatePrimitiveSegment<int16_t>;
645
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<int16_t>;
646
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<int16_t>;
647
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<int16_t>;
648
- break;
649
- }
650
- case PhysicalType::INT32: {
651
- write_data_to_segment.create_segment = CreatePrimitiveSegment<int32_t>;
652
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<int32_t>;
653
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<int32_t>;
654
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<int32_t>;
655
- break;
656
- }
657
- case PhysicalType::INT64: {
658
- write_data_to_segment.create_segment = CreatePrimitiveSegment<int64_t>;
659
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<int64_t>;
660
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<int64_t>;
661
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<int64_t>;
662
- break;
663
- }
664
- case PhysicalType::UINT8: {
665
- write_data_to_segment.create_segment = CreatePrimitiveSegment<uint8_t>;
666
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<uint8_t>;
667
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<uint8_t>;
668
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<uint8_t>;
669
- break;
670
- }
671
- case PhysicalType::UINT16: {
672
- write_data_to_segment.create_segment = CreatePrimitiveSegment<uint16_t>;
673
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<uint16_t>;
674
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<uint16_t>;
675
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<uint16_t>;
676
- break;
677
- }
678
- case PhysicalType::UINT32: {
679
- write_data_to_segment.create_segment = CreatePrimitiveSegment<uint32_t>;
680
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<uint32_t>;
681
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<uint32_t>;
682
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<uint32_t>;
683
- break;
684
- }
685
- case PhysicalType::UINT64: {
686
- write_data_to_segment.create_segment = CreatePrimitiveSegment<uint64_t>;
687
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<uint64_t>;
688
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<uint64_t>;
689
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<uint64_t>;
690
- break;
691
- }
692
- case PhysicalType::FLOAT: {
693
- write_data_to_segment.create_segment = CreatePrimitiveSegment<float>;
694
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<float>;
695
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<float>;
696
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<float>;
697
- break;
698
- }
699
- case PhysicalType::DOUBLE: {
700
- write_data_to_segment.create_segment = CreatePrimitiveSegment<double>;
701
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<double>;
702
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<double>;
703
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<double>;
704
- break;
705
- }
706
- case PhysicalType::INT128: {
707
- write_data_to_segment.create_segment = CreatePrimitiveSegment<hugeint_t>;
708
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<hugeint_t>;
709
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<hugeint_t>;
710
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<hugeint_t>;
711
- break;
712
- }
713
- case PhysicalType::INTERVAL: {
714
- write_data_to_segment.create_segment = CreatePrimitiveSegment<interval_t>;
715
- write_data_to_segment.segment_function = WriteDataToPrimitiveSegment<interval_t>;
716
- read_data_from_segment.segment_function = ReadDataFromPrimitiveSegment<interval_t>;
717
- copy_data_from_segment.segment_function = CopyDataFromPrimitiveSegment<interval_t>;
718
- break;
719
- }
720
- case PhysicalType::VARCHAR: {
721
- write_data_to_segment.create_segment = CreateListSegment;
722
- write_data_to_segment.segment_function = WriteDataToVarcharSegment;
723
- read_data_from_segment.segment_function = ReadDataFromVarcharSegment;
724
- copy_data_from_segment.segment_function = CopyDataFromListSegment;
725
-
726
- write_data_to_segment.child_functions.emplace_back(WriteDataToSegment());
727
- write_data_to_segment.child_functions.back().create_segment = CreatePrimitiveSegment<char>;
728
- copy_data_from_segment.child_functions.emplace_back(CopyDataFromSegment());
729
- copy_data_from_segment.child_functions.back().segment_function = CopyDataFromPrimitiveSegment<char>;
730
- break;
731
- }
732
- case PhysicalType::LIST: {
733
- write_data_to_segment.create_segment = CreateListSegment;
734
- write_data_to_segment.segment_function = WriteDataToListSegment;
735
- read_data_from_segment.segment_function = ReadDataFromListSegment;
736
- copy_data_from_segment.segment_function = CopyDataFromListSegment;
737
-
738
- // recurse
739
- write_data_to_segment.child_functions.emplace_back(WriteDataToSegment());
740
- read_data_from_segment.child_functions.emplace_back(ReadDataFromSegment());
741
- copy_data_from_segment.child_functions.emplace_back(CopyDataFromSegment());
742
- GetSegmentDataFunctions(write_data_to_segment.child_functions.back(),
743
- read_data_from_segment.child_functions.back(),
744
- copy_data_from_segment.child_functions.back(), ListType::GetChildType(type));
745
- break;
746
- }
747
- case PhysicalType::STRUCT: {
748
- write_data_to_segment.create_segment = CreateStructSegment;
749
- write_data_to_segment.segment_function = WriteDataToStructSegment;
750
- read_data_from_segment.segment_function = ReadDataFromStructSegment;
751
- copy_data_from_segment.segment_function = CopyDataFromStructSegment;
752
-
753
- // recurse
754
- auto child_types = StructType::GetChildTypes(type);
755
- for (idx_t i = 0; i < child_types.size(); i++) {
756
- write_data_to_segment.child_functions.emplace_back(WriteDataToSegment());
757
- read_data_from_segment.child_functions.emplace_back(ReadDataFromSegment());
758
- copy_data_from_segment.child_functions.emplace_back(CopyDataFromSegment());
759
- GetSegmentDataFunctions(write_data_to_segment.child_functions.back(),
760
- read_data_from_segment.child_functions.back(),
761
- copy_data_from_segment.child_functions.back(), child_types[i].second);
762
- }
763
- break;
764
- }
765
- default:
766
- throw InternalException("LIST aggregate not yet implemented for " + type.ToString());
767
- }
768
- }
769
-
770
64
  ListBindData::ListBindData(const LogicalType &stype_p) : stype(stype_p) {
771
65
 
772
66
  // always unnest once because the result vector is of type LIST
@@ -834,8 +128,8 @@ static void ListUpdateFunction(Vector inputs[], AggregateInputData &aggr_input_d
834
128
  state->owning_vector = new vector<AllocatedData>;
835
129
  }
836
130
  D_ASSERT(state->type);
837
- AppendRow(list_bind_data.write_data_to_segment, aggr_input_data.allocator, *state->owning_vector,
838
- state->linked_list, input, i, count);
131
+ list_bind_data.write_data_to_segment.AppendRow(aggr_input_data.allocator, *state->owning_vector,
132
+ state->linked_list, input, i, count);
839
133
  }
840
134
  }
841
135
 
@@ -865,8 +159,8 @@ static void ListCombineFunction(Vector &state, Vector &combined, AggregateInputD
865
159
 
866
160
  // copy the linked list of the state
867
161
  auto copied_linked_list = LinkedList(state->linked_list->total_capacity, nullptr, nullptr);
868
- CopyLinkedList(list_bind_data.copy_data_from_segment, state->linked_list, copied_linked_list,
869
- aggr_input_data.allocator, *owning_vector);
162
+ list_bind_data.copy_data_from_segment.CopyLinkedList(state->linked_list, copied_linked_list,
163
+ aggr_input_data.allocator, *owning_vector);
870
164
 
871
165
  // append the copied linked list to the combined state
872
166
  if (combined_ptr[i]->linked_list->last_segment) {
@@ -919,7 +213,7 @@ static void ListFinalize(Vector &state_vector, AggregateInputData &aggr_input_da
919
213
  InitializeValidities(aggr_vector, total_capacity);
920
214
 
921
215
  idx_t total_count = 0;
922
- BuildListVector(list_bind_data.read_data_from_segment, state->linked_list, aggr_vector, total_count);
216
+ list_bind_data.read_data_from_segment.BuildListVector(state->linked_list, aggr_vector, total_count);
923
217
  ListVector::Append(result, aggr_vector, total_capacity);
924
218
 
925
219
  // now destroy the state (for parallel destruction)