duckdb 0.7.2-dev0.0 → 0.7.2-dev1034.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/binding.gyp +12 -7
- package/lib/duckdb.d.ts +55 -2
- package/lib/duckdb.js +20 -1
- package/package.json +1 -1
- package/src/connection.cpp +1 -2
- package/src/database.cpp +1 -1
- package/src/duckdb/extension/icu/icu-extension.cpp +4 -0
- package/src/duckdb/extension/icu/icu-list-range.cpp +207 -0
- package/src/duckdb/extension/icu/icu-table-range.cpp +194 -0
- package/src/duckdb/extension/icu/include/icu-list-range.hpp +17 -0
- package/src/duckdb/extension/icu/include/icu-table-range.hpp +17 -0
- package/src/duckdb/extension/json/include/json_common.hpp +1 -0
- package/src/duckdb/extension/json/include/json_functions.hpp +2 -0
- package/src/duckdb/extension/json/include/json_serializer.hpp +77 -0
- package/src/duckdb/extension/json/json_functions/json_serialize_sql.cpp +147 -0
- package/src/duckdb/extension/json/json_functions/read_json.cpp +6 -5
- package/src/duckdb/extension/json/json_functions.cpp +12 -4
- package/src/duckdb/extension/json/json_scan.cpp +2 -2
- package/src/duckdb/extension/json/json_serializer.cpp +217 -0
- package/src/duckdb/extension/parquet/column_reader.cpp +94 -15
- package/src/duckdb/extension/parquet/column_writer.cpp +0 -1
- package/src/duckdb/extension/parquet/include/column_reader.hpp +1 -2
- package/src/duckdb/extension/parquet/include/decode_utils.hpp +5 -4
- package/src/duckdb/extension/parquet/include/generated_column_reader.hpp +1 -11
- package/src/duckdb/extension/parquet/include/parquet_timestamp.hpp +2 -1
- package/src/duckdb/extension/parquet/parquet-extension.cpp +12 -2
- package/src/duckdb/extension/parquet/parquet_reader.cpp +1 -1
- package/src/duckdb/extension/parquet/parquet_statistics.cpp +26 -32
- package/src/duckdb/extension/parquet/parquet_timestamp.cpp +16 -6
- package/src/duckdb/src/catalog/catalog.cpp +34 -5
- package/src/duckdb/src/catalog/catalog_entry/duck_schema_entry.cpp +4 -0
- package/src/duckdb/src/catalog/catalog_entry/duck_table_entry.cpp +2 -21
- package/src/duckdb/src/catalog/catalog_entry/scalar_function_catalog_entry.cpp +7 -6
- package/src/duckdb/src/catalog/catalog_entry/table_catalog_entry.cpp +3 -3
- package/src/duckdb/src/catalog/catalog_entry/table_function_catalog_entry.cpp +20 -1
- package/src/duckdb/src/catalog/catalog_entry/type_catalog_entry.cpp +8 -2
- package/src/duckdb/src/catalog/catalog_set.cpp +1 -0
- package/src/duckdb/src/catalog/default/default_functions.cpp +3 -0
- package/src/duckdb/src/catalog/dependency_list.cpp +12 -0
- package/src/duckdb/src/catalog/duck_catalog.cpp +34 -7
- package/src/duckdb/src/common/arrow/arrow_appender.cpp +48 -4
- package/src/duckdb/src/common/arrow/arrow_converter.cpp +1 -1
- package/src/duckdb/src/common/box_renderer.cpp +109 -23
- package/src/duckdb/src/common/enums/expression_type.cpp +8 -222
- package/src/duckdb/src/common/enums/join_type.cpp +3 -22
- package/src/duckdb/src/common/enums/logical_operator_type.cpp +2 -0
- package/src/duckdb/src/common/enums/statement_type.cpp +2 -0
- package/src/duckdb/src/common/exception.cpp +15 -1
- package/src/duckdb/src/common/field_writer.cpp +1 -0
- package/src/duckdb/src/common/operator/cast_operators.cpp +1 -1
- package/src/duckdb/src/common/preserved_error.cpp +7 -5
- package/src/duckdb/src/common/serializer/buffered_deserializer.cpp +4 -0
- package/src/duckdb/src/common/serializer/buffered_file_reader.cpp +15 -2
- package/src/duckdb/src/common/serializer/enum_serializer.cpp +1176 -0
- package/src/duckdb/src/common/sort/sort_state.cpp +5 -7
- package/src/duckdb/src/common/sort/sorted_block.cpp +0 -1
- package/src/duckdb/src/common/string_util.cpp +4 -1
- package/src/duckdb/src/common/types/bit.cpp +166 -87
- package/src/duckdb/src/common/types/blob.cpp +1 -1
- package/src/duckdb/src/common/types/chunk_collection.cpp +2 -2
- package/src/duckdb/src/common/types/column_data_collection.cpp +39 -2
- package/src/duckdb/src/common/types/column_data_collection_segment.cpp +11 -6
- package/src/duckdb/src/common/types/data_chunk.cpp +1 -1
- package/src/duckdb/src/common/types/time.cpp +13 -0
- package/src/duckdb/src/common/types/value.cpp +320 -154
- package/src/duckdb/src/common/types/vector.cpp +155 -127
- package/src/duckdb/src/common/types.cpp +313 -153
- package/src/duckdb/src/common/vector_operations/vector_cast.cpp +2 -1
- package/src/duckdb/src/execution/aggregate_hashtable.cpp +10 -5
- package/src/duckdb/src/execution/column_binding_resolver.cpp +21 -5
- package/src/duckdb/src/execution/expression_executor/execute_cast.cpp +2 -1
- package/src/duckdb/src/execution/index/art/art.cpp +6 -5
- package/src/duckdb/src/execution/operator/aggregate/physical_perfecthash_aggregate.cpp +4 -5
- package/src/duckdb/src/execution/operator/aggregate/physical_window.cpp +117 -26
- package/src/duckdb/src/execution/operator/helper/physical_limit.cpp +3 -0
- package/src/duckdb/src/execution/operator/helper/physical_vacuum.cpp +5 -3
- package/src/duckdb/src/execution/operator/join/physical_blockwise_nl_join.cpp +64 -17
- package/src/duckdb/src/execution/operator/join/physical_iejoin.cpp +2 -2
- package/src/duckdb/src/execution/operator/join/physical_index_join.cpp +12 -4
- package/src/duckdb/src/execution/operator/join/physical_piecewise_merge_join.cpp +6 -11
- package/src/duckdb/src/execution/operator/join/physical_range_join.cpp +3 -1
- package/src/duckdb/src/execution/operator/persistent/base_csv_reader.cpp +6 -3
- package/src/duckdb/src/execution/operator/persistent/buffered_csv_reader.cpp +6 -14
- package/src/duckdb/src/execution/operator/persistent/physical_copy_to_file.cpp +2 -2
- package/src/duckdb/src/execution/operator/projection/physical_projection.cpp +34 -0
- package/src/duckdb/src/execution/operator/scan/physical_positional_scan.cpp +20 -5
- package/src/duckdb/src/execution/operator/schema/physical_create_type.cpp +20 -40
- package/src/duckdb/src/execution/partitionable_hashtable.cpp +14 -2
- package/src/duckdb/src/execution/physical_plan/plan_aggregate.cpp +21 -16
- package/src/duckdb/src/execution/physical_plan/plan_asof_join.cpp +97 -0
- package/src/duckdb/src/execution/physical_plan/plan_comparison_join.cpp +95 -47
- package/src/duckdb/src/execution/physical_plan/plan_distinct.cpp +5 -8
- package/src/duckdb/src/execution/physical_plan/plan_positional_join.cpp +14 -5
- package/src/duckdb/src/execution/physical_plan_generator.cpp +3 -0
- package/src/duckdb/src/execution/window_segment_tree.cpp +173 -1
- package/src/duckdb/src/function/aggregate/algebraic/avg.cpp +0 -6
- package/src/duckdb/src/function/aggregate/distributive/bitagg.cpp +99 -95
- package/src/duckdb/src/function/aggregate/distributive/bitstring_agg.cpp +269 -0
- package/src/duckdb/src/function/aggregate/distributive/bool.cpp +2 -0
- package/src/duckdb/src/function/aggregate/distributive/count.cpp +3 -4
- package/src/duckdb/src/function/aggregate/distributive/first.cpp +1 -0
- package/src/duckdb/src/function/aggregate/distributive/minmax.cpp +2 -0
- package/src/duckdb/src/function/aggregate/distributive/sum.cpp +19 -16
- package/src/duckdb/src/function/aggregate/distributive_functions.cpp +1 -0
- package/src/duckdb/src/function/aggregate/holistic/approximate_quantile.cpp +5 -2
- package/src/duckdb/src/function/aggregate/holistic/mode.cpp +1 -1
- package/src/duckdb/src/function/aggregate/holistic/quantile.cpp +16 -1
- package/src/duckdb/src/function/aggregate/nested/list.cpp +8 -8
- package/src/duckdb/src/function/aggregate/sorted_aggregate_function.cpp +58 -16
- package/src/duckdb/src/function/cast/bit_cast.cpp +0 -2
- package/src/duckdb/src/function/cast/blob_cast.cpp +0 -1
- package/src/duckdb/src/function/cast/cast_function_set.cpp +1 -1
- package/src/duckdb/src/function/cast/enum_casts.cpp +25 -3
- package/src/duckdb/src/function/cast/list_casts.cpp +17 -4
- package/src/duckdb/src/function/cast/map_cast.cpp +5 -2
- package/src/duckdb/src/function/cast/string_cast.cpp +36 -10
- package/src/duckdb/src/function/cast/struct_cast.cpp +24 -4
- package/src/duckdb/src/function/cast/time_casts.cpp +2 -2
- package/src/duckdb/src/function/cast/union_casts.cpp +33 -7
- package/src/duckdb/src/function/function_binder.cpp +1 -8
- package/src/duckdb/src/function/scalar/bit/bitstring.cpp +100 -0
- package/src/duckdb/src/function/scalar/date/current.cpp +0 -2
- package/src/duckdb/src/function/scalar/date/date_diff.cpp +0 -1
- package/src/duckdb/src/function/scalar/date/date_part.cpp +18 -26
- package/src/duckdb/src/function/scalar/date/date_sub.cpp +0 -1
- package/src/duckdb/src/function/scalar/date/date_trunc.cpp +10 -14
- package/src/duckdb/src/function/scalar/generic/stats.cpp +2 -4
- package/src/duckdb/src/function/scalar/list/contains_or_position.cpp +4 -146
- package/src/duckdb/src/function/scalar/list/flatten.cpp +5 -12
- package/src/duckdb/src/function/scalar/list/list_aggregates.cpp +1 -1
- package/src/duckdb/src/function/scalar/list/list_concat.cpp +8 -12
- package/src/duckdb/src/function/scalar/list/list_extract.cpp +5 -12
- package/src/duckdb/src/function/scalar/list/list_lambdas.cpp +7 -3
- package/src/duckdb/src/function/scalar/list/list_value.cpp +6 -10
- package/src/duckdb/src/function/scalar/map/map.cpp +47 -1
- package/src/duckdb/src/function/scalar/map/map_entries.cpp +61 -0
- package/src/duckdb/src/function/scalar/map/map_extract.cpp +68 -26
- package/src/duckdb/src/function/scalar/map/map_keys_values.cpp +97 -0
- package/src/duckdb/src/function/scalar/math/numeric.cpp +101 -17
- package/src/duckdb/src/function/scalar/math_functions.cpp +3 -0
- package/src/duckdb/src/function/scalar/nested_functions.cpp +3 -0
- package/src/duckdb/src/function/scalar/operators/add.cpp +0 -9
- package/src/duckdb/src/function/scalar/operators/arithmetic.cpp +29 -48
- package/src/duckdb/src/function/scalar/operators/bitwise.cpp +0 -63
- package/src/duckdb/src/function/scalar/operators/multiply.cpp +5 -6
- package/src/duckdb/src/function/scalar/operators/subtract.cpp +0 -6
- package/src/duckdb/src/function/scalar/string/caseconvert.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/hex.cpp +201 -0
- package/src/duckdb/src/function/scalar/string/instr.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/length.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/like.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/regexp/regexp_extract_all.cpp +243 -0
- package/src/duckdb/src/function/scalar/string/regexp/regexp_util.cpp +79 -0
- package/src/duckdb/src/function/scalar/string/regexp.cpp +21 -80
- package/src/duckdb/src/function/scalar/string/substring.cpp +2 -6
- package/src/duckdb/src/function/scalar/string_functions.cpp +2 -0
- package/src/duckdb/src/function/scalar/struct/struct_extract.cpp +5 -10
- package/src/duckdb/src/function/scalar/struct/struct_insert.cpp +11 -14
- package/src/duckdb/src/function/scalar/struct/struct_pack.cpp +6 -7
- package/src/duckdb/src/function/table/arrow.cpp +5 -2
- package/src/duckdb/src/function/table/arrow_conversion.cpp +25 -1
- package/src/duckdb/src/function/table/checkpoint.cpp +5 -1
- package/src/duckdb/src/function/table/read_csv.cpp +55 -0
- package/src/duckdb/src/function/table/system/duckdb_constraints.cpp +2 -2
- package/src/duckdb/src/function/table/system/test_all_types.cpp +2 -2
- package/src/duckdb/src/function/table/table_scan.cpp +1 -1
- package/src/duckdb/src/function/table/version/pragma_version.cpp +2 -2
- package/src/duckdb/src/function/table_function.cpp +30 -11
- package/src/duckdb/src/include/duckdb/catalog/catalog.hpp +6 -0
- package/src/duckdb/src/include/duckdb/catalog/catalog_entry/duck_table_entry.hpp +1 -1
- package/src/duckdb/src/include/duckdb/catalog/catalog_entry/table_function_catalog_entry.hpp +6 -8
- package/src/duckdb/src/include/duckdb/catalog/dependency_list.hpp +3 -0
- package/src/duckdb/src/include/duckdb/catalog/duck_catalog.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/box_renderer.hpp +8 -2
- package/src/duckdb/src/include/duckdb/common/constants.hpp +0 -19
- package/src/duckdb/src/include/duckdb/common/enums/aggregate_handling.hpp +2 -0
- package/src/duckdb/src/include/duckdb/common/enums/expression_type.hpp +2 -3
- package/src/duckdb/src/include/duckdb/common/enums/joinref_type.hpp +7 -4
- package/src/duckdb/src/include/duckdb/common/enums/logical_operator_type.hpp +1 -0
- package/src/duckdb/src/include/duckdb/common/enums/order_type.hpp +2 -0
- package/src/duckdb/src/include/duckdb/common/enums/set_operation_type.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/enums/statement_type.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/enums/tableref_type.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/exception.hpp +69 -2
- package/src/duckdb/src/include/duckdb/common/field_writer.hpp +12 -4
- package/src/duckdb/src/include/duckdb/common/{http_stats.hpp → http_state.hpp} +18 -4
- package/src/duckdb/src/include/duckdb/common/operator/multiply.hpp +2 -0
- package/src/duckdb/src/include/duckdb/common/optional_ptr.hpp +45 -0
- package/src/duckdb/src/include/duckdb/common/preserved_error.hpp +6 -1
- package/src/duckdb/src/include/duckdb/common/serializer/buffered_deserializer.hpp +4 -2
- package/src/duckdb/src/include/duckdb/common/serializer/buffered_file_reader.hpp +8 -2
- package/src/duckdb/src/include/duckdb/common/serializer/enum_serializer.hpp +113 -0
- package/src/duckdb/src/include/duckdb/common/serializer/format_deserializer.hpp +336 -0
- package/src/duckdb/src/include/duckdb/common/serializer/format_serializer.hpp +268 -0
- package/src/duckdb/src/include/duckdb/common/serializer/serialization_traits.hpp +126 -0
- package/src/duckdb/src/include/duckdb/common/serializer.hpp +13 -0
- package/src/duckdb/src/include/duckdb/common/string_util.hpp +25 -0
- package/src/duckdb/src/include/duckdb/common/types/bit.hpp +12 -7
- package/src/duckdb/src/include/duckdb/common/types/time.hpp +3 -0
- package/src/duckdb/src/include/duckdb/common/types/value.hpp +17 -48
- package/src/duckdb/src/include/duckdb/common/types/value_map.hpp +1 -1
- package/src/duckdb/src/include/duckdb/common/types/vector.hpp +3 -1
- package/src/duckdb/src/include/duckdb/common/types.hpp +45 -8
- package/src/duckdb/src/include/duckdb/common/vector_operations/unary_executor.hpp +2 -2
- package/src/duckdb/src/include/duckdb/execution/aggregate_hashtable.hpp +1 -0
- package/src/duckdb/src/include/duckdb/execution/index/art/art.hpp +2 -2
- package/src/duckdb/src/include/duckdb/execution/operator/aggregate/physical_perfecthash_aggregate.hpp +1 -1
- package/src/duckdb/src/include/duckdb/execution/operator/join/physical_cross_product.hpp +2 -0
- package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_file_handle.hpp +1 -0
- package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_reader_options.hpp +6 -0
- package/src/duckdb/src/include/duckdb/execution/operator/projection/physical_projection.hpp +5 -0
- package/src/duckdb/src/include/duckdb/execution/partitionable_hashtable.hpp +3 -0
- package/src/duckdb/src/include/duckdb/execution/physical_plan_generator.hpp +1 -3
- package/src/duckdb/src/include/duckdb/execution/window_segment_tree.hpp +54 -0
- package/src/duckdb/src/include/duckdb/function/aggregate/distributive_functions.hpp +5 -0
- package/src/duckdb/src/include/duckdb/function/aggregate_function.hpp +18 -6
- package/src/duckdb/src/include/duckdb/function/cast/bound_cast_data.hpp +84 -0
- package/src/duckdb/src/include/duckdb/function/cast/cast_function_set.hpp +2 -2
- package/src/duckdb/src/include/duckdb/function/cast/default_casts.hpp +28 -64
- package/src/duckdb/src/include/duckdb/function/function_binder.hpp +3 -6
- package/src/duckdb/src/include/duckdb/function/scalar/bit_functions.hpp +4 -0
- package/src/duckdb/src/include/duckdb/function/scalar/list/contains_or_position.hpp +138 -0
- package/src/duckdb/src/include/duckdb/function/scalar/math_functions.hpp +8 -0
- package/src/duckdb/src/include/duckdb/function/scalar/nested_functions.hpp +59 -0
- package/src/duckdb/src/include/duckdb/function/scalar/regexp.hpp +81 -1
- package/src/duckdb/src/include/duckdb/function/scalar/string_functions.hpp +4 -0
- package/src/duckdb/src/include/duckdb/function/scalar_function.hpp +2 -2
- package/src/duckdb/src/include/duckdb/function/table/arrow.hpp +12 -1
- package/src/duckdb/src/include/duckdb/function/table_function.hpp +10 -0
- package/src/duckdb/src/include/duckdb/main/capi/capi_internal.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/client_data.hpp +3 -3
- package/src/duckdb/src/include/duckdb/main/config.hpp +3 -0
- package/src/duckdb/src/include/duckdb/main/connection_manager.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/database.hpp +1 -0
- package/src/duckdb/src/include/duckdb/main/extension_entries.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/prepared_statement.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/relation/explain_relation.hpp +2 -1
- package/src/duckdb/src/include/duckdb/main/relation.hpp +2 -1
- package/src/duckdb/src/include/duckdb/optimizer/filter_pushdown.hpp +2 -0
- package/src/duckdb/src/include/duckdb/optimizer/join_order/cardinality_estimator.hpp +2 -2
- package/src/duckdb/src/include/duckdb/optimizer/rule/list.hpp +1 -0
- package/src/duckdb/src/include/duckdb/optimizer/rule/ordered_aggregate_optimizer.hpp +24 -0
- package/src/duckdb/src/include/duckdb/parser/common_table_expression_info.hpp +4 -0
- package/src/duckdb/src/include/duckdb/parser/expression/between_expression.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/expression/bound_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/case_expression.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/expression/cast_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/collate_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/columnref_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/comparison_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/conjunction_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/constant_expression.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/expression/default_expression.hpp +1 -0
- package/src/duckdb/src/include/duckdb/parser/expression/function_expression.hpp +4 -2
- package/src/duckdb/src/include/duckdb/parser/expression/lambda_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/operator_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/parameter_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/positional_reference_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/star_expression.hpp +4 -2
- package/src/duckdb/src/include/duckdb/parser/expression/subquery_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/window_expression.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_info.hpp +5 -1
- package/src/duckdb/src/include/duckdb/parser/parsed_data/{alter_function_info.hpp → alter_scalar_function_info.hpp} +13 -13
- package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_function_info.hpp +47 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_info.hpp +6 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_data/create_table_function_info.hpp +2 -1
- package/src/duckdb/src/include/duckdb/parser/parsed_data/sample_options.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_expression.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/query_node/recursive_cte_node.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/query_node/select_node.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/query_node/set_operation_node.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/query_node.hpp +13 -2
- package/src/duckdb/src/include/duckdb/parser/result_modifier.hpp +24 -1
- package/src/duckdb/src/include/duckdb/parser/sql_statement.hpp +2 -1
- package/src/duckdb/src/include/duckdb/parser/statement/multi_statement.hpp +28 -0
- package/src/duckdb/src/include/duckdb/parser/statement/select_statement.hpp +6 -1
- package/src/duckdb/src/include/duckdb/parser/tableref/basetableref.hpp +4 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/emptytableref.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/expressionlistref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/joinref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/list.hpp +1 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/pivotref.hpp +87 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/subqueryref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/table_function_ref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref.hpp +3 -1
- package/src/duckdb/src/include/duckdb/parser/tokens.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/transformer.hpp +33 -0
- package/src/duckdb/src/include/duckdb/planner/bind_context.hpp +2 -0
- package/src/duckdb/src/include/duckdb/planner/binder.hpp +15 -4
- package/src/duckdb/src/include/duckdb/planner/bound_result_modifier.hpp +3 -0
- package/src/duckdb/src/include/duckdb/planner/expression/bound_aggregate_expression.hpp +3 -0
- package/src/duckdb/src/include/duckdb/planner/expression_binder/base_select_binder.hpp +64 -0
- package/src/duckdb/src/include/duckdb/planner/expression_binder/having_binder.hpp +2 -2
- package/src/duckdb/src/include/duckdb/planner/expression_binder/order_binder.hpp +4 -1
- package/src/duckdb/src/include/duckdb/planner/expression_binder/qualify_binder.hpp +2 -2
- package/src/duckdb/src/include/duckdb/planner/expression_binder/select_binder.hpp +9 -38
- package/src/duckdb/src/include/duckdb/planner/expression_binder.hpp +1 -1
- package/src/duckdb/src/include/duckdb/planner/logical_tokens.hpp +1 -0
- package/src/duckdb/src/include/duckdb/planner/operator/list.hpp +1 -0
- package/src/duckdb/src/include/duckdb/planner/operator/logical_asof_join.hpp +22 -0
- package/src/duckdb/src/include/duckdb/planner/operator/logical_comparison_join.hpp +5 -2
- package/src/duckdb/src/include/duckdb/planner/operator/logical_distinct.hpp +3 -0
- package/src/duckdb/src/include/duckdb/planner/query_node/bound_select_node.hpp +8 -2
- package/src/duckdb/src/include/duckdb/storage/buffer/block_handle.hpp +2 -0
- package/src/duckdb/src/include/duckdb/storage/buffer_manager.hpp +76 -44
- package/src/duckdb/src/include/duckdb/storage/checkpoint/table_data_writer.hpp +3 -2
- package/src/duckdb/src/include/duckdb/storage/checkpoint_manager.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_compress.hpp +2 -2
- package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_fetch.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_scan.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_compress.hpp +2 -2
- package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_fetch.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_scan.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/data_pointer.hpp +5 -2
- package/src/duckdb/src/include/duckdb/storage/data_table.hpp +3 -3
- package/src/duckdb/src/include/duckdb/storage/index.hpp +4 -3
- package/src/duckdb/src/include/duckdb/storage/meta_block_reader.hpp +7 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/base_statistics.hpp +93 -29
- package/src/duckdb/src/include/duckdb/storage/statistics/column_statistics.hpp +22 -3
- package/src/duckdb/src/include/duckdb/storage/statistics/distinct_statistics.hpp +8 -6
- package/src/duckdb/src/include/duckdb/storage/statistics/list_stats.hpp +41 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/node_statistics.hpp +26 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats.hpp +114 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats_union.hpp +62 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/segment_statistics.hpp +2 -7
- package/src/duckdb/src/include/duckdb/storage/statistics/string_stats.hpp +74 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/struct_stats.hpp +42 -0
- package/src/duckdb/src/include/duckdb/storage/string_uncompressed.hpp +2 -3
- package/src/duckdb/src/include/duckdb/storage/table/column_checkpoint_state.hpp +2 -1
- package/src/duckdb/src/include/duckdb/storage/table/column_data.hpp +6 -3
- package/src/duckdb/src/include/duckdb/storage/table/column_data_checkpointer.hpp +3 -2
- package/src/duckdb/src/include/duckdb/storage/table/column_segment.hpp +7 -5
- package/src/duckdb/src/include/duckdb/storage/table/list_column_data.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/table/persistent_table_data.hpp +6 -2
- package/src/duckdb/src/include/duckdb/storage/table/row_group.hpp +10 -6
- package/src/duckdb/src/include/duckdb/storage/table/row_group_collection.hpp +8 -5
- package/src/duckdb/src/include/duckdb/storage/table/row_group_segment_tree.hpp +37 -0
- package/src/duckdb/src/include/duckdb/storage/table/scan_state.hpp +10 -1
- package/src/duckdb/src/include/duckdb/storage/table/segment_base.hpp +4 -3
- package/src/duckdb/src/include/duckdb/storage/table/segment_tree.hpp +271 -26
- package/src/duckdb/src/include/duckdb/storage/table/table_statistics.hpp +5 -0
- package/src/duckdb/src/include/duckdb/storage/table/update_segment.hpp +0 -1
- package/src/duckdb/src/include/duckdb/storage/write_ahead_log.hpp +1 -1
- package/src/duckdb/src/include/duckdb/transaction/local_storage.hpp +2 -2
- package/src/duckdb/src/include/duckdb.h +50 -2
- package/src/duckdb/src/include/duckdb.hpp +0 -1
- package/src/duckdb/src/main/capi/pending-c.cpp +16 -3
- package/src/duckdb/src/main/capi/result-c.cpp +27 -1
- package/src/duckdb/src/main/capi/stream-c.cpp +25 -0
- package/src/duckdb/src/main/client_context.cpp +38 -34
- package/src/duckdb/src/main/client_data.cpp +7 -6
- package/src/duckdb/src/main/config.cpp +70 -1
- package/src/duckdb/src/main/database.cpp +19 -2
- package/src/duckdb/src/main/extension/extension_install.cpp +7 -2
- package/src/duckdb/src/main/prepared_statement.cpp +4 -0
- package/src/duckdb/src/main/query_profiler.cpp +17 -15
- package/src/duckdb/src/main/relation/explain_relation.cpp +3 -3
- package/src/duckdb/src/main/relation.cpp +3 -2
- package/src/duckdb/src/optimizer/column_lifetime_analyzer.cpp +1 -0
- package/src/duckdb/src/optimizer/deliminator.cpp +1 -1
- package/src/duckdb/src/optimizer/filter_combiner.cpp +1 -1
- package/src/duckdb/src/optimizer/filter_pullup.cpp +3 -1
- package/src/duckdb/src/optimizer/filter_pushdown.cpp +14 -8
- package/src/duckdb/src/optimizer/join_order/cardinality_estimator.cpp +105 -71
- package/src/duckdb/src/optimizer/join_order/join_order_optimizer.cpp +31 -12
- package/src/duckdb/src/optimizer/optimizer.cpp +1 -0
- package/src/duckdb/src/optimizer/pullup/pullup_from_left.cpp +2 -2
- package/src/duckdb/src/optimizer/pushdown/pushdown_aggregate.cpp +33 -5
- package/src/duckdb/src/optimizer/pushdown/pushdown_cross_product.cpp +1 -1
- package/src/duckdb/src/optimizer/pushdown/pushdown_inner_join.cpp +3 -0
- package/src/duckdb/src/optimizer/pushdown/pushdown_left_join.cpp +5 -12
- package/src/duckdb/src/optimizer/pushdown/pushdown_mark_join.cpp +2 -2
- package/src/duckdb/src/optimizer/pushdown/pushdown_single_join.cpp +1 -1
- package/src/duckdb/src/optimizer/remove_unused_columns.cpp +1 -0
- package/src/duckdb/src/optimizer/rule/move_constants.cpp +10 -4
- package/src/duckdb/src/optimizer/rule/ordered_aggregate_optimizer.cpp +30 -0
- package/src/duckdb/src/optimizer/rule/regex_optimizations.cpp +9 -2
- package/src/duckdb/src/optimizer/statistics/expression/propagate_aggregate.cpp +9 -3
- package/src/duckdb/src/optimizer/statistics/expression/propagate_and_compress.cpp +6 -7
- package/src/duckdb/src/optimizer/statistics/expression/propagate_cast.cpp +14 -11
- package/src/duckdb/src/optimizer/statistics/expression/propagate_columnref.cpp +1 -1
- package/src/duckdb/src/optimizer/statistics/expression/propagate_comparison.cpp +13 -15
- package/src/duckdb/src/optimizer/statistics/expression/propagate_conjunction.cpp +0 -1
- package/src/duckdb/src/optimizer/statistics/expression/propagate_constant.cpp +3 -75
- package/src/duckdb/src/optimizer/statistics/expression/propagate_function.cpp +7 -2
- package/src/duckdb/src/optimizer/statistics/expression/propagate_operator.cpp +10 -0
- package/src/duckdb/src/optimizer/statistics/operator/propagate_aggregate.cpp +2 -3
- package/src/duckdb/src/optimizer/statistics/operator/propagate_filter.cpp +29 -32
- package/src/duckdb/src/optimizer/statistics/operator/propagate_join.cpp +5 -5
- package/src/duckdb/src/optimizer/statistics/operator/propagate_set_operation.cpp +3 -3
- package/src/duckdb/src/optimizer/statistics_propagator.cpp +2 -1
- package/src/duckdb/src/optimizer/unnest_rewriter.cpp +2 -2
- package/src/duckdb/src/parallel/meta_pipeline.cpp +0 -4
- package/src/duckdb/src/parser/common_table_expression_info.cpp +19 -0
- package/src/duckdb/src/parser/expression/between_expression.cpp +17 -0
- package/src/duckdb/src/parser/expression/case_expression.cpp +28 -0
- package/src/duckdb/src/parser/expression/cast_expression.cpp +17 -0
- package/src/duckdb/src/parser/expression/collate_expression.cpp +16 -0
- package/src/duckdb/src/parser/expression/columnref_expression.cpp +15 -0
- package/src/duckdb/src/parser/expression/comparison_expression.cpp +16 -0
- package/src/duckdb/src/parser/expression/conjunction_expression.cpp +17 -0
- package/src/duckdb/src/parser/expression/constant_expression.cpp +14 -0
- package/src/duckdb/src/parser/expression/default_expression.cpp +7 -0
- package/src/duckdb/src/parser/expression/function_expression.cpp +35 -0
- package/src/duckdb/src/parser/expression/lambda_expression.cpp +16 -0
- package/src/duckdb/src/parser/expression/operator_expression.cpp +15 -0
- package/src/duckdb/src/parser/expression/parameter_expression.cpp +15 -0
- package/src/duckdb/src/parser/expression/positional_reference_expression.cpp +14 -0
- package/src/duckdb/src/parser/expression/star_expression.cpp +26 -6
- package/src/duckdb/src/parser/expression/subquery_expression.cpp +20 -0
- package/src/duckdb/src/parser/expression/window_expression.cpp +43 -0
- package/src/duckdb/src/parser/parsed_data/alter_info.cpp +7 -3
- package/src/duckdb/src/parser/parsed_data/alter_scalar_function_info.cpp +56 -0
- package/src/duckdb/src/parser/parsed_data/alter_table_function_info.cpp +51 -0
- package/src/duckdb/src/parser/parsed_data/create_scalar_function_info.cpp +3 -2
- package/src/duckdb/src/parser/parsed_data/create_table_function_info.cpp +6 -0
- package/src/duckdb/src/parser/parsed_data/sample_options.cpp +22 -10
- package/src/duckdb/src/parser/parsed_expression.cpp +72 -0
- package/src/duckdb/src/parser/parsed_expression_iterator.cpp +15 -1
- package/src/duckdb/src/parser/query_node/recursive_cte_node.cpp +21 -0
- package/src/duckdb/src/parser/query_node/select_node.cpp +31 -0
- package/src/duckdb/src/parser/query_node/set_operation_node.cpp +17 -0
- package/src/duckdb/src/parser/query_node.cpp +51 -1
- package/src/duckdb/src/parser/result_modifier.cpp +78 -0
- package/src/duckdb/src/parser/statement/multi_statement.cpp +18 -0
- package/src/duckdb/src/parser/statement/select_statement.cpp +12 -0
- package/src/duckdb/src/parser/tableref/basetableref.cpp +21 -0
- package/src/duckdb/src/parser/tableref/emptytableref.cpp +4 -0
- package/src/duckdb/src/parser/tableref/expressionlistref.cpp +17 -0
- package/src/duckdb/src/parser/tableref/joinref.cpp +29 -0
- package/src/duckdb/src/parser/tableref/pivotref.cpp +373 -0
- package/src/duckdb/src/parser/tableref/subqueryref.cpp +15 -0
- package/src/duckdb/src/parser/tableref/table_function.cpp +17 -0
- package/src/duckdb/src/parser/tableref.cpp +49 -0
- package/src/duckdb/src/parser/transform/expression/transform_array_access.cpp +11 -0
- package/src/duckdb/src/parser/transform/expression/transform_bool_expr.cpp +1 -1
- package/src/duckdb/src/parser/transform/expression/transform_columnref.cpp +17 -2
- package/src/duckdb/src/parser/transform/expression/transform_function.cpp +63 -42
- package/src/duckdb/src/parser/transform/expression/transform_operator.cpp +1 -1
- package/src/duckdb/src/parser/transform/expression/transform_subquery.cpp +1 -1
- package/src/duckdb/src/parser/transform/helpers/transform_alias.cpp +12 -6
- package/src/duckdb/src/parser/transform/helpers/transform_cte.cpp +24 -0
- package/src/duckdb/src/parser/transform/helpers/transform_groupby.cpp +7 -0
- package/src/duckdb/src/parser/transform/helpers/transform_orderby.cpp +0 -7
- package/src/duckdb/src/parser/transform/helpers/transform_typename.cpp +3 -2
- package/src/duckdb/src/parser/transform/statement/transform_create_function.cpp +4 -0
- package/src/duckdb/src/parser/transform/statement/transform_create_view.cpp +4 -0
- package/src/duckdb/src/parser/transform/statement/transform_pivot_stmt.cpp +179 -0
- package/src/duckdb/src/parser/transform/statement/transform_rename.cpp +3 -4
- package/src/duckdb/src/parser/transform/statement/transform_select.cpp +8 -0
- package/src/duckdb/src/parser/transform/statement/transform_select_node.cpp +2 -3
- package/src/duckdb/src/parser/transform/tableref/transform_join.cpp +12 -1
- package/src/duckdb/src/parser/transform/tableref/transform_pivot.cpp +121 -0
- package/src/duckdb/src/parser/transform/tableref/transform_tableref.cpp +2 -0
- package/src/duckdb/src/parser/transformer.cpp +15 -3
- package/src/duckdb/src/planner/bind_context.cpp +18 -25
- package/src/duckdb/src/planner/binder/expression/bind_aggregate_expression.cpp +9 -7
- package/src/duckdb/src/planner/binder/expression/bind_columnref_expression.cpp +4 -3
- package/src/duckdb/src/planner/binder/expression/bind_function_expression.cpp +23 -12
- package/src/duckdb/src/planner/binder/expression/bind_lambda.cpp +3 -2
- package/src/duckdb/src/planner/binder/expression/bind_star_expression.cpp +176 -0
- package/src/duckdb/src/planner/binder/expression/bind_subquery_expression.cpp +4 -0
- package/src/duckdb/src/planner/binder/expression/bind_unnest_expression.cpp +163 -24
- package/src/duckdb/src/planner/binder/expression/bind_window_expression.cpp +2 -2
- package/src/duckdb/src/planner/binder/query_node/bind_select_node.cpp +109 -94
- package/src/duckdb/src/planner/binder/query_node/plan_query_node.cpp +11 -0
- package/src/duckdb/src/planner/binder/query_node/plan_select_node.cpp +9 -4
- package/src/duckdb/src/planner/binder/statement/bind_copy.cpp +5 -3
- package/src/duckdb/src/planner/binder/statement/bind_create.cpp +3 -2
- package/src/duckdb/src/planner/binder/statement/bind_create_table.cpp +9 -1
- package/src/duckdb/src/planner/binder/statement/bind_delete.cpp +1 -1
- package/src/duckdb/src/planner/binder/statement/bind_insert.cpp +12 -8
- package/src/duckdb/src/planner/binder/statement/bind_logical_plan.cpp +17 -0
- package/src/duckdb/src/planner/binder/statement/bind_update.cpp +4 -2
- package/src/duckdb/src/planner/binder/tableref/bind_joinref.cpp +19 -3
- package/src/duckdb/src/planner/binder/tableref/bind_pivot.cpp +366 -0
- package/src/duckdb/src/planner/binder/tableref/bind_table_function.cpp +11 -1
- package/src/duckdb/src/planner/binder/tableref/plan_cteref.cpp +1 -0
- package/src/duckdb/src/planner/binder/tableref/plan_joinref.cpp +61 -13
- package/src/duckdb/src/planner/binder.cpp +19 -24
- package/src/duckdb/src/planner/bound_result_modifier.cpp +27 -1
- package/src/duckdb/src/planner/expression/bound_aggregate_expression.cpp +9 -2
- package/src/duckdb/src/planner/expression/bound_expression.cpp +4 -0
- package/src/duckdb/src/planner/expression/bound_window_expression.cpp +1 -1
- package/src/duckdb/src/planner/expression_binder/base_select_binder.cpp +146 -0
- package/src/duckdb/src/planner/expression_binder/having_binder.cpp +6 -3
- package/src/duckdb/src/planner/expression_binder/qualify_binder.cpp +3 -3
- package/src/duckdb/src/planner/expression_binder/select_binder.cpp +1 -132
- package/src/duckdb/src/planner/expression_binder.cpp +10 -3
- package/src/duckdb/src/planner/expression_iterator.cpp +17 -10
- package/src/duckdb/src/planner/filter/constant_filter.cpp +4 -6
- package/src/duckdb/src/planner/logical_operator.cpp +7 -2
- package/src/duckdb/src/planner/logical_operator_visitor.cpp +6 -0
- package/src/duckdb/src/planner/operator/logical_asof_join.cpp +8 -0
- package/src/duckdb/src/planner/operator/logical_distinct.cpp +3 -0
- package/src/duckdb/src/planner/planner.cpp +2 -1
- package/src/duckdb/src/planner/pragma_handler.cpp +10 -2
- package/src/duckdb/src/planner/subquery/flatten_dependent_join.cpp +3 -1
- package/src/duckdb/src/storage/buffer_manager.cpp +44 -46
- package/src/duckdb/src/storage/checkpoint/row_group_writer.cpp +1 -1
- package/src/duckdb/src/storage/checkpoint/table_data_reader.cpp +4 -15
- package/src/duckdb/src/storage/checkpoint/table_data_writer.cpp +10 -4
- package/src/duckdb/src/storage/checkpoint_manager.cpp +9 -3
- package/src/duckdb/src/storage/compression/bitpacking.cpp +28 -24
- package/src/duckdb/src/storage/compression/fixed_size_uncompressed.cpp +43 -45
- package/src/duckdb/src/storage/compression/numeric_constant.cpp +9 -10
- package/src/duckdb/src/storage/compression/patas.cpp +1 -1
- package/src/duckdb/src/storage/compression/rle.cpp +19 -15
- package/src/duckdb/src/storage/compression/validity_uncompressed.cpp +5 -5
- package/src/duckdb/src/storage/data_table.cpp +20 -20
- package/src/duckdb/src/storage/index.cpp +12 -1
- package/src/duckdb/src/storage/local_storage.cpp +20 -23
- package/src/duckdb/src/storage/meta_block_reader.cpp +22 -0
- package/src/duckdb/src/storage/statistics/base_statistics.cpp +373 -128
- package/src/duckdb/src/storage/statistics/column_statistics.cpp +57 -3
- package/src/duckdb/src/storage/statistics/distinct_statistics.cpp +8 -9
- package/src/duckdb/src/storage/statistics/list_stats.cpp +121 -0
- package/src/duckdb/src/storage/statistics/numeric_stats.cpp +591 -0
- package/src/duckdb/src/storage/statistics/numeric_stats_union.cpp +65 -0
- package/src/duckdb/src/storage/statistics/segment_statistics.cpp +2 -11
- package/src/duckdb/src/storage/statistics/string_stats.cpp +273 -0
- package/src/duckdb/src/storage/statistics/struct_stats.cpp +133 -0
- package/src/duckdb/src/storage/storage_info.cpp +2 -2
- package/src/duckdb/src/storage/table/column_checkpoint_state.cpp +4 -10
- package/src/duckdb/src/storage/table/column_data.cpp +45 -46
- package/src/duckdb/src/storage/table/column_data_checkpointer.cpp +7 -8
- package/src/duckdb/src/storage/table/column_segment.cpp +13 -14
- package/src/duckdb/src/storage/table/list_column_data.cpp +41 -59
- package/src/duckdb/src/storage/table/persistent_table_data.cpp +2 -1
- package/src/duckdb/src/storage/table/row_group.cpp +38 -32
- package/src/duckdb/src/storage/table/row_group_collection.cpp +94 -78
- package/src/duckdb/src/storage/table/scan_state.cpp +22 -3
- package/src/duckdb/src/storage/table/standard_column_data.cpp +7 -6
- package/src/duckdb/src/storage/table/struct_column_data.cpp +16 -16
- package/src/duckdb/src/storage/table/table_statistics.cpp +27 -7
- package/src/duckdb/src/storage/table/update_segment.cpp +20 -18
- package/src/duckdb/src/storage/wal_replay.cpp +8 -5
- package/src/duckdb/src/storage/write_ahead_log.cpp +2 -2
- package/src/duckdb/src/transaction/commit_state.cpp +11 -7
- package/src/duckdb/src/verification/deserialized_statement_verifier.cpp +0 -1
- package/src/duckdb/third_party/libpg_query/include/nodes/nodes.hpp +35 -0
- package/src/duckdb/third_party/libpg_query/include/nodes/parsenodes.hpp +36 -2
- package/src/duckdb/third_party/libpg_query/include/nodes/primnodes.hpp +3 -3
- package/src/duckdb/third_party/libpg_query/include/parser/gram.hpp +1022 -530
- package/src/duckdb/third_party/libpg_query/include/parser/kwlist.hpp +8 -0
- package/src/duckdb/third_party/libpg_query/src_backend_parser_gram.cpp +24462 -22828
- package/src/duckdb/third_party/re2/re2/re2.cc +9 -0
- package/src/duckdb/third_party/re2/re2/re2.h +2 -0
- package/src/duckdb/ub_extension_icu_third_party_icu_i18n.cpp +4 -4
- package/src/duckdb/ub_extension_json_json_functions.cpp +2 -0
- package/src/duckdb/ub_src_common_serializer.cpp +2 -0
- package/src/duckdb/ub_src_execution_physical_plan.cpp +2 -0
- package/src/duckdb/ub_src_function_aggregate_distributive.cpp +2 -0
- package/src/duckdb/ub_src_function_scalar_bit.cpp +2 -0
- package/src/duckdb/ub_src_function_scalar_map.cpp +4 -0
- package/src/duckdb/ub_src_function_scalar_string.cpp +2 -0
- package/src/duckdb/ub_src_function_scalar_string_regexp.cpp +4 -0
- package/src/duckdb/ub_src_main_capi.cpp +2 -0
- package/src/duckdb/ub_src_optimizer_rule.cpp +2 -0
- package/src/duckdb/ub_src_parser.cpp +2 -0
- package/src/duckdb/ub_src_parser_parsed_data.cpp +4 -2
- package/src/duckdb/ub_src_parser_statement.cpp +2 -0
- package/src/duckdb/ub_src_parser_tableref.cpp +2 -0
- package/src/duckdb/ub_src_parser_transform_statement.cpp +2 -0
- package/src/duckdb/ub_src_parser_transform_tableref.cpp +2 -0
- package/src/duckdb/ub_src_planner_binder_expression.cpp +2 -0
- package/src/duckdb/ub_src_planner_binder_tableref.cpp +2 -0
- package/src/duckdb/ub_src_planner_expression_binder.cpp +2 -0
- package/src/duckdb/ub_src_planner_operator.cpp +2 -0
- package/src/duckdb/ub_src_storage_statistics.cpp +6 -6
- package/src/duckdb/ub_src_storage_table.cpp +0 -2
- package/src/duckdb_node.hpp +2 -1
- package/src/statement.cpp +5 -5
- package/src/utils.cpp +27 -2
- package/test/extension.test.ts +44 -26
- package/test/syntax_error.test.ts +3 -1
- package/filelist.cache +0 -0
- package/src/duckdb/src/include/duckdb/main/loadable_extension.hpp +0 -59
- package/src/duckdb/src/include/duckdb/storage/statistics/list_statistics.hpp +0 -36
- package/src/duckdb/src/include/duckdb/storage/statistics/numeric_statistics.hpp +0 -75
- package/src/duckdb/src/include/duckdb/storage/statistics/string_statistics.hpp +0 -49
- package/src/duckdb/src/include/duckdb/storage/statistics/struct_statistics.hpp +0 -36
- package/src/duckdb/src/include/duckdb/storage/statistics/validity_statistics.hpp +0 -45
- package/src/duckdb/src/parser/parsed_data/alter_function_info.cpp +0 -55
- package/src/duckdb/src/storage/statistics/list_statistics.cpp +0 -94
- package/src/duckdb/src/storage/statistics/numeric_statistics.cpp +0 -307
- package/src/duckdb/src/storage/statistics/string_statistics.cpp +0 -220
- package/src/duckdb/src/storage/statistics/struct_statistics.cpp +0 -108
- package/src/duckdb/src/storage/statistics/validity_statistics.cpp +0 -91
- package/src/duckdb/src/storage/table/segment_tree.cpp +0 -179
|
@@ -30,25 +30,20 @@ PhysicalPiecewiseMergeJoin::PhysicalPiecewiseMergeJoin(LogicalOperator &op, uniq
|
|
|
30
30
|
switch (cond.comparison) {
|
|
31
31
|
case ExpressionType::COMPARE_LESSTHAN:
|
|
32
32
|
case ExpressionType::COMPARE_LESSTHANOREQUALTO:
|
|
33
|
-
lhs_orders.emplace_back(
|
|
34
|
-
|
|
35
|
-
rhs_orders.emplace_back(
|
|
36
|
-
BoundOrderByNode(OrderType::ASCENDING, OrderByNullType::NULLS_LAST, std::move(right)));
|
|
33
|
+
lhs_orders.emplace_back(OrderType::ASCENDING, OrderByNullType::NULLS_LAST, std::move(left));
|
|
34
|
+
rhs_orders.emplace_back(OrderType::ASCENDING, OrderByNullType::NULLS_LAST, std::move(right));
|
|
37
35
|
break;
|
|
38
36
|
case ExpressionType::COMPARE_GREATERTHAN:
|
|
39
37
|
case ExpressionType::COMPARE_GREATERTHANOREQUALTO:
|
|
40
|
-
lhs_orders.emplace_back(
|
|
41
|
-
|
|
42
|
-
rhs_orders.emplace_back(
|
|
43
|
-
BoundOrderByNode(OrderType::DESCENDING, OrderByNullType::NULLS_LAST, std::move(right)));
|
|
38
|
+
lhs_orders.emplace_back(OrderType::DESCENDING, OrderByNullType::NULLS_LAST, std::move(left));
|
|
39
|
+
rhs_orders.emplace_back(OrderType::DESCENDING, OrderByNullType::NULLS_LAST, std::move(right));
|
|
44
40
|
break;
|
|
45
41
|
case ExpressionType::COMPARE_NOTEQUAL:
|
|
46
42
|
case ExpressionType::COMPARE_DISTINCT_FROM:
|
|
47
43
|
// Allowed in multi-predicate joins, but can't be first/sort.
|
|
48
44
|
D_ASSERT(!lhs_orders.empty());
|
|
49
|
-
lhs_orders.emplace_back(
|
|
50
|
-
rhs_orders.emplace_back(
|
|
51
|
-
BoundOrderByNode(OrderType::INVALID, OrderByNullType::NULLS_LAST, std::move(right)));
|
|
45
|
+
lhs_orders.emplace_back(OrderType::INVALID, OrderByNullType::NULLS_LAST, std::move(left));
|
|
46
|
+
rhs_orders.emplace_back(OrderType::INVALID, OrderByNullType::NULLS_LAST, std::move(right));
|
|
52
47
|
break;
|
|
53
48
|
|
|
54
49
|
default:
|
|
@@ -46,7 +46,7 @@ void PhysicalRangeJoin::LocalSortedTable::Sink(DataChunk &input, GlobalSortState
|
|
|
46
46
|
|
|
47
47
|
// Only sort the primary key
|
|
48
48
|
DataChunk join_head;
|
|
49
|
-
join_head.data.emplace_back(
|
|
49
|
+
join_head.data.emplace_back(keys.data[0]);
|
|
50
50
|
join_head.SetCardinality(keys.size());
|
|
51
51
|
|
|
52
52
|
// Sink the data into the local sort state
|
|
@@ -334,7 +334,9 @@ idx_t PhysicalRangeJoin::SelectJoinTail(const ExpressionType &condition, Vector
|
|
|
334
334
|
case ExpressionType::COMPARE_DISTINCT_FROM:
|
|
335
335
|
return VectorOperations::DistinctFrom(left, right, sel, count, true_sel, nullptr);
|
|
336
336
|
case ExpressionType::COMPARE_NOT_DISTINCT_FROM:
|
|
337
|
+
return VectorOperations::NotDistinctFrom(left, right, sel, count, true_sel, nullptr);
|
|
337
338
|
case ExpressionType::COMPARE_EQUAL:
|
|
339
|
+
return VectorOperations::Equals(left, right, sel, count, true_sel, nullptr);
|
|
338
340
|
default:
|
|
339
341
|
throw InternalException("Unsupported comparison type for PhysicalRangeJoin");
|
|
340
342
|
}
|
|
@@ -121,6 +121,9 @@ bool TryCastFloatingValueCommaSeparated(const string_t &value_str, const Logical
|
|
|
121
121
|
}
|
|
122
122
|
|
|
123
123
|
bool BaseCSVReader::TryCastValue(const Value &value, const LogicalType &sql_type) {
|
|
124
|
+
if (value.IsNull()) {
|
|
125
|
+
return true;
|
|
126
|
+
}
|
|
124
127
|
if (options.has_format[LogicalTypeId::DATE] && sql_type.id() == LogicalTypeId::DATE) {
|
|
125
128
|
date_t result;
|
|
126
129
|
string error_message;
|
|
@@ -495,12 +498,12 @@ bool BaseCSVReader::Flush(DataChunk &insert_chunk, bool try_add_line) {
|
|
|
495
498
|
}
|
|
496
499
|
|
|
497
500
|
// figure out the exact line number
|
|
501
|
+
UnifiedVectorFormat inserted_column_data;
|
|
502
|
+
insert_chunk.data[col_idx].ToUnifiedFormat(parse_chunk.size(), inserted_column_data);
|
|
498
503
|
idx_t row_idx;
|
|
499
504
|
for (row_idx = 0; row_idx < parse_chunk.size(); row_idx++) {
|
|
500
|
-
auto &inserted_column = insert_chunk.data[col_idx];
|
|
501
505
|
auto &parsed_column = parse_chunk.data[col_idx];
|
|
502
|
-
|
|
503
|
-
if (FlatVector::IsNull(inserted_column, row_idx) && !FlatVector::IsNull(parsed_column, row_idx)) {
|
|
506
|
+
if (!inserted_column_data.validity.RowIsValid(row_idx) && !FlatVector::IsNull(parsed_column, row_idx)) {
|
|
504
507
|
break;
|
|
505
508
|
}
|
|
506
509
|
}
|
|
@@ -555,7 +555,7 @@ void BufferedCSVReader::DetectCandidateTypes(const vector<LogicalType> &type_can
|
|
|
555
555
|
// try formatting for date types if the user did not specify one and it starts with numeric values.
|
|
556
556
|
string separator;
|
|
557
557
|
if (has_format_candidates.count(sql_type.id()) && !original_options.has_format[sql_type.id()] &&
|
|
558
|
-
StartsWithNumericDate(separator, StringValue::Get(dummy_val))) {
|
|
558
|
+
!dummy_val.IsNull() && StartsWithNumericDate(separator, StringValue::Get(dummy_val))) {
|
|
559
559
|
// generate date format candidates the first time through
|
|
560
560
|
auto &type_format_candidates = format_candidates[sql_type.id()];
|
|
561
561
|
const auto had_format_candidates = has_format_candidates[sql_type.id()];
|
|
@@ -870,16 +870,7 @@ vector<LogicalType> BufferedCSVReader::SniffCSV(const vector<LogicalType> &reque
|
|
|
870
870
|
// #######
|
|
871
871
|
// ### type detection (initial)
|
|
872
872
|
// #######
|
|
873
|
-
|
|
874
|
-
vector<LogicalType> type_candidates = {
|
|
875
|
-
LogicalType::VARCHAR,
|
|
876
|
-
LogicalType::TIMESTAMP,
|
|
877
|
-
LogicalType::DATE,
|
|
878
|
-
LogicalType::TIME,
|
|
879
|
-
LogicalType::DOUBLE,
|
|
880
|
-
/* LogicalType::FLOAT,*/ LogicalType::BIGINT,
|
|
881
|
-
/*LogicalType::INTEGER,*/ /*LogicalType::SMALLINT, LogicalType::TINYINT,*/ LogicalType::BOOLEAN,
|
|
882
|
-
LogicalType::SQLNULL};
|
|
873
|
+
|
|
883
874
|
// format template candidates, ordered by descending specificity (~ from high to low)
|
|
884
875
|
std::map<LogicalTypeId, vector<const char *>> format_template_candidates = {
|
|
885
876
|
{LogicalTypeId::DATE, {"%m-%d-%Y", "%m-%d-%y", "%d-%m-%Y", "%d-%m-%y", "%Y-%m-%d", "%y-%m-%d"}},
|
|
@@ -890,8 +881,8 @@ vector<LogicalType> BufferedCSVReader::SniffCSV(const vector<LogicalType> &reque
|
|
|
890
881
|
vector<vector<LogicalType>> best_sql_types_candidates;
|
|
891
882
|
map<LogicalTypeId, vector<string>> best_format_candidates;
|
|
892
883
|
DataChunk best_header_row;
|
|
893
|
-
DetectCandidateTypes(
|
|
894
|
-
best_sql_types_candidates, best_format_candidates, best_header_row);
|
|
884
|
+
DetectCandidateTypes(options.auto_type_candidates, format_template_candidates, info_candidates, original_options,
|
|
885
|
+
best_num_cols, best_sql_types_candidates, best_format_candidates, best_header_row);
|
|
895
886
|
|
|
896
887
|
if (best_format_candidates.empty() || best_header_row.size() == 0) {
|
|
897
888
|
throw InvalidInputException(
|
|
@@ -939,7 +930,8 @@ vector<LogicalType> BufferedCSVReader::SniffCSV(const vector<LogicalType> &reque
|
|
|
939
930
|
// #######
|
|
940
931
|
// ### type detection (refining)
|
|
941
932
|
// #######
|
|
942
|
-
return RefineTypeDetection(
|
|
933
|
+
return RefineTypeDetection(options.auto_type_candidates, requested_types, best_sql_types_candidates,
|
|
934
|
+
best_format_candidates);
|
|
943
935
|
}
|
|
944
936
|
|
|
945
937
|
bool BufferedCSVReader::TryParseComplexCSV(DataChunk &insert_chunk, string &error_message) {
|
|
@@ -85,8 +85,8 @@ static string CreateDirRecursive(const vector<idx_t> &cols, const vector<string>
|
|
|
85
85
|
CreateDir(path, fs);
|
|
86
86
|
|
|
87
87
|
for (idx_t i = 0; i < cols.size(); i++) {
|
|
88
|
-
auto partition_col_name = names[cols[i]];
|
|
89
|
-
auto partition_value = values[i];
|
|
88
|
+
const auto &partition_col_name = names[cols[i]];
|
|
89
|
+
const auto &partition_value = values[i];
|
|
90
90
|
string p_dir = partition_col_name + "=" + partition_value.ToString();
|
|
91
91
|
path = fs.JoinPath(path, p_dir);
|
|
92
92
|
CreateDir(path, fs);
|
|
@@ -1,6 +1,7 @@
|
|
|
1
1
|
#include "duckdb/execution/operator/projection/physical_projection.hpp"
|
|
2
2
|
#include "duckdb/parallel/thread_context.hpp"
|
|
3
3
|
#include "duckdb/execution/expression_executor.hpp"
|
|
4
|
+
#include "duckdb/planner/expression/bound_reference_expression.hpp"
|
|
4
5
|
|
|
5
6
|
namespace duckdb {
|
|
6
7
|
|
|
@@ -35,6 +36,39 @@ unique_ptr<OperatorState> PhysicalProjection::GetOperatorState(ExecutionContext
|
|
|
35
36
|
return make_unique<ProjectionState>(context, select_list);
|
|
36
37
|
}
|
|
37
38
|
|
|
39
|
+
unique_ptr<PhysicalOperator>
|
|
40
|
+
PhysicalProjection::CreateJoinProjection(vector<LogicalType> proj_types, const vector<LogicalType> &lhs_types,
|
|
41
|
+
const vector<LogicalType> &rhs_types, const vector<idx_t> &left_projection_map,
|
|
42
|
+
const vector<idx_t> &right_projection_map, const idx_t estimated_cardinality) {
|
|
43
|
+
|
|
44
|
+
vector<unique_ptr<Expression>> proj_selects;
|
|
45
|
+
proj_selects.reserve(proj_types.size());
|
|
46
|
+
|
|
47
|
+
if (left_projection_map.empty()) {
|
|
48
|
+
for (storage_t i = 0; i < lhs_types.size(); ++i) {
|
|
49
|
+
proj_selects.emplace_back(make_unique<BoundReferenceExpression>(lhs_types[i], i));
|
|
50
|
+
}
|
|
51
|
+
} else {
|
|
52
|
+
for (auto i : left_projection_map) {
|
|
53
|
+
proj_selects.emplace_back(make_unique<BoundReferenceExpression>(lhs_types[i], i));
|
|
54
|
+
}
|
|
55
|
+
}
|
|
56
|
+
const auto left_cols = lhs_types.size();
|
|
57
|
+
|
|
58
|
+
if (right_projection_map.empty()) {
|
|
59
|
+
for (storage_t i = 0; i < rhs_types.size(); ++i) {
|
|
60
|
+
proj_selects.emplace_back(make_unique<BoundReferenceExpression>(rhs_types[i], left_cols + i));
|
|
61
|
+
}
|
|
62
|
+
|
|
63
|
+
} else {
|
|
64
|
+
for (auto i : right_projection_map) {
|
|
65
|
+
proj_selects.emplace_back(make_unique<BoundReferenceExpression>(rhs_types[i], left_cols + i));
|
|
66
|
+
}
|
|
67
|
+
}
|
|
68
|
+
|
|
69
|
+
return make_unique<PhysicalProjection>(std::move(proj_types), std::move(proj_selects), estimated_cardinality);
|
|
70
|
+
}
|
|
71
|
+
|
|
38
72
|
string PhysicalProjection::ParamsToString() const {
|
|
39
73
|
string extra_info;
|
|
40
74
|
for (auto &expr : select_list) {
|
|
@@ -12,13 +12,28 @@ namespace duckdb {
|
|
|
12
12
|
PhysicalPositionalScan::PhysicalPositionalScan(vector<LogicalType> types, unique_ptr<PhysicalOperator> left,
|
|
13
13
|
unique_ptr<PhysicalOperator> right)
|
|
14
14
|
: PhysicalOperator(PhysicalOperatorType::POSITIONAL_SCAN, std::move(types),
|
|
15
|
-
|
|
15
|
+
MaxValue(left->estimated_cardinality, right->estimated_cardinality)) {
|
|
16
16
|
|
|
17
17
|
// Manage the children ourselves
|
|
18
|
-
|
|
19
|
-
|
|
20
|
-
|
|
21
|
-
|
|
18
|
+
if (left->type == PhysicalOperatorType::TABLE_SCAN) {
|
|
19
|
+
child_tables.emplace_back(std::move(left));
|
|
20
|
+
} else if (left->type == PhysicalOperatorType::POSITIONAL_SCAN) {
|
|
21
|
+
auto &left_scan = (PhysicalPositionalScan &)*left;
|
|
22
|
+
child_tables = std::move(left_scan.child_tables);
|
|
23
|
+
} else {
|
|
24
|
+
throw InternalException("Invalid left input for PhysicalPositionalScan");
|
|
25
|
+
}
|
|
26
|
+
|
|
27
|
+
if (right->type == PhysicalOperatorType::TABLE_SCAN) {
|
|
28
|
+
child_tables.emplace_back(std::move(right));
|
|
29
|
+
} else if (right->type == PhysicalOperatorType::POSITIONAL_SCAN) {
|
|
30
|
+
auto &right_scan = (PhysicalPositionalScan &)*right;
|
|
31
|
+
auto &right_tables = right_scan.child_tables;
|
|
32
|
+
child_tables.reserve(child_tables.size() + right_tables.size());
|
|
33
|
+
std::move(right_tables.begin(), right_tables.end(), std::back_inserter(child_tables));
|
|
34
|
+
} else {
|
|
35
|
+
throw InternalException("Invalid right input for PhysicalPositionalScan");
|
|
36
|
+
}
|
|
22
37
|
}
|
|
23
38
|
|
|
24
39
|
class PositionalScanGlobalSourceState : public GlobalSourceState {
|
|
@@ -15,10 +15,11 @@ PhysicalCreateType::PhysicalCreateType(unique_ptr<CreateTypeInfo> info, idx_t es
|
|
|
15
15
|
//===--------------------------------------------------------------------===//
|
|
16
16
|
class CreateTypeGlobalState : public GlobalSinkState {
|
|
17
17
|
public:
|
|
18
|
-
explicit CreateTypeGlobalState(ClientContext &context) :
|
|
18
|
+
explicit CreateTypeGlobalState(ClientContext &context) : result(LogicalType::VARCHAR) {
|
|
19
19
|
}
|
|
20
|
-
|
|
21
|
-
|
|
20
|
+
Vector result;
|
|
21
|
+
idx_t size = 0;
|
|
22
|
+
idx_t capacity = STANDARD_VECTOR_SIZE;
|
|
22
23
|
};
|
|
23
24
|
|
|
24
25
|
unique_ptr<GlobalSinkState> PhysicalCreateType::GetGlobalSinkState(ClientContext &context) const {
|
|
@@ -28,7 +29,7 @@ unique_ptr<GlobalSinkState> PhysicalCreateType::GetGlobalSinkState(ClientContext
|
|
|
28
29
|
SinkResultType PhysicalCreateType::Sink(ExecutionContext &context, GlobalSinkState &gstate_p, LocalSinkState &lstate_p,
|
|
29
30
|
DataChunk &input) const {
|
|
30
31
|
auto &gstate = (CreateTypeGlobalState &)gstate_p;
|
|
31
|
-
idx_t total_row_count = gstate.
|
|
32
|
+
idx_t total_row_count = gstate.size + input.size();
|
|
32
33
|
if (total_row_count > NumericLimits<uint32_t>::Maximum()) {
|
|
33
34
|
throw InvalidInputException("Attempted to create ENUM of size %llu, which exceeds the maximum size of %llu",
|
|
34
35
|
total_row_count, NumericLimits<uint32_t>::Maximum());
|
|
@@ -36,15 +37,23 @@ SinkResultType PhysicalCreateType::Sink(ExecutionContext &context, GlobalSinkSta
|
|
|
36
37
|
UnifiedVectorFormat sdata;
|
|
37
38
|
input.data[0].ToUnifiedFormat(input.size(), sdata);
|
|
38
39
|
|
|
40
|
+
if (total_row_count > gstate.capacity) {
|
|
41
|
+
// We must resize our result vector
|
|
42
|
+
gstate.result.Resize(gstate.capacity, gstate.capacity * 2);
|
|
43
|
+
gstate.capacity *= 2;
|
|
44
|
+
}
|
|
45
|
+
|
|
46
|
+
auto src_ptr = (string_t *)sdata.data;
|
|
47
|
+
auto result_ptr = FlatVector::GetData<string_t>(gstate.result);
|
|
39
48
|
// Input vector has NULL value, we just throw an exception
|
|
40
49
|
for (idx_t i = 0; i < input.size(); i++) {
|
|
41
50
|
idx_t idx = sdata.sel->get_index(i);
|
|
42
51
|
if (!sdata.validity.RowIsValid(idx)) {
|
|
43
52
|
throw InvalidInputException("Attempted to create ENUM type with NULL value!");
|
|
44
53
|
}
|
|
54
|
+
result_ptr[gstate.size++] =
|
|
55
|
+
StringVector::AddStringOrBlob(gstate.result, src_ptr[idx].GetDataUnsafe(), src_ptr[idx].GetSize());
|
|
45
56
|
}
|
|
46
|
-
|
|
47
|
-
gstate.collection.Append(input);
|
|
48
57
|
return SinkResultType::NEED_MORE_INPUT;
|
|
49
58
|
}
|
|
50
59
|
|
|
@@ -72,44 +81,15 @@ void PhysicalCreateType::GetData(ExecutionContext &context, DataChunk &chunk, Gl
|
|
|
72
81
|
|
|
73
82
|
if (IsSink()) {
|
|
74
83
|
D_ASSERT(info->type == LogicalType::INVALID);
|
|
75
|
-
|
|
76
84
|
auto &g_sink_state = (CreateTypeGlobalState &)*sink_state;
|
|
77
|
-
|
|
78
|
-
|
|
79
|
-
idx_t total_row_count = collection.Count();
|
|
80
|
-
|
|
81
|
-
ColumnDataScanState scan_state;
|
|
82
|
-
collection.InitializeScan(scan_state);
|
|
83
|
-
|
|
84
|
-
DataChunk scan_chunk;
|
|
85
|
-
collection.InitializeScanChunk(scan_chunk);
|
|
86
|
-
|
|
87
|
-
Vector result(LogicalType::VARCHAR, total_row_count);
|
|
88
|
-
auto result_ptr = FlatVector::GetData<string_t>(result);
|
|
89
|
-
|
|
90
|
-
idx_t offset = 0;
|
|
91
|
-
while (collection.Scan(scan_state, scan_chunk)) {
|
|
92
|
-
idx_t src_row_count = scan_chunk.size();
|
|
93
|
-
auto &src_vec = scan_chunk.data[0];
|
|
94
|
-
D_ASSERT(src_vec.GetVectorType() == VectorType::FLAT_VECTOR);
|
|
95
|
-
D_ASSERT(src_vec.GetType().id() == LogicalType::VARCHAR);
|
|
96
|
-
|
|
97
|
-
auto src_ptr = FlatVector::GetData<string_t>(src_vec);
|
|
98
|
-
|
|
99
|
-
for (idx_t i = 0; i < src_row_count; i++) {
|
|
100
|
-
idx_t target_index = offset + i;
|
|
101
|
-
result_ptr[target_index] =
|
|
102
|
-
StringVector::AddStringOrBlob(result, src_ptr[i].GetDataUnsafe(), src_ptr[i].GetSize());
|
|
103
|
-
}
|
|
104
|
-
|
|
105
|
-
offset += src_row_count;
|
|
106
|
-
}
|
|
107
|
-
|
|
108
|
-
info->type = LogicalType::ENUM(info->name, result, total_row_count);
|
|
85
|
+
info->type = LogicalType::ENUM(info->name, g_sink_state.result, g_sink_state.size);
|
|
109
86
|
}
|
|
110
87
|
|
|
111
88
|
auto &catalog = Catalog::GetCatalog(context.client, info->catalog);
|
|
112
|
-
catalog.CreateType(context.client, info.get());
|
|
89
|
+
auto catalog_entry = catalog.CreateType(context.client, info.get());
|
|
90
|
+
D_ASSERT(catalog_entry->type == CatalogType::TYPE_ENTRY);
|
|
91
|
+
auto catalog_type = (TypeCatalogEntry *)catalog_entry;
|
|
92
|
+
LogicalType::SetCatalog(info->type, catalog_type);
|
|
113
93
|
state.finished = true;
|
|
114
94
|
}
|
|
115
95
|
|
|
@@ -62,6 +62,18 @@ PartitionableHashTable::PartitionableHashTable(ClientContext &context, Allocator
|
|
|
62
62
|
for (hash_t r = 0; r < partition_info.n_partitions; r++) {
|
|
63
63
|
sel_vectors[r].Initialize();
|
|
64
64
|
}
|
|
65
|
+
|
|
66
|
+
RowLayout layout;
|
|
67
|
+
layout.Initialize(group_types, AggregateObject::CreateAggregateObjects(bindings));
|
|
68
|
+
tuple_size = layout.GetRowWidth();
|
|
69
|
+
}
|
|
70
|
+
|
|
71
|
+
HtEntryType PartitionableHashTable::GetHTEntrySize() {
|
|
72
|
+
// we need at least STANDARD_VECTOR_SIZE entries to fit in the hash table
|
|
73
|
+
if (GroupedAggregateHashTable::GetMaxCapacity(HtEntryType::HT_WIDTH_32, tuple_size) < STANDARD_VECTOR_SIZE) {
|
|
74
|
+
return HtEntryType::HT_WIDTH_64;
|
|
75
|
+
}
|
|
76
|
+
return HtEntryType::HT_WIDTH_32;
|
|
65
77
|
}
|
|
66
78
|
|
|
67
79
|
idx_t PartitionableHashTable::ListAddChunk(HashTableList &list, DataChunk &groups, Vector &group_hashes,
|
|
@@ -74,7 +86,7 @@ idx_t PartitionableHashTable::ListAddChunk(HashTableList &list, DataChunk &group
|
|
|
74
86
|
list.back()->Finalize();
|
|
75
87
|
}
|
|
76
88
|
list.push_back(make_unique<GroupedAggregateHashTable>(context, allocator, group_types, payload_types, bindings,
|
|
77
|
-
|
|
89
|
+
GetHTEntrySize()));
|
|
78
90
|
}
|
|
79
91
|
return list.back()->AddChunk(groups, group_hashes, payload, filter);
|
|
80
92
|
}
|
|
@@ -141,7 +153,7 @@ void PartitionableHashTable::Partition() {
|
|
|
141
153
|
for (auto &unpartitioned_ht : unpartitioned_hts) {
|
|
142
154
|
for (idx_t r = 0; r < partition_info.n_partitions; r++) {
|
|
143
155
|
radix_partitioned_hts[r].push_back(make_unique<GroupedAggregateHashTable>(
|
|
144
|
-
context, allocator, group_types, payload_types, bindings,
|
|
156
|
+
context, allocator, group_types, payload_types, bindings, GetHTEntrySize()));
|
|
145
157
|
partition_hts[r] = radix_partitioned_hts[r].back().get();
|
|
146
158
|
}
|
|
147
159
|
unpartitioned_ht->Partition(partition_hts, partition_info.radix_mask, partition_info.RADIX_SHIFT);
|
|
@@ -9,7 +9,8 @@
|
|
|
9
9
|
#include "duckdb/parser/expression/comparison_expression.hpp"
|
|
10
10
|
#include "duckdb/planner/expression/bound_aggregate_expression.hpp"
|
|
11
11
|
#include "duckdb/planner/operator/logical_aggregate.hpp"
|
|
12
|
-
#include "duckdb/
|
|
12
|
+
#include "duckdb/function/function_binder.hpp"
|
|
13
|
+
|
|
13
14
|
namespace duckdb {
|
|
14
15
|
|
|
15
16
|
static uint32_t RequiredBitsForValue(uint32_t n) {
|
|
@@ -50,23 +51,20 @@ static bool CanUsePerfectHashAggregate(ClientContext &context, LogicalAggregate
|
|
|
50
51
|
// for small types we can just set the stats to [type_min, type_max]
|
|
51
52
|
switch (group_type.InternalType()) {
|
|
52
53
|
case PhysicalType::INT8:
|
|
53
|
-
stats = make_unique<NumericStatistics>(group_type, Value::MinimumValue(group_type),
|
|
54
|
-
Value::MaximumValue(group_type), StatisticsType::LOCAL_STATS);
|
|
55
|
-
break;
|
|
56
54
|
case PhysicalType::INT16:
|
|
57
|
-
stats = make_unique<NumericStatistics>(group_type, Value::MinimumValue(group_type),
|
|
58
|
-
Value::MaximumValue(group_type), StatisticsType::LOCAL_STATS);
|
|
59
55
|
break;
|
|
60
56
|
default:
|
|
61
57
|
// type is too large and there are no stats: skip perfect hashing
|
|
62
58
|
return false;
|
|
63
59
|
}
|
|
64
|
-
//
|
|
65
|
-
stats
|
|
60
|
+
// construct stats with the min and max value of the type
|
|
61
|
+
stats = NumericStats::CreateUnknown(group_type).ToUnique();
|
|
62
|
+
NumericStats::SetMin(*stats, Value::MinimumValue(group_type));
|
|
63
|
+
NumericStats::SetMax(*stats, Value::MaximumValue(group_type));
|
|
66
64
|
}
|
|
67
|
-
auto &nstats =
|
|
65
|
+
auto &nstats = *stats;
|
|
68
66
|
|
|
69
|
-
if (
|
|
67
|
+
if (!NumericStats::HasMinMax(nstats)) {
|
|
70
68
|
return false;
|
|
71
69
|
}
|
|
72
70
|
// we have a min and a max value for the stats: use that to figure out how many bits we have
|
|
@@ -75,17 +73,17 @@ static bool CanUsePerfectHashAggregate(ClientContext &context, LogicalAggregate
|
|
|
75
73
|
int64_t range;
|
|
76
74
|
switch (group_type.InternalType()) {
|
|
77
75
|
case PhysicalType::INT8:
|
|
78
|
-
range = int64_t(
|
|
76
|
+
range = int64_t(NumericStats::GetMax<int8_t>(nstats)) - int64_t(NumericStats::GetMin<int8_t>(nstats));
|
|
79
77
|
break;
|
|
80
78
|
case PhysicalType::INT16:
|
|
81
|
-
range = int64_t(
|
|
79
|
+
range = int64_t(NumericStats::GetMax<int16_t>(nstats)) - int64_t(NumericStats::GetMin<int16_t>(nstats));
|
|
82
80
|
break;
|
|
83
81
|
case PhysicalType::INT32:
|
|
84
|
-
range = int64_t(
|
|
82
|
+
range = int64_t(NumericStats::GetMax<int32_t>(nstats)) - int64_t(NumericStats::GetMin<int32_t>(nstats));
|
|
85
83
|
break;
|
|
86
84
|
case PhysicalType::INT64:
|
|
87
|
-
if (!TrySubtractOperator::Operation(
|
|
88
|
-
|
|
85
|
+
if (!TrySubtractOperator::Operation(NumericStats::GetMax<int64_t>(nstats),
|
|
86
|
+
NumericStats::GetMin<int64_t>(nstats), range)) {
|
|
89
87
|
return false;
|
|
90
88
|
}
|
|
91
89
|
break;
|
|
@@ -169,13 +167,20 @@ PhysicalPlanGenerator::ExtractAggregateExpressions(unique_ptr<PhysicalOperator>
|
|
|
169
167
|
vector<unique_ptr<Expression>> expressions;
|
|
170
168
|
vector<LogicalType> types;
|
|
171
169
|
|
|
170
|
+
// bind sorted aggregates
|
|
171
|
+
for (auto &aggr : aggregates) {
|
|
172
|
+
auto &bound_aggr = (BoundAggregateExpression &)*aggr;
|
|
173
|
+
if (bound_aggr.order_bys) {
|
|
174
|
+
// sorted aggregate!
|
|
175
|
+
FunctionBinder::BindSortedAggregate(context, bound_aggr, groups);
|
|
176
|
+
}
|
|
177
|
+
}
|
|
172
178
|
for (auto &group : groups) {
|
|
173
179
|
auto ref = make_unique<BoundReferenceExpression>(group->return_type, expressions.size());
|
|
174
180
|
types.push_back(group->return_type);
|
|
175
181
|
expressions.push_back(std::move(group));
|
|
176
182
|
group = std::move(ref);
|
|
177
183
|
}
|
|
178
|
-
|
|
179
184
|
for (auto &aggr : aggregates) {
|
|
180
185
|
auto &bound_aggr = (BoundAggregateExpression &)*aggr;
|
|
181
186
|
for (auto &child : bound_aggr.children) {
|
|
@@ -0,0 +1,97 @@
|
|
|
1
|
+
#include "duckdb/execution/operator/aggregate/physical_window.hpp"
|
|
2
|
+
#include "duckdb/execution/operator/join/physical_iejoin.hpp"
|
|
3
|
+
#include "duckdb/execution/operator/projection/physical_projection.hpp"
|
|
4
|
+
#include "duckdb/execution/physical_plan_generator.hpp"
|
|
5
|
+
#include "duckdb/main/client_context.hpp"
|
|
6
|
+
#include "duckdb/planner/expression/bound_constant_expression.hpp"
|
|
7
|
+
#include "duckdb/planner/expression/bound_reference_expression.hpp"
|
|
8
|
+
#include "duckdb/planner/expression/bound_window_expression.hpp"
|
|
9
|
+
#include "duckdb/planner/operator/logical_asof_join.hpp"
|
|
10
|
+
|
|
11
|
+
namespace duckdb {
|
|
12
|
+
|
|
13
|
+
unique_ptr<PhysicalOperator> PhysicalPlanGenerator::CreatePlan(LogicalAsOfJoin &op) {
|
|
14
|
+
// now visit the children
|
|
15
|
+
D_ASSERT(op.children.size() == 2);
|
|
16
|
+
idx_t lhs_cardinality = op.children[0]->EstimateCardinality(context);
|
|
17
|
+
idx_t rhs_cardinality = op.children[1]->EstimateCardinality(context);
|
|
18
|
+
auto left = CreatePlan(*op.children[0]);
|
|
19
|
+
auto right = CreatePlan(*op.children[1]);
|
|
20
|
+
D_ASSERT(left && right);
|
|
21
|
+
|
|
22
|
+
// Validate
|
|
23
|
+
vector<idx_t> equi_indexes;
|
|
24
|
+
auto asof_idx = op.conditions.size();
|
|
25
|
+
for (size_t c = 0; c < op.conditions.size(); ++c) {
|
|
26
|
+
auto &cond = op.conditions[c];
|
|
27
|
+
switch (cond.comparison) {
|
|
28
|
+
case ExpressionType::COMPARE_EQUAL:
|
|
29
|
+
case ExpressionType::COMPARE_NOT_DISTINCT_FROM:
|
|
30
|
+
equi_indexes.emplace_back(c);
|
|
31
|
+
break;
|
|
32
|
+
case ExpressionType::COMPARE_GREATERTHANOREQUALTO:
|
|
33
|
+
D_ASSERT(asof_idx == op.conditions.size());
|
|
34
|
+
asof_idx = c;
|
|
35
|
+
break;
|
|
36
|
+
default:
|
|
37
|
+
throw InternalException("Invalid ASOF JOIN comparison");
|
|
38
|
+
}
|
|
39
|
+
}
|
|
40
|
+
D_ASSERT(asof_idx < op.conditions.size());
|
|
41
|
+
|
|
42
|
+
// Temporary implementation: IEJoin of Window
|
|
43
|
+
// LEAD(asof_column, 1, infinity) OVER (PARTITION BY equi_column... ORDER BY asof_column) AS asof_temp
|
|
44
|
+
auto &asof_comp = op.conditions[asof_idx];
|
|
45
|
+
auto &asof_column = asof_comp.right;
|
|
46
|
+
auto asof_type = asof_column->return_type;
|
|
47
|
+
auto asof_temp = make_unique<BoundWindowExpression>(ExpressionType::WINDOW_LEAD, asof_type, nullptr, nullptr);
|
|
48
|
+
asof_temp->children.emplace_back(asof_column->Copy());
|
|
49
|
+
asof_temp->offset_expr = make_unique<BoundConstantExpression>(Value::BIGINT(1));
|
|
50
|
+
asof_temp->default_expr = make_unique<BoundConstantExpression>(Value::Infinity(asof_type));
|
|
51
|
+
for (auto equi_idx : equi_indexes) {
|
|
52
|
+
asof_temp->partitions.emplace_back(op.conditions[equi_idx].right->Copy());
|
|
53
|
+
}
|
|
54
|
+
asof_temp->orders.emplace_back(OrderType::ASCENDING, OrderByNullType::NULLS_FIRST, asof_column->Copy());
|
|
55
|
+
asof_temp->start = WindowBoundary::UNBOUNDED_PRECEDING;
|
|
56
|
+
asof_temp->end = WindowBoundary::CURRENT_ROW_ROWS;
|
|
57
|
+
|
|
58
|
+
vector<unique_ptr<Expression>> window_select;
|
|
59
|
+
window_select.emplace_back(std::move(asof_temp));
|
|
60
|
+
|
|
61
|
+
auto window_types = right->types;
|
|
62
|
+
window_types.emplace_back(asof_type);
|
|
63
|
+
|
|
64
|
+
auto window = make_unique<PhysicalWindow>(window_types, std::move(window_select), rhs_cardinality);
|
|
65
|
+
window->children.emplace_back(std::move(right));
|
|
66
|
+
|
|
67
|
+
// IEJoin(left, window, conditions || asof_column < asof_temp)
|
|
68
|
+
JoinCondition asof_upper;
|
|
69
|
+
asof_upper.left = asof_comp.left->Copy();
|
|
70
|
+
asof_upper.right = make_unique<BoundReferenceExpression>(asof_type, window_types.size() - 1);
|
|
71
|
+
asof_upper.comparison = ExpressionType::COMPARE_LESSTHAN;
|
|
72
|
+
|
|
73
|
+
// We have an equality condition, so we may have to deal with projection maps.
|
|
74
|
+
// IEJoin does not (currently) support them, so we have to do it manually
|
|
75
|
+
auto proj_types = op.types;
|
|
76
|
+
op.types.clear();
|
|
77
|
+
|
|
78
|
+
auto lhs_types = op.children[0]->types;
|
|
79
|
+
op.types = lhs_types;
|
|
80
|
+
|
|
81
|
+
auto rhs_types = op.children[1]->types;
|
|
82
|
+
op.types.insert(op.types.end(), rhs_types.begin(), rhs_types.end());
|
|
83
|
+
|
|
84
|
+
op.types.emplace_back(asof_type);
|
|
85
|
+
op.conditions.emplace_back(std::move(asof_upper));
|
|
86
|
+
auto iejoin = make_unique<PhysicalIEJoin>(op, std::move(left), std::move(window), std::move(op.conditions),
|
|
87
|
+
op.join_type, op.estimated_cardinality);
|
|
88
|
+
|
|
89
|
+
// Project away asof_temp and anything from the projection maps
|
|
90
|
+
auto proj = PhysicalProjection::CreateJoinProjection(proj_types, lhs_types, rhs_types, op.left_projection_map,
|
|
91
|
+
op.right_projection_map, lhs_cardinality);
|
|
92
|
+
proj->children.push_back(std::move(iejoin));
|
|
93
|
+
|
|
94
|
+
return proj;
|
|
95
|
+
}
|
|
96
|
+
|
|
97
|
+
} // namespace duckdb
|