duckdb 0.7.1 → 0.7.2-dev1034.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/binding.gyp +12 -7
- package/lib/duckdb.d.ts +55 -2
- package/lib/duckdb.js +20 -1
- package/package.json +1 -1
- package/src/connection.cpp +1 -2
- package/src/database.cpp +1 -1
- package/src/duckdb/extension/icu/icu-extension.cpp +4 -0
- package/src/duckdb/extension/icu/icu-list-range.cpp +207 -0
- package/src/duckdb/extension/icu/icu-table-range.cpp +194 -0
- package/src/duckdb/extension/icu/include/icu-list-range.hpp +17 -0
- package/src/duckdb/extension/icu/include/icu-table-range.hpp +17 -0
- package/src/duckdb/extension/json/include/json_common.hpp +1 -0
- package/src/duckdb/extension/json/include/json_functions.hpp +2 -0
- package/src/duckdb/extension/json/include/json_serializer.hpp +77 -0
- package/src/duckdb/extension/json/json_functions/json_serialize_sql.cpp +147 -0
- package/src/duckdb/extension/json/json_functions/read_json.cpp +6 -5
- package/src/duckdb/extension/json/json_functions.cpp +12 -4
- package/src/duckdb/extension/json/json_scan.cpp +2 -2
- package/src/duckdb/extension/json/json_serializer.cpp +217 -0
- package/src/duckdb/extension/parquet/column_reader.cpp +94 -15
- package/src/duckdb/extension/parquet/column_writer.cpp +0 -1
- package/src/duckdb/extension/parquet/include/column_reader.hpp +1 -2
- package/src/duckdb/extension/parquet/include/decode_utils.hpp +5 -4
- package/src/duckdb/extension/parquet/include/generated_column_reader.hpp +1 -11
- package/src/duckdb/extension/parquet/include/parquet_timestamp.hpp +2 -1
- package/src/duckdb/extension/parquet/parquet-extension.cpp +12 -2
- package/src/duckdb/extension/parquet/parquet_reader.cpp +1 -1
- package/src/duckdb/extension/parquet/parquet_statistics.cpp +26 -32
- package/src/duckdb/extension/parquet/parquet_timestamp.cpp +16 -6
- package/src/duckdb/src/catalog/catalog.cpp +34 -5
- package/src/duckdb/src/catalog/catalog_entry/duck_schema_entry.cpp +4 -0
- package/src/duckdb/src/catalog/catalog_entry/duck_table_entry.cpp +2 -21
- package/src/duckdb/src/catalog/catalog_entry/scalar_function_catalog_entry.cpp +7 -6
- package/src/duckdb/src/catalog/catalog_entry/table_catalog_entry.cpp +3 -3
- package/src/duckdb/src/catalog/catalog_entry/table_function_catalog_entry.cpp +20 -1
- package/src/duckdb/src/catalog/catalog_entry/type_catalog_entry.cpp +8 -2
- package/src/duckdb/src/catalog/catalog_set.cpp +1 -0
- package/src/duckdb/src/catalog/default/default_functions.cpp +3 -0
- package/src/duckdb/src/catalog/dependency_list.cpp +12 -0
- package/src/duckdb/src/catalog/duck_catalog.cpp +34 -7
- package/src/duckdb/src/common/arrow/arrow_appender.cpp +48 -4
- package/src/duckdb/src/common/arrow/arrow_converter.cpp +1 -1
- package/src/duckdb/src/common/box_renderer.cpp +109 -23
- package/src/duckdb/src/common/enums/expression_type.cpp +8 -222
- package/src/duckdb/src/common/enums/join_type.cpp +3 -22
- package/src/duckdb/src/common/enums/logical_operator_type.cpp +2 -0
- package/src/duckdb/src/common/enums/statement_type.cpp +2 -0
- package/src/duckdb/src/common/exception.cpp +15 -1
- package/src/duckdb/src/common/field_writer.cpp +1 -0
- package/src/duckdb/src/common/operator/cast_operators.cpp +1 -1
- package/src/duckdb/src/common/preserved_error.cpp +7 -5
- package/src/duckdb/src/common/serializer/buffered_deserializer.cpp +4 -0
- package/src/duckdb/src/common/serializer/buffered_file_reader.cpp +15 -2
- package/src/duckdb/src/common/serializer/enum_serializer.cpp +1176 -0
- package/src/duckdb/src/common/sort/sort_state.cpp +5 -7
- package/src/duckdb/src/common/sort/sorted_block.cpp +0 -1
- package/src/duckdb/src/common/string_util.cpp +4 -1
- package/src/duckdb/src/common/types/bit.cpp +166 -87
- package/src/duckdb/src/common/types/blob.cpp +1 -1
- package/src/duckdb/src/common/types/chunk_collection.cpp +2 -2
- package/src/duckdb/src/common/types/column_data_collection.cpp +39 -2
- package/src/duckdb/src/common/types/column_data_collection_segment.cpp +11 -6
- package/src/duckdb/src/common/types/data_chunk.cpp +1 -1
- package/src/duckdb/src/common/types/time.cpp +13 -0
- package/src/duckdb/src/common/types/value.cpp +320 -154
- package/src/duckdb/src/common/types/vector.cpp +155 -127
- package/src/duckdb/src/common/types.cpp +313 -153
- package/src/duckdb/src/common/vector_operations/vector_cast.cpp +2 -1
- package/src/duckdb/src/execution/aggregate_hashtable.cpp +10 -5
- package/src/duckdb/src/execution/column_binding_resolver.cpp +21 -5
- package/src/duckdb/src/execution/expression_executor/execute_cast.cpp +2 -1
- package/src/duckdb/src/execution/index/art/art.cpp +6 -5
- package/src/duckdb/src/execution/operator/aggregate/physical_perfecthash_aggregate.cpp +4 -5
- package/src/duckdb/src/execution/operator/aggregate/physical_window.cpp +117 -26
- package/src/duckdb/src/execution/operator/helper/physical_limit.cpp +3 -0
- package/src/duckdb/src/execution/operator/helper/physical_vacuum.cpp +5 -3
- package/src/duckdb/src/execution/operator/join/physical_blockwise_nl_join.cpp +64 -17
- package/src/duckdb/src/execution/operator/join/physical_iejoin.cpp +2 -2
- package/src/duckdb/src/execution/operator/join/physical_index_join.cpp +12 -4
- package/src/duckdb/src/execution/operator/join/physical_piecewise_merge_join.cpp +6 -11
- package/src/duckdb/src/execution/operator/join/physical_range_join.cpp +3 -1
- package/src/duckdb/src/execution/operator/persistent/base_csv_reader.cpp +6 -3
- package/src/duckdb/src/execution/operator/persistent/buffered_csv_reader.cpp +6 -14
- package/src/duckdb/src/execution/operator/persistent/physical_copy_to_file.cpp +2 -2
- package/src/duckdb/src/execution/operator/projection/physical_projection.cpp +34 -0
- package/src/duckdb/src/execution/operator/scan/physical_positional_scan.cpp +20 -5
- package/src/duckdb/src/execution/operator/schema/physical_create_type.cpp +20 -40
- package/src/duckdb/src/execution/partitionable_hashtable.cpp +14 -2
- package/src/duckdb/src/execution/physical_plan/plan_aggregate.cpp +21 -16
- package/src/duckdb/src/execution/physical_plan/plan_asof_join.cpp +97 -0
- package/src/duckdb/src/execution/physical_plan/plan_comparison_join.cpp +95 -47
- package/src/duckdb/src/execution/physical_plan/plan_distinct.cpp +5 -8
- package/src/duckdb/src/execution/physical_plan/plan_positional_join.cpp +14 -5
- package/src/duckdb/src/execution/physical_plan_generator.cpp +3 -0
- package/src/duckdb/src/execution/window_segment_tree.cpp +173 -1
- package/src/duckdb/src/function/aggregate/algebraic/avg.cpp +0 -6
- package/src/duckdb/src/function/aggregate/distributive/bitagg.cpp +99 -95
- package/src/duckdb/src/function/aggregate/distributive/bitstring_agg.cpp +269 -0
- package/src/duckdb/src/function/aggregate/distributive/bool.cpp +2 -0
- package/src/duckdb/src/function/aggregate/distributive/count.cpp +3 -4
- package/src/duckdb/src/function/aggregate/distributive/first.cpp +1 -0
- package/src/duckdb/src/function/aggregate/distributive/minmax.cpp +2 -0
- package/src/duckdb/src/function/aggregate/distributive/sum.cpp +19 -16
- package/src/duckdb/src/function/aggregate/distributive_functions.cpp +1 -0
- package/src/duckdb/src/function/aggregate/holistic/approximate_quantile.cpp +5 -2
- package/src/duckdb/src/function/aggregate/holistic/mode.cpp +1 -1
- package/src/duckdb/src/function/aggregate/holistic/quantile.cpp +16 -1
- package/src/duckdb/src/function/aggregate/nested/list.cpp +8 -8
- package/src/duckdb/src/function/aggregate/sorted_aggregate_function.cpp +58 -16
- package/src/duckdb/src/function/cast/bit_cast.cpp +0 -2
- package/src/duckdb/src/function/cast/blob_cast.cpp +0 -1
- package/src/duckdb/src/function/cast/cast_function_set.cpp +1 -1
- package/src/duckdb/src/function/cast/enum_casts.cpp +25 -3
- package/src/duckdb/src/function/cast/list_casts.cpp +17 -4
- package/src/duckdb/src/function/cast/map_cast.cpp +5 -2
- package/src/duckdb/src/function/cast/string_cast.cpp +36 -10
- package/src/duckdb/src/function/cast/struct_cast.cpp +24 -4
- package/src/duckdb/src/function/cast/time_casts.cpp +2 -2
- package/src/duckdb/src/function/cast/union_casts.cpp +33 -7
- package/src/duckdb/src/function/function_binder.cpp +1 -8
- package/src/duckdb/src/function/scalar/bit/bitstring.cpp +100 -0
- package/src/duckdb/src/function/scalar/date/current.cpp +0 -2
- package/src/duckdb/src/function/scalar/date/date_diff.cpp +0 -1
- package/src/duckdb/src/function/scalar/date/date_part.cpp +18 -26
- package/src/duckdb/src/function/scalar/date/date_sub.cpp +0 -1
- package/src/duckdb/src/function/scalar/date/date_trunc.cpp +10 -14
- package/src/duckdb/src/function/scalar/generic/stats.cpp +2 -4
- package/src/duckdb/src/function/scalar/list/contains_or_position.cpp +4 -146
- package/src/duckdb/src/function/scalar/list/flatten.cpp +5 -12
- package/src/duckdb/src/function/scalar/list/list_aggregates.cpp +1 -1
- package/src/duckdb/src/function/scalar/list/list_concat.cpp +8 -12
- package/src/duckdb/src/function/scalar/list/list_extract.cpp +5 -12
- package/src/duckdb/src/function/scalar/list/list_lambdas.cpp +7 -3
- package/src/duckdb/src/function/scalar/list/list_value.cpp +6 -10
- package/src/duckdb/src/function/scalar/map/map.cpp +47 -1
- package/src/duckdb/src/function/scalar/map/map_entries.cpp +61 -0
- package/src/duckdb/src/function/scalar/map/map_extract.cpp +68 -26
- package/src/duckdb/src/function/scalar/map/map_keys_values.cpp +97 -0
- package/src/duckdb/src/function/scalar/math/numeric.cpp +101 -17
- package/src/duckdb/src/function/scalar/math_functions.cpp +3 -0
- package/src/duckdb/src/function/scalar/nested_functions.cpp +3 -0
- package/src/duckdb/src/function/scalar/operators/add.cpp +0 -9
- package/src/duckdb/src/function/scalar/operators/arithmetic.cpp +29 -48
- package/src/duckdb/src/function/scalar/operators/bitwise.cpp +0 -63
- package/src/duckdb/src/function/scalar/operators/multiply.cpp +5 -6
- package/src/duckdb/src/function/scalar/operators/subtract.cpp +0 -6
- package/src/duckdb/src/function/scalar/string/caseconvert.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/hex.cpp +201 -0
- package/src/duckdb/src/function/scalar/string/instr.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/length.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/like.cpp +2 -6
- package/src/duckdb/src/function/scalar/string/regexp/regexp_extract_all.cpp +243 -0
- package/src/duckdb/src/function/scalar/string/regexp/regexp_util.cpp +79 -0
- package/src/duckdb/src/function/scalar/string/regexp.cpp +21 -80
- package/src/duckdb/src/function/scalar/string/substring.cpp +2 -6
- package/src/duckdb/src/function/scalar/string_functions.cpp +2 -0
- package/src/duckdb/src/function/scalar/struct/struct_extract.cpp +5 -10
- package/src/duckdb/src/function/scalar/struct/struct_insert.cpp +11 -14
- package/src/duckdb/src/function/scalar/struct/struct_pack.cpp +6 -7
- package/src/duckdb/src/function/table/arrow.cpp +5 -2
- package/src/duckdb/src/function/table/arrow_conversion.cpp +25 -1
- package/src/duckdb/src/function/table/checkpoint.cpp +5 -1
- package/src/duckdb/src/function/table/read_csv.cpp +55 -0
- package/src/duckdb/src/function/table/system/duckdb_constraints.cpp +2 -2
- package/src/duckdb/src/function/table/system/test_all_types.cpp +2 -2
- package/src/duckdb/src/function/table/table_scan.cpp +1 -1
- package/src/duckdb/src/function/table/version/pragma_version.cpp +2 -2
- package/src/duckdb/src/function/table_function.cpp +30 -11
- package/src/duckdb/src/include/duckdb/catalog/catalog.hpp +6 -0
- package/src/duckdb/src/include/duckdb/catalog/catalog_entry/duck_table_entry.hpp +1 -1
- package/src/duckdb/src/include/duckdb/catalog/catalog_entry/table_function_catalog_entry.hpp +6 -8
- package/src/duckdb/src/include/duckdb/catalog/dependency_list.hpp +3 -0
- package/src/duckdb/src/include/duckdb/catalog/duck_catalog.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/box_renderer.hpp +8 -2
- package/src/duckdb/src/include/duckdb/common/constants.hpp +0 -19
- package/src/duckdb/src/include/duckdb/common/enums/aggregate_handling.hpp +2 -0
- package/src/duckdb/src/include/duckdb/common/enums/expression_type.hpp +2 -3
- package/src/duckdb/src/include/duckdb/common/enums/joinref_type.hpp +7 -4
- package/src/duckdb/src/include/duckdb/common/enums/logical_operator_type.hpp +1 -0
- package/src/duckdb/src/include/duckdb/common/enums/order_type.hpp +2 -0
- package/src/duckdb/src/include/duckdb/common/enums/set_operation_type.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/enums/statement_type.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/enums/tableref_type.hpp +2 -1
- package/src/duckdb/src/include/duckdb/common/exception.hpp +69 -2
- package/src/duckdb/src/include/duckdb/common/field_writer.hpp +12 -4
- package/src/duckdb/src/include/duckdb/common/{http_stats.hpp → http_state.hpp} +18 -4
- package/src/duckdb/src/include/duckdb/common/operator/multiply.hpp +2 -0
- package/src/duckdb/src/include/duckdb/common/optional_ptr.hpp +45 -0
- package/src/duckdb/src/include/duckdb/common/preserved_error.hpp +6 -1
- package/src/duckdb/src/include/duckdb/common/serializer/buffered_deserializer.hpp +4 -2
- package/src/duckdb/src/include/duckdb/common/serializer/buffered_file_reader.hpp +8 -2
- package/src/duckdb/src/include/duckdb/common/serializer/enum_serializer.hpp +113 -0
- package/src/duckdb/src/include/duckdb/common/serializer/format_deserializer.hpp +336 -0
- package/src/duckdb/src/include/duckdb/common/serializer/format_serializer.hpp +268 -0
- package/src/duckdb/src/include/duckdb/common/serializer/serialization_traits.hpp +126 -0
- package/src/duckdb/src/include/duckdb/common/serializer.hpp +13 -0
- package/src/duckdb/src/include/duckdb/common/string_util.hpp +25 -0
- package/src/duckdb/src/include/duckdb/common/types/bit.hpp +12 -7
- package/src/duckdb/src/include/duckdb/common/types/time.hpp +3 -0
- package/src/duckdb/src/include/duckdb/common/types/value.hpp +17 -48
- package/src/duckdb/src/include/duckdb/common/types/value_map.hpp +1 -1
- package/src/duckdb/src/include/duckdb/common/types/vector.hpp +3 -1
- package/src/duckdb/src/include/duckdb/common/types.hpp +45 -8
- package/src/duckdb/src/include/duckdb/common/vector_operations/unary_executor.hpp +2 -2
- package/src/duckdb/src/include/duckdb/execution/aggregate_hashtable.hpp +1 -0
- package/src/duckdb/src/include/duckdb/execution/index/art/art.hpp +2 -2
- package/src/duckdb/src/include/duckdb/execution/operator/aggregate/physical_perfecthash_aggregate.hpp +1 -1
- package/src/duckdb/src/include/duckdb/execution/operator/join/physical_cross_product.hpp +2 -0
- package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_file_handle.hpp +1 -0
- package/src/duckdb/src/include/duckdb/execution/operator/persistent/csv_reader_options.hpp +6 -0
- package/src/duckdb/src/include/duckdb/execution/operator/projection/physical_projection.hpp +5 -0
- package/src/duckdb/src/include/duckdb/execution/partitionable_hashtable.hpp +3 -0
- package/src/duckdb/src/include/duckdb/execution/physical_plan_generator.hpp +1 -3
- package/src/duckdb/src/include/duckdb/execution/window_segment_tree.hpp +54 -0
- package/src/duckdb/src/include/duckdb/function/aggregate/distributive_functions.hpp +5 -0
- package/src/duckdb/src/include/duckdb/function/aggregate_function.hpp +18 -6
- package/src/duckdb/src/include/duckdb/function/cast/bound_cast_data.hpp +84 -0
- package/src/duckdb/src/include/duckdb/function/cast/cast_function_set.hpp +2 -2
- package/src/duckdb/src/include/duckdb/function/cast/default_casts.hpp +28 -64
- package/src/duckdb/src/include/duckdb/function/function_binder.hpp +3 -6
- package/src/duckdb/src/include/duckdb/function/scalar/bit_functions.hpp +4 -0
- package/src/duckdb/src/include/duckdb/function/scalar/list/contains_or_position.hpp +138 -0
- package/src/duckdb/src/include/duckdb/function/scalar/math_functions.hpp +8 -0
- package/src/duckdb/src/include/duckdb/function/scalar/nested_functions.hpp +59 -0
- package/src/duckdb/src/include/duckdb/function/scalar/regexp.hpp +81 -1
- package/src/duckdb/src/include/duckdb/function/scalar/string_functions.hpp +4 -0
- package/src/duckdb/src/include/duckdb/function/scalar_function.hpp +2 -2
- package/src/duckdb/src/include/duckdb/function/table/arrow.hpp +12 -1
- package/src/duckdb/src/include/duckdb/function/table_function.hpp +10 -0
- package/src/duckdb/src/include/duckdb/main/capi/capi_internal.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/client_data.hpp +3 -3
- package/src/duckdb/src/include/duckdb/main/config.hpp +3 -0
- package/src/duckdb/src/include/duckdb/main/connection_manager.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/database.hpp +1 -0
- package/src/duckdb/src/include/duckdb/main/extension_entries.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/prepared_statement.hpp +2 -0
- package/src/duckdb/src/include/duckdb/main/relation/explain_relation.hpp +2 -1
- package/src/duckdb/src/include/duckdb/main/relation.hpp +2 -1
- package/src/duckdb/src/include/duckdb/optimizer/filter_pushdown.hpp +2 -0
- package/src/duckdb/src/include/duckdb/optimizer/join_order/cardinality_estimator.hpp +2 -2
- package/src/duckdb/src/include/duckdb/optimizer/rule/list.hpp +1 -0
- package/src/duckdb/src/include/duckdb/optimizer/rule/ordered_aggregate_optimizer.hpp +24 -0
- package/src/duckdb/src/include/duckdb/parser/common_table_expression_info.hpp +4 -0
- package/src/duckdb/src/include/duckdb/parser/expression/between_expression.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/expression/bound_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/case_expression.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/expression/cast_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/collate_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/columnref_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/comparison_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/conjunction_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/constant_expression.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/expression/default_expression.hpp +1 -0
- package/src/duckdb/src/include/duckdb/parser/expression/function_expression.hpp +4 -2
- package/src/duckdb/src/include/duckdb/parser/expression/lambda_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/operator_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/parameter_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/positional_reference_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/star_expression.hpp +4 -2
- package/src/duckdb/src/include/duckdb/parser/expression/subquery_expression.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/expression/window_expression.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_info.hpp +5 -1
- package/src/duckdb/src/include/duckdb/parser/parsed_data/{alter_function_info.hpp → alter_scalar_function_info.hpp} +13 -13
- package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_function_info.hpp +47 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_data/alter_table_info.hpp +6 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_data/create_table_function_info.hpp +2 -1
- package/src/duckdb/src/include/duckdb/parser/parsed_data/sample_options.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/parsed_expression.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/query_node/recursive_cte_node.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/query_node/select_node.hpp +5 -0
- package/src/duckdb/src/include/duckdb/parser/query_node/set_operation_node.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/query_node.hpp +13 -2
- package/src/duckdb/src/include/duckdb/parser/result_modifier.hpp +24 -1
- package/src/duckdb/src/include/duckdb/parser/sql_statement.hpp +2 -1
- package/src/duckdb/src/include/duckdb/parser/statement/multi_statement.hpp +28 -0
- package/src/duckdb/src/include/duckdb/parser/statement/select_statement.hpp +6 -1
- package/src/duckdb/src/include/duckdb/parser/tableref/basetableref.hpp +4 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/emptytableref.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/expressionlistref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/joinref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/list.hpp +1 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/pivotref.hpp +87 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/subqueryref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref/table_function_ref.hpp +3 -0
- package/src/duckdb/src/include/duckdb/parser/tableref.hpp +3 -1
- package/src/duckdb/src/include/duckdb/parser/tokens.hpp +2 -0
- package/src/duckdb/src/include/duckdb/parser/transformer.hpp +33 -0
- package/src/duckdb/src/include/duckdb/planner/bind_context.hpp +2 -0
- package/src/duckdb/src/include/duckdb/planner/binder.hpp +15 -4
- package/src/duckdb/src/include/duckdb/planner/bound_result_modifier.hpp +3 -0
- package/src/duckdb/src/include/duckdb/planner/expression/bound_aggregate_expression.hpp +3 -0
- package/src/duckdb/src/include/duckdb/planner/expression_binder/base_select_binder.hpp +64 -0
- package/src/duckdb/src/include/duckdb/planner/expression_binder/having_binder.hpp +2 -2
- package/src/duckdb/src/include/duckdb/planner/expression_binder/order_binder.hpp +4 -1
- package/src/duckdb/src/include/duckdb/planner/expression_binder/qualify_binder.hpp +2 -2
- package/src/duckdb/src/include/duckdb/planner/expression_binder/select_binder.hpp +9 -38
- package/src/duckdb/src/include/duckdb/planner/expression_binder.hpp +1 -1
- package/src/duckdb/src/include/duckdb/planner/logical_tokens.hpp +1 -0
- package/src/duckdb/src/include/duckdb/planner/operator/list.hpp +1 -0
- package/src/duckdb/src/include/duckdb/planner/operator/logical_asof_join.hpp +22 -0
- package/src/duckdb/src/include/duckdb/planner/operator/logical_comparison_join.hpp +5 -2
- package/src/duckdb/src/include/duckdb/planner/operator/logical_distinct.hpp +3 -0
- package/src/duckdb/src/include/duckdb/planner/query_node/bound_select_node.hpp +8 -2
- package/src/duckdb/src/include/duckdb/storage/buffer/block_handle.hpp +2 -0
- package/src/duckdb/src/include/duckdb/storage/buffer_manager.hpp +76 -44
- package/src/duckdb/src/include/duckdb/storage/checkpoint/table_data_writer.hpp +3 -2
- package/src/duckdb/src/include/duckdb/storage/checkpoint_manager.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_compress.hpp +2 -2
- package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_fetch.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/chimp/chimp_scan.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_compress.hpp +2 -2
- package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_fetch.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/compression/patas/patas_scan.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/data_pointer.hpp +5 -2
- package/src/duckdb/src/include/duckdb/storage/data_table.hpp +3 -3
- package/src/duckdb/src/include/duckdb/storage/index.hpp +4 -3
- package/src/duckdb/src/include/duckdb/storage/meta_block_reader.hpp +7 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/base_statistics.hpp +93 -29
- package/src/duckdb/src/include/duckdb/storage/statistics/column_statistics.hpp +22 -3
- package/src/duckdb/src/include/duckdb/storage/statistics/distinct_statistics.hpp +8 -6
- package/src/duckdb/src/include/duckdb/storage/statistics/list_stats.hpp +41 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/node_statistics.hpp +26 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats.hpp +114 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/numeric_stats_union.hpp +62 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/segment_statistics.hpp +2 -7
- package/src/duckdb/src/include/duckdb/storage/statistics/string_stats.hpp +74 -0
- package/src/duckdb/src/include/duckdb/storage/statistics/struct_stats.hpp +42 -0
- package/src/duckdb/src/include/duckdb/storage/string_uncompressed.hpp +2 -3
- package/src/duckdb/src/include/duckdb/storage/table/column_checkpoint_state.hpp +2 -1
- package/src/duckdb/src/include/duckdb/storage/table/column_data.hpp +6 -3
- package/src/duckdb/src/include/duckdb/storage/table/column_data_checkpointer.hpp +3 -2
- package/src/duckdb/src/include/duckdb/storage/table/column_segment.hpp +7 -5
- package/src/duckdb/src/include/duckdb/storage/table/list_column_data.hpp +1 -1
- package/src/duckdb/src/include/duckdb/storage/table/persistent_table_data.hpp +6 -2
- package/src/duckdb/src/include/duckdb/storage/table/row_group.hpp +10 -6
- package/src/duckdb/src/include/duckdb/storage/table/row_group_collection.hpp +8 -5
- package/src/duckdb/src/include/duckdb/storage/table/row_group_segment_tree.hpp +37 -0
- package/src/duckdb/src/include/duckdb/storage/table/scan_state.hpp +10 -1
- package/src/duckdb/src/include/duckdb/storage/table/segment_base.hpp +4 -3
- package/src/duckdb/src/include/duckdb/storage/table/segment_tree.hpp +271 -26
- package/src/duckdb/src/include/duckdb/storage/table/table_statistics.hpp +5 -0
- package/src/duckdb/src/include/duckdb/storage/table/update_segment.hpp +0 -1
- package/src/duckdb/src/include/duckdb/storage/write_ahead_log.hpp +1 -1
- package/src/duckdb/src/include/duckdb/transaction/local_storage.hpp +2 -2
- package/src/duckdb/src/include/duckdb.h +50 -2
- package/src/duckdb/src/include/duckdb.hpp +0 -1
- package/src/duckdb/src/main/capi/pending-c.cpp +16 -3
- package/src/duckdb/src/main/capi/result-c.cpp +27 -1
- package/src/duckdb/src/main/capi/stream-c.cpp +25 -0
- package/src/duckdb/src/main/client_context.cpp +38 -34
- package/src/duckdb/src/main/client_data.cpp +7 -6
- package/src/duckdb/src/main/config.cpp +70 -1
- package/src/duckdb/src/main/database.cpp +19 -2
- package/src/duckdb/src/main/extension/extension_install.cpp +7 -2
- package/src/duckdb/src/main/prepared_statement.cpp +4 -0
- package/src/duckdb/src/main/query_profiler.cpp +17 -15
- package/src/duckdb/src/main/relation/explain_relation.cpp +3 -3
- package/src/duckdb/src/main/relation.cpp +3 -2
- package/src/duckdb/src/optimizer/column_lifetime_analyzer.cpp +1 -0
- package/src/duckdb/src/optimizer/deliminator.cpp +1 -1
- package/src/duckdb/src/optimizer/filter_combiner.cpp +1 -1
- package/src/duckdb/src/optimizer/filter_pullup.cpp +3 -1
- package/src/duckdb/src/optimizer/filter_pushdown.cpp +14 -8
- package/src/duckdb/src/optimizer/join_order/cardinality_estimator.cpp +105 -71
- package/src/duckdb/src/optimizer/join_order/join_order_optimizer.cpp +31 -12
- package/src/duckdb/src/optimizer/optimizer.cpp +1 -0
- package/src/duckdb/src/optimizer/pullup/pullup_from_left.cpp +2 -2
- package/src/duckdb/src/optimizer/pushdown/pushdown_aggregate.cpp +33 -5
- package/src/duckdb/src/optimizer/pushdown/pushdown_cross_product.cpp +1 -1
- package/src/duckdb/src/optimizer/pushdown/pushdown_inner_join.cpp +3 -0
- package/src/duckdb/src/optimizer/pushdown/pushdown_left_join.cpp +5 -12
- package/src/duckdb/src/optimizer/pushdown/pushdown_mark_join.cpp +2 -2
- package/src/duckdb/src/optimizer/pushdown/pushdown_single_join.cpp +1 -1
- package/src/duckdb/src/optimizer/remove_unused_columns.cpp +1 -0
- package/src/duckdb/src/optimizer/rule/move_constants.cpp +10 -4
- package/src/duckdb/src/optimizer/rule/ordered_aggregate_optimizer.cpp +30 -0
- package/src/duckdb/src/optimizer/rule/regex_optimizations.cpp +9 -2
- package/src/duckdb/src/optimizer/statistics/expression/propagate_aggregate.cpp +9 -3
- package/src/duckdb/src/optimizer/statistics/expression/propagate_and_compress.cpp +6 -7
- package/src/duckdb/src/optimizer/statistics/expression/propagate_cast.cpp +14 -11
- package/src/duckdb/src/optimizer/statistics/expression/propagate_columnref.cpp +1 -1
- package/src/duckdb/src/optimizer/statistics/expression/propagate_comparison.cpp +13 -15
- package/src/duckdb/src/optimizer/statistics/expression/propagate_conjunction.cpp +0 -1
- package/src/duckdb/src/optimizer/statistics/expression/propagate_constant.cpp +3 -75
- package/src/duckdb/src/optimizer/statistics/expression/propagate_function.cpp +7 -2
- package/src/duckdb/src/optimizer/statistics/expression/propagate_operator.cpp +10 -0
- package/src/duckdb/src/optimizer/statistics/operator/propagate_aggregate.cpp +2 -3
- package/src/duckdb/src/optimizer/statistics/operator/propagate_filter.cpp +29 -32
- package/src/duckdb/src/optimizer/statistics/operator/propagate_join.cpp +5 -5
- package/src/duckdb/src/optimizer/statistics/operator/propagate_set_operation.cpp +3 -3
- package/src/duckdb/src/optimizer/statistics_propagator.cpp +2 -1
- package/src/duckdb/src/optimizer/unnest_rewriter.cpp +2 -2
- package/src/duckdb/src/parallel/meta_pipeline.cpp +0 -4
- package/src/duckdb/src/parser/common_table_expression_info.cpp +19 -0
- package/src/duckdb/src/parser/expression/between_expression.cpp +17 -0
- package/src/duckdb/src/parser/expression/case_expression.cpp +28 -0
- package/src/duckdb/src/parser/expression/cast_expression.cpp +17 -0
- package/src/duckdb/src/parser/expression/collate_expression.cpp +16 -0
- package/src/duckdb/src/parser/expression/columnref_expression.cpp +15 -0
- package/src/duckdb/src/parser/expression/comparison_expression.cpp +16 -0
- package/src/duckdb/src/parser/expression/conjunction_expression.cpp +17 -0
- package/src/duckdb/src/parser/expression/constant_expression.cpp +14 -0
- package/src/duckdb/src/parser/expression/default_expression.cpp +7 -0
- package/src/duckdb/src/parser/expression/function_expression.cpp +35 -0
- package/src/duckdb/src/parser/expression/lambda_expression.cpp +16 -0
- package/src/duckdb/src/parser/expression/operator_expression.cpp +15 -0
- package/src/duckdb/src/parser/expression/parameter_expression.cpp +15 -0
- package/src/duckdb/src/parser/expression/positional_reference_expression.cpp +14 -0
- package/src/duckdb/src/parser/expression/star_expression.cpp +26 -6
- package/src/duckdb/src/parser/expression/subquery_expression.cpp +20 -0
- package/src/duckdb/src/parser/expression/window_expression.cpp +43 -0
- package/src/duckdb/src/parser/parsed_data/alter_info.cpp +7 -3
- package/src/duckdb/src/parser/parsed_data/alter_scalar_function_info.cpp +56 -0
- package/src/duckdb/src/parser/parsed_data/alter_table_function_info.cpp +51 -0
- package/src/duckdb/src/parser/parsed_data/create_scalar_function_info.cpp +3 -2
- package/src/duckdb/src/parser/parsed_data/create_table_function_info.cpp +6 -0
- package/src/duckdb/src/parser/parsed_data/sample_options.cpp +22 -10
- package/src/duckdb/src/parser/parsed_expression.cpp +72 -0
- package/src/duckdb/src/parser/parsed_expression_iterator.cpp +15 -1
- package/src/duckdb/src/parser/query_node/recursive_cte_node.cpp +21 -0
- package/src/duckdb/src/parser/query_node/select_node.cpp +31 -0
- package/src/duckdb/src/parser/query_node/set_operation_node.cpp +17 -0
- package/src/duckdb/src/parser/query_node.cpp +51 -1
- package/src/duckdb/src/parser/result_modifier.cpp +78 -0
- package/src/duckdb/src/parser/statement/multi_statement.cpp +18 -0
- package/src/duckdb/src/parser/statement/select_statement.cpp +12 -0
- package/src/duckdb/src/parser/tableref/basetableref.cpp +21 -0
- package/src/duckdb/src/parser/tableref/emptytableref.cpp +4 -0
- package/src/duckdb/src/parser/tableref/expressionlistref.cpp +17 -0
- package/src/duckdb/src/parser/tableref/joinref.cpp +29 -0
- package/src/duckdb/src/parser/tableref/pivotref.cpp +373 -0
- package/src/duckdb/src/parser/tableref/subqueryref.cpp +15 -0
- package/src/duckdb/src/parser/tableref/table_function.cpp +17 -0
- package/src/duckdb/src/parser/tableref.cpp +49 -0
- package/src/duckdb/src/parser/transform/expression/transform_array_access.cpp +11 -0
- package/src/duckdb/src/parser/transform/expression/transform_bool_expr.cpp +1 -1
- package/src/duckdb/src/parser/transform/expression/transform_columnref.cpp +17 -2
- package/src/duckdb/src/parser/transform/expression/transform_function.cpp +63 -42
- package/src/duckdb/src/parser/transform/expression/transform_operator.cpp +1 -1
- package/src/duckdb/src/parser/transform/expression/transform_subquery.cpp +1 -1
- package/src/duckdb/src/parser/transform/helpers/transform_alias.cpp +12 -6
- package/src/duckdb/src/parser/transform/helpers/transform_cte.cpp +24 -0
- package/src/duckdb/src/parser/transform/helpers/transform_groupby.cpp +7 -0
- package/src/duckdb/src/parser/transform/helpers/transform_orderby.cpp +0 -7
- package/src/duckdb/src/parser/transform/helpers/transform_typename.cpp +3 -2
- package/src/duckdb/src/parser/transform/statement/transform_create_function.cpp +4 -0
- package/src/duckdb/src/parser/transform/statement/transform_create_view.cpp +4 -0
- package/src/duckdb/src/parser/transform/statement/transform_pivot_stmt.cpp +179 -0
- package/src/duckdb/src/parser/transform/statement/transform_rename.cpp +3 -4
- package/src/duckdb/src/parser/transform/statement/transform_select.cpp +8 -0
- package/src/duckdb/src/parser/transform/statement/transform_select_node.cpp +2 -3
- package/src/duckdb/src/parser/transform/tableref/transform_join.cpp +12 -1
- package/src/duckdb/src/parser/transform/tableref/transform_pivot.cpp +121 -0
- package/src/duckdb/src/parser/transform/tableref/transform_tableref.cpp +2 -0
- package/src/duckdb/src/parser/transformer.cpp +15 -3
- package/src/duckdb/src/planner/bind_context.cpp +18 -25
- package/src/duckdb/src/planner/binder/expression/bind_aggregate_expression.cpp +9 -7
- package/src/duckdb/src/planner/binder/expression/bind_columnref_expression.cpp +4 -3
- package/src/duckdb/src/planner/binder/expression/bind_function_expression.cpp +23 -12
- package/src/duckdb/src/planner/binder/expression/bind_lambda.cpp +3 -2
- package/src/duckdb/src/planner/binder/expression/bind_star_expression.cpp +176 -0
- package/src/duckdb/src/planner/binder/expression/bind_subquery_expression.cpp +4 -0
- package/src/duckdb/src/planner/binder/expression/bind_unnest_expression.cpp +163 -24
- package/src/duckdb/src/planner/binder/expression/bind_window_expression.cpp +2 -2
- package/src/duckdb/src/planner/binder/query_node/bind_select_node.cpp +109 -94
- package/src/duckdb/src/planner/binder/query_node/plan_query_node.cpp +11 -0
- package/src/duckdb/src/planner/binder/query_node/plan_select_node.cpp +9 -4
- package/src/duckdb/src/planner/binder/statement/bind_copy.cpp +5 -3
- package/src/duckdb/src/planner/binder/statement/bind_create.cpp +3 -2
- package/src/duckdb/src/planner/binder/statement/bind_create_table.cpp +9 -1
- package/src/duckdb/src/planner/binder/statement/bind_delete.cpp +1 -1
- package/src/duckdb/src/planner/binder/statement/bind_insert.cpp +12 -8
- package/src/duckdb/src/planner/binder/statement/bind_logical_plan.cpp +17 -0
- package/src/duckdb/src/planner/binder/statement/bind_update.cpp +4 -2
- package/src/duckdb/src/planner/binder/tableref/bind_joinref.cpp +19 -3
- package/src/duckdb/src/planner/binder/tableref/bind_pivot.cpp +366 -0
- package/src/duckdb/src/planner/binder/tableref/bind_table_function.cpp +11 -1
- package/src/duckdb/src/planner/binder/tableref/plan_cteref.cpp +1 -0
- package/src/duckdb/src/planner/binder/tableref/plan_joinref.cpp +61 -13
- package/src/duckdb/src/planner/binder.cpp +19 -24
- package/src/duckdb/src/planner/bound_result_modifier.cpp +27 -1
- package/src/duckdb/src/planner/expression/bound_aggregate_expression.cpp +9 -2
- package/src/duckdb/src/planner/expression/bound_expression.cpp +4 -0
- package/src/duckdb/src/planner/expression/bound_window_expression.cpp +1 -1
- package/src/duckdb/src/planner/expression_binder/base_select_binder.cpp +146 -0
- package/src/duckdb/src/planner/expression_binder/having_binder.cpp +6 -3
- package/src/duckdb/src/planner/expression_binder/qualify_binder.cpp +3 -3
- package/src/duckdb/src/planner/expression_binder/select_binder.cpp +1 -132
- package/src/duckdb/src/planner/expression_binder.cpp +10 -3
- package/src/duckdb/src/planner/expression_iterator.cpp +17 -10
- package/src/duckdb/src/planner/filter/constant_filter.cpp +4 -6
- package/src/duckdb/src/planner/logical_operator.cpp +7 -2
- package/src/duckdb/src/planner/logical_operator_visitor.cpp +6 -0
- package/src/duckdb/src/planner/operator/logical_asof_join.cpp +8 -0
- package/src/duckdb/src/planner/operator/logical_distinct.cpp +3 -0
- package/src/duckdb/src/planner/planner.cpp +2 -1
- package/src/duckdb/src/planner/pragma_handler.cpp +10 -2
- package/src/duckdb/src/planner/subquery/flatten_dependent_join.cpp +3 -1
- package/src/duckdb/src/storage/buffer_manager.cpp +44 -46
- package/src/duckdb/src/storage/checkpoint/row_group_writer.cpp +1 -1
- package/src/duckdb/src/storage/checkpoint/table_data_reader.cpp +4 -15
- package/src/duckdb/src/storage/checkpoint/table_data_writer.cpp +10 -4
- package/src/duckdb/src/storage/checkpoint_manager.cpp +9 -3
- package/src/duckdb/src/storage/compression/bitpacking.cpp +28 -24
- package/src/duckdb/src/storage/compression/fixed_size_uncompressed.cpp +43 -45
- package/src/duckdb/src/storage/compression/numeric_constant.cpp +9 -10
- package/src/duckdb/src/storage/compression/patas.cpp +1 -1
- package/src/duckdb/src/storage/compression/rle.cpp +19 -15
- package/src/duckdb/src/storage/compression/validity_uncompressed.cpp +5 -5
- package/src/duckdb/src/storage/data_table.cpp +20 -20
- package/src/duckdb/src/storage/index.cpp +12 -1
- package/src/duckdb/src/storage/local_storage.cpp +20 -23
- package/src/duckdb/src/storage/meta_block_reader.cpp +22 -0
- package/src/duckdb/src/storage/statistics/base_statistics.cpp +373 -128
- package/src/duckdb/src/storage/statistics/column_statistics.cpp +57 -3
- package/src/duckdb/src/storage/statistics/distinct_statistics.cpp +8 -9
- package/src/duckdb/src/storage/statistics/list_stats.cpp +121 -0
- package/src/duckdb/src/storage/statistics/numeric_stats.cpp +591 -0
- package/src/duckdb/src/storage/statistics/numeric_stats_union.cpp +65 -0
- package/src/duckdb/src/storage/statistics/segment_statistics.cpp +2 -11
- package/src/duckdb/src/storage/statistics/string_stats.cpp +273 -0
- package/src/duckdb/src/storage/statistics/struct_stats.cpp +133 -0
- package/src/duckdb/src/storage/storage_info.cpp +2 -2
- package/src/duckdb/src/storage/table/column_checkpoint_state.cpp +4 -10
- package/src/duckdb/src/storage/table/column_data.cpp +45 -46
- package/src/duckdb/src/storage/table/column_data_checkpointer.cpp +7 -8
- package/src/duckdb/src/storage/table/column_segment.cpp +13 -14
- package/src/duckdb/src/storage/table/list_column_data.cpp +41 -59
- package/src/duckdb/src/storage/table/persistent_table_data.cpp +2 -1
- package/src/duckdb/src/storage/table/row_group.cpp +38 -32
- package/src/duckdb/src/storage/table/row_group_collection.cpp +94 -78
- package/src/duckdb/src/storage/table/scan_state.cpp +22 -3
- package/src/duckdb/src/storage/table/standard_column_data.cpp +7 -6
- package/src/duckdb/src/storage/table/struct_column_data.cpp +16 -16
- package/src/duckdb/src/storage/table/table_statistics.cpp +27 -7
- package/src/duckdb/src/storage/table/update_segment.cpp +20 -18
- package/src/duckdb/src/storage/wal_replay.cpp +8 -5
- package/src/duckdb/src/storage/write_ahead_log.cpp +2 -2
- package/src/duckdb/src/transaction/commit_state.cpp +11 -7
- package/src/duckdb/src/verification/deserialized_statement_verifier.cpp +0 -1
- package/src/duckdb/third_party/libpg_query/include/nodes/nodes.hpp +35 -0
- package/src/duckdb/third_party/libpg_query/include/nodes/parsenodes.hpp +36 -2
- package/src/duckdb/third_party/libpg_query/include/nodes/primnodes.hpp +3 -3
- package/src/duckdb/third_party/libpg_query/include/parser/gram.hpp +1022 -530
- package/src/duckdb/third_party/libpg_query/include/parser/kwlist.hpp +8 -0
- package/src/duckdb/third_party/libpg_query/src_backend_parser_gram.cpp +24462 -22828
- package/src/duckdb/third_party/re2/re2/re2.cc +9 -0
- package/src/duckdb/third_party/re2/re2/re2.h +2 -0
- package/src/duckdb/ub_extension_icu_third_party_icu_i18n.cpp +4 -4
- package/src/duckdb/ub_extension_json_json_functions.cpp +2 -0
- package/src/duckdb/ub_src_common_serializer.cpp +2 -0
- package/src/duckdb/ub_src_execution_physical_plan.cpp +2 -0
- package/src/duckdb/ub_src_function_aggregate_distributive.cpp +2 -0
- package/src/duckdb/ub_src_function_scalar_bit.cpp +2 -0
- package/src/duckdb/ub_src_function_scalar_map.cpp +4 -0
- package/src/duckdb/ub_src_function_scalar_string.cpp +2 -0
- package/src/duckdb/ub_src_function_scalar_string_regexp.cpp +4 -0
- package/src/duckdb/ub_src_main_capi.cpp +2 -0
- package/src/duckdb/ub_src_optimizer_rule.cpp +2 -0
- package/src/duckdb/ub_src_parser.cpp +2 -0
- package/src/duckdb/ub_src_parser_parsed_data.cpp +4 -2
- package/src/duckdb/ub_src_parser_statement.cpp +2 -0
- package/src/duckdb/ub_src_parser_tableref.cpp +2 -0
- package/src/duckdb/ub_src_parser_transform_statement.cpp +2 -0
- package/src/duckdb/ub_src_parser_transform_tableref.cpp +2 -0
- package/src/duckdb/ub_src_planner_binder_expression.cpp +2 -0
- package/src/duckdb/ub_src_planner_binder_tableref.cpp +2 -0
- package/src/duckdb/ub_src_planner_expression_binder.cpp +2 -0
- package/src/duckdb/ub_src_planner_operator.cpp +2 -0
- package/src/duckdb/ub_src_storage_statistics.cpp +6 -6
- package/src/duckdb/ub_src_storage_table.cpp +0 -2
- package/src/duckdb_node.hpp +2 -1
- package/src/statement.cpp +5 -5
- package/src/utils.cpp +27 -2
- package/test/extension.test.ts +44 -26
- package/test/syntax_error.test.ts +3 -1
- package/filelist.cache +0 -0
- package/src/duckdb/src/include/duckdb/main/loadable_extension.hpp +0 -59
- package/src/duckdb/src/include/duckdb/storage/statistics/list_statistics.hpp +0 -36
- package/src/duckdb/src/include/duckdb/storage/statistics/numeric_statistics.hpp +0 -75
- package/src/duckdb/src/include/duckdb/storage/statistics/string_statistics.hpp +0 -49
- package/src/duckdb/src/include/duckdb/storage/statistics/struct_statistics.hpp +0 -36
- package/src/duckdb/src/include/duckdb/storage/statistics/validity_statistics.hpp +0 -45
- package/src/duckdb/src/parser/parsed_data/alter_function_info.cpp +0 -55
- package/src/duckdb/src/storage/statistics/list_statistics.cpp +0 -94
- package/src/duckdb/src/storage/statistics/numeric_statistics.cpp +0 -307
- package/src/duckdb/src/storage/statistics/string_statistics.cpp +0 -220
- package/src/duckdb/src/storage/statistics/struct_statistics.cpp +0 -108
- package/src/duckdb/src/storage/statistics/validity_statistics.cpp +0 -91
- package/src/duckdb/src/storage/table/segment_tree.cpp +0 -179
| @@ -2,86 +2,61 @@ | |
| 2 2 | 
             
            #include "duckdb/common/field_writer.hpp"
         | 
| 3 3 | 
             
            #include "duckdb/common/string_util.hpp"
         | 
| 4 4 | 
             
            #include "duckdb/common/types/vector.hpp"
         | 
| 5 | 
            -
            #include "duckdb/storage/statistics/ | 
| 6 | 
            -
            #include "duckdb/storage/statistics/ | 
| 7 | 
            -
            #include "duckdb/storage/statistics/ | 
| 8 | 
            -
            #include "duckdb/storage/statistics/string_statistics.hpp"
         | 
| 9 | 
            -
            #include "duckdb/storage/statistics/struct_statistics.hpp"
         | 
| 10 | 
            -
            #include "duckdb/storage/statistics/validity_statistics.hpp"
         | 
| 5 | 
            +
            #include "duckdb/storage/statistics/base_statistics.hpp"
         | 
| 6 | 
            +
            #include "duckdb/storage/statistics/list_stats.hpp"
         | 
| 7 | 
            +
            #include "duckdb/storage/statistics/struct_stats.hpp"
         | 
| 11 8 |  | 
| 12 9 | 
             
            namespace duckdb {
         | 
| 13 10 |  | 
| 14 | 
            -
            BaseStatistics::BaseStatistics( | 
| 15 | 
            -
                : type(std::move(type)), stats_type(stats_type) {
         | 
| 11 | 
            +
            BaseStatistics::BaseStatistics() : type(LogicalType::INVALID) {
         | 
| 16 12 | 
             
            }
         | 
| 17 13 |  | 
| 18 | 
            -
            BaseStatistics | 
| 19 | 
            -
             | 
| 20 | 
            -
             | 
| 21 | 
            -
            void BaseStatistics::InitializeBase() {
         | 
| 22 | 
            -
            	validity_stats = make_unique<ValidityStatistics>(false);
         | 
| 23 | 
            -
            	if (stats_type == GLOBAL_STATS) {
         | 
| 24 | 
            -
            		distinct_stats = make_unique<DistinctStatistics>();
         | 
| 25 | 
            -
            	}
         | 
| 14 | 
            +
            BaseStatistics::BaseStatistics(LogicalType type) {
         | 
| 15 | 
            +
            	Construct(*this, std::move(type));
         | 
| 26 16 | 
             
            }
         | 
| 27 17 |  | 
| 28 | 
            -
             | 
| 29 | 
            -
            	 | 
| 30 | 
            -
             | 
| 31 | 
            -
             | 
| 32 | 
            -
             | 
| 18 | 
            +
            void BaseStatistics::Construct(BaseStatistics &stats, LogicalType type) {
         | 
| 19 | 
            +
            	stats.distinct_count = 0;
         | 
| 20 | 
            +
            	stats.type = std::move(type);
         | 
| 21 | 
            +
            	switch (GetStatsType(stats.type)) {
         | 
| 22 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 23 | 
            +
            		ListStats::Construct(stats);
         | 
| 24 | 
            +
            		break;
         | 
| 25 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 26 | 
            +
            		StructStats::Construct(stats);
         | 
| 27 | 
            +
            		break;
         | 
| 28 | 
            +
            	default:
         | 
| 29 | 
            +
            		break;
         | 
| 33 30 | 
             
            	}
         | 
| 34 | 
            -
            	return ((ValidityStatistics &)*validity_stats).has_null;
         | 
| 35 31 | 
             
            }
         | 
| 36 32 |  | 
| 37 | 
            -
             | 
| 38 | 
            -
            	if (!validity_stats) {
         | 
| 39 | 
            -
            		// we don't know
         | 
| 40 | 
            -
            		// solid maybe
         | 
| 41 | 
            -
            		return true;
         | 
| 42 | 
            -
            	}
         | 
| 43 | 
            -
            	return ((ValidityStatistics &)*validity_stats).has_no_null;
         | 
| 33 | 
            +
            BaseStatistics::~BaseStatistics() {
         | 
| 44 34 | 
             
            }
         | 
| 45 35 |  | 
| 46 | 
            -
             | 
| 47 | 
            -
            	 | 
| 48 | 
            -
             | 
| 49 | 
            -
            	 | 
| 50 | 
            -
            	 | 
| 51 | 
            -
            	 | 
| 36 | 
            +
            BaseStatistics::BaseStatistics(BaseStatistics &&other) noexcept {
         | 
| 37 | 
            +
            	std::swap(type, other.type);
         | 
| 38 | 
            +
            	has_null = other.has_null;
         | 
| 39 | 
            +
            	has_no_null = other.has_no_null;
         | 
| 40 | 
            +
            	distinct_count = other.distinct_count;
         | 
| 41 | 
            +
            	stats_union = other.stats_union;
         | 
| 42 | 
            +
            	std::swap(child_stats, other.child_stats);
         | 
| 52 43 | 
             
            }
         | 
| 53 44 |  | 
| 54 | 
            -
             | 
| 55 | 
            -
            	 | 
| 56 | 
            -
             | 
| 57 | 
            -
             | 
| 58 | 
            -
             | 
| 59 | 
            -
             | 
| 60 | 
            -
             | 
| 61 | 
            -
            	 | 
| 45 | 
            +
            BaseStatistics &BaseStatistics::operator=(BaseStatistics &&other) noexcept {
         | 
| 46 | 
            +
            	std::swap(type, other.type);
         | 
| 47 | 
            +
            	has_null = other.has_null;
         | 
| 48 | 
            +
            	has_no_null = other.has_no_null;
         | 
| 49 | 
            +
            	distinct_count = other.distinct_count;
         | 
| 50 | 
            +
            	stats_union = other.stats_union;
         | 
| 51 | 
            +
            	std::swap(child_stats, other.child_stats);
         | 
| 52 | 
            +
            	return *this;
         | 
| 62 53 | 
             
            }
         | 
| 63 54 |  | 
| 64 | 
            -
             | 
| 65 | 
            -
            	 | 
| 66 | 
            -
             | 
| 67 | 
            -
            	if (stats_type == GLOBAL_STATS) {
         | 
| 68 | 
            -
            		MergeInternal(distinct_stats, other.distinct_stats);
         | 
| 55 | 
            +
            StatisticsType BaseStatistics::GetStatsType(const LogicalType &type) {
         | 
| 56 | 
            +
            	if (type.id() == LogicalTypeId::SQLNULL) {
         | 
| 57 | 
            +
            		return StatisticsType::BASE_STATS;
         | 
| 69 58 | 
             
            	}
         | 
| 70 | 
            -
            }
         | 
| 71 | 
            -
             | 
| 72 | 
            -
            idx_t BaseStatistics::GetDistinctCount() {
         | 
| 73 | 
            -
            	if (distinct_stats) {
         | 
| 74 | 
            -
            		auto &d_stats = (DistinctStatistics &)*distinct_stats;
         | 
| 75 | 
            -
            		return d_stats.GetCount();
         | 
| 76 | 
            -
            	}
         | 
| 77 | 
            -
            	return 0;
         | 
| 78 | 
            -
            }
         | 
| 79 | 
            -
             | 
| 80 | 
            -
            unique_ptr<BaseStatistics> BaseStatistics::CreateEmpty(LogicalType type, StatisticsType stats_type) {
         | 
| 81 | 
            -
            	unique_ptr<BaseStatistics> result;
         | 
| 82 59 | 
             
            	switch (type.InternalType()) {
         | 
| 83 | 
            -
            	case PhysicalType::BIT:
         | 
| 84 | 
            -
            		return make_unique<ValidityStatistics>(false, false);
         | 
| 85 60 | 
             
            	case PhysicalType::BOOL:
         | 
| 86 61 | 
             
            	case PhysicalType::INT8:
         | 
| 87 62 | 
             
            	case PhysicalType::INT16:
         | 
| @@ -94,113 +69,323 @@ unique_ptr<BaseStatistics> BaseStatistics::CreateEmpty(LogicalType type, Statist | |
| 94 69 | 
             
            	case PhysicalType::INT128:
         | 
| 95 70 | 
             
            	case PhysicalType::FLOAT:
         | 
| 96 71 | 
             
            	case PhysicalType::DOUBLE:
         | 
| 97 | 
            -
            		 | 
| 98 | 
            -
            		break;
         | 
| 72 | 
            +
            		return StatisticsType::NUMERIC_STATS;
         | 
| 99 73 | 
             
            	case PhysicalType::VARCHAR:
         | 
| 100 | 
            -
            		 | 
| 101 | 
            -
            		break;
         | 
| 74 | 
            +
            		return StatisticsType::STRING_STATS;
         | 
| 102 75 | 
             
            	case PhysicalType::STRUCT:
         | 
| 103 | 
            -
            		 | 
| 104 | 
            -
            		break;
         | 
| 76 | 
            +
            		return StatisticsType::STRUCT_STATS;
         | 
| 105 77 | 
             
            	case PhysicalType::LIST:
         | 
| 106 | 
            -
            		 | 
| 107 | 
            -
             | 
| 78 | 
            +
            		return StatisticsType::LIST_STATS;
         | 
| 79 | 
            +
            	case PhysicalType::BIT:
         | 
| 108 80 | 
             
            	case PhysicalType::INTERVAL:
         | 
| 109 81 | 
             
            	default:
         | 
| 110 | 
            -
            		 | 
| 82 | 
            +
            		return StatisticsType::BASE_STATS;
         | 
| 111 83 | 
             
            	}
         | 
| 112 | 
            -
             | 
| 84 | 
            +
            }
         | 
| 85 | 
            +
             | 
| 86 | 
            +
            StatisticsType BaseStatistics::GetStatsType() const {
         | 
| 87 | 
            +
            	return GetStatsType(GetType());
         | 
| 88 | 
            +
            }
         | 
| 89 | 
            +
             | 
| 90 | 
            +
            void BaseStatistics::InitializeUnknown() {
         | 
| 91 | 
            +
            	has_null = true;
         | 
| 92 | 
            +
            	has_no_null = true;
         | 
| 93 | 
            +
            }
         | 
| 94 | 
            +
             | 
| 95 | 
            +
            void BaseStatistics::InitializeEmpty() {
         | 
| 96 | 
            +
            	has_null = false;
         | 
| 97 | 
            +
            	has_no_null = true;
         | 
| 98 | 
            +
            }
         | 
| 99 | 
            +
             | 
| 100 | 
            +
            bool BaseStatistics::CanHaveNull() const {
         | 
| 101 | 
            +
            	return has_null;
         | 
| 102 | 
            +
            }
         | 
| 103 | 
            +
             | 
| 104 | 
            +
            bool BaseStatistics::CanHaveNoNull() const {
         | 
| 105 | 
            +
            	return has_no_null;
         | 
| 106 | 
            +
            }
         | 
| 107 | 
            +
             | 
| 108 | 
            +
            bool BaseStatistics::IsConstant() const {
         | 
| 109 | 
            +
            	if (type.id() == LogicalTypeId::VALIDITY) {
         | 
| 110 | 
            +
            		// validity mask
         | 
| 111 | 
            +
            		if (CanHaveNull() && !CanHaveNoNull()) {
         | 
| 112 | 
            +
            			return true;
         | 
| 113 | 
            +
            		}
         | 
| 114 | 
            +
            		if (!CanHaveNull() && CanHaveNoNull()) {
         | 
| 115 | 
            +
            			return true;
         | 
| 116 | 
            +
            		}
         | 
| 117 | 
            +
            		return false;
         | 
| 118 | 
            +
            	}
         | 
| 119 | 
            +
            	switch (GetStatsType()) {
         | 
| 120 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 121 | 
            +
            		return NumericStats::IsConstant(*this);
         | 
| 122 | 
            +
            	default:
         | 
| 123 | 
            +
            		break;
         | 
| 124 | 
            +
            	}
         | 
| 125 | 
            +
            	return false;
         | 
| 126 | 
            +
            }
         | 
| 127 | 
            +
             | 
| 128 | 
            +
            void BaseStatistics::Merge(const BaseStatistics &other) {
         | 
| 129 | 
            +
            	has_null = has_null || other.has_null;
         | 
| 130 | 
            +
            	has_no_null = has_no_null || other.has_no_null;
         | 
| 131 | 
            +
            	switch (GetStatsType()) {
         | 
| 132 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 133 | 
            +
            		NumericStats::Merge(*this, other);
         | 
| 134 | 
            +
            		break;
         | 
| 135 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 136 | 
            +
            		StringStats::Merge(*this, other);
         | 
| 137 | 
            +
            		break;
         | 
| 138 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 139 | 
            +
            		ListStats::Merge(*this, other);
         | 
| 140 | 
            +
            		break;
         | 
| 141 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 142 | 
            +
            		StructStats::Merge(*this, other);
         | 
| 143 | 
            +
            		break;
         | 
| 144 | 
            +
            	default:
         | 
| 145 | 
            +
            		break;
         | 
| 146 | 
            +
            	}
         | 
| 147 | 
            +
            }
         | 
| 148 | 
            +
             | 
| 149 | 
            +
            idx_t BaseStatistics::GetDistinctCount() {
         | 
| 150 | 
            +
            	return distinct_count;
         | 
| 151 | 
            +
            }
         | 
| 152 | 
            +
             | 
| 153 | 
            +
            BaseStatistics BaseStatistics::CreateUnknownType(LogicalType type) {
         | 
| 154 | 
            +
            	switch (GetStatsType(type)) {
         | 
| 155 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 156 | 
            +
            		return NumericStats::CreateUnknown(std::move(type));
         | 
| 157 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 158 | 
            +
            		return StringStats::CreateUnknown(std::move(type));
         | 
| 159 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 160 | 
            +
            		return ListStats::CreateUnknown(std::move(type));
         | 
| 161 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 162 | 
            +
            		return StructStats::CreateUnknown(std::move(type));
         | 
| 163 | 
            +
            	default:
         | 
| 164 | 
            +
            		return BaseStatistics(std::move(type));
         | 
| 165 | 
            +
            	}
         | 
| 166 | 
            +
            }
         | 
| 167 | 
            +
             | 
| 168 | 
            +
            BaseStatistics BaseStatistics::CreateEmptyType(LogicalType type) {
         | 
| 169 | 
            +
            	switch (GetStatsType(type)) {
         | 
| 170 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 171 | 
            +
            		return NumericStats::CreateEmpty(std::move(type));
         | 
| 172 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 173 | 
            +
            		return StringStats::CreateEmpty(std::move(type));
         | 
| 174 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 175 | 
            +
            		return ListStats::CreateEmpty(std::move(type));
         | 
| 176 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 177 | 
            +
            		return StructStats::CreateEmpty(std::move(type));
         | 
| 178 | 
            +
            	default:
         | 
| 179 | 
            +
            		return BaseStatistics(std::move(type));
         | 
| 180 | 
            +
            	}
         | 
| 181 | 
            +
            }
         | 
| 182 | 
            +
             | 
| 183 | 
            +
            BaseStatistics BaseStatistics::CreateUnknown(LogicalType type) {
         | 
| 184 | 
            +
            	auto result = CreateUnknownType(std::move(type));
         | 
| 185 | 
            +
            	result.InitializeUnknown();
         | 
| 113 186 | 
             
            	return result;
         | 
| 114 187 | 
             
            }
         | 
| 115 188 |  | 
| 116 | 
            -
             | 
| 117 | 
            -
            	 | 
| 118 | 
            -
             | 
| 189 | 
            +
            BaseStatistics BaseStatistics::CreateEmpty(LogicalType type) {
         | 
| 190 | 
            +
            	if (type.InternalType() == PhysicalType::BIT) {
         | 
| 191 | 
            +
            		// FIXME: this special case should not be necessary
         | 
| 192 | 
            +
            		// but currently InitializeEmpty sets StatsInfo::CAN_HAVE_VALID_VALUES
         | 
| 193 | 
            +
            		BaseStatistics result(std::move(type));
         | 
| 194 | 
            +
            		result.Set(StatsInfo::CANNOT_HAVE_NULL_VALUES);
         | 
| 195 | 
            +
            		result.Set(StatsInfo::CANNOT_HAVE_VALID_VALUES);
         | 
| 196 | 
            +
            		return result;
         | 
| 197 | 
            +
            	}
         | 
| 198 | 
            +
            	auto result = CreateEmptyType(std::move(type));
         | 
| 199 | 
            +
            	result.InitializeEmpty();
         | 
| 119 200 | 
             
            	return result;
         | 
| 120 201 | 
             
            }
         | 
| 121 202 |  | 
| 122 | 
            -
            void BaseStatistics:: | 
| 123 | 
            -
            	 | 
| 124 | 
            -
             | 
| 203 | 
            +
            void BaseStatistics::Copy(const BaseStatistics &other) {
         | 
| 204 | 
            +
            	D_ASSERT(GetType() == other.GetType());
         | 
| 205 | 
            +
            	CopyBase(other);
         | 
| 206 | 
            +
            	stats_union = other.stats_union;
         | 
| 207 | 
            +
            	switch (GetStatsType()) {
         | 
| 208 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 209 | 
            +
            		ListStats::Copy(*this, other);
         | 
| 210 | 
            +
            		break;
         | 
| 211 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 212 | 
            +
            		StructStats::Copy(*this, other);
         | 
| 213 | 
            +
            		break;
         | 
| 214 | 
            +
            	default:
         | 
| 215 | 
            +
            		break;
         | 
| 125 216 | 
             
            	}
         | 
| 126 | 
            -
             | 
| 127 | 
            -
             | 
| 217 | 
            +
            }
         | 
| 218 | 
            +
             | 
| 219 | 
            +
            BaseStatistics BaseStatistics::Copy() const {
         | 
| 220 | 
            +
            	BaseStatistics result(type);
         | 
| 221 | 
            +
            	result.Copy(*this);
         | 
| 222 | 
            +
            	return result;
         | 
| 223 | 
            +
            }
         | 
| 224 | 
            +
             | 
| 225 | 
            +
            unique_ptr<BaseStatistics> BaseStatistics::ToUnique() const {
         | 
| 226 | 
            +
            	auto result = unique_ptr<BaseStatistics>(new BaseStatistics(type));
         | 
| 227 | 
            +
            	result->Copy(*this);
         | 
| 228 | 
            +
            	return result;
         | 
| 229 | 
            +
            }
         | 
| 230 | 
            +
             | 
| 231 | 
            +
            void BaseStatistics::CopyBase(const BaseStatistics &other) {
         | 
| 232 | 
            +
            	has_null = other.has_null;
         | 
| 233 | 
            +
            	has_no_null = other.has_no_null;
         | 
| 234 | 
            +
            	distinct_count = other.distinct_count;
         | 
| 235 | 
            +
            }
         | 
| 236 | 
            +
             | 
| 237 | 
            +
            void BaseStatistics::Set(StatsInfo info) {
         | 
| 238 | 
            +
            	switch (info) {
         | 
| 239 | 
            +
            	case StatsInfo::CAN_HAVE_NULL_VALUES:
         | 
| 240 | 
            +
            		has_null = true;
         | 
| 241 | 
            +
            		break;
         | 
| 242 | 
            +
            	case StatsInfo::CANNOT_HAVE_NULL_VALUES:
         | 
| 243 | 
            +
            		has_null = false;
         | 
| 244 | 
            +
            		break;
         | 
| 245 | 
            +
            	case StatsInfo::CAN_HAVE_VALID_VALUES:
         | 
| 246 | 
            +
            		has_no_null = true;
         | 
| 247 | 
            +
            		break;
         | 
| 248 | 
            +
            	case StatsInfo::CANNOT_HAVE_VALID_VALUES:
         | 
| 249 | 
            +
            		has_no_null = false;
         | 
| 250 | 
            +
            		break;
         | 
| 251 | 
            +
            	case StatsInfo::CAN_HAVE_NULL_AND_VALID_VALUES:
         | 
| 252 | 
            +
            		has_null = true;
         | 
| 253 | 
            +
            		has_no_null = true;
         | 
| 254 | 
            +
            		break;
         | 
| 255 | 
            +
            	default:
         | 
| 256 | 
            +
            		throw InternalException("Unrecognized StatsInfo for BaseStatistics::Set");
         | 
| 128 257 | 
             
            	}
         | 
| 129 258 | 
             
            }
         | 
| 130 259 |  | 
| 260 | 
            +
            void BaseStatistics::CombineValidity(BaseStatistics &left, BaseStatistics &right) {
         | 
| 261 | 
            +
            	has_null = left.has_null || right.has_null;
         | 
| 262 | 
            +
            	has_no_null = left.has_no_null || right.has_no_null;
         | 
| 263 | 
            +
            }
         | 
| 264 | 
            +
             | 
| 265 | 
            +
            void BaseStatistics::CopyValidity(BaseStatistics &stats) {
         | 
| 266 | 
            +
            	has_null = stats.has_null;
         | 
| 267 | 
            +
            	has_no_null = stats.has_no_null;
         | 
| 268 | 
            +
            }
         | 
| 269 | 
            +
             | 
| 131 270 | 
             
            void BaseStatistics::Serialize(Serializer &serializer) const {
         | 
| 132 271 | 
             
            	FieldWriter writer(serializer);
         | 
| 133 | 
            -
            	 | 
| 272 | 
            +
            	writer.WriteField<bool>(has_null);
         | 
| 273 | 
            +
            	writer.WriteField<bool>(has_no_null);
         | 
| 134 274 | 
             
            	Serialize(writer);
         | 
| 135 | 
            -
            	auto ptype = type.InternalType();
         | 
| 136 | 
            -
            	if (ptype != PhysicalType::BIT) {
         | 
| 137 | 
            -
            		writer.WriteField<StatisticsType>(stats_type);
         | 
| 138 | 
            -
            		writer.WriteOptional<BaseStatistics>(distinct_stats);
         | 
| 139 | 
            -
            	}
         | 
| 140 275 | 
             
            	writer.Finalize();
         | 
| 141 276 | 
             
            }
         | 
| 142 277 |  | 
| 143 | 
            -
            void BaseStatistics:: | 
| 278 | 
            +
            void BaseStatistics::SetDistinctCount(idx_t count) {
         | 
| 279 | 
            +
            	this->distinct_count = count;
         | 
| 144 280 | 
             
            }
         | 
| 145 281 |  | 
| 146 | 
            -
             | 
| 147 | 
            -
            	 | 
| 148 | 
            -
            	 | 
| 149 | 
            -
             | 
| 150 | 
            -
            	auto ptype = type.InternalType();
         | 
| 151 | 
            -
            	switch (ptype) {
         | 
| 152 | 
            -
            	case PhysicalType::BIT:
         | 
| 153 | 
            -
            		result = ValidityStatistics::Deserialize(reader);
         | 
| 154 | 
            -
            		break;
         | 
| 155 | 
            -
            	case PhysicalType::BOOL:
         | 
| 156 | 
            -
            	case PhysicalType::INT8:
         | 
| 157 | 
            -
            	case PhysicalType::INT16:
         | 
| 158 | 
            -
            	case PhysicalType::INT32:
         | 
| 159 | 
            -
            	case PhysicalType::INT64:
         | 
| 160 | 
            -
            	case PhysicalType::UINT8:
         | 
| 161 | 
            -
            	case PhysicalType::UINT16:
         | 
| 162 | 
            -
            	case PhysicalType::UINT32:
         | 
| 163 | 
            -
            	case PhysicalType::UINT64:
         | 
| 164 | 
            -
            	case PhysicalType::INT128:
         | 
| 165 | 
            -
            	case PhysicalType::FLOAT:
         | 
| 166 | 
            -
            	case PhysicalType::DOUBLE:
         | 
| 167 | 
            -
            		result = NumericStatistics::Deserialize(reader, std::move(type));
         | 
| 168 | 
            -
            		break;
         | 
| 169 | 
            -
            	case PhysicalType::VARCHAR:
         | 
| 170 | 
            -
            		result = StringStatistics::Deserialize(reader, std::move(type));
         | 
| 282 | 
            +
            void BaseStatistics::Serialize(FieldWriter &writer) const {
         | 
| 283 | 
            +
            	switch (GetStatsType()) {
         | 
| 284 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 285 | 
            +
            		NumericStats::Serialize(*this, writer);
         | 
| 171 286 | 
             
            		break;
         | 
| 172 | 
            -
            	case  | 
| 173 | 
            -
            		 | 
| 287 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 288 | 
            +
            		StringStats::Serialize(*this, writer);
         | 
| 174 289 | 
             
            		break;
         | 
| 175 | 
            -
            	case  | 
| 176 | 
            -
            		 | 
| 290 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 291 | 
            +
            		ListStats::Serialize(*this, writer);
         | 
| 177 292 | 
             
            		break;
         | 
| 178 | 
            -
            	case  | 
| 179 | 
            -
            		 | 
| 293 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 294 | 
            +
            		StructStats::Serialize(*this, writer);
         | 
| 180 295 | 
             
            		break;
         | 
| 181 296 | 
             
            	default:
         | 
| 182 | 
            -
            		 | 
| 297 | 
            +
            		break;
         | 
| 183 298 | 
             
            	}
         | 
| 184 | 
            -
             | 
| 185 | 
            -
             | 
| 186 | 
            -
             | 
| 187 | 
            -
             | 
| 188 | 
            -
            		 | 
| 299 | 
            +
            }
         | 
| 300 | 
            +
            BaseStatistics BaseStatistics::DeserializeType(FieldReader &reader, LogicalType type) {
         | 
| 301 | 
            +
            	switch (GetStatsType(type)) {
         | 
| 302 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 303 | 
            +
            		return NumericStats::Deserialize(reader, std::move(type));
         | 
| 304 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 305 | 
            +
            		return StringStats::Deserialize(reader, std::move(type));
         | 
| 306 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 307 | 
            +
            		return ListStats::Deserialize(reader, std::move(type));
         | 
| 308 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 309 | 
            +
            		return StructStats::Deserialize(reader, std::move(type));
         | 
| 310 | 
            +
            	default:
         | 
| 311 | 
            +
            		return BaseStatistics(std::move(type));
         | 
| 189 312 | 
             
            	}
         | 
| 313 | 
            +
            }
         | 
| 190 314 |  | 
| 315 | 
            +
            BaseStatistics BaseStatistics::Deserialize(Deserializer &source, LogicalType type) {
         | 
| 316 | 
            +
            	FieldReader reader(source);
         | 
| 317 | 
            +
            	bool has_null = reader.ReadRequired<bool>();
         | 
| 318 | 
            +
            	bool has_no_null = reader.ReadRequired<bool>();
         | 
| 319 | 
            +
            	auto result = DeserializeType(reader, std::move(type));
         | 
| 320 | 
            +
            	result.has_null = has_null;
         | 
| 321 | 
            +
            	result.has_no_null = has_no_null;
         | 
| 191 322 | 
             
            	reader.Finalize();
         | 
| 192 323 | 
             
            	return result;
         | 
| 193 324 | 
             
            }
         | 
| 194 325 |  | 
| 195 326 | 
             
            string BaseStatistics::ToString() const {
         | 
| 196 | 
            -
            	 | 
| 197 | 
            -
             | 
| 327 | 
            +
            	auto has_n = has_null ? "true" : "false";
         | 
| 328 | 
            +
            	auto has_n_n = has_no_null ? "true" : "false";
         | 
| 329 | 
            +
            	string result =
         | 
| 330 | 
            +
            	    StringUtil::Format("%s%s", StringUtil::Format("[Has Null: %s, Has No Null: %s]", has_n, has_n_n),
         | 
| 331 | 
            +
            	                       distinct_count > 0 ? StringUtil::Format("[Approx Unique: %lld]", distinct_count) : "");
         | 
| 332 | 
            +
            	switch (GetStatsType()) {
         | 
| 333 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 334 | 
            +
            		result = NumericStats::ToString(*this) + result;
         | 
| 335 | 
            +
            		break;
         | 
| 336 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 337 | 
            +
            		result = StringStats::ToString(*this) + result;
         | 
| 338 | 
            +
            		break;
         | 
| 339 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 340 | 
            +
            		result = ListStats::ToString(*this) + result;
         | 
| 341 | 
            +
            		break;
         | 
| 342 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 343 | 
            +
            		result = StructStats::ToString(*this) + result;
         | 
| 344 | 
            +
            		break;
         | 
| 345 | 
            +
            	default:
         | 
| 346 | 
            +
            		break;
         | 
| 347 | 
            +
            	}
         | 
| 348 | 
            +
            	return result;
         | 
| 198 349 | 
             
            }
         | 
| 199 350 |  | 
| 200 351 | 
             
            void BaseStatistics::Verify(Vector &vector, const SelectionVector &sel, idx_t count) const {
         | 
| 201 352 | 
             
            	D_ASSERT(vector.GetType() == this->type);
         | 
| 202 | 
            -
            	 | 
| 203 | 
            -
             | 
| 353 | 
            +
            	switch (GetStatsType()) {
         | 
| 354 | 
            +
            	case StatisticsType::NUMERIC_STATS:
         | 
| 355 | 
            +
            		NumericStats::Verify(*this, vector, sel, count);
         | 
| 356 | 
            +
            		break;
         | 
| 357 | 
            +
            	case StatisticsType::STRING_STATS:
         | 
| 358 | 
            +
            		StringStats::Verify(*this, vector, sel, count);
         | 
| 359 | 
            +
            		break;
         | 
| 360 | 
            +
            	case StatisticsType::LIST_STATS:
         | 
| 361 | 
            +
            		ListStats::Verify(*this, vector, sel, count);
         | 
| 362 | 
            +
            		break;
         | 
| 363 | 
            +
            	case StatisticsType::STRUCT_STATS:
         | 
| 364 | 
            +
            		StructStats::Verify(*this, vector, sel, count);
         | 
| 365 | 
            +
            		break;
         | 
| 366 | 
            +
            	default:
         | 
| 367 | 
            +
            		break;
         | 
| 368 | 
            +
            	}
         | 
| 369 | 
            +
            	if (has_null && has_no_null) {
         | 
| 370 | 
            +
            		// nothing to verify
         | 
| 371 | 
            +
            		return;
         | 
| 372 | 
            +
            	}
         | 
| 373 | 
            +
            	UnifiedVectorFormat vdata;
         | 
| 374 | 
            +
            	vector.ToUnifiedFormat(count, vdata);
         | 
| 375 | 
            +
            	for (idx_t i = 0; i < count; i++) {
         | 
| 376 | 
            +
            		auto idx = sel.get_index(i);
         | 
| 377 | 
            +
            		auto index = vdata.sel->get_index(idx);
         | 
| 378 | 
            +
            		bool row_is_valid = vdata.validity.RowIsValid(index);
         | 
| 379 | 
            +
            		if (row_is_valid && !has_no_null) {
         | 
| 380 | 
            +
            			throw InternalException(
         | 
| 381 | 
            +
            			    "Statistics mismatch: vector labeled as having only NULL values, but vector contains valid values: %s",
         | 
| 382 | 
            +
            			    vector.ToString(count));
         | 
| 383 | 
            +
            		}
         | 
| 384 | 
            +
            		if (!row_is_valid && !has_null) {
         | 
| 385 | 
            +
            			throw InternalException(
         | 
| 386 | 
            +
            			    "Statistics mismatch: vector labeled as not having NULL values, but vector contains null values: %s",
         | 
| 387 | 
            +
            			    vector.ToString(count));
         | 
| 388 | 
            +
            		}
         | 
| 204 389 | 
             
            	}
         | 
| 205 390 | 
             
            }
         | 
| 206 391 |  | 
| @@ -209,4 +394,64 @@ void BaseStatistics::Verify(Vector &vector, idx_t count) const { | |
| 209 394 | 
             
            	Verify(vector, *sel, count);
         | 
| 210 395 | 
             
            }
         | 
| 211 396 |  | 
| 397 | 
            +
            BaseStatistics BaseStatistics::FromConstantType(const Value &input) {
         | 
| 398 | 
            +
            	switch (GetStatsType(input.type())) {
         | 
| 399 | 
            +
            	case StatisticsType::NUMERIC_STATS: {
         | 
| 400 | 
            +
            		auto result = NumericStats::CreateEmpty(input.type());
         | 
| 401 | 
            +
            		NumericStats::SetMin(result, input);
         | 
| 402 | 
            +
            		NumericStats::SetMax(result, input);
         | 
| 403 | 
            +
            		return result;
         | 
| 404 | 
            +
            	}
         | 
| 405 | 
            +
            	case StatisticsType::STRING_STATS: {
         | 
| 406 | 
            +
            		auto result = StringStats::CreateEmpty(input.type());
         | 
| 407 | 
            +
            		if (!input.IsNull()) {
         | 
| 408 | 
            +
            			auto &string_value = StringValue::Get(input);
         | 
| 409 | 
            +
            			StringStats::Update(result, string_t(string_value));
         | 
| 410 | 
            +
            		}
         | 
| 411 | 
            +
            		return result;
         | 
| 412 | 
            +
            	}
         | 
| 413 | 
            +
            	case StatisticsType::LIST_STATS: {
         | 
| 414 | 
            +
            		auto result = ListStats::CreateEmpty(input.type());
         | 
| 415 | 
            +
            		auto &child_stats = ListStats::GetChildStats(result);
         | 
| 416 | 
            +
            		if (!input.IsNull()) {
         | 
| 417 | 
            +
            			auto &list_children = ListValue::GetChildren(input);
         | 
| 418 | 
            +
            			for (auto &child_element : list_children) {
         | 
| 419 | 
            +
            				child_stats.Merge(FromConstant(child_element));
         | 
| 420 | 
            +
            			}
         | 
| 421 | 
            +
            		}
         | 
| 422 | 
            +
            		return result;
         | 
| 423 | 
            +
            	}
         | 
| 424 | 
            +
            	case StatisticsType::STRUCT_STATS: {
         | 
| 425 | 
            +
            		auto result = StructStats::CreateEmpty(input.type());
         | 
| 426 | 
            +
            		auto &child_types = StructType::GetChildTypes(input.type());
         | 
| 427 | 
            +
            		if (input.IsNull()) {
         | 
| 428 | 
            +
            			for (idx_t i = 0; i < child_types.size(); i++) {
         | 
| 429 | 
            +
            				StructStats::SetChildStats(result, i, FromConstant(Value(child_types[i].second)));
         | 
| 430 | 
            +
            			}
         | 
| 431 | 
            +
            		} else {
         | 
| 432 | 
            +
            			auto &struct_children = StructValue::GetChildren(input);
         | 
| 433 | 
            +
            			for (idx_t i = 0; i < child_types.size(); i++) {
         | 
| 434 | 
            +
            				StructStats::SetChildStats(result, i, FromConstant(struct_children[i]));
         | 
| 435 | 
            +
            			}
         | 
| 436 | 
            +
            		}
         | 
| 437 | 
            +
            		return result;
         | 
| 438 | 
            +
            	}
         | 
| 439 | 
            +
            	default:
         | 
| 440 | 
            +
            		return BaseStatistics(input.type());
         | 
| 441 | 
            +
            	}
         | 
| 442 | 
            +
            }
         | 
| 443 | 
            +
             | 
| 444 | 
            +
            BaseStatistics BaseStatistics::FromConstant(const Value &input) {
         | 
| 445 | 
            +
            	auto result = FromConstantType(input);
         | 
| 446 | 
            +
            	result.SetDistinctCount(1);
         | 
| 447 | 
            +
            	if (input.IsNull()) {
         | 
| 448 | 
            +
            		result.Set(StatsInfo::CAN_HAVE_NULL_VALUES);
         | 
| 449 | 
            +
            		result.Set(StatsInfo::CANNOT_HAVE_VALID_VALUES);
         | 
| 450 | 
            +
            	} else {
         | 
| 451 | 
            +
            		result.Set(StatsInfo::CANNOT_HAVE_NULL_VALUES);
         | 
| 452 | 
            +
            		result.Set(StatsInfo::CAN_HAVE_VALID_VALUES);
         | 
| 453 | 
            +
            	}
         | 
| 454 | 
            +
            	return result;
         | 
| 455 | 
            +
            }
         | 
| 456 | 
            +
             | 
| 212 457 | 
             
            } // namespace duckdb
         | 
| @@ -1,13 +1,67 @@ | |
| 1 1 | 
             
            #include "duckdb/storage/statistics/column_statistics.hpp"
         | 
| 2 | 
            +
            #include "duckdb/common/serializer.hpp"
         | 
| 2 3 |  | 
| 3 4 | 
             
            namespace duckdb {
         | 
| 4 5 |  | 
| 5 | 
            -
            ColumnStatistics::ColumnStatistics( | 
| 6 | 
            +
            ColumnStatistics::ColumnStatistics(BaseStatistics stats_p) : stats(std::move(stats_p)) {
         | 
| 7 | 
            +
            	if (DistinctStatistics::TypeIsSupported(stats.GetType())) {
         | 
| 8 | 
            +
            		distinct_stats = make_unique<DistinctStatistics>();
         | 
| 9 | 
            +
            	}
         | 
| 10 | 
            +
            }
         | 
| 11 | 
            +
            ColumnStatistics::ColumnStatistics(BaseStatistics stats_p, unique_ptr<DistinctStatistics> distinct_stats_p)
         | 
| 12 | 
            +
                : stats(std::move(stats_p)), distinct_stats(std::move(distinct_stats_p)) {
         | 
| 6 13 | 
             
            }
         | 
| 7 14 |  | 
| 8 15 | 
             
            shared_ptr<ColumnStatistics> ColumnStatistics::CreateEmptyStats(const LogicalType &type) {
         | 
| 9 | 
            -
            	 | 
| 10 | 
            -
             | 
| 16 | 
            +
            	return make_shared<ColumnStatistics>(BaseStatistics::CreateEmpty(type));
         | 
| 17 | 
            +
            }
         | 
| 18 | 
            +
             | 
| 19 | 
            +
            void ColumnStatistics::Merge(ColumnStatistics &other) {
         | 
| 20 | 
            +
            	stats.Merge(other.stats);
         | 
| 21 | 
            +
            	if (distinct_stats) {
         | 
| 22 | 
            +
            		distinct_stats->Merge(*other.distinct_stats);
         | 
| 23 | 
            +
            	}
         | 
| 24 | 
            +
            }
         | 
| 25 | 
            +
             | 
| 26 | 
            +
            BaseStatistics &ColumnStatistics::Statistics() {
         | 
| 27 | 
            +
            	return stats;
         | 
| 28 | 
            +
            }
         | 
| 29 | 
            +
             | 
| 30 | 
            +
            bool ColumnStatistics::HasDistinctStats() {
         | 
| 31 | 
            +
            	return distinct_stats.get();
         | 
| 32 | 
            +
            }
         | 
| 33 | 
            +
             | 
| 34 | 
            +
            DistinctStatistics &ColumnStatistics::DistinctStats() {
         | 
| 35 | 
            +
            	if (!distinct_stats) {
         | 
| 36 | 
            +
            		throw InternalException("DistinctStats called without distinct_stats");
         | 
| 37 | 
            +
            	}
         | 
| 38 | 
            +
            	return *distinct_stats;
         | 
| 39 | 
            +
            }
         | 
| 40 | 
            +
             | 
| 41 | 
            +
            void ColumnStatistics::SetDistinct(unique_ptr<DistinctStatistics> distinct) {
         | 
| 42 | 
            +
            	this->distinct_stats = std::move(distinct);
         | 
| 43 | 
            +
            }
         | 
| 44 | 
            +
             | 
| 45 | 
            +
            void ColumnStatistics::UpdateDistinctStatistics(Vector &v, idx_t count) {
         | 
| 46 | 
            +
            	if (!distinct_stats) {
         | 
| 47 | 
            +
            		return;
         | 
| 48 | 
            +
            	}
         | 
| 49 | 
            +
            	auto &d_stats = (DistinctStatistics &)*distinct_stats;
         | 
| 50 | 
            +
            	d_stats.Update(v, count);
         | 
| 51 | 
            +
            }
         | 
| 52 | 
            +
             | 
| 53 | 
            +
            shared_ptr<ColumnStatistics> ColumnStatistics::Copy() const {
         | 
| 54 | 
            +
            	return make_shared<ColumnStatistics>(stats.Copy(), distinct_stats ? distinct_stats->Copy() : nullptr);
         | 
| 55 | 
            +
            }
         | 
| 56 | 
            +
            void ColumnStatistics::Serialize(Serializer &serializer) const {
         | 
| 57 | 
            +
            	stats.Serialize(serializer);
         | 
| 58 | 
            +
            	serializer.WriteOptional(distinct_stats);
         | 
| 59 | 
            +
            }
         | 
| 60 | 
            +
             | 
| 61 | 
            +
            shared_ptr<ColumnStatistics> ColumnStatistics::Deserialize(Deserializer &source, const LogicalType &type) {
         | 
| 62 | 
            +
            	auto stats = BaseStatistics::Deserialize(source, type);
         | 
| 63 | 
            +
            	auto distinct_stats = source.ReadOptional<DistinctStatistics>();
         | 
| 64 | 
            +
            	return make_shared<ColumnStatistics>(stats.Copy(), std::move(distinct_stats));
         | 
| 11 65 | 
             
            }
         | 
| 12 66 |  | 
| 13 67 | 
             
            } // namespace duckdb
         | 
| @@ -7,23 +7,18 @@ | |
| 7 7 |  | 
| 8 8 | 
             
            namespace duckdb {
         | 
| 9 9 |  | 
| 10 | 
            -
            DistinctStatistics::DistinctStatistics()
         | 
| 11 | 
            -
                : BaseStatistics(LogicalType::INVALID, StatisticsType::LOCAL_STATS), log(make_unique<HyperLogLog>()),
         | 
| 12 | 
            -
                  sample_count(0), total_count(0) {
         | 
| 10 | 
            +
            DistinctStatistics::DistinctStatistics() : log(make_unique<HyperLogLog>()), sample_count(0), total_count(0) {
         | 
| 13 11 | 
             
            }
         | 
| 14 12 |  | 
| 15 13 | 
             
            DistinctStatistics::DistinctStatistics(unique_ptr<HyperLogLog> log, idx_t sample_count, idx_t total_count)
         | 
| 16 | 
            -
                :  | 
| 17 | 
            -
                  sample_count(sample_count), total_count(total_count) {
         | 
| 14 | 
            +
                : log(std::move(log)), sample_count(sample_count), total_count(total_count) {
         | 
| 18 15 | 
             
            }
         | 
| 19 16 |  | 
| 20 | 
            -
            unique_ptr< | 
| 17 | 
            +
            unique_ptr<DistinctStatistics> DistinctStatistics::Copy() const {
         | 
| 21 18 | 
             
            	return make_unique<DistinctStatistics>(log->Copy(), sample_count, total_count);
         | 
| 22 19 | 
             
            }
         | 
| 23 20 |  | 
| 24 | 
            -
            void DistinctStatistics::Merge(const  | 
| 25 | 
            -
            	BaseStatistics::Merge(other_p);
         | 
| 26 | 
            -
            	auto &other = (const DistinctStatistics &)other_p;
         | 
| 21 | 
            +
            void DistinctStatistics::Merge(const DistinctStatistics &other) {
         | 
| 27 22 | 
             
            	log = log->Merge(*other.log);
         | 
| 28 23 | 
             
            	sample_count += other.sample_count;
         | 
| 29 24 | 
             
            	total_count += other.total_count;
         | 
| @@ -99,4 +94,8 @@ idx_t DistinctStatistics::GetCount() const { | |
| 99 94 | 
             
            	return MinValue<idx_t>(estimate, total_count);
         | 
| 100 95 | 
             
            }
         | 
| 101 96 |  | 
| 97 | 
            +
            bool DistinctStatistics::TypeIsSupported(const LogicalType &type) {
         | 
| 98 | 
            +
            	return type.InternalType() != PhysicalType::LIST && type.InternalType() != PhysicalType::STRUCT;
         | 
| 99 | 
            +
            }
         | 
| 100 | 
            +
             | 
| 102 101 | 
             
            } // namespace duckdb
         |