AutoRAG 0.3.22__tar.gz → 0.3.24__tar.gz
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- {autorag-0.3.22 → autorag-0.3.24}/PKG-INFO +22 -19
- {autorag-0.3.22 → autorag-0.3.24}/README.md +4 -1
- {autorag-0.3.22 → autorag-0.3.24}/autorag/__init__.py +8 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/table_hybrid_parse.py +6 -6
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/evolve/openai_query_evolve.py +8 -8
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/filter/dontknow.py +4 -4
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/filter/passage_dependency.py +4 -4
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/generation_gt/openai_gen_gt.py +4 -4
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/query/openai_gen_query.py +8 -8
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/generation.py +26 -9
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/util.py +10 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/node_line.py +4 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/minimax_llm.py +6 -5
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/openai_llm.py +52 -35
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/nvidia.py +3 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/support.py +2 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/utils/util.py +17 -1
- autorag-0.3.24/docs/source/azure_openai.md +188 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/evaluate_metrics/generation.md +14 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/index.rst +1 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/local_model.md +1 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/generator/minimax_llm.md +13 -9
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/generator/openai_llm.md +20 -1
- {autorag-0.3.22 → autorag-0.3.24}/pyproject.toml +18 -19
- autorag-0.3.24/sample_config/rag/english/non_gpu/simple_azure_openai.yaml +44 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/simple_minimax.yaml +1 -1
- autorag-0.3.24/tests/autorag/data/qa/evolve/test_openai_query_evolve.py +68 -0
- autorag-0.3.24/tests/autorag/data/qa/filter/test_dontknow.py +159 -0
- autorag-0.3.24/tests/autorag/data/qa/filter/test_passage_dependency.py +159 -0
- autorag-0.3.24/tests/autorag/data/qa/generation_gt/test_openai_gen_gt.py +82 -0
- autorag-0.3.24/tests/autorag/data/qa/query/test_openai_gen_query.py +122 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/metric/test_generation_metric.py +40 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/generator/test_minimax.py +65 -50
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/generator/test_openai.py +70 -29
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/generator/test_vllm.py +6 -5
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/hybridretrieval/test_run_hybrid_retrieval.py +5 -4
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/lexicalretrieval/test_run_lexical_retrieval.py +6 -2
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_passage_filter_run.py +4 -5
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_nvidia_reranker.py +10 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_passage_reranker_run.py +4 -6
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/semanticretrieval/test_run_semantic_retrieval.py +5 -4
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_cli.py +7 -5
- autorag-0.3.24/tests/autorag/test_node_line.py +30 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/utils/test_util.py +24 -0
- autorag-0.3.24/tests/autorag/vectordb/test_couchbase.py +77 -0
- autorag-0.3.24/tests/autorag/vectordb/test_milvus.py +75 -0
- autorag-0.3.24/tests/autorag/vectordb/test_pinecone.py +74 -0
- autorag-0.3.24/tests/autorag/vectordb/test_qdrant.py +74 -0
- autorag-0.3.24/tests/autorag/vectordb/test_weaviate.py +73 -0
- {autorag-0.3.22 → autorag-0.3.24}/uv.lock +4197 -3953
- autorag-0.3.22/.github/FUNDING.yml +0 -2
- autorag-0.3.22/.github/ISSUE_TEMPLATE/bug_report.md +0 -35
- autorag-0.3.22/.github/ISSUE_TEMPLATE/feature_request.md +0 -20
- autorag-0.3.22/.github/copilot-instructions.md +0 -317
- autorag-0.3.22/.github/dependabot.yml +0 -11
- autorag-0.3.22/.github/workflows/publish.yml +0 -63
- autorag-0.3.22/.github/workflows/sphinx.yml +0 -41
- autorag-0.3.22/.github/workflows/test.yml +0 -53
- autorag-0.3.22/CODE_OF_CONDUCT.md +0 -39
- autorag-0.3.22/tests/autorag/data/qa/evolve/test_openai_query_evolve.py +0 -88
- autorag-0.3.22/tests/autorag/data/qa/filter/test_dontknow.py +0 -177
- autorag-0.3.22/tests/autorag/data/qa/filter/test_passage_dependency.py +0 -180
- autorag-0.3.22/tests/autorag/data/qa/generation_gt/test_openai_gen_gt.py +0 -102
- autorag-0.3.22/tests/autorag/data/qa/query/test_openai_gen_query.py +0 -155
- autorag-0.3.22/tests/autorag/vectordb/test_couchbase.py +0 -83
- autorag-0.3.22/tests/autorag/vectordb/test_milvus.py +0 -81
- autorag-0.3.22/tests/autorag/vectordb/test_pinecone.py +0 -80
- autorag-0.3.22/tests/autorag/vectordb/test_qdrant.py +0 -80
- autorag-0.3.22/tests/autorag/vectordb/test_weaviate.py +0 -79
- {autorag-0.3.22 → autorag-0.3.24}/.gitignore +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/.pre-commit-config.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/CNAME +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/CONTRIBUTING.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/LICENSE +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/chunker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/cli.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/dashboard.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/chunk/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/chunk/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/chunk/langchain_chunk.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/chunk/llama_index_chunk.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/chunk/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/corpus/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/corpus/langchain.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/corpus/llama_index.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/qacreation/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/qacreation/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/qacreation/llama_index.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/qacreation/llama_index_default_prompt.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/qacreation/ragas.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/legacy/qacreation/simple.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/clova.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/langchain_parse.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/llamaparse.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/parse/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/evolve/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/evolve/llama_index_query_evolve.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/evolve/prompt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/extract_evidence.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/filter/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/filter/prompt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/generation_gt/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/generation_gt/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/generation_gt/llama_index_gen_gt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/generation_gt/prompt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/query/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/query/llama_gen_query.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/query/prompt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/sample.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/qa/schema.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/utils/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/data/utils/util.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/deploy/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/deploy/api.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/deploy/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/deploy/gradio.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/deploy/swagger.yml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/embedding/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/embedding/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/embedding/vllm.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/generation.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/deepeval_prompt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/g_eval_prompts/coh_detailed.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/g_eval_prompts/con_detailed.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/g_eval_prompts/flu_detailed.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/g_eval_prompts/rel_detailed.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/retrieval.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/metric/retrieval_contents.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/retrieval.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/retrieval_contents.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluation/util.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/evaluator.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/llama_index_llm.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/vllm.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/generator/vllm_api.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/hybridretrieval/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/hybridretrieval/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/hybridretrieval/hybrid_cc.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/hybridretrieval/hybrid_rrf.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/hybridretrieval/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/lexicalretrieval/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/lexicalretrieval/bm25.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/lexicalretrieval/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passageaugmenter/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passageaugmenter/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passageaugmenter/pass_passage_augmenter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passageaugmenter/prev_next_augmenter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passageaugmenter/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/longllmlingua.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/pass_compressor.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/refine.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagecompressor/tree_summarize.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/pass_passage_filter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/percentile_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/recency.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/similarity_percentile_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/similarity_threshold_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagefilter/threshold_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/cohere.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/colbert.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/flag_embedding.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/flag_embedding_llm.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/flashrank.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/jina.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/koreranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/mixedbreadai.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/monot5.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/openvino.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/pass_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/rankgpt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/sentence_transformer.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/tart/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/tart/modeling_enc_t5.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/tart/tart.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/tart/tokenization_enc_t5.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/time_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/upr.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/passagereranker/voyageai.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/chat_fstring.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/fstring.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/long_context_reorder.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/promptmaker/window_replacement.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/hyde.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/multi_query_expansion.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/pass_query_expansion.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/query_decompose.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/queryexpansion/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/retrieval/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/retrieval/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/retrieval/run_util.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/semanticretrieval/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/semanticretrieval/run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/semanticretrieval/vectordb.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/nodes/util.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/parser.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/schema/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/schema/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/schema/metricinput.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/schema/module.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/schema/node.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/strategy.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/utils/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/utils/cast.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/utils/preprocess.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/validator.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/chroma.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/couchbase.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/milvus.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/pinecone.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/qdrant.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/vectordb/weaviate.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/autorag/web.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/CNAME +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/Makefile +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/make.bat +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/requirements.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/CNAME +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/data_creation.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/data_creation_pipeline.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/data_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/dcg.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/f1_score.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/integration/couchbase_search_index.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/integration/nvidia_api.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/integration/nvidia_nim.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/integration/ollama_autorag.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/map.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/mrr.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/ndcg.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/ndcg_formula.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/node_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/node_line_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/node_line_summary.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/node_lines.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/node_summary.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/normal_distribution.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/project_folder_example.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/project_folders.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/qa/data_creation_schema.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/resources_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/roadmap/RAG_paradigms.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/roadmap/advanced_RAG.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/roadmap/cycle.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/roadmap/merger.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/roadmap/node_line_modular.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/roadmap/policy.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/samsung_sundae.jpeg +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/score_fusion.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/trial_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/trial_json.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/trial_summary.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/web_interface.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/web_interface_gradio.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/compact_structured.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/detail_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/full_modules.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/full_structured.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/full_yaml_structure.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/gpu_modules.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/half_structured.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/non_gpu_modules.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/sample_yaml_folder.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/_static/yaml/simple_structured.png +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.chunk.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.corpus.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.legacy.corpus.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.legacy.qacreation.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.legacy.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.parse.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.qa.evolve.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.qa.filter.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.qa.generation_gt.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.qa.query.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.qa.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.qacreation.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.data.utils.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.deploy.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.evaluation.metric.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.evaluation.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.generator.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.passageaugmenter.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.passagecompressor.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.passagefilter.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.passagereranker.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.passagereranker.tart.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.promptmaker.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.queryexpansion.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.retrieval.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.nodes.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.schema.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.utils.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/autorag.vectordb.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/api_spec/modules.rst +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/conf.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/chunk/chunk.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/chunk/langchain_chunk.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/chunk/llama_index_chunk.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/data_creation.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/data_format.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/legacy/legacy.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/legacy/parse.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/legacy/ragas.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/legacy/tutorial.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/parse/clova.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/parse/langchain_parse.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/parse/llama_parse.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/parse/parse.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/parse/table_hybrid_parse.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/qa_creation/answer_gen.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/qa_creation/evolve.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/qa_creation/filter.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/qa_creation/qa_creation.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/qa_creation/query_gen.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/data_creation/tutorial.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/deploy/api_endpoint.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/deploy/web.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/evaluate_metrics/retrieval.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/evaluate_metrics/retrieval_contents.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/install.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/llm/aws_bedrock.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/llm/huggingface_llm.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/llm/llm.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/llm/nvidia_nim.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/llm/ollama.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/chroma.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/couchbase.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/milvus.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/pinecone.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/qdrant.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/vectordb.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/integration/vectordb/weaviate.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/migration.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/generator/generator.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/generator/llama_index_llm.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/generator/vllm.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/generator/vllm_api.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/index.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_augmenter/passage_augmenter.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_augmenter/prev_next_augmenter.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_compressor/longllmlingua.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_compressor/passage_compressor.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_compressor/refine.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_compressor/tree_summarize.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_filter/passage_filter.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_filter/percentile_cutoff.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_filter/recency_filter.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_filter/similarity_percentile_cutoff.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_filter/similarity_threshold_cutoff.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_filter/threshold_cutoff.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/cohere.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/colbert.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/flag_embedding_llm_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/flag_embedding_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/flashrank_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/jina_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/koreranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/mixedbreadai_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/monot5.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/nvidia_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/openvino_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/passage_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/rankgpt.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/sentence_transformer_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/tart.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/time_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/upr.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/passage_reranker/voyageai_reranker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/prompt_maker/chat_fstring.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/prompt_maker/fstring.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/prompt_maker/long_context_reorder.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/prompt_maker/prompt_maker.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/prompt_maker/window_replacement.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/query_expansion/hyde.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/query_expansion/multi_query_expansion.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/query_expansion/query_decompose.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/query_expansion/query_expansion.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/retrieval/bm25.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/retrieval/hybrid_cc.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/retrieval/hybrid_rrf.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/retrieval/retrieval.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/nodes/retrieval/vectordb.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/optimization/custom_config.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/optimization/folder_structure.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/optimization/optimization.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/optimization/sample_config.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/optimization/strategies.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/roadmap/modular_rag.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/structure.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/test_your_rag.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/troubleshooting.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/docs/source/tutorial.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/chunk/chunk_full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/chunk/chunk_ko.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/chunk/simple_chunk.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/all_files_full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/file_types_full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/parse_hybird.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/parse_ko.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/parse_multimodal.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/parse_ocr.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/parse/simple_parse.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu/compact_local.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu/compact_openai.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu/full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu/half.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu_api/compact.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu_api/full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/gpu_api/half.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/compact.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/half.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/simple_bedrock.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/simple_local.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/simple_ollama.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/english/non_gpu/simple_openai.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/extracted_sample.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/gpu/compact_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/gpu/full_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/gpu/half_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/gpu_api/compact_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/gpu_api/full_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/gpu_api/half_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/non_gpu/compact_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/non_gpu/full_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/non_gpu/half_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_config/rag/korean/non_gpu/simple_korean.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_dataset/README.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_dataset/eli5/load_eli5_dataset.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_dataset/hotpotqa/load_hotpotqa_dataset.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_dataset/msmarco/load_msmarco_dataset.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/sample_dataset/triviaqa/load_triviaqa_dataset.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/__init__.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/chunk/test_chunk_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/chunk/test_chunk_run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/chunk/test_langchain_chunk.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/chunk/test_llama_index_chunk.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/corpus/test_base_corpus_legacy.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/corpus/test_langchain.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/corpus/test_llama_index_corpus.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/qacreation/test_base_qacreation.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/qacreation/test_llama_index_qacreation.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/qacreation/test_ragas_qa_creation.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/legacy/qacreation/test_simple.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/parse/test_clova.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/parse/test_langchain_parse.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/parse/test_llamaparse.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/parse/test_parse_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/parse/test_parse_run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/parse/test_table_hybrid_parse.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/evolve/base_test_query_evolve.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/evolve/test_llama_index_query_evolve.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/generation_gt/base_test_generation_gt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/generation_gt/test_llama_index_gen_gt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/query/base_test_query_gen.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/query/test_llama_gen_query.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/test_data_creation_piepline.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/test_sample.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/data/qa/test_schema.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/embedding/test_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/metric/test_retrieval_contents_metric.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/metric/test_retrieval_metric.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/test_evaluate_util.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/test_generation_evaluate.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/test_retrieval_contents_evaluate.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/evaluate/test_retrieval_evaluate.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/generator/test_generator_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/generator/test_llama_index_llm.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/generator/test_run_generator_node.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/hybridretrieval/test_hybrid_cc.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/hybridretrieval/test_hybrid_rrf.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/lexicalretrieval/test_bm25.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passageaugmenter/test_base_passage_augmenter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passageaugmenter/test_pass_passage_augmenter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passageaugmenter/test_prev_next_augmenter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passageaugmenter/test_run_passage_augmenter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagecompressor/test_base_passage_compressor.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagecompressor/test_longllmlingua.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagecompressor/test_pass_compressor.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagecompressor/test_refine.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagecompressor/test_run_passage_compressor_node.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagecompressor/test_tree_summarize.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_pass_passage_filter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_passage_filter_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_percentile_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_recency_filter.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_similarity_percentile_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_similarity_threshold_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagefilter/test_threshold_cutoff.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_cohere_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_colbert_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_flag_embedding.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_flag_embedding_llm.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_flashrank_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_jina_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_koreranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_mixedbreadai_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_monot5.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_openvino_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_pass_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_passage_reranker_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_rankgpt.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_sentence_transformer.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_tart.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_time_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_upr.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/passagereranker/test_voyageai_reranker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/promptmaker/test_chat_fstring.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/promptmaker/test_fstring.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/promptmaker/test_long_context_reorder.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/promptmaker/test_prompt_maker_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/promptmaker/test_prompt_maker_run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/promptmaker/test_window_replacement.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/queryexpansion/test_hyde.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/queryexpansion/test_multi_query_expansion.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/queryexpansion/test_pass_query_expansion.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/queryexpansion/test_query_decompose.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/queryexpansion/test_query_expansion_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/queryexpansion/test_query_expansion_run.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/retrieval/test_hybrid_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/retrieval/test_retrieval_base.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/nodes/semanticretrieval/test_vectordb.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/schema/test_base_schema.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/schema/test_metricinput_schema.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/schema/test_module_schema.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/schema/test_node_schema.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_chunker.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_dashboard.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_deploy.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_evaluator.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_parser.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_strategy.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_support.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_validator.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/test_web.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/utils/test_preprocess.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/vectordb/test_base_vectordb.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/autorag/vectordb/test_chroma.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/conftest.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/delete_tests.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/mock.py +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/README.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/chunk_data/sample_parsed.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/corpus_data_sample.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/data_creation/raw_dir/sample1.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/data_creation/raw_dir/sample2.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/data_creation/raw_dir/sample3.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/dataset_sample_gen_by_autorag/corpus.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/dataset_sample_gen_by_autorag/qa.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/full.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files/baseball_1.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files/csv_sample.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files_full/baseball_1.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files_full/csv_sample.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files_full/html_sample.html +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files_full/json_sample.json +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files_full/markdown_sample.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/all_files_full/xml_sample.xml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/clova_data/result_sample.json +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/clova_data/result_table.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/clova_data/result_text.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/config/all_files.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/config/lack_full_parse.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/config/lack_simple_parse.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/config/perfect_full_parse.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/config/perfect_simple_parse.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/csv_data/csv_sample.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/eng_text/baseball_1.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/eng_text/baseball_2.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/html_data/html_sample.html +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/hybrid_data/nfl_rulebook_both.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/json_data/json_sample.json +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/korean_table/only_table/kbo_only_table.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/korean_table/table_text/kbo_table_text.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/korean_text/korean_texts_two_page.pdf +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/markdown_data/markdown_sample.md +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/parse_data/xml_data/xml_sample.xml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/qa_data_sample.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/qa_gen_prompts/prompt1.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/qa_gen_prompts/prompt2.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/qa_gen_prompts/prompt3.txt +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/qa_test_data_sample.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/config.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/generator/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/generator/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/generator/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/prompt_maker/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/prompt_maker/1.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/prompt_maker/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/prompt_maker/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/post_retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/pre_retrieve_node_line/query_expansion/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/pre_retrieve_node_line/query_expansion/1.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/pre_retrieve_node_line/query_expansion/2.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/pre_retrieve_node_line/query_expansion/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/pre_retrieve_node_line/query_expansion/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/pre_retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/hybrid_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/hybrid_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/hybrid_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/lexical_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/lexical_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/lexical_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/passage_reranker/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/passage_reranker/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/passage_reranker/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/semantic_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/semantic_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/semantic_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/0/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/config.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/pre_retrieve_node_line/query_expansion/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/pre_retrieve_node_line/query_expansion/1.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/pre_retrieve_node_line/query_expansion/2.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/pre_retrieve_node_line/query_expansion/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/pre_retrieve_node_line/query_expansion/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/pre_retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/hybrid_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/hybrid_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/hybrid_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/lexical_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/lexical_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/lexical_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/passage_reranker/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/passage_reranker/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/passage_reranker/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/semantic_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/semantic_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/semantic_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/1/retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/config.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/post_retrieve_node_line/prompt_maker/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/post_retrieve_node_line/prompt_maker/1.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/post_retrieve_node_line/prompt_maker/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/post_retrieve_node_line/prompt_maker/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/pre_retrieve_node_line/query_expansion/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/pre_retrieve_node_line/query_expansion/1.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/pre_retrieve_node_line/query_expansion/2.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/pre_retrieve_node_line/query_expansion/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/pre_retrieve_node_line/query_expansion/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/pre_retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/hybrid_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/hybrid_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/hybrid_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/lexical_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/lexical_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/lexical_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/passage_reranker/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/passage_reranker/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/passage_reranker/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/semantic_retrieval/0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/semantic_retrieval/best_0.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/semantic_retrieval/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/2/retrieve_node_line/summary.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/3/config.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/best.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/data/corpus.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/data/qa.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/bm25_porter_stemmer.pkl +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/chroma/6595d304-5270-4d9f-a267-c191366804ce/data_level0.bin +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/chroma/6595d304-5270-4d9f-a267-c191366804ce/header.bin +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/chroma/6595d304-5270-4d9f-a267-c191366804ce/length.bin +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/chroma/6595d304-5270-4d9f-a267-c191366804ce/link_lists.bin +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/chroma/chroma.sqlite3 +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/resources/vectordb.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/result_project/trial.json +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/sample_contents_nqa.csv +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/sample_project/data/corpus.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/sample_project/data/qa.parquet +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/sample_project/resources/bm25_gpt2.pkl +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/sample_project/resources/bm25_porter_stemmer.pkl +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/sample_project/resources/vectordb.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple_chunk.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple_milvus.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple_mock.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple_mock_with_llm.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple_parse.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/simple_with_llm.yaml +0 -0
- {autorag-0.3.22 → autorag-0.3.24}/tests/resources/test_bm25_retrieval.pkl +0 -0
|
@@ -1,6 +1,6 @@
|
|
|
1
1
|
Metadata-Version: 2.4
|
|
2
2
|
Name: AutoRAG
|
|
3
|
-
Version: 0.3.
|
|
3
|
+
Version: 0.3.24
|
|
4
4
|
Summary: Automatically Evaluate RAG pipelines with your own data. Find optimal structure for new RAG product.
|
|
5
5
|
Project-URL: Homepage, https://github.com/Marker-Inc-Korea/AutoRAG
|
|
6
6
|
Author-email: Marker-Inc <vkehfdl1@gmail.com>
|
|
@@ -217,8 +217,8 @@ Classifier: Topic :: Scientific/Engineering :: Artificial Intelligence
|
|
|
217
217
|
Classifier: Topic :: Software Development :: Libraries
|
|
218
218
|
Classifier: Topic :: Software Development :: Libraries :: Python Modules
|
|
219
219
|
Requires-Python: <3.13,>=3.10
|
|
220
|
-
Requires-Dist: aiohttp
|
|
221
|
-
Requires-Dist: banks>=2.
|
|
220
|
+
Requires-Dist: aiohttp<3.14,>=3.12.13
|
|
221
|
+
Requires-Dist: banks>=2.4.2
|
|
222
222
|
Requires-Dist: chromadb>=1.0.0
|
|
223
223
|
Requires-Dist: click>=8.2.1
|
|
224
224
|
Requires-Dist: cohere>=5.18.0
|
|
@@ -228,19 +228,19 @@ Requires-Dist: emoji>=2.14.1
|
|
|
228
228
|
Requires-Dist: evaluate>=0.4.4
|
|
229
229
|
Requires-Dist: fastapi>=0.115.13
|
|
230
230
|
Requires-Dist: fastparquet>=2024.11.0
|
|
231
|
-
Requires-Dist: gradio>=
|
|
232
|
-
Requires-Dist: grpcio-health-checking<1.
|
|
233
|
-
Requires-Dist: grpcio-status<1.
|
|
234
|
-
Requires-Dist: grpcio-tools<1.
|
|
235
|
-
Requires-Dist: grpcio<1.
|
|
231
|
+
Requires-Dist: gradio>=6.15.0
|
|
232
|
+
Requires-Dist: grpcio-health-checking<1.83.0,>=1.66.2
|
|
233
|
+
Requires-Dist: grpcio-status<1.83.0,>=1.66.2
|
|
234
|
+
Requires-Dist: grpcio-tools<1.83.0,>=1.66.2
|
|
235
|
+
Requires-Dist: grpcio<1.83.0,>=1.66.2
|
|
236
236
|
Requires-Dist: ipykernel>=6.29.5
|
|
237
237
|
Requires-Dist: ipywidgets-bokeh>=1.6.0
|
|
238
238
|
Requires-Dist: ipywidgets>=8.1.7
|
|
239
239
|
Requires-Dist: langchain-community>=0.4.1
|
|
240
|
-
Requires-Dist: langchain-core>=1.
|
|
241
|
-
Requires-Dist: langchain-text-splitters>=1.1.
|
|
242
|
-
Requires-Dist: langchain>=1.
|
|
243
|
-
Requires-Dist: langsmith>=0.
|
|
240
|
+
Requires-Dist: langchain-core>=1.3.3
|
|
241
|
+
Requires-Dist: langchain-text-splitters>=1.1.2
|
|
242
|
+
Requires-Dist: langchain>=1.3.9
|
|
243
|
+
Requires-Dist: langsmith>=0.8.18
|
|
244
244
|
Requires-Dist: llama-index-core>=0.12.42
|
|
245
245
|
Requires-Dist: llama-index-embeddings-ollama>=0.6.0
|
|
246
246
|
Requires-Dist: llama-index-embeddings-openai-like>=0.1.1
|
|
@@ -252,14 +252,14 @@ Requires-Dist: llama-index-readers-file>=0.4.9
|
|
|
252
252
|
Requires-Dist: llama-index-retrievers-bm25>=0.5.2
|
|
253
253
|
Requires-Dist: llama-index>=0.12.42
|
|
254
254
|
Requires-Dist: mixedbread-ai>=2.2.6
|
|
255
|
-
Requires-Dist: numpy
|
|
255
|
+
Requires-Dist: numpy>=2.2.6
|
|
256
256
|
Requires-Dist: openai>=2.8.0
|
|
257
257
|
Requires-Dist: pandas>=2.2.3
|
|
258
258
|
Requires-Dist: panel>=1.7.1
|
|
259
259
|
Requires-Dist: pinecone[grpc]
|
|
260
|
-
Requires-Dist: pyarrow>=
|
|
260
|
+
Requires-Dist: pyarrow>=23.0.1
|
|
261
261
|
Requires-Dist: pydantic>=2.9.2
|
|
262
|
-
Requires-Dist: pymilvus
|
|
262
|
+
Requires-Dist: pymilvus<3,>=2.6.0b0
|
|
263
263
|
Requires-Dist: pyngrok>=7.2.11
|
|
264
264
|
Requires-Dist: pyyaml>=6.0.2
|
|
265
265
|
Requires-Dist: qdrant-client>=1.12.1
|
|
@@ -270,7 +270,7 @@ Requires-Dist: rouge-score>=0.1.2
|
|
|
270
270
|
Requires-Dist: sacrebleu>=2.5.1
|
|
271
271
|
Requires-Dist: scikit-learn>=1.7.0
|
|
272
272
|
Requires-Dist: seaborn>=0.13.2
|
|
273
|
-
Requires-Dist: streamlit>=1.
|
|
273
|
+
Requires-Dist: streamlit>=1.54.0
|
|
274
274
|
Requires-Dist: tiktoken>=0.9.0
|
|
275
275
|
Requires-Dist: tokenlog>=0.0.3
|
|
276
276
|
Requires-Dist: tqdm>=4.67.1
|
|
@@ -292,7 +292,7 @@ Requires-Dist: pdfminer-six>=20250506; extra == 'all'
|
|
|
292
292
|
Requires-Dist: pdfplumber>=0.11.7; extra == 'all'
|
|
293
293
|
Requires-Dist: peft>=0.15.2; extra == 'all'
|
|
294
294
|
Requires-Dist: pymupdf>=1.26.1; extra == 'all'
|
|
295
|
-
Requires-Dist: pypdf2
|
|
295
|
+
Requires-Dist: pypdf2>=3.0.1; extra == 'all'
|
|
296
296
|
Requires-Dist: sentence-transformers>=4.1.0; extra == 'all'
|
|
297
297
|
Requires-Dist: sentencepiece>=0.2.0; extra == 'all'
|
|
298
298
|
Requires-Dist: sudachidict-core; extra == 'all'
|
|
@@ -327,11 +327,14 @@ Requires-Dist: pdf2image>=1.17.0; extra == 'parse'
|
|
|
327
327
|
Requires-Dist: pdfminer-six>=20250506; extra == 'parse'
|
|
328
328
|
Requires-Dist: pdfplumber>=0.11.7; extra == 'parse'
|
|
329
329
|
Requires-Dist: pymupdf>=1.26.1; extra == 'parse'
|
|
330
|
-
Requires-Dist: pypdf2
|
|
330
|
+
Requires-Dist: pypdf2>=3.0.1; extra == 'parse'
|
|
331
331
|
Requires-Dist: unstructured[pdf]>=0.17.2; extra == 'parse'
|
|
332
332
|
Description-Content-Type: text/markdown
|
|
333
333
|
|
|
334
|
-
# AutoRAG
|
|
334
|
+
# AutoRAG (Legacy)
|
|
335
|
+
|
|
336
|
+
> [!NOTE]
|
|
337
|
+
> This is the original Python-based AutoRAG — the RAG AutoML tool for automatically finding an optimal RAG pipeline for your data. It now lives in the `legacy/` directory of the AutoRAG monorepo and is in **maintenance mode**: it continues to receive bug fixes, dependency updates, and PyPI releases (`pip install AutoRAG`), but new feature development happens in [AutoRAG 2.0](../README.md) at the repository root.
|
|
335
338
|
|
|
336
339
|
RAG AutoML tool for automatically finding an optimal RAG pipeline for your data.
|
|
337
340
|
|
|
@@ -1,4 +1,7 @@
|
|
|
1
|
-
# AutoRAG
|
|
1
|
+
# AutoRAG (Legacy)
|
|
2
|
+
|
|
3
|
+
> [!NOTE]
|
|
4
|
+
> This is the original Python-based AutoRAG — the RAG AutoML tool for automatically finding an optimal RAG pipeline for your data. It now lives in the `legacy/` directory of the AutoRAG monorepo and is in **maintenance mode**: it continues to receive bug fixes, dependency updates, and PyPI releases (`pip install AutoRAG`), but new feature development happens in [AutoRAG 2.0](../README.md) at the repository root.
|
|
2
5
|
|
|
3
6
|
RAG AutoML tool for automatically finding an optimal RAG pipeline for your data.
|
|
4
7
|
|
|
@@ -10,6 +10,11 @@ from llama_index.llms.openai import OpenAI
|
|
|
10
10
|
from llama_index.llms.openai_like import OpenAILike
|
|
11
11
|
from rich.logging import RichHandler
|
|
12
12
|
|
|
13
|
+
try:
|
|
14
|
+
from llama_index.llms.azure_openai import AzureOpenAI
|
|
15
|
+
except ImportError:
|
|
16
|
+
AzureOpenAI = None
|
|
17
|
+
|
|
13
18
|
|
|
14
19
|
class LazyInit:
|
|
15
20
|
def __init__(self, factory, *args, **kwargs):
|
|
@@ -58,6 +63,9 @@ generator_models = {
|
|
|
58
63
|
"bedrock": AutoRAGBedrock,
|
|
59
64
|
}
|
|
60
65
|
|
|
66
|
+
if AzureOpenAI is not None:
|
|
67
|
+
generator_models["azure_openai"] = AzureOpenAI
|
|
68
|
+
|
|
61
69
|
try:
|
|
62
70
|
from llama_index.llms.huggingface import HuggingFaceLLM
|
|
63
71
|
from llama_index.llms.ollama import Ollama
|
|
@@ -3,7 +3,7 @@ import tempfile
|
|
|
3
3
|
from glob import glob
|
|
4
4
|
from typing import List, Tuple, Dict
|
|
5
5
|
|
|
6
|
-
from PyPDF2 import
|
|
6
|
+
from PyPDF2 import PdfReader, PdfWriter
|
|
7
7
|
import pdfplumber
|
|
8
8
|
|
|
9
9
|
from autorag.support import get_support_modules
|
|
@@ -76,8 +76,8 @@ def save_page_by_table(data_path: str, text_dir: str, table_dir: str) -> Dict[st
|
|
|
76
76
|
file_name = os.path.basename(data_path).split(".pdf")[0]
|
|
77
77
|
|
|
78
78
|
with open(data_path, "rb") as input_data:
|
|
79
|
-
pdf_reader =
|
|
80
|
-
num_pages = pdf_reader.
|
|
79
|
+
pdf_reader = PdfReader(input_data)
|
|
80
|
+
num_pages = len(pdf_reader.pages)
|
|
81
81
|
|
|
82
82
|
path_map_dict = {}
|
|
83
83
|
for page_num in range(num_pages):
|
|
@@ -100,9 +100,9 @@ def _get_output_path(
|
|
|
100
100
|
return os.path.join(directory, f"{file_name}_page_{page_num + 1}.pdf")
|
|
101
101
|
|
|
102
102
|
|
|
103
|
-
def _save_single_page(pdf_reader:
|
|
104
|
-
pdf_writer =
|
|
105
|
-
pdf_writer.
|
|
103
|
+
def _save_single_page(pdf_reader: PdfReader, page_num: int, output_pdf_path: str):
|
|
104
|
+
pdf_writer = PdfWriter()
|
|
105
|
+
pdf_writer.add_page(pdf_reader.pages[page_num])
|
|
106
106
|
|
|
107
107
|
with open(output_pdf_path, "wb") as output_file:
|
|
108
108
|
pdf_writer.write(output_file)
|
|
@@ -30,12 +30,12 @@ async def query_evolve_openai_base(
|
|
|
30
30
|
user_prompt = f"Question: {original_query}\nContext: {context_str}\nOutput: "
|
|
31
31
|
messages.append(ChatMessage(role=MessageRole.USER, content=user_prompt))
|
|
32
32
|
|
|
33
|
-
completion = await client.
|
|
33
|
+
completion = await client.responses.parse(
|
|
34
34
|
model=model_name,
|
|
35
|
-
|
|
36
|
-
|
|
35
|
+
input=to_openai_message_dicts(messages),
|
|
36
|
+
text_format=Response,
|
|
37
37
|
)
|
|
38
|
-
row["query"] = completion.
|
|
38
|
+
row["query"] = completion.output_parsed.evolved_query
|
|
39
39
|
return row
|
|
40
40
|
|
|
41
41
|
|
|
@@ -72,10 +72,10 @@ async def compress_ragas(
|
|
|
72
72
|
user_prompt = f"Question: {original_query}\nOutput: "
|
|
73
73
|
messages.append(ChatMessage(role=MessageRole.USER, content=user_prompt))
|
|
74
74
|
|
|
75
|
-
completion = await client.
|
|
75
|
+
completion = await client.responses.parse(
|
|
76
76
|
model=model_name,
|
|
77
|
-
|
|
78
|
-
|
|
77
|
+
input=to_openai_message_dicts(messages),
|
|
78
|
+
text_format=Response,
|
|
79
79
|
)
|
|
80
|
-
row["query"] = completion.
|
|
80
|
+
row["query"] = completion.output_parsed.evolved_query
|
|
81
81
|
return row
|
|
@@ -77,14 +77,14 @@ async def dontknow_filter_openai(
|
|
|
77
77
|
system_prompt: List[ChatMessage] = FILTER_PROMPT["dontknow_filter"][lang]
|
|
78
78
|
result = []
|
|
79
79
|
for gen_gt in row["generation_gt"]:
|
|
80
|
-
completion = await client.
|
|
80
|
+
completion = await client.responses.parse(
|
|
81
81
|
model=model_name,
|
|
82
|
-
|
|
82
|
+
input=to_openai_message_dicts(
|
|
83
83
|
system_prompt + [ChatMessage(role=MessageRole.USER, content=gen_gt)]
|
|
84
84
|
),
|
|
85
|
-
|
|
85
|
+
text_format=Response,
|
|
86
86
|
)
|
|
87
|
-
result.append(completion.
|
|
87
|
+
result.append(completion.output_parsed.is_dont_know)
|
|
88
88
|
return not any(result)
|
|
89
89
|
|
|
90
90
|
|
|
@@ -37,9 +37,9 @@ async def passage_dependency_filter_openai(
|
|
|
37
37
|
assert "query" in row.keys(), "query column is not in the row."
|
|
38
38
|
system_prompt: List[ChatMessage] = FILTER_PROMPT["passage_dependency"][lang]
|
|
39
39
|
query = row["query"]
|
|
40
|
-
completion = await client.
|
|
40
|
+
completion = await client.responses.parse(
|
|
41
41
|
model=model_name,
|
|
42
|
-
|
|
42
|
+
input=to_openai_message_dicts(
|
|
43
43
|
system_prompt
|
|
44
44
|
+ [
|
|
45
45
|
ChatMessage(
|
|
@@ -48,9 +48,9 @@ async def passage_dependency_filter_openai(
|
|
|
48
48
|
)
|
|
49
49
|
]
|
|
50
50
|
),
|
|
51
|
-
|
|
51
|
+
text_format=Response,
|
|
52
52
|
)
|
|
53
|
-
return not completion.
|
|
53
|
+
return not completion.output_parsed.is_passage_dependent
|
|
54
54
|
|
|
55
55
|
|
|
56
56
|
async def passage_dependency_filter_llama_index(
|
|
@@ -25,16 +25,16 @@ async def make_gen_gt_openai(
|
|
|
25
25
|
passage_str = "\n".join(retrieval_gt_contents)
|
|
26
26
|
user_prompt = f"Text:\n<|text_start|>\n{passage_str}\n<|text_end|>\n\nQuestion:\n{query}\n\nAnswer:"
|
|
27
27
|
|
|
28
|
-
completion = await client.
|
|
28
|
+
completion = await client.responses.parse(
|
|
29
29
|
model=model_name,
|
|
30
|
-
|
|
30
|
+
input=[
|
|
31
31
|
{"role": "system", "content": system_prompt},
|
|
32
32
|
{"role": "user", "content": user_prompt},
|
|
33
33
|
],
|
|
34
34
|
temperature=0.0,
|
|
35
|
-
|
|
35
|
+
text_format=Response,
|
|
36
36
|
)
|
|
37
|
-
response: Response = completion.
|
|
37
|
+
response: Response = completion.output_parsed
|
|
38
38
|
return add_gen_gt(row, response.answer)
|
|
39
39
|
|
|
40
40
|
|
|
@@ -27,12 +27,12 @@ async def query_gen_openai_base(
|
|
|
27
27
|
user_prompt = f"{context_str}\n\nGenerated Question from the Text:\n"
|
|
28
28
|
messages.append(ChatMessage(role=MessageRole.USER, content=user_prompt))
|
|
29
29
|
|
|
30
|
-
completion = await client.
|
|
30
|
+
completion = await client.responses.parse(
|
|
31
31
|
model=model_name,
|
|
32
|
-
|
|
33
|
-
|
|
32
|
+
input=to_openai_message_dicts(messages),
|
|
33
|
+
text_format=Response,
|
|
34
34
|
)
|
|
35
|
-
row["query"] = completion.
|
|
35
|
+
row["query"] = completion.output_parsed.query
|
|
36
36
|
return row
|
|
37
37
|
|
|
38
38
|
|
|
@@ -86,10 +86,10 @@ async def two_hop_incremental(
|
|
|
86
86
|
user_prompt = f"{context_str}\n\nGenerated two-hop Question from two Documents:\n"
|
|
87
87
|
messages.append(ChatMessage(role=MessageRole.USER, content=user_prompt))
|
|
88
88
|
|
|
89
|
-
completion = await client.
|
|
89
|
+
completion = await client.responses.parse(
|
|
90
90
|
model=model_name,
|
|
91
|
-
|
|
92
|
-
|
|
91
|
+
input=to_openai_message_dicts(messages),
|
|
92
|
+
text_format=TwoHopIncrementalResponse,
|
|
93
93
|
)
|
|
94
|
-
row["query"] = completion.
|
|
94
|
+
row["query"] = completion.output_parsed.two_hop_question
|
|
95
95
|
return row
|
|
@@ -19,6 +19,7 @@ from autorag.evaluation.metric.deepeval_prompt import FaithfulnessTemplate
|
|
|
19
19
|
from autorag.evaluation.metric.util import (
|
|
20
20
|
autorag_metric_loop,
|
|
21
21
|
calculate_cosine_similarity,
|
|
22
|
+
remove_think_tags as _remove_think_tags,
|
|
22
23
|
)
|
|
23
24
|
from autorag.nodes.generator import OpenAILLM
|
|
24
25
|
from autorag.nodes.generator.base import BaseGenerator
|
|
@@ -36,7 +37,11 @@ from autorag.utils.util import (
|
|
|
36
37
|
|
|
37
38
|
@convert_inputs_to_list
|
|
38
39
|
def huggingface_evaluate(
|
|
39
|
-
instance,
|
|
40
|
+
instance,
|
|
41
|
+
key: str,
|
|
42
|
+
metric_inputs: List[MetricInput],
|
|
43
|
+
remove_think_tags: bool = True,
|
|
44
|
+
**kwargs,
|
|
40
45
|
) -> List[float]:
|
|
41
46
|
"""
|
|
42
47
|
Compute huggingface evaluate metric.
|
|
@@ -49,6 +54,8 @@ def huggingface_evaluate(
|
|
|
49
54
|
"""
|
|
50
55
|
|
|
51
56
|
def compute_score(gt: List[str], pred: str) -> float:
|
|
57
|
+
if remove_think_tags:
|
|
58
|
+
pred = _remove_think_tags(pred)
|
|
52
59
|
return max(
|
|
53
60
|
list(
|
|
54
61
|
map(
|
|
@@ -189,6 +196,7 @@ def bleu(
|
|
|
189
196
|
max_ngram_order: int = 4,
|
|
190
197
|
trg_lang: str = "",
|
|
191
198
|
effective_order: bool = True,
|
|
199
|
+
remove_think_tags: bool = True,
|
|
192
200
|
**kwargs,
|
|
193
201
|
) -> List[float]:
|
|
194
202
|
"""
|
|
@@ -213,14 +221,15 @@ def bleu(
|
|
|
213
221
|
**kwargs,
|
|
214
222
|
)
|
|
215
223
|
|
|
216
|
-
|
|
217
|
-
|
|
218
|
-
|
|
219
|
-
|
|
220
|
-
|
|
221
|
-
|
|
222
|
-
)
|
|
223
|
-
|
|
224
|
+
def compute_score(metric_input: MetricInput) -> float:
|
|
225
|
+
prediction = metric_input.generated_texts
|
|
226
|
+
if remove_think_tags:
|
|
227
|
+
prediction = _remove_think_tags(prediction)
|
|
228
|
+
return bleu_instance.sentence_score(
|
|
229
|
+
prediction, metric_input.generation_gt
|
|
230
|
+
).score
|
|
231
|
+
|
|
232
|
+
result = list(map(compute_score, metric_inputs))
|
|
224
233
|
return result
|
|
225
234
|
|
|
226
235
|
|
|
@@ -230,6 +239,7 @@ def meteor(
|
|
|
230
239
|
alpha: float = 0.9,
|
|
231
240
|
beta: float = 3.0,
|
|
232
241
|
gamma: float = 0.5,
|
|
242
|
+
remove_think_tags: bool = True,
|
|
233
243
|
) -> List[float]:
|
|
234
244
|
"""
|
|
235
245
|
Compute meteor score for generation.
|
|
@@ -250,6 +260,7 @@ def meteor(
|
|
|
250
260
|
meteor_instance,
|
|
251
261
|
"meteor",
|
|
252
262
|
metric_inputs,
|
|
263
|
+
remove_think_tags=remove_think_tags,
|
|
253
264
|
alpha=alpha,
|
|
254
265
|
beta=beta,
|
|
255
266
|
gamma=gamma,
|
|
@@ -265,6 +276,7 @@ def rouge(
|
|
|
265
276
|
use_stemmer: bool = False,
|
|
266
277
|
split_summaries: bool = False,
|
|
267
278
|
batch: int = os.cpu_count(),
|
|
279
|
+
remove_think_tags: bool = True,
|
|
268
280
|
) -> List[float]:
|
|
269
281
|
"""
|
|
270
282
|
Compute rouge score for generation.
|
|
@@ -295,6 +307,8 @@ def rouge(
|
|
|
295
307
|
)
|
|
296
308
|
|
|
297
309
|
async def compute(gt: List[str], pred: str) -> float:
|
|
310
|
+
if remove_think_tags:
|
|
311
|
+
pred = _remove_think_tags(pred)
|
|
298
312
|
return rouge_instance.score_multi(targets=gt, prediction=pred)[
|
|
299
313
|
rouge_type
|
|
300
314
|
].fmeasure
|
|
@@ -476,8 +490,11 @@ def bert_score(
|
|
|
476
490
|
lang: str = "en",
|
|
477
491
|
batch: int = 128,
|
|
478
492
|
n_threads: int = os.cpu_count(),
|
|
493
|
+
remove_think_tags: bool = True,
|
|
479
494
|
) -> List[float]:
|
|
480
495
|
generations = [metric_input.generated_texts for metric_input in metric_inputs]
|
|
496
|
+
if remove_think_tags:
|
|
497
|
+
generations = [_remove_think_tags(generation) for generation in generations]
|
|
481
498
|
generation_gt = [metric_input.generation_gt for metric_input in metric_inputs]
|
|
482
499
|
evaluator = evaluate.load("bertscore")
|
|
483
500
|
|
|
@@ -1,4 +1,5 @@
|
|
|
1
1
|
import functools
|
|
2
|
+
import re
|
|
2
3
|
from typing import List
|
|
3
4
|
|
|
4
5
|
import numpy as np
|
|
@@ -6,6 +7,15 @@ import numpy as np
|
|
|
6
7
|
from autorag.schema.metricinput import MetricInput
|
|
7
8
|
from autorag.utils.util import convert_inputs_to_list
|
|
8
9
|
|
|
10
|
+
_THINK_TAG_PATTERN = re.compile(
|
|
11
|
+
r"<think(?:ing)?>[\s\S]*?</think(?:ing)?>\s*", re.IGNORECASE
|
|
12
|
+
)
|
|
13
|
+
|
|
14
|
+
|
|
15
|
+
def remove_think_tags(text: str) -> str:
|
|
16
|
+
"""Remove complete ``<think>`` or ``<thinking>`` reasoning blocks."""
|
|
17
|
+
return _THINK_TAG_PATTERN.sub("", text)
|
|
18
|
+
|
|
9
19
|
|
|
10
20
|
def calculate_cosine_similarity(a, b):
|
|
11
21
|
dot_product = np.dot(a, b)
|
|
@@ -50,6 +50,10 @@ def run_node_line(
|
|
|
50
50
|
os.path.join(node_line_dir, node.node_type, "summary.csv")
|
|
51
51
|
)
|
|
52
52
|
best_node_row = node_summary_df.loc[node_summary_df["is_best"]]
|
|
53
|
+
if best_node_row.empty:
|
|
54
|
+
raise ValueError(
|
|
55
|
+
f"No best module found for node type {node.node_type} in summary file."
|
|
56
|
+
)
|
|
53
57
|
summary_lst.append(
|
|
54
58
|
{
|
|
55
59
|
"node_type": node.node_type,
|
|
@@ -22,9 +22,10 @@ THINK_OPEN_TAG = "<think>"
|
|
|
22
22
|
THINK_CLOSE_TAG = "</think>"
|
|
23
23
|
|
|
24
24
|
MAX_TOKEN_DICT = {
|
|
25
|
-
"MiniMax-
|
|
26
|
-
"MiniMax-M2.7
|
|
27
|
-
"MiniMax-M2.
|
|
25
|
+
"MiniMax-M3": 1_000_000,
|
|
26
|
+
"MiniMax-M2.7": 204_800,
|
|
27
|
+
"MiniMax-M2.7-highspeed": 204_800,
|
|
28
|
+
"MiniMax-M2.5": 204_800,
|
|
28
29
|
"MiniMax-M2.5-highspeed": 204_800,
|
|
29
30
|
}
|
|
30
31
|
|
|
@@ -38,7 +39,7 @@ class MiniMaxLLM(BaseGenerator):
|
|
|
38
39
|
It returns pseudo token ids and log probs since MiniMax does not support logprobs.
|
|
39
40
|
|
|
40
41
|
:param project_dir: The project directory.
|
|
41
|
-
:param llm: A model name for MiniMax. For example, ``MiniMax-
|
|
42
|
+
:param llm: A model name for MiniMax. For example, ``MiniMax-M3`` or ``MiniMax-M2.7``.
|
|
42
43
|
:param batch: Batch size for API calls. Default is 16.
|
|
43
44
|
:param api_key: MiniMax API key. You can also set this to env variable ``MINIMAX_API_KEY``.
|
|
44
45
|
:param kwargs: Extra parameters for the MiniMax chat completion API.
|
|
@@ -172,7 +173,7 @@ class MiniMaxLLM(BaseGenerator):
|
|
|
172
173
|
)
|
|
173
174
|
answer = response.choices[0].message.content
|
|
174
175
|
|
|
175
|
-
# Strip thinking tags if present (MiniMax M2.
|
|
176
|
+
# Strip thinking tags if present (MiniMax M2.7+ may include them)
|
|
176
177
|
if answer and THINK_OPEN_TAG in answer:
|
|
177
178
|
answer = strip_think_tags(answer)
|
|
178
179
|
|
|
@@ -1,5 +1,6 @@
|
|
|
1
1
|
import logging
|
|
2
|
-
|
|
2
|
+
import re
|
|
3
|
+
from typing import List, Optional, Tuple, Union, Dict
|
|
3
4
|
|
|
4
5
|
import pandas as pd
|
|
5
6
|
import tiktoken
|
|
@@ -16,16 +17,36 @@ from autorag.utils.util import (
|
|
|
16
17
|
|
|
17
18
|
logger = logging.getLogger("AutoRAG")
|
|
18
19
|
|
|
20
|
+
GPT_5_LONG_CONTEXT = (
|
|
21
|
+
1_050_000 # gpt-5.4 / gpt-5.5 / gpt-5.6 families share a 1.05M context window
|
|
22
|
+
)
|
|
23
|
+
|
|
24
|
+
# gpt-5 family prefixes that support the 1.05M context window.
|
|
25
|
+
# Used as a fallback for dated snapshots (e.g. gpt-5.6-sol-2026-07-09)
|
|
26
|
+
# that are not listed explicitly in MAX_TOKEN_DICT.
|
|
27
|
+
GPT_5_LONG_CONTEXT_PREFIXES = ("gpt-5.4", "gpt-5.5", "gpt-5.6")
|
|
28
|
+
|
|
29
|
+
# Only models still served by the OpenAI API are listed. Models retired per
|
|
30
|
+
# https://developers.openai.com/api/docs/deprecations (gpt-5, gpt-5.1,
|
|
31
|
+
# gpt-5-mini/nano, chatgpt-4o-latest, o1-preview, o1-mini, gpt-4-32k,
|
|
32
|
+
# gpt-4 vision/turbo previews, gpt-3.5-turbo-0613/16k snapshots, ...)
|
|
33
|
+
# are intentionally excluded so they fail fast with a clear error.
|
|
19
34
|
MAX_TOKEN_DICT = { # model name : token limit
|
|
20
|
-
|
|
21
|
-
"gpt-5.
|
|
22
|
-
"gpt-5":
|
|
23
|
-
"gpt-5-
|
|
24
|
-
"gpt-5-
|
|
25
|
-
"gpt-5-
|
|
26
|
-
|
|
27
|
-
"gpt-5
|
|
28
|
-
"gpt-5-
|
|
35
|
+
# gpt-5.6 family (July 2026) - 1.05M context window
|
|
36
|
+
"gpt-5.6": GPT_5_LONG_CONTEXT, # alias routing to gpt-5.6-sol
|
|
37
|
+
"gpt-5.6-sol": GPT_5_LONG_CONTEXT,
|
|
38
|
+
"gpt-5.6-terra": GPT_5_LONG_CONTEXT,
|
|
39
|
+
"gpt-5.6-luna": GPT_5_LONG_CONTEXT,
|
|
40
|
+
"gpt-5.6-pro": GPT_5_LONG_CONTEXT,
|
|
41
|
+
# gpt-5.5 (April 2026)
|
|
42
|
+
"gpt-5.5": GPT_5_LONG_CONTEXT,
|
|
43
|
+
"gpt-5.5-pro": GPT_5_LONG_CONTEXT,
|
|
44
|
+
"gpt-5.5-chat-latest": GPT_5_LONG_CONTEXT,
|
|
45
|
+
# gpt-5.4
|
|
46
|
+
"gpt-5.4": GPT_5_LONG_CONTEXT,
|
|
47
|
+
"gpt-5.4-pro": GPT_5_LONG_CONTEXT,
|
|
48
|
+
"gpt-5.4-mini": GPT_5_LONG_CONTEXT,
|
|
49
|
+
"gpt-5.4-nano": GPT_5_LONG_CONTEXT,
|
|
29
50
|
"gpt-4.1": 1_000_000,
|
|
30
51
|
"gpt-4.1-2025-04-14": 1_000_000,
|
|
31
52
|
"gpt-4.1-mini": 1_047_576,
|
|
@@ -33,10 +54,6 @@ MAX_TOKEN_DICT = { # model name : token limit
|
|
|
33
54
|
"gpt-4.1-nano": 1_000_000,
|
|
34
55
|
"gpt-4.1-nano-2025-04-14": 1_000_000,
|
|
35
56
|
"o1": 200_000,
|
|
36
|
-
"o1-preview": 128_000,
|
|
37
|
-
"o1-preview-2024-09-12": 128_000,
|
|
38
|
-
"o1-mini": 128_000,
|
|
39
|
-
"o1-mini-2024-09-12": 128_000,
|
|
40
57
|
"o1-pro": 200_000,
|
|
41
58
|
"o1-pro-2025-03-19": 200_000,
|
|
42
59
|
"o3": 200_000,
|
|
@@ -49,28 +66,35 @@ MAX_TOKEN_DICT = { # model name : token limit
|
|
|
49
66
|
"gpt-4o": 128_000,
|
|
50
67
|
"gpt-4o-2024-08-06": 128_000,
|
|
51
68
|
"gpt-4o-2024-05-13": 128_000,
|
|
52
|
-
"chatgpt-4o-latest": 128_000,
|
|
53
69
|
"gpt-4-turbo": 128_000,
|
|
54
70
|
"gpt-4-turbo-2024-04-09": 128_000,
|
|
55
|
-
"gpt-4-turbo-preview": 128_000,
|
|
56
|
-
"gpt-4-0125-preview": 128_000,
|
|
57
|
-
"gpt-4-1106-preview": 128_000,
|
|
58
|
-
"gpt-4-vision-preview": 128_000,
|
|
59
|
-
"gpt-4-1106-vision-preview": 128_000,
|
|
60
71
|
"gpt-4": 8_192,
|
|
61
72
|
"gpt-4-0613": 8_192,
|
|
62
|
-
"gpt-4-32k": 32_768,
|
|
63
|
-
"gpt-4-32k-0613": 32_768,
|
|
64
73
|
"gpt-3.5-turbo-0125": 16_385,
|
|
65
74
|
"gpt-3.5-turbo": 16_385,
|
|
66
75
|
"gpt-3.5-turbo-1106": 16_385,
|
|
67
76
|
"gpt-3.5-turbo-instruct": 4_096,
|
|
68
|
-
"gpt-3.5-turbo-16k": 16_385,
|
|
69
|
-
"gpt-3.5-turbo-0613": 4_096,
|
|
70
|
-
"gpt-3.5-turbo-16k-0613": 16_385,
|
|
71
77
|
}
|
|
72
78
|
|
|
73
79
|
|
|
80
|
+
def get_max_token_size(llm: str) -> Optional[int]:
|
|
81
|
+
"""Resolve the context window size for an OpenAI model name.
|
|
82
|
+
|
|
83
|
+
Exact matches in MAX_TOKEN_DICT win first. Dated snapshots
|
|
84
|
+
(e.g. ``gpt-5.6-sol-2026-07-09``) fall back to their base model entry,
|
|
85
|
+
and unknown gpt-5 family variants fall back to the family context window
|
|
86
|
+
so that newly released snapshots keep working without a code change.
|
|
87
|
+
"""
|
|
88
|
+
if llm in MAX_TOKEN_DICT:
|
|
89
|
+
return MAX_TOKEN_DICT[llm]
|
|
90
|
+
base = re.sub(r"-\d{4}-\d{2}-\d{2}$", "", llm)
|
|
91
|
+
if base in MAX_TOKEN_DICT:
|
|
92
|
+
return MAX_TOKEN_DICT[base]
|
|
93
|
+
if base.startswith(GPT_5_LONG_CONTEXT_PREFIXES):
|
|
94
|
+
return GPT_5_LONG_CONTEXT
|
|
95
|
+
return None
|
|
96
|
+
|
|
97
|
+
|
|
74
98
|
class OpenAILLM(BaseGenerator):
|
|
75
99
|
def __init__(self, project_dir, llm: str, batch: int = 16, *args, **kwargs):
|
|
76
100
|
super().__init__(project_dir, llm, *args, **kwargs)
|
|
@@ -84,14 +108,13 @@ class OpenAILLM(BaseGenerator):
|
|
|
84
108
|
except KeyError:
|
|
85
109
|
self.tokenizer = tiktoken.get_encoding("o200k_base")
|
|
86
110
|
|
|
87
|
-
|
|
88
|
-
|
|
89
|
-
) # because of chat token usage
|
|
90
|
-
if self.max_token_size is None:
|
|
111
|
+
max_tokens = get_max_token_size(self.llm)
|
|
112
|
+
if max_tokens is None:
|
|
91
113
|
raise ValueError(
|
|
92
114
|
f"Model {self.llm} does not supported. "
|
|
93
115
|
f"Please select the model between {list(MAX_TOKEN_DICT.keys())}"
|
|
94
116
|
)
|
|
117
|
+
self.max_token_size = max_tokens - 7 # because of chat token usage
|
|
95
118
|
|
|
96
119
|
@result_to_dataframe(["generated_texts", "generated_tokens", "generated_log_probs"])
|
|
97
120
|
def pure(self, previous_result: pd.DataFrame, *args, **kwargs):
|
|
@@ -309,12 +332,6 @@ class OpenAILLM(BaseGenerator):
|
|
|
309
332
|
async def get_result_gpt_5(self, prompt: Union[str, List[dict]], **kwargs):
|
|
310
333
|
if not self.llm.startswith("gpt-5"):
|
|
311
334
|
raise ValueError("get_result_gpt_5 is only for gpt-5 models.")
|
|
312
|
-
api_key = getattr(self.client, "api_key", None)
|
|
313
|
-
if isinstance(api_key, str) and api_key.startswith("mock_"):
|
|
314
|
-
answer = "Why not"
|
|
315
|
-
tokens = self.tokenizer.encode(answer, allowed_special="all")
|
|
316
|
-
pseudo_log_probs = [0.5] * len(tokens)
|
|
317
|
-
return answer, tokens, pseudo_log_probs
|
|
318
335
|
messages = parse_prompt(prompt)
|
|
319
336
|
instruction = "\n\n".join(
|
|
320
337
|
[msg["content"] for msg in messages if msg["role"] == "system"]
|
|
@@ -131,6 +131,8 @@ def get_support_modules(module_name: str) -> Callable:
|
|
|
131
131
|
),
|
|
132
132
|
"flashrank_reranker": ("autorag.nodes.passagereranker", "FlashRankReranker"),
|
|
133
133
|
"FlashRankReranker": ("autorag.nodes.passagereranker", "FlashRankReranker"),
|
|
134
|
+
"nvidia_reranker": ("autorag.nodes.passagereranker", "NvidiaReranker"),
|
|
135
|
+
"NvidiaReranker": ("autorag.nodes.passagereranker", "NvidiaReranker"),
|
|
134
136
|
# passage_filter
|
|
135
137
|
"pass_passage_filter": ("autorag.nodes.passagefilter", "PassPassageFilter"),
|
|
136
138
|
"similarity_threshold_cutoff": (
|
|
@@ -302,7 +302,23 @@ async def process_batch(tasks, batch_size: int = 64) -> List[Any]:
|
|
|
302
302
|
|
|
303
303
|
for i in range(0, len(tasks), batch_size):
|
|
304
304
|
batch = tasks[i : i + batch_size]
|
|
305
|
-
batch_results = await asyncio.gather(*batch)
|
|
305
|
+
batch_results = await asyncio.gather(*batch, return_exceptions=True)
|
|
306
|
+
|
|
307
|
+
first_exception = None
|
|
308
|
+
for offset, item in enumerate(batch_results):
|
|
309
|
+
if isinstance(item, BaseException):
|
|
310
|
+
logger.error(
|
|
311
|
+
"process_batch task %d (batch offset %d) raised %s: %s",
|
|
312
|
+
i + offset,
|
|
313
|
+
offset,
|
|
314
|
+
type(item).__name__,
|
|
315
|
+
item,
|
|
316
|
+
)
|
|
317
|
+
if first_exception is None:
|
|
318
|
+
first_exception = item
|
|
319
|
+
if first_exception is not None:
|
|
320
|
+
raise first_exception
|
|
321
|
+
|
|
306
322
|
results.extend(batch_results)
|
|
307
323
|
|
|
308
324
|
return results
|