edsl 0.1.39.dev3__py3-none-any.whl → 0.1.39.dev4__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (344) hide show
  1. edsl/Base.py +413 -332
  2. edsl/BaseDiff.py +260 -260
  3. edsl/TemplateLoader.py +24 -24
  4. edsl/__init__.py +57 -49
  5. edsl/__version__.py +1 -1
  6. edsl/agents/Agent.py +1071 -867
  7. edsl/agents/AgentList.py +551 -413
  8. edsl/agents/Invigilator.py +284 -233
  9. edsl/agents/InvigilatorBase.py +257 -270
  10. edsl/agents/PromptConstructor.py +272 -354
  11. edsl/agents/QuestionInstructionPromptBuilder.py +128 -0
  12. edsl/agents/QuestionTemplateReplacementsBuilder.py +137 -0
  13. edsl/agents/__init__.py +2 -3
  14. edsl/agents/descriptors.py +99 -99
  15. edsl/agents/prompt_helpers.py +129 -129
  16. edsl/agents/question_option_processor.py +172 -0
  17. edsl/auto/AutoStudy.py +130 -117
  18. edsl/auto/StageBase.py +243 -230
  19. edsl/auto/StageGenerateSurvey.py +178 -178
  20. edsl/auto/StageLabelQuestions.py +125 -125
  21. edsl/auto/StagePersona.py +61 -61
  22. edsl/auto/StagePersonaDimensionValueRanges.py +88 -88
  23. edsl/auto/StagePersonaDimensionValues.py +74 -74
  24. edsl/auto/StagePersonaDimensions.py +69 -69
  25. edsl/auto/StageQuestions.py +74 -73
  26. edsl/auto/SurveyCreatorPipeline.py +21 -21
  27. edsl/auto/utilities.py +218 -224
  28. edsl/base/Base.py +279 -279
  29. edsl/config.py +177 -157
  30. edsl/conversation/Conversation.py +290 -290
  31. edsl/conversation/car_buying.py +59 -58
  32. edsl/conversation/chips.py +95 -95
  33. edsl/conversation/mug_negotiation.py +81 -81
  34. edsl/conversation/next_speaker_utilities.py +93 -93
  35. edsl/coop/CoopFunctionsMixin.py +15 -0
  36. edsl/coop/ExpectedParrotKeyHandler.py +125 -0
  37. edsl/coop/PriceFetcher.py +54 -54
  38. edsl/coop/__init__.py +2 -2
  39. edsl/coop/coop.py +1106 -1028
  40. edsl/coop/utils.py +131 -131
  41. edsl/data/Cache.py +573 -555
  42. edsl/data/CacheEntry.py +230 -233
  43. edsl/data/CacheHandler.py +168 -149
  44. edsl/data/RemoteCacheSync.py +186 -78
  45. edsl/data/SQLiteDict.py +292 -292
  46. edsl/data/__init__.py +5 -4
  47. edsl/data/hack.py +10 -0
  48. edsl/data/orm.py +10 -10
  49. edsl/data_transfer_models.py +74 -73
  50. edsl/enums.py +202 -175
  51. edsl/exceptions/BaseException.py +21 -21
  52. edsl/exceptions/__init__.py +54 -54
  53. edsl/exceptions/agents.py +54 -42
  54. edsl/exceptions/cache.py +5 -5
  55. edsl/exceptions/configuration.py +16 -16
  56. edsl/exceptions/coop.py +10 -10
  57. edsl/exceptions/data.py +14 -14
  58. edsl/exceptions/general.py +34 -34
  59. edsl/exceptions/inference_services.py +5 -0
  60. edsl/exceptions/jobs.py +33 -33
  61. edsl/exceptions/language_models.py +63 -63
  62. edsl/exceptions/prompts.py +15 -15
  63. edsl/exceptions/questions.py +109 -91
  64. edsl/exceptions/results.py +29 -29
  65. edsl/exceptions/scenarios.py +29 -22
  66. edsl/exceptions/surveys.py +37 -37
  67. edsl/inference_services/AnthropicService.py +106 -87
  68. edsl/inference_services/AvailableModelCacheHandler.py +184 -0
  69. edsl/inference_services/AvailableModelFetcher.py +215 -0
  70. edsl/inference_services/AwsBedrock.py +118 -120
  71. edsl/inference_services/AzureAI.py +215 -217
  72. edsl/inference_services/DeepInfraService.py +18 -18
  73. edsl/inference_services/GoogleService.py +143 -148
  74. edsl/inference_services/GroqService.py +20 -20
  75. edsl/inference_services/InferenceServiceABC.py +80 -147
  76. edsl/inference_services/InferenceServicesCollection.py +138 -97
  77. edsl/inference_services/MistralAIService.py +120 -123
  78. edsl/inference_services/OllamaService.py +18 -18
  79. edsl/inference_services/OpenAIService.py +236 -224
  80. edsl/inference_services/PerplexityService.py +160 -163
  81. edsl/inference_services/ServiceAvailability.py +135 -0
  82. edsl/inference_services/TestService.py +90 -89
  83. edsl/inference_services/TogetherAIService.py +172 -170
  84. edsl/inference_services/data_structures.py +134 -0
  85. edsl/inference_services/models_available_cache.py +118 -118
  86. edsl/inference_services/rate_limits_cache.py +25 -25
  87. edsl/inference_services/registry.py +41 -41
  88. edsl/inference_services/write_available.py +10 -10
  89. edsl/jobs/AnswerQuestionFunctionConstructor.py +223 -0
  90. edsl/jobs/Answers.py +43 -56
  91. edsl/jobs/FetchInvigilator.py +47 -0
  92. edsl/jobs/InterviewTaskManager.py +98 -0
  93. edsl/jobs/InterviewsConstructor.py +50 -0
  94. edsl/jobs/Jobs.py +823 -898
  95. edsl/jobs/JobsChecks.py +172 -147
  96. edsl/jobs/JobsComponentConstructor.py +189 -0
  97. edsl/jobs/JobsPrompts.py +270 -268
  98. edsl/jobs/JobsRemoteInferenceHandler.py +311 -239
  99. edsl/jobs/JobsRemoteInferenceLogger.py +239 -0
  100. edsl/jobs/RequestTokenEstimator.py +30 -0
  101. edsl/jobs/__init__.py +1 -1
  102. edsl/jobs/async_interview_runner.py +138 -0
  103. edsl/jobs/buckets/BucketCollection.py +104 -63
  104. edsl/jobs/buckets/ModelBuckets.py +65 -65
  105. edsl/jobs/buckets/TokenBucket.py +283 -251
  106. edsl/jobs/buckets/TokenBucketAPI.py +211 -0
  107. edsl/jobs/buckets/TokenBucketClient.py +191 -0
  108. edsl/jobs/check_survey_scenario_compatibility.py +85 -0
  109. edsl/jobs/data_structures.py +120 -0
  110. edsl/jobs/decorators.py +35 -0
  111. edsl/jobs/interviews/Interview.py +396 -661
  112. edsl/jobs/interviews/InterviewExceptionCollection.py +99 -99
  113. edsl/jobs/interviews/InterviewExceptionEntry.py +186 -186
  114. edsl/jobs/interviews/InterviewStatistic.py +63 -63
  115. edsl/jobs/interviews/InterviewStatisticsCollection.py +25 -25
  116. edsl/jobs/interviews/InterviewStatusDictionary.py +78 -78
  117. edsl/jobs/interviews/InterviewStatusLog.py +92 -92
  118. edsl/jobs/interviews/ReportErrors.py +66 -66
  119. edsl/jobs/interviews/interview_status_enum.py +9 -9
  120. edsl/jobs/jobs_status_enums.py +9 -0
  121. edsl/jobs/loggers/HTMLTableJobLogger.py +304 -0
  122. edsl/jobs/results_exceptions_handler.py +98 -0
  123. edsl/jobs/runners/JobsRunnerAsyncio.py +151 -466
  124. edsl/jobs/runners/JobsRunnerStatus.py +297 -330
  125. edsl/jobs/tasks/QuestionTaskCreator.py +244 -242
  126. edsl/jobs/tasks/TaskCreators.py +64 -64
  127. edsl/jobs/tasks/TaskHistory.py +470 -450
  128. edsl/jobs/tasks/TaskStatusLog.py +23 -23
  129. edsl/jobs/tasks/task_status_enum.py +161 -163
  130. edsl/jobs/tokens/InterviewTokenUsage.py +27 -27
  131. edsl/jobs/tokens/TokenUsage.py +34 -34
  132. edsl/language_models/ComputeCost.py +63 -0
  133. edsl/language_models/LanguageModel.py +626 -668
  134. edsl/language_models/ModelList.py +164 -155
  135. edsl/language_models/PriceManager.py +127 -0
  136. edsl/language_models/RawResponseHandler.py +106 -0
  137. edsl/language_models/RegisterLanguageModelsMeta.py +184 -184
  138. edsl/language_models/ServiceDataSources.py +0 -0
  139. edsl/language_models/__init__.py +2 -3
  140. edsl/language_models/fake_openai_call.py +15 -15
  141. edsl/language_models/fake_openai_service.py +61 -61
  142. edsl/language_models/key_management/KeyLookup.py +63 -0
  143. edsl/language_models/key_management/KeyLookupBuilder.py +273 -0
  144. edsl/language_models/key_management/KeyLookupCollection.py +38 -0
  145. edsl/language_models/key_management/__init__.py +0 -0
  146. edsl/language_models/key_management/models.py +131 -0
  147. edsl/language_models/model.py +256 -0
  148. edsl/language_models/repair.py +156 -156
  149. edsl/language_models/utilities.py +65 -64
  150. edsl/notebooks/Notebook.py +263 -258
  151. edsl/notebooks/NotebookToLaTeX.py +142 -0
  152. edsl/notebooks/__init__.py +1 -1
  153. edsl/prompts/Prompt.py +352 -362
  154. edsl/prompts/__init__.py +2 -2
  155. edsl/questions/ExceptionExplainer.py +77 -0
  156. edsl/questions/HTMLQuestion.py +103 -0
  157. edsl/questions/QuestionBase.py +518 -664
  158. edsl/questions/QuestionBasePromptsMixin.py +221 -217
  159. edsl/questions/QuestionBudget.py +227 -227
  160. edsl/questions/QuestionCheckBox.py +359 -359
  161. edsl/questions/QuestionExtract.py +180 -182
  162. edsl/questions/QuestionFreeText.py +113 -114
  163. edsl/questions/QuestionFunctional.py +166 -166
  164. edsl/questions/QuestionList.py +223 -231
  165. edsl/questions/QuestionMatrix.py +265 -0
  166. edsl/questions/QuestionMultipleChoice.py +330 -286
  167. edsl/questions/QuestionNumerical.py +151 -153
  168. edsl/questions/QuestionRank.py +314 -324
  169. edsl/questions/Quick.py +41 -41
  170. edsl/questions/SimpleAskMixin.py +74 -73
  171. edsl/questions/__init__.py +27 -26
  172. edsl/questions/{AnswerValidatorMixin.py → answer_validator_mixin.py} +334 -289
  173. edsl/questions/compose_questions.py +98 -98
  174. edsl/questions/data_structures.py +20 -0
  175. edsl/questions/decorators.py +21 -21
  176. edsl/questions/derived/QuestionLikertFive.py +76 -76
  177. edsl/questions/derived/QuestionLinearScale.py +90 -87
  178. edsl/questions/derived/QuestionTopK.py +93 -93
  179. edsl/questions/derived/QuestionYesNo.py +82 -82
  180. edsl/questions/descriptors.py +427 -413
  181. edsl/questions/loop_processor.py +149 -0
  182. edsl/questions/prompt_templates/question_budget.jinja +13 -13
  183. edsl/questions/prompt_templates/question_checkbox.jinja +32 -32
  184. edsl/questions/prompt_templates/question_extract.jinja +11 -11
  185. edsl/questions/prompt_templates/question_free_text.jinja +3 -3
  186. edsl/questions/prompt_templates/question_linear_scale.jinja +11 -11
  187. edsl/questions/prompt_templates/question_list.jinja +17 -17
  188. edsl/questions/prompt_templates/question_multiple_choice.jinja +33 -33
  189. edsl/questions/prompt_templates/question_numerical.jinja +36 -36
  190. edsl/questions/{QuestionBaseGenMixin.py → question_base_gen_mixin.py} +168 -161
  191. edsl/questions/question_registry.py +177 -177
  192. edsl/questions/{RegisterQuestionsMeta.py → register_questions_meta.py} +71 -71
  193. edsl/questions/{ResponseValidatorABC.py → response_validator_abc.py} +188 -174
  194. edsl/questions/response_validator_factory.py +34 -0
  195. edsl/questions/settings.py +12 -12
  196. edsl/questions/templates/budget/answering_instructions.jinja +7 -7
  197. edsl/questions/templates/budget/question_presentation.jinja +7 -7
  198. edsl/questions/templates/checkbox/answering_instructions.jinja +10 -10
  199. edsl/questions/templates/checkbox/question_presentation.jinja +22 -22
  200. edsl/questions/templates/extract/answering_instructions.jinja +7 -7
  201. edsl/questions/templates/likert_five/answering_instructions.jinja +10 -10
  202. edsl/questions/templates/likert_five/question_presentation.jinja +11 -11
  203. edsl/questions/templates/linear_scale/answering_instructions.jinja +5 -5
  204. edsl/questions/templates/linear_scale/question_presentation.jinja +5 -5
  205. edsl/questions/templates/list/answering_instructions.jinja +3 -3
  206. edsl/questions/templates/list/question_presentation.jinja +5 -5
  207. edsl/questions/templates/matrix/__init__.py +1 -0
  208. edsl/questions/templates/matrix/answering_instructions.jinja +5 -0
  209. edsl/questions/templates/matrix/question_presentation.jinja +20 -0
  210. edsl/questions/templates/multiple_choice/answering_instructions.jinja +9 -9
  211. edsl/questions/templates/multiple_choice/question_presentation.jinja +11 -11
  212. edsl/questions/templates/numerical/answering_instructions.jinja +6 -6
  213. edsl/questions/templates/numerical/question_presentation.jinja +6 -6
  214. edsl/questions/templates/rank/answering_instructions.jinja +11 -11
  215. edsl/questions/templates/rank/question_presentation.jinja +15 -15
  216. edsl/questions/templates/top_k/answering_instructions.jinja +8 -8
  217. edsl/questions/templates/top_k/question_presentation.jinja +22 -22
  218. edsl/questions/templates/yes_no/answering_instructions.jinja +6 -6
  219. edsl/questions/templates/yes_no/question_presentation.jinja +11 -11
  220. edsl/results/CSSParameterizer.py +108 -108
  221. edsl/results/Dataset.py +587 -424
  222. edsl/results/DatasetExportMixin.py +594 -731
  223. edsl/results/DatasetTree.py +295 -275
  224. edsl/results/MarkdownToDocx.py +122 -0
  225. edsl/results/MarkdownToPDF.py +111 -0
  226. edsl/results/Result.py +557 -465
  227. edsl/results/Results.py +1183 -1165
  228. edsl/results/ResultsExportMixin.py +45 -43
  229. edsl/results/ResultsGGMixin.py +121 -121
  230. edsl/results/TableDisplay.py +125 -198
  231. edsl/results/TextEditor.py +50 -0
  232. edsl/results/__init__.py +2 -2
  233. edsl/results/file_exports.py +252 -0
  234. edsl/results/{ResultsFetchMixin.py → results_fetch_mixin.py} +33 -33
  235. edsl/results/{Selector.py → results_selector.py} +145 -135
  236. edsl/results/{ResultsToolsMixin.py → results_tools_mixin.py} +98 -98
  237. edsl/results/smart_objects.py +96 -0
  238. edsl/results/table_data_class.py +12 -0
  239. edsl/results/table_display.css +77 -77
  240. edsl/results/table_renderers.py +118 -0
  241. edsl/results/tree_explore.py +115 -115
  242. edsl/scenarios/ConstructDownloadLink.py +109 -0
  243. edsl/scenarios/DocumentChunker.py +102 -0
  244. edsl/scenarios/DocxScenario.py +16 -0
  245. edsl/scenarios/FileStore.py +511 -632
  246. edsl/scenarios/PdfExtractor.py +40 -0
  247. edsl/scenarios/Scenario.py +498 -601
  248. edsl/scenarios/ScenarioHtmlMixin.py +65 -64
  249. edsl/scenarios/ScenarioList.py +1458 -1287
  250. edsl/scenarios/ScenarioListExportMixin.py +45 -52
  251. edsl/scenarios/ScenarioListPdfMixin.py +239 -261
  252. edsl/scenarios/__init__.py +3 -4
  253. edsl/scenarios/directory_scanner.py +96 -0
  254. edsl/scenarios/file_methods.py +85 -0
  255. edsl/scenarios/handlers/__init__.py +13 -0
  256. edsl/scenarios/handlers/csv.py +38 -0
  257. edsl/scenarios/handlers/docx.py +76 -0
  258. edsl/scenarios/handlers/html.py +37 -0
  259. edsl/scenarios/handlers/json.py +111 -0
  260. edsl/scenarios/handlers/latex.py +5 -0
  261. edsl/scenarios/handlers/md.py +51 -0
  262. edsl/scenarios/handlers/pdf.py +68 -0
  263. edsl/scenarios/handlers/png.py +39 -0
  264. edsl/scenarios/handlers/pptx.py +105 -0
  265. edsl/scenarios/handlers/py.py +294 -0
  266. edsl/scenarios/handlers/sql.py +313 -0
  267. edsl/scenarios/handlers/sqlite.py +149 -0
  268. edsl/scenarios/handlers/txt.py +33 -0
  269. edsl/scenarios/{ScenarioJoin.py → scenario_join.py} +131 -127
  270. edsl/scenarios/scenario_selector.py +156 -0
  271. edsl/shared.py +1 -1
  272. edsl/study/ObjectEntry.py +173 -173
  273. edsl/study/ProofOfWork.py +113 -113
  274. edsl/study/SnapShot.py +80 -80
  275. edsl/study/Study.py +521 -528
  276. edsl/study/__init__.py +4 -4
  277. edsl/surveys/ConstructDAG.py +92 -0
  278. edsl/surveys/DAG.py +148 -148
  279. edsl/surveys/EditSurvey.py +221 -0
  280. edsl/surveys/InstructionHandler.py +100 -0
  281. edsl/surveys/Memory.py +31 -31
  282. edsl/surveys/MemoryManagement.py +72 -0
  283. edsl/surveys/MemoryPlan.py +244 -244
  284. edsl/surveys/Rule.py +327 -326
  285. edsl/surveys/RuleCollection.py +385 -387
  286. edsl/surveys/RuleManager.py +172 -0
  287. edsl/surveys/Simulator.py +75 -0
  288. edsl/surveys/Survey.py +1280 -1801
  289. edsl/surveys/SurveyCSS.py +273 -261
  290. edsl/surveys/SurveyExportMixin.py +259 -259
  291. edsl/surveys/{SurveyFlowVisualizationMixin.py → SurveyFlowVisualization.py} +181 -179
  292. edsl/surveys/SurveyQualtricsImport.py +284 -284
  293. edsl/surveys/SurveyToApp.py +141 -0
  294. edsl/surveys/__init__.py +5 -3
  295. edsl/surveys/base.py +53 -53
  296. edsl/surveys/descriptors.py +60 -56
  297. edsl/surveys/instructions/ChangeInstruction.py +48 -49
  298. edsl/surveys/instructions/Instruction.py +56 -65
  299. edsl/surveys/instructions/InstructionCollection.py +82 -77
  300. edsl/templates/error_reporting/base.html +23 -23
  301. edsl/templates/error_reporting/exceptions_by_model.html +34 -34
  302. edsl/templates/error_reporting/exceptions_by_question_name.html +16 -16
  303. edsl/templates/error_reporting/exceptions_by_type.html +16 -16
  304. edsl/templates/error_reporting/interview_details.html +115 -115
  305. edsl/templates/error_reporting/interviews.html +19 -19
  306. edsl/templates/error_reporting/overview.html +4 -4
  307. edsl/templates/error_reporting/performance_plot.html +1 -1
  308. edsl/templates/error_reporting/report.css +73 -73
  309. edsl/templates/error_reporting/report.html +117 -117
  310. edsl/templates/error_reporting/report.js +25 -25
  311. edsl/test_h +1 -0
  312. edsl/tools/__init__.py +1 -1
  313. edsl/tools/clusters.py +192 -192
  314. edsl/tools/embeddings.py +27 -27
  315. edsl/tools/embeddings_plotting.py +118 -118
  316. edsl/tools/plotting.py +112 -112
  317. edsl/tools/summarize.py +18 -18
  318. edsl/utilities/PrettyList.py +56 -0
  319. edsl/utilities/SystemInfo.py +28 -28
  320. edsl/utilities/__init__.py +22 -22
  321. edsl/utilities/ast_utilities.py +25 -25
  322. edsl/utilities/data/Registry.py +6 -6
  323. edsl/utilities/data/__init__.py +1 -1
  324. edsl/utilities/data/scooter_results.json +1 -1
  325. edsl/utilities/decorators.py +77 -77
  326. edsl/utilities/gcp_bucket/cloud_storage.py +96 -96
  327. edsl/utilities/gcp_bucket/example.py +50 -0
  328. edsl/utilities/interface.py +627 -627
  329. edsl/utilities/is_notebook.py +18 -0
  330. edsl/utilities/is_valid_variable_name.py +11 -0
  331. edsl/utilities/naming_utilities.py +263 -263
  332. edsl/utilities/remove_edsl_version.py +24 -0
  333. edsl/utilities/repair_functions.py +28 -28
  334. edsl/utilities/restricted_python.py +70 -70
  335. edsl/utilities/utilities.py +436 -424
  336. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev4.dist-info}/LICENSE +21 -21
  337. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev4.dist-info}/METADATA +13 -11
  338. edsl-0.1.39.dev4.dist-info/RECORD +361 -0
  339. edsl/language_models/KeyLookup.py +0 -30
  340. edsl/language_models/registry.py +0 -190
  341. edsl/language_models/unused/ReplicateBase.py +0 -83
  342. edsl/results/ResultsDBMixin.py +0 -238
  343. edsl-0.1.39.dev3.dist-info/RECORD +0 -277
  344. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev4.dist-info}/WHEEL +0 -0
edsl/config.py CHANGED
@@ -1,157 +1,177 @@
1
- """This module provides a Config class that loads environment variables from a .env file and sets them as class attributes."""
2
-
3
- import os
4
- from dotenv import load_dotenv, find_dotenv
5
- from edsl.exceptions import (
6
- InvalidEnvironmentVariableError,
7
- MissingEnvironmentVariableError,
8
- )
9
-
10
- # valid values for EDSL_RUN_MODE
11
- EDSL_RUN_MODES = [
12
- "development",
13
- "development-testrun",
14
- "production",
15
- ]
16
-
17
- # `default` is used to impute values only in "production" mode
18
- # `info` gives a brief description of the env var
19
- CONFIG_MAP = {
20
- "EDSL_RUN_MODE": {
21
- "default": "production",
22
- "info": "This config var determines the run mode of the application.",
23
- },
24
- "EDSL_API_TIMEOUT": {
25
- "default": "60",
26
- "info": "This config var determines the maximum number of seconds to wait for an API call to return.",
27
- },
28
- "EDSL_BACKOFF_START_SEC": {
29
- "default": "1",
30
- "info": "This config var determines the number of seconds to wait before retrying a failed API call.",
31
- },
32
- "EDSL_BACKOFF_MAX_SEC": {
33
- "default": "60",
34
- "info": "This config var determines the maximum number of seconds to wait before retrying a failed API call.",
35
- },
36
- "EDSL_DATABASE_PATH": {
37
- "default": f"sqlite:///{os.path.join(os.getcwd(), '.edsl_cache/data.db')}",
38
- "info": "This config var determines the path to the cache file.",
39
- },
40
- "EDSL_DEFAULT_MODEL": {
41
- "default": "gpt-4o",
42
- "info": "This config var holds the default model that will be used if a model is not explicitly passed.",
43
- },
44
- "EDSL_FETCH_TOKEN_PRICES": {
45
- "default": "True",
46
- "info": "This config var determines whether to fetch prices for tokens used in remote inference",
47
- },
48
- "EDSL_MAX_ATTEMPTS": {
49
- "default": "5",
50
- "info": "This config var determines the maximum number of times to retry a failed API call.",
51
- },
52
- "EDSL_SERVICE_RPM_BASELINE": {
53
- "default": "100",
54
- "info": "This config var holds the maximum number of requests per minute. Model-specific values provided in env vars such as EDSL_SERVICE_RPM_OPENAI will override this. value for the corresponding model",
55
- },
56
- "EDSL_SERVICE_TPM_BASELINE": {
57
- "default": "2000000",
58
- "info": "This config var holds the maximum number of tokens per minute for all models. Model-specific values provided in env vars such as EDSL_SERVICE_TPM_OPENAI will override this value for the corresponding model.",
59
- },
60
- "EXPECTED_PARROT_URL": {
61
- "default": "https://www.expectedparrot.com",
62
- "info": "This config var holds the URL of the Expected Parrot API.",
63
- },
64
- "EDSL_MAX_CONCURRENT_TASKS": {
65
- "default": "500",
66
- "info": "This config var determines the maximum number of concurrent tasks that can be run by the async job-runner",
67
- },
68
- "EDSL_OPEN_EXCEPTION_REPORT_URL": {
69
- "default": "False",
70
- "info": "This config var determines whether to open the exception report URL in the browser",
71
- },
72
- }
73
-
74
-
75
- class Config:
76
- """A class that loads environment variables from a .env file and sets them as class attributes."""
77
-
78
- def __init__(self):
79
- """Initialize the Config class."""
80
- self._set_run_mode()
81
- self._load_dotenv()
82
- self._set_env_vars()
83
-
84
- def _set_run_mode(self) -> None:
85
- """
86
- Sets EDSL_RUN_MODE as a class attribute.
87
- """
88
- run_mode = os.getenv("EDSL_RUN_MODE")
89
- default = CONFIG_MAP.get("EDSL_RUN_MODE").get("default")
90
- if run_mode is None:
91
- run_mode = default
92
- if run_mode not in EDSL_RUN_MODES:
93
- raise InvalidEnvironmentVariableError(
94
- f"Value `{run_mode}` is not allowed for EDSL_RUN_MODE."
95
- )
96
- self.EDSL_RUN_MODE = run_mode
97
-
98
- def _load_dotenv(self) -> None:
99
- """
100
- Loads the .env
101
- - The .env will override existing env vars **unless** EDSL_RUN_MODE=="development-testrun"
102
- """
103
-
104
- if self.EDSL_RUN_MODE == "development-testrun":
105
- override = False
106
- else:
107
- override = True
108
- _ = load_dotenv(dotenv_path=find_dotenv(usecwd=True), override=override)
109
-
110
- def __contains__(self, env_var: str) -> bool:
111
- """
112
- Checks if an env var is set as a class attribute.
113
- """
114
- return env_var in self.__dict__
115
-
116
- def _set_env_vars(self) -> None:
117
- """
118
- Sets env vars as class attributes.
119
- - EDSL_RUN_MODE is not set my this method, but by _set_run_mode
120
- - If an env var is not set and has a default value in the CONFIG_MAP, sets it to the default value.
121
- """
122
- # for each env var in the CONFIG_MAP
123
- for env_var, config in CONFIG_MAP.items():
124
- # EDSL_RUN_MODE is already set by _set_run_mode
125
- if env_var == "EDSL_RUN_MODE":
126
- continue
127
- value = os.getenv(env_var)
128
- default_value = config.get("default")
129
- # if an env var exists, set it as a class attribute
130
- if value:
131
- setattr(self, env_var, value)
132
- # otherwise, if EDSL_RUN_MODE == "production" set it to its default value
133
- elif self.EDSL_RUN_MODE == "production":
134
- setattr(self, env_var, default_value)
135
-
136
- def get(self, env_var: str) -> str:
137
- """
138
- Returns the value of an environment variable.
139
- """
140
- if env_var not in CONFIG_MAP:
141
- raise InvalidEnvironmentVariableError(f"{env_var} is not a valid env var. ")
142
- elif env_var not in self.__dict__:
143
- info = CONFIG_MAP[env_var].get("info")
144
- raise MissingEnvironmentVariableError(f"{env_var} is not set. {info}")
145
- return self.__dict__.get(env_var)
146
-
147
- def show(self) -> str:
148
- """Print the currently set environment vars."""
149
- max_env_var_length = max(len(env_var) for env_var in self.__dict__)
150
- print("Here are the current configuration settings:")
151
- for env_var, value in self.__dict__.items():
152
- print(f"{env_var:<{max_env_var_length}} : {value}")
153
-
154
-
155
- # Note: Python modules are singletons. As such, once this module is imported
156
- # the same instance of it is reused across the application.
157
- CONFIG = Config()
1
+ """This module provides a Config class that loads environment variables from a .env file and sets them as class attributes."""
2
+
3
+ import os
4
+ import platformdirs
5
+ from dotenv import load_dotenv, find_dotenv
6
+ from edsl.exceptions.configuration import (
7
+ InvalidEnvironmentVariableError,
8
+ MissingEnvironmentVariableError,
9
+ )
10
+
11
+ cache_dir = platformdirs.user_cache_dir("edsl")
12
+ os.makedirs(cache_dir, exist_ok=True)
13
+
14
+ # valid values for EDSL_RUN_MODE
15
+ EDSL_RUN_MODES = [
16
+ "development",
17
+ "development-testrun",
18
+ "production",
19
+ ]
20
+
21
+ # `default` is used to impute values only in "production" mode
22
+ # `info` gives a brief description of the env var
23
+ CONFIG_MAP = {
24
+ "EDSL_RUN_MODE": {
25
+ "default": "production",
26
+ "info": "This config var determines the run mode of the application.",
27
+ },
28
+ "EDSL_API_TIMEOUT": {
29
+ "default": "60",
30
+ "info": "This config var determines the maximum number of seconds to wait for an API call to return.",
31
+ },
32
+ "EDSL_BACKOFF_START_SEC": {
33
+ "default": "1",
34
+ "info": "This config var determines the number of seconds to wait before retrying a failed API call.",
35
+ },
36
+ "EDSL_BACKOFF_MAX_SEC": {
37
+ "default": "60",
38
+ "info": "This config var determines the maximum number of seconds to wait before retrying a failed API call.",
39
+ },
40
+ "EDSL_DATABASE_PATH": {
41
+ # "default": f"sqlite:///{os.path.join(os.getcwd(), '.edsl_cache/data.db')}",
42
+ "default": f"sqlite:///{os.path.join(platformdirs.user_cache_dir('edsl'), 'lm_model_calls.db')}",
43
+ "info": "This config var determines the path to the cache file.",
44
+ },
45
+ "EDSL_DEFAULT_MODEL": {
46
+ "default": "gpt-4o",
47
+ "info": "This config var holds the default model that will be used if a model is not explicitly passed.",
48
+ },
49
+ "EDSL_FETCH_TOKEN_PRICES": {
50
+ "default": "True",
51
+ "info": "This config var determines whether to fetch prices for tokens used in remote inference",
52
+ },
53
+ "EDSL_MAX_ATTEMPTS": {
54
+ "default": "5",
55
+ "info": "This config var determines the maximum number of times to retry a failed API call.",
56
+ },
57
+ "EDSL_SERVICE_RPM_BASELINE": {
58
+ "default": "100",
59
+ "info": "This config var holds the maximum number of requests per minute. Model-specific values provided in env vars such as EDSL_SERVICE_RPM_OPENAI will override this. value for the corresponding model",
60
+ },
61
+ "EDSL_SERVICE_TPM_BASELINE": {
62
+ "default": "2000000",
63
+ "info": "This config var holds the maximum number of tokens per minute for all models. Model-specific values provided in env vars such as EDSL_SERVICE_TPM_OPENAI will override this value for the corresponding model.",
64
+ },
65
+ "EXPECTED_PARROT_URL": {
66
+ "default": "https://www.expectedparrot.com",
67
+ "info": "This config var holds the URL of the Expected Parrot API.",
68
+ },
69
+ "EDSL_MAX_CONCURRENT_TASKS": {
70
+ "default": "500",
71
+ "info": "This config var determines the maximum number of concurrent tasks that can be run by the async job-runner",
72
+ },
73
+ "EDSL_OPEN_EXCEPTION_REPORT_URL": {
74
+ "default": "False",
75
+ "info": "This config var determines whether to open the exception report URL in the browser",
76
+ },
77
+ "EDSL_REMOTE_TOKEN_BUCKET_URL": {
78
+ "default": "None",
79
+ "info": "This config var holds the URL of the remote token bucket server.",
80
+ },
81
+ }
82
+
83
+
84
+ class Config:
85
+ """A class that loads environment variables from a .env file and sets them as class attributes."""
86
+
87
+ def __init__(self):
88
+ """Initialize the Config class."""
89
+ self._set_run_mode()
90
+ self._load_dotenv()
91
+ self._set_env_vars()
92
+
93
+ def show_path_to_dot_env(self):
94
+ print(find_dotenv(usecwd=True))
95
+
96
+ def _set_run_mode(self) -> None:
97
+ """
98
+ Sets EDSL_RUN_MODE as a class attribute.
99
+ """
100
+ run_mode = os.getenv("EDSL_RUN_MODE")
101
+ default = CONFIG_MAP.get("EDSL_RUN_MODE").get("default")
102
+ if run_mode is None:
103
+ run_mode = default
104
+ if run_mode not in EDSL_RUN_MODES:
105
+ raise InvalidEnvironmentVariableError(
106
+ f"Value `{run_mode}` is not allowed for EDSL_RUN_MODE."
107
+ )
108
+ self.EDSL_RUN_MODE = run_mode
109
+
110
+ def _load_dotenv(self) -> None:
111
+ """
112
+ Loads the .env
113
+ - The .env will override existing env vars **unless** EDSL_RUN_MODE=="development-testrun"
114
+ """
115
+
116
+ if self.EDSL_RUN_MODE == "development-testrun":
117
+ override = False
118
+ else:
119
+ override = True
120
+ _ = load_dotenv(dotenv_path=find_dotenv(usecwd=True), override=override)
121
+
122
+ def __contains__(self, env_var: str) -> bool:
123
+ """
124
+ Checks if an env var is set as a class attribute.
125
+ """
126
+ return env_var in self.__dict__
127
+
128
+ def _set_env_vars(self) -> None:
129
+ """
130
+ Sets env vars as class attributes.
131
+ - EDSL_RUN_MODE is not set my this method, but by _set_run_mode
132
+ - If an env var is not set and has a default value in the CONFIG_MAP, sets it to the default value.
133
+ """
134
+ # for each env var in the CONFIG_MAP
135
+ for env_var, config in CONFIG_MAP.items():
136
+ # EDSL_RUN_MODE is already set by _set_run_mode
137
+ if env_var == "EDSL_RUN_MODE":
138
+ continue
139
+ value = os.getenv(env_var)
140
+ default_value = config.get("default")
141
+ # if an env var exists, set it as a class attribute
142
+ if value:
143
+ setattr(self, env_var, value)
144
+ # otherwise, if EDSL_RUN_MODE == "production" set it to its default value
145
+ elif self.EDSL_RUN_MODE == "production":
146
+ setattr(self, env_var, default_value)
147
+
148
+ def get(self, env_var: str) -> str:
149
+ """
150
+ Returns the value of an environment variable.
151
+ """
152
+ if env_var not in CONFIG_MAP:
153
+ raise InvalidEnvironmentVariableError(f"{env_var} is not a valid env var. ")
154
+ elif env_var not in self.__dict__:
155
+ info = CONFIG_MAP[env_var].get("info")
156
+ raise MissingEnvironmentVariableError(f"{env_var} is not set. {info}")
157
+ return self.__dict__.get(env_var)
158
+
159
+ def __iter__(self):
160
+ """Iterate over the environment variables."""
161
+ return iter(self.__dict__)
162
+
163
+ def items(self):
164
+ """Iterate over the environment variables and their values."""
165
+ return self.__dict__.items()
166
+
167
+ def show(self) -> str:
168
+ """Print the currently set environment vars."""
169
+ max_env_var_length = max(len(env_var) for env_var in self.__dict__)
170
+ print("Here are the current configuration settings:")
171
+ for env_var, value in self.__dict__.items():
172
+ print(f"{env_var:<{max_env_var_length}} : {value}")
173
+
174
+
175
+ # Note: Python modules are singletons. As such, once this module is imported
176
+ # the same instance of it is reused across the application.
177
+ CONFIG = Config()