edsl 0.1.39.dev3__py3-none-any.whl → 0.1.39.dev5__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (341) hide show
  1. edsl/Base.py +413 -332
  2. edsl/BaseDiff.py +260 -260
  3. edsl/TemplateLoader.py +24 -24
  4. edsl/__init__.py +57 -49
  5. edsl/__version__.py +1 -1
  6. edsl/agents/Agent.py +1071 -867
  7. edsl/agents/AgentList.py +551 -413
  8. edsl/agents/Invigilator.py +284 -233
  9. edsl/agents/InvigilatorBase.py +257 -270
  10. edsl/agents/PromptConstructor.py +272 -354
  11. edsl/agents/QuestionInstructionPromptBuilder.py +128 -0
  12. edsl/agents/QuestionTemplateReplacementsBuilder.py +137 -0
  13. edsl/agents/__init__.py +2 -3
  14. edsl/agents/descriptors.py +99 -99
  15. edsl/agents/prompt_helpers.py +129 -129
  16. edsl/agents/question_option_processor.py +172 -0
  17. edsl/auto/AutoStudy.py +130 -117
  18. edsl/auto/StageBase.py +243 -230
  19. edsl/auto/StageGenerateSurvey.py +178 -178
  20. edsl/auto/StageLabelQuestions.py +125 -125
  21. edsl/auto/StagePersona.py +61 -61
  22. edsl/auto/StagePersonaDimensionValueRanges.py +88 -88
  23. edsl/auto/StagePersonaDimensionValues.py +74 -74
  24. edsl/auto/StagePersonaDimensions.py +69 -69
  25. edsl/auto/StageQuestions.py +74 -73
  26. edsl/auto/SurveyCreatorPipeline.py +21 -21
  27. edsl/auto/utilities.py +218 -224
  28. edsl/base/Base.py +279 -279
  29. edsl/config.py +177 -157
  30. edsl/conversation/Conversation.py +290 -290
  31. edsl/conversation/car_buying.py +59 -58
  32. edsl/conversation/chips.py +95 -95
  33. edsl/conversation/mug_negotiation.py +81 -81
  34. edsl/conversation/next_speaker_utilities.py +93 -93
  35. edsl/coop/CoopFunctionsMixin.py +15 -0
  36. edsl/coop/ExpectedParrotKeyHandler.py +125 -0
  37. edsl/coop/PriceFetcher.py +54 -54
  38. edsl/coop/__init__.py +2 -2
  39. edsl/coop/coop.py +1106 -1028
  40. edsl/coop/utils.py +131 -131
  41. edsl/data/Cache.py +573 -555
  42. edsl/data/CacheEntry.py +230 -233
  43. edsl/data/CacheHandler.py +168 -149
  44. edsl/data/RemoteCacheSync.py +186 -78
  45. edsl/data/SQLiteDict.py +292 -292
  46. edsl/data/__init__.py +5 -4
  47. edsl/data/orm.py +10 -10
  48. edsl/data_transfer_models.py +74 -73
  49. edsl/enums.py +202 -175
  50. edsl/exceptions/BaseException.py +21 -21
  51. edsl/exceptions/__init__.py +54 -54
  52. edsl/exceptions/agents.py +54 -42
  53. edsl/exceptions/cache.py +5 -5
  54. edsl/exceptions/configuration.py +16 -16
  55. edsl/exceptions/coop.py +10 -10
  56. edsl/exceptions/data.py +14 -14
  57. edsl/exceptions/general.py +34 -34
  58. edsl/exceptions/inference_services.py +5 -0
  59. edsl/exceptions/jobs.py +33 -33
  60. edsl/exceptions/language_models.py +63 -63
  61. edsl/exceptions/prompts.py +15 -15
  62. edsl/exceptions/questions.py +109 -91
  63. edsl/exceptions/results.py +29 -29
  64. edsl/exceptions/scenarios.py +29 -22
  65. edsl/exceptions/surveys.py +37 -37
  66. edsl/inference_services/AnthropicService.py +106 -87
  67. edsl/inference_services/AvailableModelCacheHandler.py +184 -0
  68. edsl/inference_services/AvailableModelFetcher.py +215 -0
  69. edsl/inference_services/AwsBedrock.py +118 -120
  70. edsl/inference_services/AzureAI.py +215 -217
  71. edsl/inference_services/DeepInfraService.py +18 -18
  72. edsl/inference_services/GoogleService.py +143 -148
  73. edsl/inference_services/GroqService.py +20 -20
  74. edsl/inference_services/InferenceServiceABC.py +80 -147
  75. edsl/inference_services/InferenceServicesCollection.py +138 -97
  76. edsl/inference_services/MistralAIService.py +120 -123
  77. edsl/inference_services/OllamaService.py +18 -18
  78. edsl/inference_services/OpenAIService.py +236 -224
  79. edsl/inference_services/PerplexityService.py +160 -163
  80. edsl/inference_services/ServiceAvailability.py +135 -0
  81. edsl/inference_services/TestService.py +90 -89
  82. edsl/inference_services/TogetherAIService.py +172 -170
  83. edsl/inference_services/data_structures.py +134 -0
  84. edsl/inference_services/models_available_cache.py +118 -118
  85. edsl/inference_services/rate_limits_cache.py +25 -25
  86. edsl/inference_services/registry.py +41 -41
  87. edsl/inference_services/write_available.py +10 -10
  88. edsl/jobs/AnswerQuestionFunctionConstructor.py +223 -0
  89. edsl/jobs/Answers.py +43 -56
  90. edsl/jobs/FetchInvigilator.py +47 -0
  91. edsl/jobs/InterviewTaskManager.py +98 -0
  92. edsl/jobs/InterviewsConstructor.py +50 -0
  93. edsl/jobs/Jobs.py +823 -898
  94. edsl/jobs/JobsChecks.py +172 -147
  95. edsl/jobs/JobsComponentConstructor.py +189 -0
  96. edsl/jobs/JobsPrompts.py +270 -268
  97. edsl/jobs/JobsRemoteInferenceHandler.py +311 -239
  98. edsl/jobs/JobsRemoteInferenceLogger.py +239 -0
  99. edsl/jobs/RequestTokenEstimator.py +30 -0
  100. edsl/jobs/__init__.py +1 -1
  101. edsl/jobs/async_interview_runner.py +138 -0
  102. edsl/jobs/buckets/BucketCollection.py +104 -63
  103. edsl/jobs/buckets/ModelBuckets.py +65 -65
  104. edsl/jobs/buckets/TokenBucket.py +283 -251
  105. edsl/jobs/buckets/TokenBucketAPI.py +211 -0
  106. edsl/jobs/buckets/TokenBucketClient.py +191 -0
  107. edsl/jobs/check_survey_scenario_compatibility.py +85 -0
  108. edsl/jobs/data_structures.py +120 -0
  109. edsl/jobs/decorators.py +35 -0
  110. edsl/jobs/interviews/Interview.py +396 -661
  111. edsl/jobs/interviews/InterviewExceptionCollection.py +99 -99
  112. edsl/jobs/interviews/InterviewExceptionEntry.py +186 -186
  113. edsl/jobs/interviews/InterviewStatistic.py +63 -63
  114. edsl/jobs/interviews/InterviewStatisticsCollection.py +25 -25
  115. edsl/jobs/interviews/InterviewStatusDictionary.py +78 -78
  116. edsl/jobs/interviews/InterviewStatusLog.py +92 -92
  117. edsl/jobs/interviews/ReportErrors.py +66 -66
  118. edsl/jobs/interviews/interview_status_enum.py +9 -9
  119. edsl/jobs/jobs_status_enums.py +9 -0
  120. edsl/jobs/loggers/HTMLTableJobLogger.py +304 -0
  121. edsl/jobs/results_exceptions_handler.py +98 -0
  122. edsl/jobs/runners/JobsRunnerAsyncio.py +151 -466
  123. edsl/jobs/runners/JobsRunnerStatus.py +297 -330
  124. edsl/jobs/tasks/QuestionTaskCreator.py +244 -242
  125. edsl/jobs/tasks/TaskCreators.py +64 -64
  126. edsl/jobs/tasks/TaskHistory.py +470 -450
  127. edsl/jobs/tasks/TaskStatusLog.py +23 -23
  128. edsl/jobs/tasks/task_status_enum.py +161 -163
  129. edsl/jobs/tokens/InterviewTokenUsage.py +27 -27
  130. edsl/jobs/tokens/TokenUsage.py +34 -34
  131. edsl/language_models/ComputeCost.py +63 -0
  132. edsl/language_models/LanguageModel.py +626 -668
  133. edsl/language_models/ModelList.py +164 -155
  134. edsl/language_models/PriceManager.py +127 -0
  135. edsl/language_models/RawResponseHandler.py +106 -0
  136. edsl/language_models/RegisterLanguageModelsMeta.py +184 -184
  137. edsl/language_models/ServiceDataSources.py +0 -0
  138. edsl/language_models/__init__.py +2 -3
  139. edsl/language_models/fake_openai_call.py +15 -15
  140. edsl/language_models/fake_openai_service.py +61 -61
  141. edsl/language_models/key_management/KeyLookup.py +63 -0
  142. edsl/language_models/key_management/KeyLookupBuilder.py +273 -0
  143. edsl/language_models/key_management/KeyLookupCollection.py +38 -0
  144. edsl/language_models/key_management/__init__.py +0 -0
  145. edsl/language_models/key_management/models.py +131 -0
  146. edsl/language_models/model.py +256 -0
  147. edsl/language_models/repair.py +156 -156
  148. edsl/language_models/utilities.py +65 -64
  149. edsl/notebooks/Notebook.py +263 -258
  150. edsl/notebooks/NotebookToLaTeX.py +142 -0
  151. edsl/notebooks/__init__.py +1 -1
  152. edsl/prompts/Prompt.py +352 -362
  153. edsl/prompts/__init__.py +2 -2
  154. edsl/questions/ExceptionExplainer.py +77 -0
  155. edsl/questions/HTMLQuestion.py +103 -0
  156. edsl/questions/QuestionBase.py +518 -664
  157. edsl/questions/QuestionBasePromptsMixin.py +221 -217
  158. edsl/questions/QuestionBudget.py +227 -227
  159. edsl/questions/QuestionCheckBox.py +359 -359
  160. edsl/questions/QuestionExtract.py +180 -182
  161. edsl/questions/QuestionFreeText.py +113 -114
  162. edsl/questions/QuestionFunctional.py +166 -166
  163. edsl/questions/QuestionList.py +223 -231
  164. edsl/questions/QuestionMatrix.py +265 -0
  165. edsl/questions/QuestionMultipleChoice.py +330 -286
  166. edsl/questions/QuestionNumerical.py +151 -153
  167. edsl/questions/QuestionRank.py +314 -324
  168. edsl/questions/Quick.py +41 -41
  169. edsl/questions/SimpleAskMixin.py +74 -73
  170. edsl/questions/__init__.py +27 -26
  171. edsl/questions/{AnswerValidatorMixin.py → answer_validator_mixin.py} +334 -289
  172. edsl/questions/compose_questions.py +98 -98
  173. edsl/questions/data_structures.py +20 -0
  174. edsl/questions/decorators.py +21 -21
  175. edsl/questions/derived/QuestionLikertFive.py +76 -76
  176. edsl/questions/derived/QuestionLinearScale.py +90 -87
  177. edsl/questions/derived/QuestionTopK.py +93 -93
  178. edsl/questions/derived/QuestionYesNo.py +82 -82
  179. edsl/questions/descriptors.py +427 -413
  180. edsl/questions/loop_processor.py +149 -0
  181. edsl/questions/prompt_templates/question_budget.jinja +13 -13
  182. edsl/questions/prompt_templates/question_checkbox.jinja +32 -32
  183. edsl/questions/prompt_templates/question_extract.jinja +11 -11
  184. edsl/questions/prompt_templates/question_free_text.jinja +3 -3
  185. edsl/questions/prompt_templates/question_linear_scale.jinja +11 -11
  186. edsl/questions/prompt_templates/question_list.jinja +17 -17
  187. edsl/questions/prompt_templates/question_multiple_choice.jinja +33 -33
  188. edsl/questions/prompt_templates/question_numerical.jinja +36 -36
  189. edsl/questions/{QuestionBaseGenMixin.py → question_base_gen_mixin.py} +168 -161
  190. edsl/questions/question_registry.py +177 -177
  191. edsl/questions/{RegisterQuestionsMeta.py → register_questions_meta.py} +71 -71
  192. edsl/questions/{ResponseValidatorABC.py → response_validator_abc.py} +188 -174
  193. edsl/questions/response_validator_factory.py +34 -0
  194. edsl/questions/settings.py +12 -12
  195. edsl/questions/templates/budget/answering_instructions.jinja +7 -7
  196. edsl/questions/templates/budget/question_presentation.jinja +7 -7
  197. edsl/questions/templates/checkbox/answering_instructions.jinja +10 -10
  198. edsl/questions/templates/checkbox/question_presentation.jinja +22 -22
  199. edsl/questions/templates/extract/answering_instructions.jinja +7 -7
  200. edsl/questions/templates/likert_five/answering_instructions.jinja +10 -10
  201. edsl/questions/templates/likert_five/question_presentation.jinja +11 -11
  202. edsl/questions/templates/linear_scale/answering_instructions.jinja +5 -5
  203. edsl/questions/templates/linear_scale/question_presentation.jinja +5 -5
  204. edsl/questions/templates/list/answering_instructions.jinja +3 -3
  205. edsl/questions/templates/list/question_presentation.jinja +5 -5
  206. edsl/questions/templates/matrix/__init__.py +1 -0
  207. edsl/questions/templates/matrix/answering_instructions.jinja +5 -0
  208. edsl/questions/templates/matrix/question_presentation.jinja +20 -0
  209. edsl/questions/templates/multiple_choice/answering_instructions.jinja +9 -9
  210. edsl/questions/templates/multiple_choice/question_presentation.jinja +11 -11
  211. edsl/questions/templates/numerical/answering_instructions.jinja +6 -6
  212. edsl/questions/templates/numerical/question_presentation.jinja +6 -6
  213. edsl/questions/templates/rank/answering_instructions.jinja +11 -11
  214. edsl/questions/templates/rank/question_presentation.jinja +15 -15
  215. edsl/questions/templates/top_k/answering_instructions.jinja +8 -8
  216. edsl/questions/templates/top_k/question_presentation.jinja +22 -22
  217. edsl/questions/templates/yes_no/answering_instructions.jinja +6 -6
  218. edsl/questions/templates/yes_no/question_presentation.jinja +11 -11
  219. edsl/results/CSSParameterizer.py +108 -108
  220. edsl/results/Dataset.py +587 -424
  221. edsl/results/DatasetExportMixin.py +594 -731
  222. edsl/results/DatasetTree.py +295 -275
  223. edsl/results/MarkdownToDocx.py +122 -0
  224. edsl/results/MarkdownToPDF.py +111 -0
  225. edsl/results/Result.py +557 -465
  226. edsl/results/Results.py +1183 -1165
  227. edsl/results/ResultsExportMixin.py +45 -43
  228. edsl/results/ResultsGGMixin.py +121 -121
  229. edsl/results/TableDisplay.py +125 -198
  230. edsl/results/TextEditor.py +50 -0
  231. edsl/results/__init__.py +2 -2
  232. edsl/results/file_exports.py +252 -0
  233. edsl/results/{ResultsFetchMixin.py → results_fetch_mixin.py} +33 -33
  234. edsl/results/{Selector.py → results_selector.py} +145 -135
  235. edsl/results/{ResultsToolsMixin.py → results_tools_mixin.py} +98 -98
  236. edsl/results/smart_objects.py +96 -0
  237. edsl/results/table_data_class.py +12 -0
  238. edsl/results/table_display.css +77 -77
  239. edsl/results/table_renderers.py +118 -0
  240. edsl/results/tree_explore.py +115 -115
  241. edsl/scenarios/ConstructDownloadLink.py +109 -0
  242. edsl/scenarios/DocumentChunker.py +102 -0
  243. edsl/scenarios/DocxScenario.py +16 -0
  244. edsl/scenarios/FileStore.py +511 -632
  245. edsl/scenarios/PdfExtractor.py +40 -0
  246. edsl/scenarios/Scenario.py +498 -601
  247. edsl/scenarios/ScenarioHtmlMixin.py +65 -64
  248. edsl/scenarios/ScenarioList.py +1458 -1287
  249. edsl/scenarios/ScenarioListExportMixin.py +45 -52
  250. edsl/scenarios/ScenarioListPdfMixin.py +239 -261
  251. edsl/scenarios/__init__.py +3 -4
  252. edsl/scenarios/directory_scanner.py +96 -0
  253. edsl/scenarios/file_methods.py +85 -0
  254. edsl/scenarios/handlers/__init__.py +13 -0
  255. edsl/scenarios/handlers/csv.py +38 -0
  256. edsl/scenarios/handlers/docx.py +76 -0
  257. edsl/scenarios/handlers/html.py +37 -0
  258. edsl/scenarios/handlers/json.py +111 -0
  259. edsl/scenarios/handlers/latex.py +5 -0
  260. edsl/scenarios/handlers/md.py +51 -0
  261. edsl/scenarios/handlers/pdf.py +68 -0
  262. edsl/scenarios/handlers/png.py +39 -0
  263. edsl/scenarios/handlers/pptx.py +105 -0
  264. edsl/scenarios/handlers/py.py +294 -0
  265. edsl/scenarios/handlers/sql.py +313 -0
  266. edsl/scenarios/handlers/sqlite.py +149 -0
  267. edsl/scenarios/handlers/txt.py +33 -0
  268. edsl/scenarios/{ScenarioJoin.py → scenario_join.py} +131 -127
  269. edsl/scenarios/scenario_selector.py +156 -0
  270. edsl/shared.py +1 -1
  271. edsl/study/ObjectEntry.py +173 -173
  272. edsl/study/ProofOfWork.py +113 -113
  273. edsl/study/SnapShot.py +80 -80
  274. edsl/study/Study.py +521 -528
  275. edsl/study/__init__.py +4 -4
  276. edsl/surveys/ConstructDAG.py +92 -0
  277. edsl/surveys/DAG.py +148 -148
  278. edsl/surveys/EditSurvey.py +221 -0
  279. edsl/surveys/InstructionHandler.py +100 -0
  280. edsl/surveys/Memory.py +31 -31
  281. edsl/surveys/MemoryManagement.py +72 -0
  282. edsl/surveys/MemoryPlan.py +244 -244
  283. edsl/surveys/Rule.py +327 -326
  284. edsl/surveys/RuleCollection.py +385 -387
  285. edsl/surveys/RuleManager.py +172 -0
  286. edsl/surveys/Simulator.py +75 -0
  287. edsl/surveys/Survey.py +1280 -1801
  288. edsl/surveys/SurveyCSS.py +273 -261
  289. edsl/surveys/SurveyExportMixin.py +259 -259
  290. edsl/surveys/{SurveyFlowVisualizationMixin.py → SurveyFlowVisualization.py} +181 -179
  291. edsl/surveys/SurveyQualtricsImport.py +284 -284
  292. edsl/surveys/SurveyToApp.py +141 -0
  293. edsl/surveys/__init__.py +5 -3
  294. edsl/surveys/base.py +53 -53
  295. edsl/surveys/descriptors.py +60 -56
  296. edsl/surveys/instructions/ChangeInstruction.py +48 -49
  297. edsl/surveys/instructions/Instruction.py +56 -65
  298. edsl/surveys/instructions/InstructionCollection.py +82 -77
  299. edsl/templates/error_reporting/base.html +23 -23
  300. edsl/templates/error_reporting/exceptions_by_model.html +34 -34
  301. edsl/templates/error_reporting/exceptions_by_question_name.html +16 -16
  302. edsl/templates/error_reporting/exceptions_by_type.html +16 -16
  303. edsl/templates/error_reporting/interview_details.html +115 -115
  304. edsl/templates/error_reporting/interviews.html +19 -19
  305. edsl/templates/error_reporting/overview.html +4 -4
  306. edsl/templates/error_reporting/performance_plot.html +1 -1
  307. edsl/templates/error_reporting/report.css +73 -73
  308. edsl/templates/error_reporting/report.html +117 -117
  309. edsl/templates/error_reporting/report.js +25 -25
  310. edsl/tools/__init__.py +1 -1
  311. edsl/tools/clusters.py +192 -192
  312. edsl/tools/embeddings.py +27 -27
  313. edsl/tools/embeddings_plotting.py +118 -118
  314. edsl/tools/plotting.py +112 -112
  315. edsl/tools/summarize.py +18 -18
  316. edsl/utilities/PrettyList.py +56 -0
  317. edsl/utilities/SystemInfo.py +28 -28
  318. edsl/utilities/__init__.py +22 -22
  319. edsl/utilities/ast_utilities.py +25 -25
  320. edsl/utilities/data/Registry.py +6 -6
  321. edsl/utilities/data/__init__.py +1 -1
  322. edsl/utilities/data/scooter_results.json +1 -1
  323. edsl/utilities/decorators.py +77 -77
  324. edsl/utilities/gcp_bucket/cloud_storage.py +96 -96
  325. edsl/utilities/interface.py +627 -627
  326. edsl/utilities/is_notebook.py +18 -0
  327. edsl/utilities/is_valid_variable_name.py +11 -0
  328. edsl/utilities/naming_utilities.py +263 -263
  329. edsl/utilities/remove_edsl_version.py +24 -0
  330. edsl/utilities/repair_functions.py +28 -28
  331. edsl/utilities/restricted_python.py +70 -70
  332. edsl/utilities/utilities.py +436 -424
  333. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev5.dist-info}/LICENSE +21 -21
  334. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev5.dist-info}/METADATA +13 -11
  335. edsl-0.1.39.dev5.dist-info/RECORD +358 -0
  336. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev5.dist-info}/WHEEL +1 -1
  337. edsl/language_models/KeyLookup.py +0 -30
  338. edsl/language_models/registry.py +0 -190
  339. edsl/language_models/unused/ReplicateBase.py +0 -83
  340. edsl/results/ResultsDBMixin.py +0 -238
  341. edsl-0.1.39.dev3.dist-info/RECORD +0 -277
edsl/config.py CHANGED
@@ -1,157 +1,177 @@
1
- """This module provides a Config class that loads environment variables from a .env file and sets them as class attributes."""
2
-
3
- import os
4
- from dotenv import load_dotenv, find_dotenv
5
- from edsl.exceptions import (
6
- InvalidEnvironmentVariableError,
7
- MissingEnvironmentVariableError,
8
- )
9
-
10
- # valid values for EDSL_RUN_MODE
11
- EDSL_RUN_MODES = [
12
- "development",
13
- "development-testrun",
14
- "production",
15
- ]
16
-
17
- # `default` is used to impute values only in "production" mode
18
- # `info` gives a brief description of the env var
19
- CONFIG_MAP = {
20
- "EDSL_RUN_MODE": {
21
- "default": "production",
22
- "info": "This config var determines the run mode of the application.",
23
- },
24
- "EDSL_API_TIMEOUT": {
25
- "default": "60",
26
- "info": "This config var determines the maximum number of seconds to wait for an API call to return.",
27
- },
28
- "EDSL_BACKOFF_START_SEC": {
29
- "default": "1",
30
- "info": "This config var determines the number of seconds to wait before retrying a failed API call.",
31
- },
32
- "EDSL_BACKOFF_MAX_SEC": {
33
- "default": "60",
34
- "info": "This config var determines the maximum number of seconds to wait before retrying a failed API call.",
35
- },
36
- "EDSL_DATABASE_PATH": {
37
- "default": f"sqlite:///{os.path.join(os.getcwd(), '.edsl_cache/data.db')}",
38
- "info": "This config var determines the path to the cache file.",
39
- },
40
- "EDSL_DEFAULT_MODEL": {
41
- "default": "gpt-4o",
42
- "info": "This config var holds the default model that will be used if a model is not explicitly passed.",
43
- },
44
- "EDSL_FETCH_TOKEN_PRICES": {
45
- "default": "True",
46
- "info": "This config var determines whether to fetch prices for tokens used in remote inference",
47
- },
48
- "EDSL_MAX_ATTEMPTS": {
49
- "default": "5",
50
- "info": "This config var determines the maximum number of times to retry a failed API call.",
51
- },
52
- "EDSL_SERVICE_RPM_BASELINE": {
53
- "default": "100",
54
- "info": "This config var holds the maximum number of requests per minute. Model-specific values provided in env vars such as EDSL_SERVICE_RPM_OPENAI will override this. value for the corresponding model",
55
- },
56
- "EDSL_SERVICE_TPM_BASELINE": {
57
- "default": "2000000",
58
- "info": "This config var holds the maximum number of tokens per minute for all models. Model-specific values provided in env vars such as EDSL_SERVICE_TPM_OPENAI will override this value for the corresponding model.",
59
- },
60
- "EXPECTED_PARROT_URL": {
61
- "default": "https://www.expectedparrot.com",
62
- "info": "This config var holds the URL of the Expected Parrot API.",
63
- },
64
- "EDSL_MAX_CONCURRENT_TASKS": {
65
- "default": "500",
66
- "info": "This config var determines the maximum number of concurrent tasks that can be run by the async job-runner",
67
- },
68
- "EDSL_OPEN_EXCEPTION_REPORT_URL": {
69
- "default": "False",
70
- "info": "This config var determines whether to open the exception report URL in the browser",
71
- },
72
- }
73
-
74
-
75
- class Config:
76
- """A class that loads environment variables from a .env file and sets them as class attributes."""
77
-
78
- def __init__(self):
79
- """Initialize the Config class."""
80
- self._set_run_mode()
81
- self._load_dotenv()
82
- self._set_env_vars()
83
-
84
- def _set_run_mode(self) -> None:
85
- """
86
- Sets EDSL_RUN_MODE as a class attribute.
87
- """
88
- run_mode = os.getenv("EDSL_RUN_MODE")
89
- default = CONFIG_MAP.get("EDSL_RUN_MODE").get("default")
90
- if run_mode is None:
91
- run_mode = default
92
- if run_mode not in EDSL_RUN_MODES:
93
- raise InvalidEnvironmentVariableError(
94
- f"Value `{run_mode}` is not allowed for EDSL_RUN_MODE."
95
- )
96
- self.EDSL_RUN_MODE = run_mode
97
-
98
- def _load_dotenv(self) -> None:
99
- """
100
- Loads the .env
101
- - The .env will override existing env vars **unless** EDSL_RUN_MODE=="development-testrun"
102
- """
103
-
104
- if self.EDSL_RUN_MODE == "development-testrun":
105
- override = False
106
- else:
107
- override = True
108
- _ = load_dotenv(dotenv_path=find_dotenv(usecwd=True), override=override)
109
-
110
- def __contains__(self, env_var: str) -> bool:
111
- """
112
- Checks if an env var is set as a class attribute.
113
- """
114
- return env_var in self.__dict__
115
-
116
- def _set_env_vars(self) -> None:
117
- """
118
- Sets env vars as class attributes.
119
- - EDSL_RUN_MODE is not set my this method, but by _set_run_mode
120
- - If an env var is not set and has a default value in the CONFIG_MAP, sets it to the default value.
121
- """
122
- # for each env var in the CONFIG_MAP
123
- for env_var, config in CONFIG_MAP.items():
124
- # EDSL_RUN_MODE is already set by _set_run_mode
125
- if env_var == "EDSL_RUN_MODE":
126
- continue
127
- value = os.getenv(env_var)
128
- default_value = config.get("default")
129
- # if an env var exists, set it as a class attribute
130
- if value:
131
- setattr(self, env_var, value)
132
- # otherwise, if EDSL_RUN_MODE == "production" set it to its default value
133
- elif self.EDSL_RUN_MODE == "production":
134
- setattr(self, env_var, default_value)
135
-
136
- def get(self, env_var: str) -> str:
137
- """
138
- Returns the value of an environment variable.
139
- """
140
- if env_var not in CONFIG_MAP:
141
- raise InvalidEnvironmentVariableError(f"{env_var} is not a valid env var. ")
142
- elif env_var not in self.__dict__:
143
- info = CONFIG_MAP[env_var].get("info")
144
- raise MissingEnvironmentVariableError(f"{env_var} is not set. {info}")
145
- return self.__dict__.get(env_var)
146
-
147
- def show(self) -> str:
148
- """Print the currently set environment vars."""
149
- max_env_var_length = max(len(env_var) for env_var in self.__dict__)
150
- print("Here are the current configuration settings:")
151
- for env_var, value in self.__dict__.items():
152
- print(f"{env_var:<{max_env_var_length}} : {value}")
153
-
154
-
155
- # Note: Python modules are singletons. As such, once this module is imported
156
- # the same instance of it is reused across the application.
157
- CONFIG = Config()
1
+ """This module provides a Config class that loads environment variables from a .env file and sets them as class attributes."""
2
+
3
+ import os
4
+ import platformdirs
5
+ from dotenv import load_dotenv, find_dotenv
6
+ from edsl.exceptions.configuration import (
7
+ InvalidEnvironmentVariableError,
8
+ MissingEnvironmentVariableError,
9
+ )
10
+
11
+ cache_dir = platformdirs.user_cache_dir("edsl")
12
+ os.makedirs(cache_dir, exist_ok=True)
13
+
14
+ # valid values for EDSL_RUN_MODE
15
+ EDSL_RUN_MODES = [
16
+ "development",
17
+ "development-testrun",
18
+ "production",
19
+ ]
20
+
21
+ # `default` is used to impute values only in "production" mode
22
+ # `info` gives a brief description of the env var
23
+ CONFIG_MAP = {
24
+ "EDSL_RUN_MODE": {
25
+ "default": "production",
26
+ "info": "This config var determines the run mode of the application.",
27
+ },
28
+ "EDSL_API_TIMEOUT": {
29
+ "default": "60",
30
+ "info": "This config var determines the maximum number of seconds to wait for an API call to return.",
31
+ },
32
+ "EDSL_BACKOFF_START_SEC": {
33
+ "default": "1",
34
+ "info": "This config var determines the number of seconds to wait before retrying a failed API call.",
35
+ },
36
+ "EDSL_BACKOFF_MAX_SEC": {
37
+ "default": "60",
38
+ "info": "This config var determines the maximum number of seconds to wait before retrying a failed API call.",
39
+ },
40
+ "EDSL_DATABASE_PATH": {
41
+ # "default": f"sqlite:///{os.path.join(os.getcwd(), '.edsl_cache/data.db')}",
42
+ "default": f"sqlite:///{os.path.join(platformdirs.user_cache_dir('edsl'), 'lm_model_calls.db')}",
43
+ "info": "This config var determines the path to the cache file.",
44
+ },
45
+ "EDSL_DEFAULT_MODEL": {
46
+ "default": "gpt-4o",
47
+ "info": "This config var holds the default model that will be used if a model is not explicitly passed.",
48
+ },
49
+ "EDSL_FETCH_TOKEN_PRICES": {
50
+ "default": "True",
51
+ "info": "This config var determines whether to fetch prices for tokens used in remote inference",
52
+ },
53
+ "EDSL_MAX_ATTEMPTS": {
54
+ "default": "5",
55
+ "info": "This config var determines the maximum number of times to retry a failed API call.",
56
+ },
57
+ "EDSL_SERVICE_RPM_BASELINE": {
58
+ "default": "100",
59
+ "info": "This config var holds the maximum number of requests per minute. Model-specific values provided in env vars such as EDSL_SERVICE_RPM_OPENAI will override this. value for the corresponding model",
60
+ },
61
+ "EDSL_SERVICE_TPM_BASELINE": {
62
+ "default": "2000000",
63
+ "info": "This config var holds the maximum number of tokens per minute for all models. Model-specific values provided in env vars such as EDSL_SERVICE_TPM_OPENAI will override this value for the corresponding model.",
64
+ },
65
+ "EXPECTED_PARROT_URL": {
66
+ "default": "https://www.expectedparrot.com",
67
+ "info": "This config var holds the URL of the Expected Parrot API.",
68
+ },
69
+ "EDSL_MAX_CONCURRENT_TASKS": {
70
+ "default": "500",
71
+ "info": "This config var determines the maximum number of concurrent tasks that can be run by the async job-runner",
72
+ },
73
+ "EDSL_OPEN_EXCEPTION_REPORT_URL": {
74
+ "default": "False",
75
+ "info": "This config var determines whether to open the exception report URL in the browser",
76
+ },
77
+ "EDSL_REMOTE_TOKEN_BUCKET_URL": {
78
+ "default": "None",
79
+ "info": "This config var holds the URL of the remote token bucket server.",
80
+ },
81
+ }
82
+
83
+
84
+ class Config:
85
+ """A class that loads environment variables from a .env file and sets them as class attributes."""
86
+
87
+ def __init__(self):
88
+ """Initialize the Config class."""
89
+ self._set_run_mode()
90
+ self._load_dotenv()
91
+ self._set_env_vars()
92
+
93
+ def show_path_to_dot_env(self):
94
+ print(find_dotenv(usecwd=True))
95
+
96
+ def _set_run_mode(self) -> None:
97
+ """
98
+ Sets EDSL_RUN_MODE as a class attribute.
99
+ """
100
+ run_mode = os.getenv("EDSL_RUN_MODE")
101
+ default = CONFIG_MAP.get("EDSL_RUN_MODE").get("default")
102
+ if run_mode is None:
103
+ run_mode = default
104
+ if run_mode not in EDSL_RUN_MODES:
105
+ raise InvalidEnvironmentVariableError(
106
+ f"Value `{run_mode}` is not allowed for EDSL_RUN_MODE."
107
+ )
108
+ self.EDSL_RUN_MODE = run_mode
109
+
110
+ def _load_dotenv(self) -> None:
111
+ """
112
+ Loads the .env
113
+ - The .env will override existing env vars **unless** EDSL_RUN_MODE=="development-testrun"
114
+ """
115
+
116
+ if self.EDSL_RUN_MODE == "development-testrun":
117
+ override = False
118
+ else:
119
+ override = True
120
+ _ = load_dotenv(dotenv_path=find_dotenv(usecwd=True), override=override)
121
+
122
+ def __contains__(self, env_var: str) -> bool:
123
+ """
124
+ Checks if an env var is set as a class attribute.
125
+ """
126
+ return env_var in self.__dict__
127
+
128
+ def _set_env_vars(self) -> None:
129
+ """
130
+ Sets env vars as class attributes.
131
+ - EDSL_RUN_MODE is not set my this method, but by _set_run_mode
132
+ - If an env var is not set and has a default value in the CONFIG_MAP, sets it to the default value.
133
+ """
134
+ # for each env var in the CONFIG_MAP
135
+ for env_var, config in CONFIG_MAP.items():
136
+ # EDSL_RUN_MODE is already set by _set_run_mode
137
+ if env_var == "EDSL_RUN_MODE":
138
+ continue
139
+ value = os.getenv(env_var)
140
+ default_value = config.get("default")
141
+ # if an env var exists, set it as a class attribute
142
+ if value:
143
+ setattr(self, env_var, value)
144
+ # otherwise, if EDSL_RUN_MODE == "production" set it to its default value
145
+ elif self.EDSL_RUN_MODE == "production":
146
+ setattr(self, env_var, default_value)
147
+
148
+ def get(self, env_var: str) -> str:
149
+ """
150
+ Returns the value of an environment variable.
151
+ """
152
+ if env_var not in CONFIG_MAP:
153
+ raise InvalidEnvironmentVariableError(f"{env_var} is not a valid env var. ")
154
+ elif env_var not in self.__dict__:
155
+ info = CONFIG_MAP[env_var].get("info")
156
+ raise MissingEnvironmentVariableError(f"{env_var} is not set. {info}")
157
+ return self.__dict__.get(env_var)
158
+
159
+ def __iter__(self):
160
+ """Iterate over the environment variables."""
161
+ return iter(self.__dict__)
162
+
163
+ def items(self):
164
+ """Iterate over the environment variables and their values."""
165
+ return self.__dict__.items()
166
+
167
+ def show(self) -> str:
168
+ """Print the currently set environment vars."""
169
+ max_env_var_length = max(len(env_var) for env_var in self.__dict__)
170
+ print("Here are the current configuration settings:")
171
+ for env_var, value in self.__dict__.items():
172
+ print(f"{env_var:<{max_env_var_length}} : {value}")
173
+
174
+
175
+ # Note: Python modules are singletons. As such, once this module is imported
176
+ # the same instance of it is reused across the application.
177
+ CONFIG = Config()