edsl 0.1.39.dev3__py3-none-any.whl → 0.1.39.dev5__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (341) hide show
  1. edsl/Base.py +413 -332
  2. edsl/BaseDiff.py +260 -260
  3. edsl/TemplateLoader.py +24 -24
  4. edsl/__init__.py +57 -49
  5. edsl/__version__.py +1 -1
  6. edsl/agents/Agent.py +1071 -867
  7. edsl/agents/AgentList.py +551 -413
  8. edsl/agents/Invigilator.py +284 -233
  9. edsl/agents/InvigilatorBase.py +257 -270
  10. edsl/agents/PromptConstructor.py +272 -354
  11. edsl/agents/QuestionInstructionPromptBuilder.py +128 -0
  12. edsl/agents/QuestionTemplateReplacementsBuilder.py +137 -0
  13. edsl/agents/__init__.py +2 -3
  14. edsl/agents/descriptors.py +99 -99
  15. edsl/agents/prompt_helpers.py +129 -129
  16. edsl/agents/question_option_processor.py +172 -0
  17. edsl/auto/AutoStudy.py +130 -117
  18. edsl/auto/StageBase.py +243 -230
  19. edsl/auto/StageGenerateSurvey.py +178 -178
  20. edsl/auto/StageLabelQuestions.py +125 -125
  21. edsl/auto/StagePersona.py +61 -61
  22. edsl/auto/StagePersonaDimensionValueRanges.py +88 -88
  23. edsl/auto/StagePersonaDimensionValues.py +74 -74
  24. edsl/auto/StagePersonaDimensions.py +69 -69
  25. edsl/auto/StageQuestions.py +74 -73
  26. edsl/auto/SurveyCreatorPipeline.py +21 -21
  27. edsl/auto/utilities.py +218 -224
  28. edsl/base/Base.py +279 -279
  29. edsl/config.py +177 -157
  30. edsl/conversation/Conversation.py +290 -290
  31. edsl/conversation/car_buying.py +59 -58
  32. edsl/conversation/chips.py +95 -95
  33. edsl/conversation/mug_negotiation.py +81 -81
  34. edsl/conversation/next_speaker_utilities.py +93 -93
  35. edsl/coop/CoopFunctionsMixin.py +15 -0
  36. edsl/coop/ExpectedParrotKeyHandler.py +125 -0
  37. edsl/coop/PriceFetcher.py +54 -54
  38. edsl/coop/__init__.py +2 -2
  39. edsl/coop/coop.py +1106 -1028
  40. edsl/coop/utils.py +131 -131
  41. edsl/data/Cache.py +573 -555
  42. edsl/data/CacheEntry.py +230 -233
  43. edsl/data/CacheHandler.py +168 -149
  44. edsl/data/RemoteCacheSync.py +186 -78
  45. edsl/data/SQLiteDict.py +292 -292
  46. edsl/data/__init__.py +5 -4
  47. edsl/data/orm.py +10 -10
  48. edsl/data_transfer_models.py +74 -73
  49. edsl/enums.py +202 -175
  50. edsl/exceptions/BaseException.py +21 -21
  51. edsl/exceptions/__init__.py +54 -54
  52. edsl/exceptions/agents.py +54 -42
  53. edsl/exceptions/cache.py +5 -5
  54. edsl/exceptions/configuration.py +16 -16
  55. edsl/exceptions/coop.py +10 -10
  56. edsl/exceptions/data.py +14 -14
  57. edsl/exceptions/general.py +34 -34
  58. edsl/exceptions/inference_services.py +5 -0
  59. edsl/exceptions/jobs.py +33 -33
  60. edsl/exceptions/language_models.py +63 -63
  61. edsl/exceptions/prompts.py +15 -15
  62. edsl/exceptions/questions.py +109 -91
  63. edsl/exceptions/results.py +29 -29
  64. edsl/exceptions/scenarios.py +29 -22
  65. edsl/exceptions/surveys.py +37 -37
  66. edsl/inference_services/AnthropicService.py +106 -87
  67. edsl/inference_services/AvailableModelCacheHandler.py +184 -0
  68. edsl/inference_services/AvailableModelFetcher.py +215 -0
  69. edsl/inference_services/AwsBedrock.py +118 -120
  70. edsl/inference_services/AzureAI.py +215 -217
  71. edsl/inference_services/DeepInfraService.py +18 -18
  72. edsl/inference_services/GoogleService.py +143 -148
  73. edsl/inference_services/GroqService.py +20 -20
  74. edsl/inference_services/InferenceServiceABC.py +80 -147
  75. edsl/inference_services/InferenceServicesCollection.py +138 -97
  76. edsl/inference_services/MistralAIService.py +120 -123
  77. edsl/inference_services/OllamaService.py +18 -18
  78. edsl/inference_services/OpenAIService.py +236 -224
  79. edsl/inference_services/PerplexityService.py +160 -163
  80. edsl/inference_services/ServiceAvailability.py +135 -0
  81. edsl/inference_services/TestService.py +90 -89
  82. edsl/inference_services/TogetherAIService.py +172 -170
  83. edsl/inference_services/data_structures.py +134 -0
  84. edsl/inference_services/models_available_cache.py +118 -118
  85. edsl/inference_services/rate_limits_cache.py +25 -25
  86. edsl/inference_services/registry.py +41 -41
  87. edsl/inference_services/write_available.py +10 -10
  88. edsl/jobs/AnswerQuestionFunctionConstructor.py +223 -0
  89. edsl/jobs/Answers.py +43 -56
  90. edsl/jobs/FetchInvigilator.py +47 -0
  91. edsl/jobs/InterviewTaskManager.py +98 -0
  92. edsl/jobs/InterviewsConstructor.py +50 -0
  93. edsl/jobs/Jobs.py +823 -898
  94. edsl/jobs/JobsChecks.py +172 -147
  95. edsl/jobs/JobsComponentConstructor.py +189 -0
  96. edsl/jobs/JobsPrompts.py +270 -268
  97. edsl/jobs/JobsRemoteInferenceHandler.py +311 -239
  98. edsl/jobs/JobsRemoteInferenceLogger.py +239 -0
  99. edsl/jobs/RequestTokenEstimator.py +30 -0
  100. edsl/jobs/__init__.py +1 -1
  101. edsl/jobs/async_interview_runner.py +138 -0
  102. edsl/jobs/buckets/BucketCollection.py +104 -63
  103. edsl/jobs/buckets/ModelBuckets.py +65 -65
  104. edsl/jobs/buckets/TokenBucket.py +283 -251
  105. edsl/jobs/buckets/TokenBucketAPI.py +211 -0
  106. edsl/jobs/buckets/TokenBucketClient.py +191 -0
  107. edsl/jobs/check_survey_scenario_compatibility.py +85 -0
  108. edsl/jobs/data_structures.py +120 -0
  109. edsl/jobs/decorators.py +35 -0
  110. edsl/jobs/interviews/Interview.py +396 -661
  111. edsl/jobs/interviews/InterviewExceptionCollection.py +99 -99
  112. edsl/jobs/interviews/InterviewExceptionEntry.py +186 -186
  113. edsl/jobs/interviews/InterviewStatistic.py +63 -63
  114. edsl/jobs/interviews/InterviewStatisticsCollection.py +25 -25
  115. edsl/jobs/interviews/InterviewStatusDictionary.py +78 -78
  116. edsl/jobs/interviews/InterviewStatusLog.py +92 -92
  117. edsl/jobs/interviews/ReportErrors.py +66 -66
  118. edsl/jobs/interviews/interview_status_enum.py +9 -9
  119. edsl/jobs/jobs_status_enums.py +9 -0
  120. edsl/jobs/loggers/HTMLTableJobLogger.py +304 -0
  121. edsl/jobs/results_exceptions_handler.py +98 -0
  122. edsl/jobs/runners/JobsRunnerAsyncio.py +151 -466
  123. edsl/jobs/runners/JobsRunnerStatus.py +297 -330
  124. edsl/jobs/tasks/QuestionTaskCreator.py +244 -242
  125. edsl/jobs/tasks/TaskCreators.py +64 -64
  126. edsl/jobs/tasks/TaskHistory.py +470 -450
  127. edsl/jobs/tasks/TaskStatusLog.py +23 -23
  128. edsl/jobs/tasks/task_status_enum.py +161 -163
  129. edsl/jobs/tokens/InterviewTokenUsage.py +27 -27
  130. edsl/jobs/tokens/TokenUsage.py +34 -34
  131. edsl/language_models/ComputeCost.py +63 -0
  132. edsl/language_models/LanguageModel.py +626 -668
  133. edsl/language_models/ModelList.py +164 -155
  134. edsl/language_models/PriceManager.py +127 -0
  135. edsl/language_models/RawResponseHandler.py +106 -0
  136. edsl/language_models/RegisterLanguageModelsMeta.py +184 -184
  137. edsl/language_models/ServiceDataSources.py +0 -0
  138. edsl/language_models/__init__.py +2 -3
  139. edsl/language_models/fake_openai_call.py +15 -15
  140. edsl/language_models/fake_openai_service.py +61 -61
  141. edsl/language_models/key_management/KeyLookup.py +63 -0
  142. edsl/language_models/key_management/KeyLookupBuilder.py +273 -0
  143. edsl/language_models/key_management/KeyLookupCollection.py +38 -0
  144. edsl/language_models/key_management/__init__.py +0 -0
  145. edsl/language_models/key_management/models.py +131 -0
  146. edsl/language_models/model.py +256 -0
  147. edsl/language_models/repair.py +156 -156
  148. edsl/language_models/utilities.py +65 -64
  149. edsl/notebooks/Notebook.py +263 -258
  150. edsl/notebooks/NotebookToLaTeX.py +142 -0
  151. edsl/notebooks/__init__.py +1 -1
  152. edsl/prompts/Prompt.py +352 -362
  153. edsl/prompts/__init__.py +2 -2
  154. edsl/questions/ExceptionExplainer.py +77 -0
  155. edsl/questions/HTMLQuestion.py +103 -0
  156. edsl/questions/QuestionBase.py +518 -664
  157. edsl/questions/QuestionBasePromptsMixin.py +221 -217
  158. edsl/questions/QuestionBudget.py +227 -227
  159. edsl/questions/QuestionCheckBox.py +359 -359
  160. edsl/questions/QuestionExtract.py +180 -182
  161. edsl/questions/QuestionFreeText.py +113 -114
  162. edsl/questions/QuestionFunctional.py +166 -166
  163. edsl/questions/QuestionList.py +223 -231
  164. edsl/questions/QuestionMatrix.py +265 -0
  165. edsl/questions/QuestionMultipleChoice.py +330 -286
  166. edsl/questions/QuestionNumerical.py +151 -153
  167. edsl/questions/QuestionRank.py +314 -324
  168. edsl/questions/Quick.py +41 -41
  169. edsl/questions/SimpleAskMixin.py +74 -73
  170. edsl/questions/__init__.py +27 -26
  171. edsl/questions/{AnswerValidatorMixin.py → answer_validator_mixin.py} +334 -289
  172. edsl/questions/compose_questions.py +98 -98
  173. edsl/questions/data_structures.py +20 -0
  174. edsl/questions/decorators.py +21 -21
  175. edsl/questions/derived/QuestionLikertFive.py +76 -76
  176. edsl/questions/derived/QuestionLinearScale.py +90 -87
  177. edsl/questions/derived/QuestionTopK.py +93 -93
  178. edsl/questions/derived/QuestionYesNo.py +82 -82
  179. edsl/questions/descriptors.py +427 -413
  180. edsl/questions/loop_processor.py +149 -0
  181. edsl/questions/prompt_templates/question_budget.jinja +13 -13
  182. edsl/questions/prompt_templates/question_checkbox.jinja +32 -32
  183. edsl/questions/prompt_templates/question_extract.jinja +11 -11
  184. edsl/questions/prompt_templates/question_free_text.jinja +3 -3
  185. edsl/questions/prompt_templates/question_linear_scale.jinja +11 -11
  186. edsl/questions/prompt_templates/question_list.jinja +17 -17
  187. edsl/questions/prompt_templates/question_multiple_choice.jinja +33 -33
  188. edsl/questions/prompt_templates/question_numerical.jinja +36 -36
  189. edsl/questions/{QuestionBaseGenMixin.py → question_base_gen_mixin.py} +168 -161
  190. edsl/questions/question_registry.py +177 -177
  191. edsl/questions/{RegisterQuestionsMeta.py → register_questions_meta.py} +71 -71
  192. edsl/questions/{ResponseValidatorABC.py → response_validator_abc.py} +188 -174
  193. edsl/questions/response_validator_factory.py +34 -0
  194. edsl/questions/settings.py +12 -12
  195. edsl/questions/templates/budget/answering_instructions.jinja +7 -7
  196. edsl/questions/templates/budget/question_presentation.jinja +7 -7
  197. edsl/questions/templates/checkbox/answering_instructions.jinja +10 -10
  198. edsl/questions/templates/checkbox/question_presentation.jinja +22 -22
  199. edsl/questions/templates/extract/answering_instructions.jinja +7 -7
  200. edsl/questions/templates/likert_five/answering_instructions.jinja +10 -10
  201. edsl/questions/templates/likert_five/question_presentation.jinja +11 -11
  202. edsl/questions/templates/linear_scale/answering_instructions.jinja +5 -5
  203. edsl/questions/templates/linear_scale/question_presentation.jinja +5 -5
  204. edsl/questions/templates/list/answering_instructions.jinja +3 -3
  205. edsl/questions/templates/list/question_presentation.jinja +5 -5
  206. edsl/questions/templates/matrix/__init__.py +1 -0
  207. edsl/questions/templates/matrix/answering_instructions.jinja +5 -0
  208. edsl/questions/templates/matrix/question_presentation.jinja +20 -0
  209. edsl/questions/templates/multiple_choice/answering_instructions.jinja +9 -9
  210. edsl/questions/templates/multiple_choice/question_presentation.jinja +11 -11
  211. edsl/questions/templates/numerical/answering_instructions.jinja +6 -6
  212. edsl/questions/templates/numerical/question_presentation.jinja +6 -6
  213. edsl/questions/templates/rank/answering_instructions.jinja +11 -11
  214. edsl/questions/templates/rank/question_presentation.jinja +15 -15
  215. edsl/questions/templates/top_k/answering_instructions.jinja +8 -8
  216. edsl/questions/templates/top_k/question_presentation.jinja +22 -22
  217. edsl/questions/templates/yes_no/answering_instructions.jinja +6 -6
  218. edsl/questions/templates/yes_no/question_presentation.jinja +11 -11
  219. edsl/results/CSSParameterizer.py +108 -108
  220. edsl/results/Dataset.py +587 -424
  221. edsl/results/DatasetExportMixin.py +594 -731
  222. edsl/results/DatasetTree.py +295 -275
  223. edsl/results/MarkdownToDocx.py +122 -0
  224. edsl/results/MarkdownToPDF.py +111 -0
  225. edsl/results/Result.py +557 -465
  226. edsl/results/Results.py +1183 -1165
  227. edsl/results/ResultsExportMixin.py +45 -43
  228. edsl/results/ResultsGGMixin.py +121 -121
  229. edsl/results/TableDisplay.py +125 -198
  230. edsl/results/TextEditor.py +50 -0
  231. edsl/results/__init__.py +2 -2
  232. edsl/results/file_exports.py +252 -0
  233. edsl/results/{ResultsFetchMixin.py → results_fetch_mixin.py} +33 -33
  234. edsl/results/{Selector.py → results_selector.py} +145 -135
  235. edsl/results/{ResultsToolsMixin.py → results_tools_mixin.py} +98 -98
  236. edsl/results/smart_objects.py +96 -0
  237. edsl/results/table_data_class.py +12 -0
  238. edsl/results/table_display.css +77 -77
  239. edsl/results/table_renderers.py +118 -0
  240. edsl/results/tree_explore.py +115 -115
  241. edsl/scenarios/ConstructDownloadLink.py +109 -0
  242. edsl/scenarios/DocumentChunker.py +102 -0
  243. edsl/scenarios/DocxScenario.py +16 -0
  244. edsl/scenarios/FileStore.py +511 -632
  245. edsl/scenarios/PdfExtractor.py +40 -0
  246. edsl/scenarios/Scenario.py +498 -601
  247. edsl/scenarios/ScenarioHtmlMixin.py +65 -64
  248. edsl/scenarios/ScenarioList.py +1458 -1287
  249. edsl/scenarios/ScenarioListExportMixin.py +45 -52
  250. edsl/scenarios/ScenarioListPdfMixin.py +239 -261
  251. edsl/scenarios/__init__.py +3 -4
  252. edsl/scenarios/directory_scanner.py +96 -0
  253. edsl/scenarios/file_methods.py +85 -0
  254. edsl/scenarios/handlers/__init__.py +13 -0
  255. edsl/scenarios/handlers/csv.py +38 -0
  256. edsl/scenarios/handlers/docx.py +76 -0
  257. edsl/scenarios/handlers/html.py +37 -0
  258. edsl/scenarios/handlers/json.py +111 -0
  259. edsl/scenarios/handlers/latex.py +5 -0
  260. edsl/scenarios/handlers/md.py +51 -0
  261. edsl/scenarios/handlers/pdf.py +68 -0
  262. edsl/scenarios/handlers/png.py +39 -0
  263. edsl/scenarios/handlers/pptx.py +105 -0
  264. edsl/scenarios/handlers/py.py +294 -0
  265. edsl/scenarios/handlers/sql.py +313 -0
  266. edsl/scenarios/handlers/sqlite.py +149 -0
  267. edsl/scenarios/handlers/txt.py +33 -0
  268. edsl/scenarios/{ScenarioJoin.py → scenario_join.py} +131 -127
  269. edsl/scenarios/scenario_selector.py +156 -0
  270. edsl/shared.py +1 -1
  271. edsl/study/ObjectEntry.py +173 -173
  272. edsl/study/ProofOfWork.py +113 -113
  273. edsl/study/SnapShot.py +80 -80
  274. edsl/study/Study.py +521 -528
  275. edsl/study/__init__.py +4 -4
  276. edsl/surveys/ConstructDAG.py +92 -0
  277. edsl/surveys/DAG.py +148 -148
  278. edsl/surveys/EditSurvey.py +221 -0
  279. edsl/surveys/InstructionHandler.py +100 -0
  280. edsl/surveys/Memory.py +31 -31
  281. edsl/surveys/MemoryManagement.py +72 -0
  282. edsl/surveys/MemoryPlan.py +244 -244
  283. edsl/surveys/Rule.py +327 -326
  284. edsl/surveys/RuleCollection.py +385 -387
  285. edsl/surveys/RuleManager.py +172 -0
  286. edsl/surveys/Simulator.py +75 -0
  287. edsl/surveys/Survey.py +1280 -1801
  288. edsl/surveys/SurveyCSS.py +273 -261
  289. edsl/surveys/SurveyExportMixin.py +259 -259
  290. edsl/surveys/{SurveyFlowVisualizationMixin.py → SurveyFlowVisualization.py} +181 -179
  291. edsl/surveys/SurveyQualtricsImport.py +284 -284
  292. edsl/surveys/SurveyToApp.py +141 -0
  293. edsl/surveys/__init__.py +5 -3
  294. edsl/surveys/base.py +53 -53
  295. edsl/surveys/descriptors.py +60 -56
  296. edsl/surveys/instructions/ChangeInstruction.py +48 -49
  297. edsl/surveys/instructions/Instruction.py +56 -65
  298. edsl/surveys/instructions/InstructionCollection.py +82 -77
  299. edsl/templates/error_reporting/base.html +23 -23
  300. edsl/templates/error_reporting/exceptions_by_model.html +34 -34
  301. edsl/templates/error_reporting/exceptions_by_question_name.html +16 -16
  302. edsl/templates/error_reporting/exceptions_by_type.html +16 -16
  303. edsl/templates/error_reporting/interview_details.html +115 -115
  304. edsl/templates/error_reporting/interviews.html +19 -19
  305. edsl/templates/error_reporting/overview.html +4 -4
  306. edsl/templates/error_reporting/performance_plot.html +1 -1
  307. edsl/templates/error_reporting/report.css +73 -73
  308. edsl/templates/error_reporting/report.html +117 -117
  309. edsl/templates/error_reporting/report.js +25 -25
  310. edsl/tools/__init__.py +1 -1
  311. edsl/tools/clusters.py +192 -192
  312. edsl/tools/embeddings.py +27 -27
  313. edsl/tools/embeddings_plotting.py +118 -118
  314. edsl/tools/plotting.py +112 -112
  315. edsl/tools/summarize.py +18 -18
  316. edsl/utilities/PrettyList.py +56 -0
  317. edsl/utilities/SystemInfo.py +28 -28
  318. edsl/utilities/__init__.py +22 -22
  319. edsl/utilities/ast_utilities.py +25 -25
  320. edsl/utilities/data/Registry.py +6 -6
  321. edsl/utilities/data/__init__.py +1 -1
  322. edsl/utilities/data/scooter_results.json +1 -1
  323. edsl/utilities/decorators.py +77 -77
  324. edsl/utilities/gcp_bucket/cloud_storage.py +96 -96
  325. edsl/utilities/interface.py +627 -627
  326. edsl/utilities/is_notebook.py +18 -0
  327. edsl/utilities/is_valid_variable_name.py +11 -0
  328. edsl/utilities/naming_utilities.py +263 -263
  329. edsl/utilities/remove_edsl_version.py +24 -0
  330. edsl/utilities/repair_functions.py +28 -28
  331. edsl/utilities/restricted_python.py +70 -70
  332. edsl/utilities/utilities.py +436 -424
  333. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev5.dist-info}/LICENSE +21 -21
  334. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev5.dist-info}/METADATA +13 -11
  335. edsl-0.1.39.dev5.dist-info/RECORD +358 -0
  336. {edsl-0.1.39.dev3.dist-info → edsl-0.1.39.dev5.dist-info}/WHEEL +1 -1
  337. edsl/language_models/KeyLookup.py +0 -30
  338. edsl/language_models/registry.py +0 -190
  339. edsl/language_models/unused/ReplicateBase.py +0 -83
  340. edsl/results/ResultsDBMixin.py +0 -238
  341. edsl-0.1.39.dev3.dist-info/RECORD +0 -277
@@ -1,270 +1,257 @@
1
- from abc import ABC, abstractmethod
2
- import asyncio
3
- from typing import Coroutine, Dict, Any, Optional
4
-
5
- from edsl.prompts.Prompt import Prompt
6
- from edsl.utilities.decorators import jupyter_nb_handler
7
- from edsl.data_transfer_models import AgentResponseDict
8
-
9
- from edsl.data.Cache import Cache
10
-
11
- from edsl.questions.QuestionBase import QuestionBase
12
- from edsl.scenarios.Scenario import Scenario
13
- from edsl.surveys.MemoryPlan import MemoryPlan
14
- from edsl.language_models.LanguageModel import LanguageModel
15
-
16
- from edsl.data_transfer_models import EDSLResultObjectInput
17
- from edsl.agents.PromptConstructor import PromptConstructor
18
-
19
- from edsl.agents.prompt_helpers import PromptPlan
20
-
21
-
22
- class InvigilatorBase(ABC):
23
- """An invigiator (someone who administers an exam) is a class that is responsible for administering a question to an agent.
24
-
25
- >>> InvigilatorBase.example().answer_question()
26
- {'message': [{'text': 'SPAM!'}], 'usage': {'prompt_tokens': 1, 'completion_tokens': 1}}
27
-
28
- >>> InvigilatorBase.example().get_failed_task_result(failure_reason="Failed to get response").comment
29
- 'Failed to get response'
30
-
31
- This returns an empty prompt because there is no memory the agent needs to have at q0.
32
-
33
-
34
- """
35
-
36
- def __init__(
37
- self,
38
- agent: "Agent",
39
- question: QuestionBase,
40
- scenario: Scenario,
41
- model: LanguageModel,
42
- memory_plan: MemoryPlan,
43
- current_answers: dict,
44
- survey: Optional["Survey"],
45
- cache: Optional[Cache] = None,
46
- iteration: Optional[int] = 1,
47
- additional_prompt_data: Optional[dict] = None,
48
- sidecar_model: Optional[LanguageModel] = None,
49
- raise_validation_errors: Optional[bool] = True,
50
- prompt_plan: Optional["PromptPlan"] = None,
51
- ):
52
- """Initialize a new Invigilator."""
53
- self.agent = agent
54
- self.question = question
55
- self.scenario = scenario
56
- self.model = model
57
- self.memory_plan = memory_plan
58
- self.current_answers = current_answers or {}
59
- self.iteration = iteration
60
- self.additional_prompt_data = additional_prompt_data
61
- self.cache = cache
62
- self.sidecar_model = sidecar_model
63
- self.survey = survey
64
- self.raise_validation_errors = raise_validation_errors
65
- if prompt_plan is None:
66
- self.prompt_plan = PromptPlan()
67
- else:
68
- self.prompt_plan = prompt_plan
69
-
70
- self.raw_model_response = (
71
- None # placeholder for the raw response from the model
72
- )
73
-
74
- @property
75
- def prompt_constructor(self) -> PromptConstructor:
76
- """Return the prompt constructor."""
77
- return PromptConstructor(self, prompt_plan=self.prompt_plan)
78
-
79
- def to_dict(self, include_cache=False):
80
- attributes = [
81
- "agent",
82
- "question",
83
- "scenario",
84
- "model",
85
- "memory_plan",
86
- "current_answers",
87
- "iteration",
88
- "additional_prompt_data",
89
- "sidecar_model",
90
- "survey",
91
- ]
92
- if include_cache:
93
- attributes.append("cache")
94
-
95
- def serialize_attribute(attr):
96
- value = getattr(self, attr)
97
- if value is None:
98
- return None
99
- if hasattr(value, "to_dict"):
100
- return value.to_dict()
101
- if isinstance(value, (int, float, str, bool, dict, list)):
102
- return value
103
- return str(value)
104
-
105
- return {attr: serialize_attribute(attr) for attr in attributes}
106
-
107
- @classmethod
108
- def from_dict(cls, data):
109
- from edsl.agents.Agent import Agent
110
- from edsl.questions import QuestionBase
111
- from edsl.scenarios.Scenario import Scenario
112
- from edsl.surveys.MemoryPlan import MemoryPlan
113
- from edsl.language_models.LanguageModel import LanguageModel
114
- from edsl.surveys.Survey import Survey
115
- from edsl.data.Cache import Cache
116
-
117
- agent = Agent.from_dict(data["agent"])
118
- question = QuestionBase.from_dict(data["question"])
119
- scenario = Scenario.from_dict(data["scenario"])
120
- model = LanguageModel.from_dict(data["model"])
121
- memory_plan = MemoryPlan.from_dict(data["memory_plan"])
122
- survey = Survey.from_dict(data["survey"])
123
- current_answers = data["current_answers"]
124
- iteration = data["iteration"]
125
- additional_prompt_data = data["additional_prompt_data"]
126
- if "cache" not in data:
127
- cache = {}
128
- else:
129
- cache = Cache.from_dict(data["cache"])
130
-
131
- if data["sidecar_model"] is None:
132
- sidecar_model = None
133
- else:
134
- sidecar_model = LanguageModel.from_dict(data["sidecar_model"])
135
-
136
- return cls(
137
- agent=agent,
138
- question=question,
139
- scenario=scenario,
140
- model=model,
141
- memory_plan=memory_plan,
142
- current_answers=current_answers,
143
- survey=survey,
144
- iteration=iteration,
145
- additional_prompt_data=additional_prompt_data,
146
- cache=cache,
147
- sidecar_model=sidecar_model,
148
- )
149
-
150
- def __repr__(self) -> str:
151
- """Return a string representation of the Invigilator.
152
-
153
- >>> InvigilatorBase.example().__repr__()
154
- 'InvigilatorExample(...)'
155
-
156
- """
157
- return f"{self.__class__.__name__}(agent={repr(self.agent)}, question={repr(self.question)}, scneario={repr(self.scenario)}, model={repr(self.model)}, memory_plan={repr(self.memory_plan)}, current_answers={repr(self.current_answers)}, iteration{repr(self.iteration)}, additional_prompt_data={repr(self.additional_prompt_data)}, cache={repr(self.cache)}, sidecarmodel={repr(self.sidecar_model)})"
158
-
159
- def get_failed_task_result(self, failure_reason) -> EDSLResultObjectInput:
160
- """Return an AgentResponseDict used in case the question-asking fails.
161
-
162
- Possible reasons include:
163
- - Legimately skipped because of skip logic
164
- - Failed to get response from the model
165
-
166
- """
167
- data = {
168
- "answer": None,
169
- "generated_tokens": None,
170
- "comment": failure_reason,
171
- "question_name": self.question.question_name,
172
- "prompts": self.get_prompts(),
173
- "cached_response": None,
174
- "raw_model_response": None,
175
- "cache_used": None,
176
- "cache_key": None,
177
- }
178
- return EDSLResultObjectInput(**data)
179
-
180
- def get_prompts(self) -> Dict[str, Prompt]:
181
- """Return the prompt used."""
182
-
183
- return {
184
- "user_prompt": Prompt("NA"),
185
- "system_prompt": Prompt("NA"),
186
- }
187
-
188
- @abstractmethod
189
- async def async_answer_question(self):
190
- """Asnwer a question."""
191
- pass
192
-
193
- @jupyter_nb_handler
194
- def answer_question(self) -> Coroutine:
195
- """Return a function that gets the answers to the question."""
196
-
197
- async def main():
198
- """Return the answer to the question."""
199
- results = await asyncio.gather(self.async_answer_question())
200
- return results[0] # Since there's only one task, return its result
201
-
202
- return main()
203
-
204
- @classmethod
205
- def example(
206
- cls, throw_an_exception=False, question=None, scenario=None, survey=None
207
- ) -> "InvigilatorBase":
208
- """Return an example invigilator.
209
-
210
- >>> InvigilatorBase.example()
211
- InvigilatorExample(...)
212
-
213
- """
214
- from edsl.agents.Agent import Agent
215
- from edsl.questions import QuestionMultipleChoice
216
- from edsl.scenarios.Scenario import Scenario
217
- from edsl.language_models import LanguageModel
218
- from edsl.surveys.MemoryPlan import MemoryPlan
219
-
220
- from edsl.enums import InferenceServiceType
221
-
222
- from edsl import Model
223
-
224
- model = Model("test", canned_response="SPAM!")
225
-
226
- if throw_an_exception:
227
- model.throw_an_exception = True
228
- agent = Agent.example()
229
- # question = QuestionMultipleChoice.example()
230
- from edsl.surveys import Survey
231
-
232
- if not survey:
233
- survey = Survey.example()
234
-
235
- if question not in survey.questions and question is not None:
236
- survey.add_question(question)
237
-
238
- question = question or survey.questions[0]
239
- scenario = scenario or Scenario.example()
240
- # memory_plan = None #memory_plan = MemoryPlan()
241
- from edsl import Survey
242
-
243
- memory_plan = MemoryPlan(survey=survey)
244
- current_answers = None
245
- from edsl.agents.PromptConstructor import PromptConstructor
246
-
247
- class InvigilatorExample(InvigilatorBase):
248
- """An example invigilator."""
249
-
250
- async def async_answer_question(self):
251
- """Answer a question."""
252
- return await self.model.async_execute_model_call(
253
- user_prompt="Hello", system_prompt="Hi"
254
- )
255
-
256
- return InvigilatorExample(
257
- agent=agent,
258
- question=question,
259
- scenario=scenario,
260
- survey=survey,
261
- model=model,
262
- memory_plan=memory_plan,
263
- current_answers=current_answers,
264
- )
265
-
266
-
267
- if __name__ == "__main__":
268
- import doctest
269
-
270
- doctest.testmod(optionflags=doctest.ELLIPSIS)
1
+ from abc import ABC, abstractmethod
2
+ import asyncio
3
+ from typing import Coroutine, Dict, Any, Optional, TYPE_CHECKING
4
+
5
+ from edsl.utilities.decorators import jupyter_nb_handler
6
+ from edsl.data_transfer_models import AgentResponseDict
7
+
8
+ if TYPE_CHECKING:
9
+ from edsl.prompts.Prompt import Prompt
10
+ from edsl.data.Cache import Cache
11
+ from edsl.questions.QuestionBase import QuestionBase
12
+ from edsl.scenarios.Scenario import Scenario
13
+ from edsl.surveys.MemoryPlan import MemoryPlan
14
+ from edsl.language_models.LanguageModel import LanguageModel
15
+ from edsl.surveys.Survey import Survey
16
+ from edsl.agents.Agent import Agent
17
+ from edsl.language_models.key_management.KeyLookup import KeyLookup
18
+
19
+ from edsl.data_transfer_models import EDSLResultObjectInput
20
+ from edsl.agents.PromptConstructor import PromptConstructor
21
+ from edsl.agents.prompt_helpers import PromptPlan
22
+
23
+
24
+ class InvigilatorBase(ABC):
25
+ """An invigiator (someone who administers an exam) is a class that is responsible for administering a question to an agent.
26
+
27
+ >>> InvigilatorBase.example().answer_question()
28
+ {'message': [{'text': 'SPAM!'}], 'usage': {'prompt_tokens': 1, 'completion_tokens': 1}}
29
+
30
+ >>> InvigilatorBase.example().get_failed_task_result(failure_reason="Failed to get response").comment
31
+ 'Failed to get response'
32
+
33
+ This returns an empty prompt because there is no memory the agent needs to have at q0.
34
+ """
35
+
36
+ def __init__(
37
+ self,
38
+ agent: "Agent",
39
+ question: "QuestionBase",
40
+ scenario: "Scenario",
41
+ model: "LanguageModel",
42
+ memory_plan: "MemoryPlan",
43
+ current_answers: dict,
44
+ survey: Optional["Survey"],
45
+ cache: Optional["Cache"] = None,
46
+ iteration: Optional[int] = 1,
47
+ additional_prompt_data: Optional[dict] = None,
48
+ raise_validation_errors: Optional[bool] = True,
49
+ prompt_plan: Optional["PromptPlan"] = None,
50
+ key_lookup: Optional["KeyLookup"] = None,
51
+ ):
52
+ """Initialize a new Invigilator."""
53
+ self.agent = agent
54
+ self.question = question
55
+ self.scenario = scenario
56
+ self.model = model
57
+ self.memory_plan = memory_plan
58
+ self.current_answers = current_answers or {}
59
+ self.iteration = iteration
60
+ self.additional_prompt_data = additional_prompt_data
61
+ self.cache = cache
62
+ self.survey = survey
63
+ self.raise_validation_errors = raise_validation_errors
64
+ self.key_lookup = key_lookup
65
+
66
+ if prompt_plan is None:
67
+ self.prompt_plan = PromptPlan()
68
+ else:
69
+ self.prompt_plan = prompt_plan
70
+
71
+ # placeholder to store the raw model response
72
+ self.raw_model_response = None
73
+
74
+ @property
75
+ def prompt_constructor(self) -> PromptConstructor:
76
+ """Return the prompt constructor."""
77
+ return PromptConstructor(self, prompt_plan=self.prompt_plan)
78
+
79
+ def to_dict(self, include_cache=False) -> Dict[str, Any]:
80
+ attributes = [
81
+ "agent",
82
+ "question",
83
+ "scenario",
84
+ "model",
85
+ "memory_plan",
86
+ "current_answers",
87
+ "iteration",
88
+ "additional_prompt_data",
89
+ "survey",
90
+ ]
91
+ if include_cache:
92
+ attributes.append("cache")
93
+
94
+ def serialize_attribute(attr):
95
+ value = getattr(self, attr)
96
+ if value is None:
97
+ return None
98
+ if hasattr(value, "to_dict"):
99
+ return value.to_dict()
100
+ if isinstance(value, (int, float, str, bool, dict, list)):
101
+ return value
102
+ return str(value)
103
+
104
+ return {attr: serialize_attribute(attr) for attr in attributes}
105
+
106
+ @classmethod
107
+ def from_dict(cls, data) -> "InvigilatorBase":
108
+ from edsl.agents.Agent import Agent
109
+ from edsl.questions import QuestionBase
110
+ from edsl.scenarios.Scenario import Scenario
111
+ from edsl.surveys.MemoryPlan import MemoryPlan
112
+ from edsl.language_models.LanguageModel import LanguageModel
113
+ from edsl.surveys.Survey import Survey
114
+ from edsl.data.Cache import Cache
115
+
116
+ attributes_to_classes = {
117
+ "agent": Agent,
118
+ "question": QuestionBase,
119
+ "scenario": Scenario,
120
+ "model": LanguageModel,
121
+ "memory_plan": MemoryPlan,
122
+ "survey": Survey,
123
+ "cache": Cache,
124
+ }
125
+ d = {}
126
+ for attr, cls_ in attributes_to_classes.items():
127
+ if attr in data and data[attr] is not None:
128
+ if attr not in data:
129
+ d[attr] = {}
130
+ else:
131
+ d[attr] = cls_.from_dict(data[attr])
132
+
133
+ d["current_answers"] = data["current_answers"]
134
+ d["iteration"] = data["iteration"]
135
+ d["additional_prompt_data"] = data["additional_prompt_data"]
136
+
137
+ d = cls(**d)
138
+
139
+ def __repr__(self) -> str:
140
+ """Return a string representation of the Invigilator.
141
+
142
+ >>> InvigilatorBase.example().__repr__()
143
+ 'InvigilatorExample(...)'
144
+
145
+ """
146
+ return f"{self.__class__.__name__}(agent={repr(self.agent)}, question={repr(self.question)}, scneario={repr(self.scenario)}, model={repr(self.model)}, memory_plan={repr(self.memory_plan)}, current_answers={repr(self.current_answers)}, iteration{repr(self.iteration)}, additional_prompt_data={repr(self.additional_prompt_data)}, cache={repr(self.cache)})"
147
+
148
+ def get_failed_task_result(self, failure_reason: str) -> EDSLResultObjectInput:
149
+ """Return an AgentResponseDict used in case the question-asking fails.
150
+
151
+ Possible reasons include:
152
+ - Legimately skipped because of skip logic
153
+ - Failed to get response from the model
154
+
155
+ """
156
+ data = {
157
+ "answer": None,
158
+ "generated_tokens": None,
159
+ "comment": failure_reason,
160
+ "question_name": self.question.question_name,
161
+ "prompts": self.get_prompts(),
162
+ "cached_response": None,
163
+ "raw_model_response": None,
164
+ "cache_used": None,
165
+ "cache_key": None,
166
+ }
167
+ return EDSLResultObjectInput(**data)
168
+
169
+ def get_prompts(self) -> Dict[str, "Prompt"]:
170
+ """Return the prompt used."""
171
+ from edsl.prompts.Prompt import Prompt
172
+
173
+ return {
174
+ "user_prompt": Prompt("NA"),
175
+ "system_prompt": Prompt("NA"),
176
+ }
177
+
178
+ @abstractmethod
179
+ async def async_answer_question(self):
180
+ """Asnwer a question."""
181
+ pass
182
+
183
+ @jupyter_nb_handler
184
+ def answer_question(self) -> Coroutine:
185
+ """Return a function that gets the answers to the question."""
186
+
187
+ async def main():
188
+ """Return the answer to the question."""
189
+ results = await asyncio.gather(self.async_answer_question())
190
+ return results[0] # Since there's only one task, return its result
191
+
192
+ return main()
193
+
194
+ @classmethod
195
+ def example(
196
+ cls, throw_an_exception=False, question=None, scenario=None, survey=None
197
+ ) -> "InvigilatorBase":
198
+ """Return an example invigilator.
199
+
200
+ >>> InvigilatorBase.example()
201
+ InvigilatorExample(...)
202
+
203
+ >>> InvigilatorBase.example().answer_question()
204
+ {'message': [{'text': 'SPAM!'}], 'usage': {'prompt_tokens': 1, 'completion_tokens': 1}}
205
+
206
+ >>> InvigilatorBase.example(throw_an_exception=True).answer_question()
207
+ Traceback (most recent call last):
208
+ ...
209
+ Exception: This is a test error
210
+ """
211
+ from edsl.agents.Agent import Agent
212
+ from edsl.scenarios.Scenario import Scenario
213
+ from edsl.surveys.MemoryPlan import MemoryPlan
214
+ from edsl.language_models.model import Model
215
+ from edsl.surveys.Survey import Survey
216
+
217
+ model = Model("test", canned_response="SPAM!")
218
+
219
+ if throw_an_exception:
220
+ model.throw_exception = True
221
+ agent = Agent.example()
222
+
223
+ if not survey:
224
+ survey = Survey.example()
225
+
226
+ if question not in survey.questions and question is not None:
227
+ survey.add_question(question)
228
+
229
+ question = question or survey.questions[0]
230
+ scenario = scenario or Scenario.example()
231
+ memory_plan = MemoryPlan(survey=survey)
232
+ current_answers = None
233
+
234
+ class InvigilatorExample(cls):
235
+ """An example invigilator."""
236
+
237
+ async def async_answer_question(self):
238
+ """Answer a question."""
239
+ return await self.model.async_execute_model_call(
240
+ user_prompt="Hello", system_prompt="Hi"
241
+ )
242
+
243
+ return InvigilatorExample(
244
+ agent=agent,
245
+ question=question,
246
+ scenario=scenario,
247
+ survey=survey,
248
+ model=model,
249
+ memory_plan=memory_plan,
250
+ current_answers=current_answers,
251
+ )
252
+
253
+
254
+ if __name__ == "__main__":
255
+ import doctest
256
+
257
+ doctest.testmod(optionflags=doctest.ELLIPSIS)