dcatoolkit 0.1.7__tar.gz → 0.1.9__tar.gz

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -1,21 +1,21 @@
1
- MIT License
2
-
3
- Copyright (c) 2024 Raheel Syed Ahmed
4
-
5
- Permission is hereby granted, free of charge, to any person obtaining a copy
6
- of this software and associated documentation files (the "Software"), to deal
7
- in the Software without restriction, including without limitation the rights
8
- to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
9
- copies of the Software, and to permit persons to whom the Software is
10
- furnished to do so, subject to the following conditions:
11
-
12
- The above copyright notice and this permission notice shall be included in all
13
- copies or substantial portions of the Software.
14
-
15
- THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
16
- IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
17
- FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
18
- AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
19
- LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
20
- OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE
21
- SOFTWARE.
1
+ MIT License
2
+
3
+ Copyright (c) 2024 Raheel Syed Ahmed
4
+
5
+ Permission is hereby granted, free of charge, to any person obtaining a copy
6
+ of this software and associated documentation files (the "Software"), to deal
7
+ in the Software without restriction, including without limitation the rights
8
+ to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
9
+ copies of the Software, and to permit persons to whom the Software is
10
+ furnished to do so, subject to the following conditions:
11
+
12
+ The above copyright notice and this permission notice shall be included in all
13
+ copies or substantial portions of the Software.
14
+
15
+ THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
16
+ IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
17
+ FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
18
+ AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
19
+ LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
20
+ OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE
21
+ SOFTWARE.
@@ -1,72 +1,72 @@
1
- Metadata-Version: 2.1
2
- Name: dcatoolkit
3
- Version: 0.1.7
4
- Summary: Collection of useful modules and representations for managing DCA output data.
5
- Author-email: Raheel Syed Ahmed <raheelsyedahmed@gmail.com>
6
- Maintainer-email: Raheel Syed Ahmed <raheelsyedahmed@gmail.com>
7
- License: MIT License
8
-
9
- Copyright (c) 2024 Raheel Syed Ahmed
10
-
11
- Permission is hereby granted, free of charge, to any person obtaining a copy
12
- of this software and associated documentation files (the "Software"), to deal
13
- in the Software without restriction, including without limitation the rights
14
- to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
15
- copies of the Software, and to permit persons to whom the Software is
16
- furnished to do so, subject to the following conditions:
17
-
18
- The above copyright notice and this permission notice shall be included in all
19
- copies or substantial portions of the Software.
20
-
21
- THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
22
- IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
23
- FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
24
- AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
25
- LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
26
- OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE
27
- SOFTWARE.
28
-
29
- Keywords: dca,toolkit,DI,coevolution
30
- Classifier: Development Status :: 4 - Beta
31
- Classifier: Intended Audience :: Science/Research
32
- Classifier: Topic :: Scientific/Engineering :: Bio-Informatics
33
- Classifier: License :: OSI Approved :: MIT License
34
- Classifier: Programming Language :: Python :: 3
35
- Classifier: Programming Language :: Python :: 3.10
36
- Classifier: Programming Language :: Python :: 3.11
37
- Classifier: Programming Language :: Python :: 3.12
38
- Requires-Python: >=3.10
39
- Description-Content-Type: text/markdown
40
- License-File: LICENSE
41
- Requires-Dist: biotite
42
- Requires-Dist: matplotlib>=3.8.0
43
- Requires-Dist: numpy>=1.26.0
44
- Requires-Dist: pandas>=2.1.0
45
- Requires-Dist: scikit-learn>=1.3
46
- Requires-Dist: scipy>=1.11.0
47
- Provides-Extra: tests
48
- Requires-Dist: pytest; extra == "tests"
49
- Provides-Extra: docs
50
- Requires-Dist: sphinx; extra == "docs"
51
- Requires-Dist: pdoc; extra == "docs"
52
- Requires-Dist: numpydoc; extra == "docs"
53
- Provides-Extra: lint
54
- Requires-Dist: ruffle; extra == "lint"
55
-
56
- # dcatoolkit
57
- Collection of useful modules and representations for managing DCA output data.
58
-
59
- ## Major Sections
60
- ### Representations
61
- * Use Pairs to load lists, tuples, sets, and ndarrays with the correct orientation of elements. This will allow you to yield integer pairs that can be mirrored (where y becomes x and vice versa) and to subset various pairs.
62
- * Use DirectInformationData to create 3-column structured ndarrays that can be sorted by "DI", mapped to a protein with a ResidueAlignment, and used to generate output for other programs (including UCSF Chimera)
63
- * Use ResidueAlignment to generate a reference map. Indices of one sequence of characters can be linked to their corresponding indices of the other sequence of characters. The dictionaries produced, domain-to-protein and protein-to-domain, allow for forward mapping and backmapping.
64
- * Use StructureInformation to find contacts in a protein structure and find atomic information related to specific pairs of interest.
65
- ### Analytics
66
- * Use MSATools to load in Multiple Sequence Alignment (MSA) data and provide functionality including generating frequency statistics on "gappiness" in the MSA and filtering and cleaning MSAs.
67
-
68
-
69
- ## Diagram of Hidden Markov Machine & Direct Coupling Analysis Pipeline
70
- <p align="center">
71
- <img src="https://github.com/user-attachments/assets/4768e08f-d513-4dbf-abc5-c80c1b3d42aa"/>
72
- </p>
1
+ Metadata-Version: 2.1
2
+ Name: dcatoolkit
3
+ Version: 0.1.9
4
+ Summary: Collection of useful modules and representations for managing DCA output data.
5
+ Author-email: Raheel Syed Ahmed <raheelsyedahmed@gmail.com>
6
+ Maintainer-email: Raheel Syed Ahmed <raheelsyedahmed@gmail.com>
7
+ License: MIT License
8
+
9
+ Copyright (c) 2024 Raheel Syed Ahmed
10
+
11
+ Permission is hereby granted, free of charge, to any person obtaining a copy
12
+ of this software and associated documentation files (the "Software"), to deal
13
+ in the Software without restriction, including without limitation the rights
14
+ to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
15
+ copies of the Software, and to permit persons to whom the Software is
16
+ furnished to do so, subject to the following conditions:
17
+
18
+ The above copyright notice and this permission notice shall be included in all
19
+ copies or substantial portions of the Software.
20
+
21
+ THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
22
+ IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
23
+ FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
24
+ AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
25
+ LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
26
+ OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE
27
+ SOFTWARE.
28
+
29
+ Keywords: dca,toolkit,DI,coevolution
30
+ Classifier: Development Status :: 4 - Beta
31
+ Classifier: Intended Audience :: Science/Research
32
+ Classifier: Topic :: Scientific/Engineering :: Bio-Informatics
33
+ Classifier: License :: OSI Approved :: MIT License
34
+ Classifier: Programming Language :: Python :: 3
35
+ Classifier: Programming Language :: Python :: 3.10
36
+ Classifier: Programming Language :: Python :: 3.11
37
+ Classifier: Programming Language :: Python :: 3.12
38
+ Requires-Python: >=3.10
39
+ Description-Content-Type: text/markdown
40
+ License-File: LICENSE
41
+ Requires-Dist: biotite
42
+ Requires-Dist: matplotlib>=3.8.0
43
+ Requires-Dist: numpy>=1.26.0
44
+ Requires-Dist: pandas>=2.1.0
45
+ Requires-Dist: scikit-learn>=1.3
46
+ Requires-Dist: scipy>=1.11.0
47
+ Provides-Extra: tests
48
+ Requires-Dist: pytest; extra == "tests"
49
+ Provides-Extra: docs
50
+ Requires-Dist: sphinx; extra == "docs"
51
+ Requires-Dist: pdoc; extra == "docs"
52
+ Requires-Dist: numpydoc; extra == "docs"
53
+ Provides-Extra: lint
54
+ Requires-Dist: ruffle; extra == "lint"
55
+
56
+ # dcatoolkit
57
+ Collection of useful modules and representations for managing DCA output data.
58
+
59
+ ## Major Sections
60
+ ### Representations
61
+ * Use Pairs to load lists, tuples, sets, and ndarrays with the correct orientation of elements. This will allow you to yield integer pairs that can be mirrored (where y becomes x and vice versa) and to subset various pairs.
62
+ * Use DirectInformationData to create 3-column structured ndarrays that can be sorted by "DI", mapped to a protein with a ResidueAlignment, and used to generate output for other programs (including UCSF Chimera)
63
+ * Use ResidueAlignment to generate a reference map. Indices of one sequence of characters can be linked to their corresponding indices of the other sequence of characters. The dictionaries produced, domain-to-protein and protein-to-domain, allow for forward mapping and backmapping.
64
+ * Use StructureInformation to find contacts in a protein structure and find atomic information related to specific pairs of interest.
65
+ ### Analytics
66
+ * Use MSATools to load in Multiple Sequence Alignment (MSA) data and provide functionality including generating frequency statistics on "gappiness" in the MSA and filtering and cleaning MSAs.
67
+
68
+
69
+ ## Diagram of Hidden Markov Machine & Direct Coupling Analysis Pipeline
70
+ <p align="center">
71
+ <img src="https://github.com/user-attachments/assets/4768e08f-d513-4dbf-abc5-c80c1b3d42aa"/>
72
+ </p>
@@ -1,17 +1,17 @@
1
- # dcatoolkit
2
- Collection of useful modules and representations for managing DCA output data.
3
-
4
- ## Major Sections
5
- ### Representations
6
- * Use Pairs to load lists, tuples, sets, and ndarrays with the correct orientation of elements. This will allow you to yield integer pairs that can be mirrored (where y becomes x and vice versa) and to subset various pairs.
7
- * Use DirectInformationData to create 3-column structured ndarrays that can be sorted by "DI", mapped to a protein with a ResidueAlignment, and used to generate output for other programs (including UCSF Chimera)
8
- * Use ResidueAlignment to generate a reference map. Indices of one sequence of characters can be linked to their corresponding indices of the other sequence of characters. The dictionaries produced, domain-to-protein and protein-to-domain, allow for forward mapping and backmapping.
9
- * Use StructureInformation to find contacts in a protein structure and find atomic information related to specific pairs of interest.
10
- ### Analytics
11
- * Use MSATools to load in Multiple Sequence Alignment (MSA) data and provide functionality including generating frequency statistics on "gappiness" in the MSA and filtering and cleaning MSAs.
12
-
13
-
14
- ## Diagram of Hidden Markov Machine & Direct Coupling Analysis Pipeline
15
- <p align="center">
16
- <img src="https://github.com/user-attachments/assets/4768e08f-d513-4dbf-abc5-c80c1b3d42aa"/>
17
- </p>
1
+ # dcatoolkit
2
+ Collection of useful modules and representations for managing DCA output data.
3
+
4
+ ## Major Sections
5
+ ### Representations
6
+ * Use Pairs to load lists, tuples, sets, and ndarrays with the correct orientation of elements. This will allow you to yield integer pairs that can be mirrored (where y becomes x and vice versa) and to subset various pairs.
7
+ * Use DirectInformationData to create 3-column structured ndarrays that can be sorted by "DI", mapped to a protein with a ResidueAlignment, and used to generate output for other programs (including UCSF Chimera)
8
+ * Use ResidueAlignment to generate a reference map. Indices of one sequence of characters can be linked to their corresponding indices of the other sequence of characters. The dictionaries produced, domain-to-protein and protein-to-domain, allow for forward mapping and backmapping.
9
+ * Use StructureInformation to find contacts in a protein structure and find atomic information related to specific pairs of interest.
10
+ ### Analytics
11
+ * Use MSATools to load in Multiple Sequence Alignment (MSA) data and provide functionality including generating frequency statistics on "gappiness" in the MSA and filtering and cleaning MSAs.
12
+
13
+
14
+ ## Diagram of Hidden Markov Machine & Direct Coupling Analysis Pipeline
15
+ <p align="center">
16
+ <img src="https://github.com/user-attachments/assets/4768e08f-d513-4dbf-abc5-c80c1b3d42aa"/>
17
+ </p>
@@ -1,54 +1,54 @@
1
- [build-system]
2
- requires = ["setuptools >= 61.0"]
3
- build-backend = "setuptools.build_meta"
4
-
5
- [project]
6
- name = "dcatoolkit"
7
- version = "0.1.7"
8
- description = "Collection of useful modules and representations for managing DCA output data."
9
- keywords = ["dca", "toolkit", "DI", "coevolution"]
10
-
11
- readme = "README.md"
12
- license = {file = "LICENSE"}
13
-
14
- requires-python = ">=3.10"
15
-
16
- authors = [
17
- {name = "Raheel Syed Ahmed", email = "raheelsyedahmed@gmail.com"}
18
- ]
19
- maintainers = [
20
- {name = "Raheel Syed Ahmed", email = "raheelsyedahmed@gmail.com"}
21
- ]
22
-
23
- dependencies = [
24
- "biotite",
25
- "matplotlib>=3.8.0",
26
- "numpy>=1.26.0",
27
- "pandas>=2.1.0",
28
- "scikit-learn>=1.3",
29
- "scipy>=1.11.0",
30
- ]
31
-
32
- classifiers = [
33
- "Development Status :: 4 - Beta",
34
- "Intended Audience :: Science/Research",
35
- "Topic :: Scientific/Engineering :: Bio-Informatics",
36
- "License :: OSI Approved :: MIT License",
37
- "Programming Language :: Python :: 3",
38
- "Programming Language :: Python :: 3.10",
39
- "Programming Language :: Python :: 3.11",
40
- "Programming Language :: Python :: 3.12",
41
- ]
42
-
43
- [project.optional-dependencies]
44
- tests = [
45
- "pytest",
46
- ]
47
- docs = [
48
- "sphinx",
49
- "pdoc",
50
- "numpydoc"
51
- ]
52
- lint = [
53
- "ruffle",
1
+ [build-system]
2
+ requires = ["setuptools >= 61.0"]
3
+ build-backend = "setuptools.build_meta"
4
+
5
+ [project]
6
+ name = "dcatoolkit"
7
+ version = "0.1.9"
8
+ description = "Collection of useful modules and representations for managing DCA output data."
9
+ keywords = ["dca", "toolkit", "DI", "coevolution"]
10
+
11
+ readme = "README.md"
12
+ license = {file = "LICENSE"}
13
+
14
+ requires-python = ">=3.10"
15
+
16
+ authors = [
17
+ {name = "Raheel Syed Ahmed", email = "raheelsyedahmed@gmail.com"}
18
+ ]
19
+ maintainers = [
20
+ {name = "Raheel Syed Ahmed", email = "raheelsyedahmed@gmail.com"}
21
+ ]
22
+
23
+ dependencies = [
24
+ "biotite",
25
+ "matplotlib>=3.8.0",
26
+ "numpy>=1.26.0",
27
+ "pandas>=2.1.0",
28
+ "scikit-learn>=1.3",
29
+ "scipy>=1.11.0",
30
+ ]
31
+
32
+ classifiers = [
33
+ "Development Status :: 4 - Beta",
34
+ "Intended Audience :: Science/Research",
35
+ "Topic :: Scientific/Engineering :: Bio-Informatics",
36
+ "License :: OSI Approved :: MIT License",
37
+ "Programming Language :: Python :: 3",
38
+ "Programming Language :: Python :: 3.10",
39
+ "Programming Language :: Python :: 3.11",
40
+ "Programming Language :: Python :: 3.12",
41
+ ]
42
+
43
+ [project.optional-dependencies]
44
+ tests = [
45
+ "pytest",
46
+ ]
47
+ docs = [
48
+ "sphinx",
49
+ "pdoc",
50
+ "numpydoc"
51
+ ]
52
+ lint = [
53
+ "ruffle",
54
54
  ]
@@ -1,4 +1,4 @@
1
- [egg_info]
2
- tag_build =
3
- tag_date = 0
4
-
1
+ [egg_info]
2
+ tag_build =
3
+ tag_date = 0
4
+
@@ -1,4 +1,4 @@
1
-
2
- __version__ = "0.1.7"
3
- from .representation import Pairs, DirectInformationData, StructureInformation, ResidueAlignment
1
+
2
+ __version__ = "0.1.9"
3
+ from .representation import Pairs, DirectInformationData, StructureInformation, ResidueAlignment
4
4
  from .analytics import MSATools