simple_text_extract 3.0.3 → 3.0.5
Sign up to get free protection for your applications and to get access to all the features.
- checksums.yaml +4 -4
- data/.ruby-version +1 -1
- data/Gemfile.lock +22 -16
- data/bin/setup +1 -1
- data/lib/simple_text_extract/extract.rb +1 -1
- data/lib/simple_text_extract/version.rb +1 -1
- data/lib/simple_text_extract.rb +1 -1
- metadata +3 -3
checksums.yaml
CHANGED
@@ -1,7 +1,7 @@
|
|
1
1
|
---
|
2
2
|
SHA256:
|
3
|
-
metadata.gz:
|
4
|
-
data.tar.gz:
|
3
|
+
metadata.gz: 48a3c8805698e6f4af386789e4b4f11e9a3e38425bb1c2ce467ec7cefa99f2ba
|
4
|
+
data.tar.gz: f66a75c47984d63cf41b4f09f8c874f94da1f2928dddf268b4aac5ce48413f85
|
5
5
|
SHA512:
|
6
|
-
metadata.gz:
|
7
|
-
data.tar.gz:
|
6
|
+
metadata.gz: 5c8a892a0916945062f298c1b78728f62962ccca8404c9aabc7724db73b0806f00c1074bac7e13a94d390c7b67d19b7f5398ef96ea6ef41439ea66ba4bf78d3b
|
7
|
+
data.tar.gz: f0de6277f88c6debeaaac4bea385b3ba6e4d9180f53e5a2cef288ce5f80a30e92f67ad7e46c08aa908fb5ee30ccaffb1aaf0ea5ac4560cb2aa8262c8a610ce05
|
data/.ruby-version
CHANGED
@@ -1 +1 @@
|
|
1
|
-
3.
|
1
|
+
3.2.2
|
data/Gemfile.lock
CHANGED
@@ -1,7 +1,7 @@
|
|
1
1
|
PATH
|
2
2
|
remote: .
|
3
3
|
specs:
|
4
|
-
simple_text_extract (3.0.
|
4
|
+
simple_text_extract (3.0.5)
|
5
5
|
roo (~> 2.10.0)
|
6
6
|
rubyzip (~> 2.3.2)
|
7
7
|
spreadsheet (~> 1.3.0)
|
@@ -10,42 +10,47 @@ GEM
|
|
10
10
|
remote: https://rubygems.org/
|
11
11
|
specs:
|
12
12
|
ast (2.4.2)
|
13
|
+
base64 (0.1.1)
|
13
14
|
coderay (1.1.3)
|
14
15
|
json (2.6.3)
|
16
|
+
language_server-protocol (3.17.0.3)
|
15
17
|
memory_profiler (1.0.1)
|
16
18
|
method_source (1.0.0)
|
17
|
-
minitest (5.
|
18
|
-
mocha (2.0
|
19
|
+
minitest (5.20.0)
|
20
|
+
mocha (2.1.0)
|
19
21
|
ruby2_keywords (>= 0.0.5)
|
20
|
-
nokogiri (1.
|
22
|
+
nokogiri (1.15.4-arm64-darwin)
|
21
23
|
racc (~> 1.4)
|
22
|
-
nokogiri (1.
|
24
|
+
nokogiri (1.15.4-x86_64-linux)
|
23
25
|
racc (~> 1.4)
|
24
|
-
parallel (1.
|
25
|
-
parser (3.2.2.
|
26
|
+
parallel (1.23.0)
|
27
|
+
parser (3.2.2.4)
|
26
28
|
ast (~> 2.4.1)
|
29
|
+
racc
|
27
30
|
pry (0.14.2)
|
28
31
|
coderay (~> 1.1)
|
29
32
|
method_source (~> 1.0)
|
30
|
-
racc (1.
|
33
|
+
racc (1.7.1)
|
31
34
|
rainbow (3.1.1)
|
32
35
|
rake (13.0.6)
|
33
|
-
regexp_parser (2.
|
34
|
-
rexml (3.2.
|
36
|
+
regexp_parser (2.8.2)
|
37
|
+
rexml (3.2.6)
|
35
38
|
roo (2.10.0)
|
36
39
|
nokogiri (~> 1)
|
37
40
|
rubyzip (>= 1.3.0, < 3.0.0)
|
38
|
-
rubocop (1.
|
41
|
+
rubocop (1.57.1)
|
42
|
+
base64 (~> 0.1.1)
|
39
43
|
json (~> 2.3)
|
44
|
+
language_server-protocol (>= 3.17.0)
|
40
45
|
parallel (~> 1.10)
|
41
|
-
parser (>= 3.2.
|
46
|
+
parser (>= 3.2.2.4)
|
42
47
|
rainbow (>= 2.2.2, < 4.0)
|
43
48
|
regexp_parser (>= 1.8, < 3.0)
|
44
49
|
rexml (>= 3.2.5, < 4.0)
|
45
|
-
rubocop-ast (>= 1.28.
|
50
|
+
rubocop-ast (>= 1.28.1, < 2.0)
|
46
51
|
ruby-progressbar (~> 1.7)
|
47
52
|
unicode-display_width (>= 2.4.0, < 3.0)
|
48
|
-
rubocop-ast (1.
|
53
|
+
rubocop-ast (1.29.0)
|
49
54
|
parser (>= 3.2.1.0)
|
50
55
|
ruby-ole (1.2.12.2)
|
51
56
|
ruby-progressbar (1.13.0)
|
@@ -53,10 +58,11 @@ GEM
|
|
53
58
|
rubyzip (2.3.2)
|
54
59
|
spreadsheet (1.3.0)
|
55
60
|
ruby-ole
|
56
|
-
unicode-display_width (2.
|
61
|
+
unicode-display_width (2.5.0)
|
57
62
|
|
58
63
|
PLATFORMS
|
59
64
|
arm64-darwin-21
|
65
|
+
arm64-darwin-23
|
60
66
|
x86_64-linux
|
61
67
|
|
62
68
|
DEPENDENCIES
|
@@ -69,4 +75,4 @@ DEPENDENCIES
|
|
69
75
|
simple_text_extract!
|
70
76
|
|
71
77
|
BUNDLED WITH
|
72
|
-
2.4.
|
78
|
+
2.4.10
|
data/bin/setup
CHANGED
@@ -28,7 +28,7 @@ class SimpleTextExtract::Extract # rubocop:disable Metrics/ClassLength
|
|
28
28
|
end
|
29
29
|
|
30
30
|
def to_s
|
31
|
-
@to_s ||= extract.to_s.gsub(/[^\S\n]+/, " ").gsub(/\s?\n\s+/, "\n").strip
|
31
|
+
@to_s ||= extract.to_s.scrub.gsub(/[^\S\n]+/, " ").gsub(/\s?\n\s+/, "\n").strip
|
32
32
|
end
|
33
33
|
|
34
34
|
private
|
data/lib/simple_text_extract.rb
CHANGED
metadata
CHANGED
@@ -1,14 +1,14 @@
|
|
1
1
|
--- !ruby/object:Gem::Specification
|
2
2
|
name: simple_text_extract
|
3
3
|
version: !ruby/object:Gem::Version
|
4
|
-
version: 3.0.
|
4
|
+
version: 3.0.5
|
5
5
|
platform: ruby
|
6
6
|
authors:
|
7
7
|
- Nick Weiland
|
8
8
|
autorequire:
|
9
9
|
bindir: exe
|
10
10
|
cert_chain: []
|
11
|
-
date: 2023-
|
11
|
+
date: 2023-10-23 00:00:00.000000000 Z
|
12
12
|
dependencies:
|
13
13
|
- !ruby/object:Gem::Dependency
|
14
14
|
name: roo
|
@@ -97,7 +97,7 @@ required_rubygems_version: !ruby/object:Gem::Requirement
|
|
97
97
|
requirements:
|
98
98
|
- antiword
|
99
99
|
- pdftotext/poppler
|
100
|
-
rubygems_version: 3.
|
100
|
+
rubygems_version: 3.4.10
|
101
101
|
signing_key:
|
102
102
|
specification_version: 4
|
103
103
|
summary: Extract text from various file types before resorting to an OCR solution.
|