graphed_fuzzy_search 0.1.0

Sign up to get free protection for your applications and to get access to all the features.
checksums.yaml ADDED
@@ -0,0 +1,7 @@
1
+ ---
2
+ SHA256:
3
+ metadata.gz: cd030d3d31c4838402d85fb1cb7ec8e65f641afca3f9c395294e14192bb74938
4
+ data.tar.gz: 454a974ca7e1c1c7a479b881024227128d3ab32470be050ef1777bc30d02e653
5
+ SHA512:
6
+ metadata.gz: 220d763cff6b305ff3627d591c8f99f1242f3ac761c15651bbfbaa21d2fac568777ede19c9d37b4bf4f76884a13ea4e3576076a1d65da9806c829cd6d3e7006d
7
+ data.tar.gz: 52a98f177ea3bcf4688a390c922b8babedb732b154ae06355d847bf35ad9a54e805cc10ff262e86aa04fa13325d5d01051a7ee8896728c134a560ff1e7483d88
data/.gitignore ADDED
@@ -0,0 +1,11 @@
1
+ /.bundle/
2
+ /.yardoc
3
+ /_yardoc/
4
+ /coverage/
5
+ /doc/
6
+ /pkg/
7
+ /spec/reports/
8
+ /tmp/
9
+
10
+ # rspec failure tracking
11
+ .rspec_status
data/.rspec ADDED
@@ -0,0 +1,3 @@
1
+ --format documentation
2
+ --color
3
+ --require spec_helper
data/.travis.yml ADDED
@@ -0,0 +1,5 @@
1
+ sudo: false
2
+ language: ruby
3
+ rvm:
4
+ - 2.6.0
5
+ before_install: gem install bundler -v 1.16.1
data/Gemfile ADDED
@@ -0,0 +1,6 @@
1
+ source "https://rubygems.org"
2
+
3
+ git_source(:github) {|repo_name| "https://github.com/#{repo_name}" }
4
+
5
+ # Specify your gem's dependencies in graphed_fuzzy_search.gemspec
6
+ gemspec
data/Gemfile.lock ADDED
@@ -0,0 +1,35 @@
1
+ PATH
2
+ remote: .
3
+ specs:
4
+ graphed_fuzzy_search (0.1.0)
5
+
6
+ GEM
7
+ remote: https://rubygems.org/
8
+ specs:
9
+ diff-lcs (1.3)
10
+ rake (12.3.0)
11
+ rspec (3.7.0)
12
+ rspec-core (~> 3.7.0)
13
+ rspec-expectations (~> 3.7.0)
14
+ rspec-mocks (~> 3.7.0)
15
+ rspec-core (3.7.0)
16
+ rspec-support (~> 3.7.0)
17
+ rspec-expectations (3.7.0)
18
+ diff-lcs (>= 1.2.0, < 2.0)
19
+ rspec-support (~> 3.7.0)
20
+ rspec-mocks (3.7.0)
21
+ diff-lcs (>= 1.2.0, < 2.0)
22
+ rspec-support (~> 3.7.0)
23
+ rspec-support (3.7.0)
24
+
25
+ PLATFORMS
26
+ ruby
27
+
28
+ DEPENDENCIES
29
+ bundler
30
+ graphed_fuzzy_search!
31
+ rake
32
+ rspec (~> 3.0)
33
+
34
+ BUNDLED WITH
35
+ 1.16.1
data/LICENSE.txt ADDED
@@ -0,0 +1,21 @@
1
+ The MIT License (MIT)
2
+
3
+ Copyright (c) 2018 Sorah Fukumori
4
+
5
+ Permission is hereby granted, free of charge, to any person obtaining a copy
6
+ of this software and associated documentation files (the "Software"), to deal
7
+ in the Software without restriction, including without limitation the rights
8
+ to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
9
+ copies of the Software, and to permit persons to whom the Software is
10
+ furnished to do so, subject to the following conditions:
11
+
12
+ The above copyright notice and this permission notice shall be included in
13
+ all copies or substantial portions of the Software.
14
+
15
+ THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
16
+ IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
17
+ FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
18
+ AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
19
+ LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
20
+ OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN
21
+ THE SOFTWARE.
data/README.md ADDED
@@ -0,0 +1,103 @@
1
+ # GraphedFuzzySearch: filter items like Slack switcher, Atom command palette can
2
+
3
+
4
+ ``` ruby
5
+ s = GraphedFuzzySearch.new(%w(john-appleseed john-doe jonathan-doe alice-eve eve-doe))
6
+ p s.query('d') #=> ["john-doe", "jonathan-doe", "eve-doe"]
7
+ p s.query('dj') #=> ["john-doe", "jonathan-doe"]
8
+ p s.query('djoh') #=> ["john-doe"]
9
+ p s.query('a') #=> ["alice-eve", "john-appleseed"]
10
+ ```
11
+
12
+ ## Installation
13
+
14
+ Add this line to your application's Gemfile:
15
+
16
+ ```ruby
17
+ gem 'graphed_fuzzy_search'
18
+ ```
19
+
20
+ And then execute:
21
+
22
+ $ bundle
23
+
24
+ Or install it yourself as:
25
+
26
+ $ gem install graphed_fuzzy_search
27
+
28
+ ## Usage
29
+
30
+ ### Basic
31
+
32
+ ``` ruby
33
+ s = GraphedFuzzySearch.new(%w(john-appleseed john-doe jonathan-doe alice-eve eve-doe))
34
+ p s.query('d') #=> ["john-doe", "jonathan-doe", "eve-doe"]
35
+ p s.query('dj') #=> ["john-doe", "jonathan-doe"]
36
+ p s.query('djoh') #=> ["john-doe"]
37
+ p s.query('a') #=> ["alice-eve", "john-appleseed"]
38
+ ```
39
+
40
+ ### Multiple attributes
41
+
42
+ ``` ruby
43
+ Item = Struct.new(:name, :email)
44
+ items = [Item.new('John Doe', 'john-doe@example.com')]
45
+ s = GraphedFuzzySearch.new(items, attributes: %i(name email))
46
+ # Hint: default +attributes:+ is %i(name)
47
+ ```
48
+
49
+ ### Custom tokenize
50
+
51
+ ``` ruby
52
+ p GraphedFuzzySearch.new(["token_one token_two"]).query('one') #=> ["token_one token_two"]
53
+ p GraphedFuzzySearch.new(["token_one token_two"], token_regex: /\w+/).query('one') #=> []
54
+ ```
55
+
56
+ ## Internal
57
+
58
+ - index
59
+ 1. Split into tokens (substring determined by `token_regex` defaults to `/[^\p{Punct}\p{Space}]+/`)
60
+ 2. construct a trie for tokens. each trie is independent for each item.
61
+ - query
62
+ 1. walk all tries by every character of a given query.
63
+
64
+ ```
65
+ Items:
66
+ ax-by
67
+ cy-ax
68
+ by
69
+ bye
70
+
71
+ Tokens:
72
+ ax by
73
+ cy ax
74
+ by
75
+ bye
76
+
77
+ Tries:
78
+ a->x b->y, a->b x->b, b->a y->a
79
+ c->y a->x, c->a y->a, a->c x->c
80
+ b->y
81
+ b->y->e
82
+ ```
83
+
84
+ ## Plans
85
+
86
+ Pull Requests are welcomed.
87
+
88
+ - Dump/Load an index (tree)
89
+ - Compressing an index
90
+
91
+ ## Development
92
+
93
+ After checking out the repo, run `bin/setup` to install dependencies. Then, run `rake spec` to run the tests. You can also run `bin/console` for an interactive prompt that will allow you to experiment.
94
+
95
+ To install this gem onto your local machine, run `bundle exec rake install`. To release a new version, update the version number in `version.rb`, and then run `bundle exec rake release`, which will create a git tag for the version, push git commits and tags, and push the `.gem` file to [rubygems.org](https://rubygems.org).
96
+
97
+ ## Contributing
98
+
99
+ Bug reports and pull requests are welcome on GitHub at https://github.com/sorah/graphed_fuzzy_search.
100
+
101
+ ## License
102
+
103
+ The gem is available as open source under the terms of the [MIT License](https://opensource.org/licenses/MIT).
data/Rakefile ADDED
@@ -0,0 +1,6 @@
1
+ require "bundler/gem_tasks"
2
+ require "rspec/core/rake_task"
3
+
4
+ RSpec::Core::RakeTask.new(:spec)
5
+
6
+ task :default => :spec
data/bin/console ADDED
@@ -0,0 +1,14 @@
1
+ #!/usr/bin/env ruby
2
+
3
+ require "bundler/setup"
4
+ require "graphed_fuzzy_search"
5
+
6
+ # You can add fixtures and/or initialization code here to make experimenting
7
+ # with your gem easier. You can also use a different console, if you like.
8
+
9
+ # (If you use this, don't forget to add pry to your Gemfile!)
10
+ # require "pry"
11
+ # Pry.start
12
+
13
+ require "irb"
14
+ IRB.start(__FILE__)
data/bin/setup ADDED
@@ -0,0 +1,8 @@
1
+ #!/usr/bin/env bash
2
+ set -euo pipefail
3
+ IFS=$'\n\t'
4
+ set -vx
5
+
6
+ bundle install
7
+
8
+ # Do any other automated setup that you need to do here
@@ -0,0 +1,26 @@
1
+
2
+ lib = File.expand_path("../lib", __FILE__)
3
+ $LOAD_PATH.unshift(lib) unless $LOAD_PATH.include?(lib)
4
+ require "graphed_fuzzy_search/version"
5
+
6
+ Gem::Specification.new do |spec|
7
+ spec.name = "graphed_fuzzy_search"
8
+ spec.version = GraphedFuzzySearch::VERSION
9
+ spec.authors = ["Sorah Fukumori"]
10
+ spec.email = ["sorah@cookpad.com"]
11
+
12
+ spec.summary = %q{filter items like Slack switcher, Atom command palette}
13
+ spec.homepage = "https://github.com/sorah/graphed_fuzzy_search"
14
+ spec.license = "MIT"
15
+
16
+ spec.files = `git ls-files -z`.split("\x0").reject do |f|
17
+ f.match(%r{^(test|spec|features)/})
18
+ end
19
+ spec.bindir = "exe"
20
+ spec.executables = spec.files.grep(%r{^exe/}) { |f| File.basename(f) }
21
+ spec.require_paths = ["lib"]
22
+
23
+ spec.add_development_dependency "bundler"
24
+ spec.add_development_dependency "rake"
25
+ spec.add_development_dependency "rspec", "~> 3.0"
26
+ end
@@ -0,0 +1,131 @@
1
+ #require "graphed_fuzzy_search/version"
2
+
3
+ module GraphedFuzzySearch
4
+ def self.new(*args, **kwargs)
5
+ Collection.new(*args, **kwargs)
6
+ end
7
+
8
+ DEFAULT_TOKEN_REGEX = /[^\p{Punct}\p{Space}]+/
9
+
10
+ class Collection
11
+ def initialize(objects, attributes: [:to_s], token_regex: DEFAULT_TOKEN_REGEX, normal_weight: 1, different_token_weight: 5)
12
+ @objects = objects
13
+ @attributes = attributes
14
+ @token_regex = token_regex
15
+ @normal_weight = normal_weight
16
+ @different_token_weight = different_token_weight
17
+ trees
18
+ end
19
+
20
+ def inspect
21
+ "#<#{self.class.name}>"
22
+ end
23
+
24
+ attr_reader :objects, :attributes, :token_regex
25
+ attr_reader :normal_weight, :different_token_weight
26
+
27
+ def query(*args, **kwargs)
28
+ query_raw(*args, **kwargs).map{ |(node, _)| node.item.object }
29
+ end
30
+
31
+ def query_raw(str, max_scan: str.size)
32
+ str = str.downcase
33
+ chars = str.chars
34
+ i = 0
35
+ needles = trees.map { |root| [root, 0] }
36
+ while i < chars.size && i < max_scan && !needles.empty?
37
+ char = chars[i]
38
+ needles.map! { |(node, weight)| [node[char], weight] }
39
+ needles.select!(&:first)
40
+ needles.map! { |(conn, weight)| [conn.node, weight + conn.weight] }
41
+ i += 1
42
+ end
43
+ needles.map! { |(conn, weight)|
44
+ if conn.item.key.start_with?(str)
45
+ weight *= 0.7
46
+ end
47
+ [conn, weight]
48
+ }
49
+ needles.sort_by(&:last)
50
+ end
51
+
52
+ def items
53
+ @items ||= objects.map do |obj|
54
+ attrs = attributes.flat_map do |k|
55
+ [*obj.send(k)]
56
+ end
57
+ tokens = attrs.flat_map do |attr|
58
+ attr.to_s.downcase.scan(token_regex)
59
+ end
60
+ Item.new(obj, attrs[0], tokens)
61
+ end
62
+ end
63
+
64
+ def trees
65
+ @trees ||= items.map do |item|
66
+ root = Node.new(nil, [], item)
67
+ token_and_heads = item.tokens.map do |token|
68
+ root.mine(normal_weight, token)
69
+ [token, root[token[0]].node]
70
+ end
71
+ token_and_heads.each do |(_, ah)|
72
+ token_and_heads.each do |(bt, _)|
73
+ root.walk(bt.each_char) do |n|
74
+ n.connect(different_token_weight, ah)
75
+ end
76
+ end
77
+ end
78
+ root
79
+ end
80
+ end
81
+ end
82
+
83
+ Item = Struct.new(:object, :key, :tokens)
84
+ Connection = Struct.new(:weight, :node)
85
+
86
+ Node = Struct.new(:str, :adjacents, :item) do
87
+ def inspect
88
+ "#<#{self.class.name} @item=#{item.inspect} @str=#{str.inspect} adjacents=#{adjacents.map(&:first).join}>"
89
+ end
90
+ def head
91
+ str[0]
92
+ end
93
+
94
+ def length
95
+ str.size
96
+ end
97
+
98
+ def [](k)
99
+ self.adjacents ||= []
100
+ _, adjacent = adjacents.find { |c, _| c == k }
101
+ adjacent
102
+ end
103
+
104
+ def walk(enum)
105
+ weight = 0
106
+ enum.inject(self) do |r, i|
107
+ c = r[i]
108
+ raise unless c
109
+ weight += c.weight
110
+ yield c.node
111
+ c.node
112
+ end
113
+ weight
114
+ end
115
+
116
+ def mine(weight, str)
117
+ str.each_char.inject(self) do |r, i|
118
+ r.connect(weight, Node.new(i, [], self.item))
119
+ end
120
+ end
121
+
122
+ def connect(weight, node)
123
+ return nil if node.__id__ == self.__id__
124
+ self.adjacents ||= []
125
+ connection = self[node.str]
126
+ return connection.node if connection
127
+ adjacents << [node.head, Connection.new(weight, node)]
128
+ node
129
+ end
130
+ end
131
+ end
@@ -0,0 +1,3 @@
1
+ module GraphedFuzzySearch
2
+ VERSION = "0.1.0"
3
+ end
metadata ADDED
@@ -0,0 +1,99 @@
1
+ --- !ruby/object:Gem::Specification
2
+ name: graphed_fuzzy_search
3
+ version: !ruby/object:Gem::Version
4
+ version: 0.1.0
5
+ platform: ruby
6
+ authors:
7
+ - Sorah Fukumori
8
+ autorequire:
9
+ bindir: exe
10
+ cert_chain: []
11
+ date: 2018-01-04 00:00:00.000000000 Z
12
+ dependencies:
13
+ - !ruby/object:Gem::Dependency
14
+ name: bundler
15
+ requirement: !ruby/object:Gem::Requirement
16
+ requirements:
17
+ - - ">="
18
+ - !ruby/object:Gem::Version
19
+ version: '0'
20
+ type: :development
21
+ prerelease: false
22
+ version_requirements: !ruby/object:Gem::Requirement
23
+ requirements:
24
+ - - ">="
25
+ - !ruby/object:Gem::Version
26
+ version: '0'
27
+ - !ruby/object:Gem::Dependency
28
+ name: rake
29
+ requirement: !ruby/object:Gem::Requirement
30
+ requirements:
31
+ - - ">="
32
+ - !ruby/object:Gem::Version
33
+ version: '0'
34
+ type: :development
35
+ prerelease: false
36
+ version_requirements: !ruby/object:Gem::Requirement
37
+ requirements:
38
+ - - ">="
39
+ - !ruby/object:Gem::Version
40
+ version: '0'
41
+ - !ruby/object:Gem::Dependency
42
+ name: rspec
43
+ requirement: !ruby/object:Gem::Requirement
44
+ requirements:
45
+ - - "~>"
46
+ - !ruby/object:Gem::Version
47
+ version: '3.0'
48
+ type: :development
49
+ prerelease: false
50
+ version_requirements: !ruby/object:Gem::Requirement
51
+ requirements:
52
+ - - "~>"
53
+ - !ruby/object:Gem::Version
54
+ version: '3.0'
55
+ description:
56
+ email:
57
+ - sorah@cookpad.com
58
+ executables: []
59
+ extensions: []
60
+ extra_rdoc_files: []
61
+ files:
62
+ - ".gitignore"
63
+ - ".rspec"
64
+ - ".travis.yml"
65
+ - Gemfile
66
+ - Gemfile.lock
67
+ - LICENSE.txt
68
+ - README.md
69
+ - Rakefile
70
+ - bin/console
71
+ - bin/setup
72
+ - graphed_fuzzy_search.gemspec
73
+ - lib/graphed_fuzzy_search.rb
74
+ - lib/graphed_fuzzy_search/version.rb
75
+ homepage: https://github.com/sorah/graphed_fuzzy_search
76
+ licenses:
77
+ - MIT
78
+ metadata: {}
79
+ post_install_message:
80
+ rdoc_options: []
81
+ require_paths:
82
+ - lib
83
+ required_ruby_version: !ruby/object:Gem::Requirement
84
+ requirements:
85
+ - - ">="
86
+ - !ruby/object:Gem::Version
87
+ version: '0'
88
+ required_rubygems_version: !ruby/object:Gem::Requirement
89
+ requirements:
90
+ - - ">="
91
+ - !ruby/object:Gem::Version
92
+ version: '0'
93
+ requirements: []
94
+ rubyforge_project:
95
+ rubygems_version: 2.7.3
96
+ signing_key:
97
+ specification_version: 4
98
+ summary: filter items like Slack switcher, Atom command palette
99
+ test_files: []