httpd_log_parser 0.1.0

Sign up to get free protection for your applications and to get access to all the features.
checksums.yaml ADDED
@@ -0,0 +1,7 @@
1
+ ---
2
+ SHA1:
3
+ metadata.gz: 671703113f0a090506818a054d2d188e3feb6ead
4
+ data.tar.gz: 282022bc528dde26dc3697a1d82a7c0b339d1f0c
5
+ SHA512:
6
+ metadata.gz: baf2397f3ca8beb6c70e5e6202eacce9e65d29e1742b75f2dd1a0c209951131cddae45c80f05ad1b8b21e95d172377a6d21d4791dafc570a0a125fb0c760f137
7
+ data.tar.gz: 8d1a738434f85ba5ab28528d0ea357ea153c31e9a5c92a979a3861a01980197550cbd55228de631c100fb69abf00fdaf1e6b7245c0cc3e116f64ce60df21c248
data/.gitignore ADDED
@@ -0,0 +1,9 @@
1
+ /.bundle/
2
+ /.yardoc
3
+ /Gemfile.lock
4
+ /_yardoc/
5
+ /coverage/
6
+ /doc/
7
+ /pkg/
8
+ /spec/reports/
9
+ /tmp/
data/.travis.yml ADDED
@@ -0,0 +1,4 @@
1
+ language: ruby
2
+ rvm:
3
+ - 2.2.3
4
+ before_install: gem install bundler -v 1.10.6
data/Gemfile ADDED
@@ -0,0 +1,4 @@
1
+ source 'https://rubygems.org'
2
+
3
+ # Specify your gem's dependencies in httpd_log_parser.gemspec
4
+ gemspec
data/LICENSE.txt ADDED
@@ -0,0 +1,21 @@
1
+ The MIT License (MIT)
2
+
3
+ Copyright (c) 2016 inpwjp
4
+
5
+ Permission is hereby granted, free of charge, to any person obtaining a copy
6
+ of this software and associated documentation files (the "Software"), to deal
7
+ in the Software without restriction, including without limitation the rights
8
+ to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
9
+ copies of the Software, and to permit persons to whom the Software is
10
+ furnished to do so, subject to the following conditions:
11
+
12
+ The above copyright notice and this permission notice shall be included in
13
+ all copies or substantial portions of the Software.
14
+
15
+ THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
16
+ IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
17
+ FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE
18
+ AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
19
+ LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
20
+ OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN
21
+ THE SOFTWARE.
data/README.md ADDED
@@ -0,0 +1,41 @@
1
+ # HttpdLogParser
2
+
3
+ Welcome to your new gem! In this directory, you'll find the files you need to be able to package up your Ruby library into a gem. Put your Ruby code in the file `lib/httpd_log_parser`. To experiment with that code, run `bin/console` for an interactive prompt.
4
+
5
+ TODO: Delete this and the text above, and describe your gem
6
+
7
+ ## Installation
8
+
9
+ Add this line to your application's Gemfile:
10
+
11
+ ```ruby
12
+ gem 'httpd_log_parser'
13
+ ```
14
+
15
+ And then execute:
16
+
17
+ $ bundle
18
+
19
+ Or install it yourself as:
20
+
21
+ $ gem install httpd_log_parser
22
+
23
+ ## Usage
24
+
25
+ TODO: Write usage instructions here
26
+
27
+ ## Development
28
+
29
+ After checking out the repo, run `bin/setup` to install dependencies. Then, run `rake false` to run the tests. You can also run `bin/console` for an interactive prompt that will allow you to experiment.
30
+
31
+ To install this gem onto your local machine, run `bundle exec rake install`. To release a new version, update the version number in `version.rb`, and then run `bundle exec rake release`, which will create a git tag for the version, push git commits and tags, and push the `.gem` file to [rubygems.org](https://rubygems.org).
32
+
33
+ ## Contributing
34
+
35
+ Bug reports and pull requests are welcome on GitHub at https://github.com/[USERNAME]/httpd_log_parser.
36
+
37
+
38
+ ## License
39
+
40
+ The gem is available as open source under the terms of the [MIT License](http://opensource.org/licenses/MIT).
41
+
data/Rakefile ADDED
@@ -0,0 +1 @@
1
+ require "bundler/gem_tasks"
data/bin/console ADDED
@@ -0,0 +1,14 @@
1
+ #!/usr/bin/env ruby
2
+
3
+ require "bundler/setup"
4
+ require "httpd_log_parser"
5
+
6
+ # You can add fixtures and/or initialization code here to make experimenting
7
+ # with your gem easier. You can also use a different console, if you like.
8
+
9
+ # (If you use this, don't forget to add pry to your Gemfile!)
10
+ # require "pry"
11
+ # Pry.start
12
+
13
+ require "irb"
14
+ IRB.start
data/bin/setup ADDED
@@ -0,0 +1,7 @@
1
+ #!/bin/bash
2
+ set -euo pipefail
3
+ IFS=$'\n\t'
4
+
5
+ bundle install
6
+
7
+ # Do any other automated setup that you need to do here
@@ -0,0 +1,32 @@
1
+ # coding: utf-8
2
+ lib = File.expand_path('../lib', __FILE__)
3
+ $LOAD_PATH.unshift(lib) unless $LOAD_PATH.include?(lib)
4
+ require 'httpd_log_parser/version'
5
+
6
+ Gem::Specification.new do |spec|
7
+ spec.name = "httpd_log_parser"
8
+ spec.version = HttpdLogParser::VERSION
9
+ spec.authors = ["inpwjp"]
10
+ spec.email = ["inpw@mua.biglobe.ne.jp"]
11
+
12
+ spec.summary = %q{This library is, convert ruby object from Apache log(common, combined).}
13
+ spec.description = %q{This library is, convert ruby object from Apache log(common, combined).}
14
+ spec.homepage = "http://github.com./inpwjp"
15
+ spec.license = "MIT"
16
+
17
+ # Prevent pushing this gem to RubyGems.org by setting 'allowed_push_host', or
18
+ # delete this section to allow pushing this gem to any host.
19
+ # if spec.respond_to?(:metadata)
20
+ # spec.metadata['allowed_push_host'] = "TODO: Set to 'http://mygemserver.com'"
21
+ # else
22
+ # raise "RubyGems 2.0 or newer is required to protect against public gem pushes."
23
+ # end
24
+
25
+ spec.files = `git ls-files -z`.split("\x0").reject { |f| f.match(%r{^(test|spec|features)/}) }
26
+ spec.bindir = "exe"
27
+ spec.executables = spec.files.grep(%r{^exe/}) { |f| File.basename(f) }
28
+ spec.require_paths = ["lib"]
29
+
30
+ spec.add_development_dependency "bundler", "~> 1.10"
31
+ spec.add_development_dependency "rake", "~> 10.0"
32
+ end
@@ -0,0 +1,28 @@
1
+
2
+ module HttpdLogParser
3
+ class Base
4
+ attr_accessor :remote_host, :client_type, :user_name, :access_time, :request_code, :response_status, :object_bytes, :referer_uri, :user_agent, :request_method, :uri, :url, :query
5
+ @@format = '%d/%b/%Y:%H:%M:%S %Z'
6
+
7
+ def filter_log
8
+ if access_time.nil?
9
+ return nil
10
+ end
11
+
12
+ yield
13
+
14
+ log
15
+ end
16
+
17
+ def delete_images()
18
+ filter_log do
19
+ images = ["js", "css", "json", "png", "gif", "jpg","woff", "ico", "pdf", "mp3", "jsp", "xml"]
20
+ images.each do |image|
21
+ if /#{image}$/ =~ self.url
22
+ return nil
23
+ end
24
+ end
25
+ end
26
+ end
27
+ end
28
+ end
@@ -0,0 +1,52 @@
1
+ require 'csv'
2
+ require 'date'
3
+ require 'digest/md5'
4
+
5
+ module HttpdLogParser
6
+ class Combined < Base
7
+ def set_string(str)
8
+ str.gsub!(/\[(\d+\/[a-zA-Z]{3}\/\d+.+?)\]/, "\"\\1\"")
9
+ if str.include?(' ')
10
+ begin
11
+ line = CSV::parse_line(str, {col_sep: ' ', skip_blanks: true})
12
+ self.remote_host = line[0]
13
+ self.client_type = line[1]
14
+ self.user_name = line[2]
15
+ self.access_time = DateTime.strptime(line[3], @@format)
16
+ self.request_code = line[4]
17
+ self.response_status = line[5]
18
+ self.object_bytes = line[6]
19
+ if line[7] == "_"
20
+ self.referer_uri = nil
21
+ else
22
+ self.referer_uri = line[7]
23
+ end
24
+ self.user_agent = line[8]
25
+ request_datum = self.request_code.split(" ")
26
+ self.request_method = request_datum[0]
27
+ self.uri = request_datum[1]
28
+ if ! self.uri.nil?
29
+ self.url, self.query = self.uri.split('?')
30
+ end
31
+ rescue => e
32
+ $stderr.puts e
33
+ $stderr.puts line
34
+ return self
35
+ end
36
+ end
37
+ self
38
+ end
39
+
40
+ def uniq_hash
41
+ Digest::MD5.hexdigest(self.user_name.to_s + self.user_agent.to_s + self.remote_host)
42
+ end
43
+
44
+ def referer_hash
45
+ Digest::MD5.hexdigest(self.referer_uri)
46
+ end
47
+
48
+ def log
49
+ "#{self.remote_host} #{self.client_type} #{self.user_name} [#{self.access_time.strftime(@@format)}] \"#{self.request_code}\" #{self.response_status} #{self.object_bytes} \"#{self.referer_uri}\" \"#{self.user_agent}\""
50
+ end
51
+ end
52
+ end
@@ -0,0 +1,42 @@
1
+ require 'csv'
2
+ require 'date'
3
+ require 'digest/md5'
4
+
5
+ module HttpdLogParser
6
+ class Common < Base
7
+ def set_string(str)
8
+ str.gsub!(/\[(\d+\/[a-zA-Z]{3}\/\d+.+?)\]/, "\"\\1\"")
9
+ if str.include?(' ')
10
+ begin
11
+ line = CSV::parse_line(str, {col_sep: ' ', skip_blanks: true})
12
+ self.remote_host = line[0]
13
+ self.client_type = line[1]
14
+ self.user_name = line[2]
15
+ self.access_time = DateTime.strptime(line[3], @@format)
16
+ self.request_code = line[4]
17
+ self.response_status = line[5]
18
+ self.object_bytes = line[6]
19
+ request_datum = self.request_code.split(" ")
20
+ self.request_method = request_datum[0]
21
+ self.uri = request_datum[1]
22
+ if ! self.uri.nil?
23
+ self.url, self.query = self.uri.split('?')
24
+ end
25
+ rescue => e
26
+ $stderr.puts e
27
+ $stderr.puts line
28
+ return self
29
+ end
30
+ end
31
+ self
32
+ end
33
+
34
+ def uniq_hash
35
+ Digest::MD5.hexdigest(self.user_name.to_s + self.remote_host)
36
+ end
37
+
38
+ def log
39
+ "#{self.remote_host} #{self.client_type} #{self.user_name} [#{self.access_time.strftime(@@format)}] \"#{self.request_code}\" #{self.response_status} #{self.object_bytes}"
40
+ end
41
+ end
42
+ end
@@ -0,0 +1,3 @@
1
+ module HttpdLogParser
2
+ VERSION = "0.1.0"
3
+ end
@@ -0,0 +1,3 @@
1
+ require 'httpd_log_parser/base'
2
+ require 'httpd_log_parser/common'
3
+ require 'httpd_log_parser/combined'
metadata ADDED
@@ -0,0 +1,86 @@
1
+ --- !ruby/object:Gem::Specification
2
+ name: httpd_log_parser
3
+ version: !ruby/object:Gem::Version
4
+ version: 0.1.0
5
+ platform: ruby
6
+ authors:
7
+ - inpwjp
8
+ autorequire:
9
+ bindir: exe
10
+ cert_chain: []
11
+ date: 2016-04-20 00:00:00.000000000 Z
12
+ dependencies:
13
+ - !ruby/object:Gem::Dependency
14
+ name: bundler
15
+ requirement: !ruby/object:Gem::Requirement
16
+ requirements:
17
+ - - "~>"
18
+ - !ruby/object:Gem::Version
19
+ version: '1.10'
20
+ type: :development
21
+ prerelease: false
22
+ version_requirements: !ruby/object:Gem::Requirement
23
+ requirements:
24
+ - - "~>"
25
+ - !ruby/object:Gem::Version
26
+ version: '1.10'
27
+ - !ruby/object:Gem::Dependency
28
+ name: rake
29
+ requirement: !ruby/object:Gem::Requirement
30
+ requirements:
31
+ - - "~>"
32
+ - !ruby/object:Gem::Version
33
+ version: '10.0'
34
+ type: :development
35
+ prerelease: false
36
+ version_requirements: !ruby/object:Gem::Requirement
37
+ requirements:
38
+ - - "~>"
39
+ - !ruby/object:Gem::Version
40
+ version: '10.0'
41
+ description: This library is, convert ruby object from Apache log(common, combined).
42
+ email:
43
+ - inpw@mua.biglobe.ne.jp
44
+ executables: []
45
+ extensions: []
46
+ extra_rdoc_files: []
47
+ files:
48
+ - ".gitignore"
49
+ - ".travis.yml"
50
+ - Gemfile
51
+ - LICENSE.txt
52
+ - README.md
53
+ - Rakefile
54
+ - bin/console
55
+ - bin/setup
56
+ - httpd_log_parser.gemspec
57
+ - lib/httpd_log_parser.rb
58
+ - lib/httpd_log_parser/base.rb
59
+ - lib/httpd_log_parser/combined.rb
60
+ - lib/httpd_log_parser/common.rb
61
+ - lib/httpd_log_parser/version.rb
62
+ homepage: http://github.com./inpwjp
63
+ licenses:
64
+ - MIT
65
+ metadata: {}
66
+ post_install_message:
67
+ rdoc_options: []
68
+ require_paths:
69
+ - lib
70
+ required_ruby_version: !ruby/object:Gem::Requirement
71
+ requirements:
72
+ - - ">="
73
+ - !ruby/object:Gem::Version
74
+ version: '0'
75
+ required_rubygems_version: !ruby/object:Gem::Requirement
76
+ requirements:
77
+ - - ">="
78
+ - !ruby/object:Gem::Version
79
+ version: '0'
80
+ requirements: []
81
+ rubyforge_project:
82
+ rubygems_version: 2.4.5.1
83
+ signing_key:
84
+ specification_version: 4
85
+ summary: This library is, convert ruby object from Apache log(common, combined).
86
+ test_files: []