python-voiceio 0.9.6__tar.gz → 0.9.8__tar.gz

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (130) hide show
  1. {python_voiceio-0.9.6/python_voiceio.egg-info → python_voiceio-0.9.8}/PKG-INFO +1 -1
  2. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/pyproject.toml +1 -1
  3. {python_voiceio-0.9.6 → python_voiceio-0.9.8/python_voiceio.egg-info}/PKG-INFO +1 -1
  4. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_platform.py +31 -0
  5. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_postcorrect.py +19 -0
  6. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_postprocess.py +36 -3
  7. python_voiceio-0.9.8/voiceio/__init__.py +1 -0
  8. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/config.py +2 -1
  9. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/platform.py +21 -6
  10. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/postcorrect.py +34 -10
  11. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/postprocess.py +39 -13
  12. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/__init__.py +1 -1
  13. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/_pystray.py +1 -1
  14. python_voiceio-0.9.6/voiceio/__init__.py +0 -1
  15. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/LICENSE +0 -0
  16. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/README.md +0 -0
  17. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/SOURCES.txt +0 -0
  18. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/dependency_links.txt +0 -0
  19. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/entry_points.txt +0 -0
  20. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/requires.txt +0 -0
  21. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/top_level.txt +0 -0
  22. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/setup.cfg +0 -0
  23. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_adjudicate.py +0 -0
  24. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_app_wiring.py +0 -0
  25. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_audio_quality.py +0 -0
  26. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_audit.py +0 -0
  27. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_autocorrect.py +0 -0
  28. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_autocorrect_state.py +0 -0
  29. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_backend_probes.py +0 -0
  30. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_cli.py +0 -0
  31. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_clipboard_read.py +0 -0
  32. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_commands.py +0 -0
  33. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_concurrency_lockdown.py +0 -0
  34. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_config.py +0 -0
  35. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_correct_batch.py +0 -0
  36. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_corrections.py +0 -0
  37. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_evaluate.py +0 -0
  38. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_fallback.py +0 -0
  39. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_health.py +0 -0
  40. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_hints.py +0 -0
  41. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_history.py +0 -0
  42. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_ibus_pending.py +0 -0
  43. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_ibus_ping.py +0 -0
  44. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_ibus_typer.py +0 -0
  45. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_llm.py +0 -0
  46. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_llm_api.py +0 -0
  47. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_numbers.py +0 -0
  48. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_prebuffer.py +0 -0
  49. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_prompt.py +0 -0
  50. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_recorder_integration.py +0 -0
  51. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_retention.py +0 -0
  52. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_robustness.py +0 -0
  53. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_security_hardening.py +0 -0
  54. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_snapshots.py +0 -0
  55. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_streaming.py +0 -0
  56. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_tokens.py +0 -0
  57. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_transcriber.py +0 -0
  58. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_tts.py +0 -0
  59. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_vad.py +0 -0
  60. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_vocabulary.py +0 -0
  61. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_wizard.py +0 -0
  62. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_wordfreq.py +0 -0
  63. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/__main__.py +0 -0
  64. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/app.py +0 -0
  65. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/audit.py +0 -0
  66. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/autocorrect.py +0 -0
  67. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/autocorrect_state.py +0 -0
  68. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/backends.py +0 -0
  69. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/cli.py +0 -0
  70. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/clipboard_read.py +0 -0
  71. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/commands.py +0 -0
  72. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/consent.py +0 -0
  73. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/corrections.py +0 -0
  74. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/demo.py +0 -0
  75. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/evaluate.py +0 -0
  76. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/feedback.py +0 -0
  77. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/health.py +0 -0
  78. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hints.py +0 -0
  79. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/history.py +0 -0
  80. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/__init__.py +0 -0
  81. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/base.py +0 -0
  82. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/chain.py +0 -0
  83. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/evdev.py +0 -0
  84. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/pynput_backend.py +0 -0
  85. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/socket_backend.py +0 -0
  86. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/ibus/__init__.py +0 -0
  87. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/ibus/engine.py +0 -0
  88. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/ibus/pending.py +0 -0
  89. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/llm.py +0 -0
  90. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/llm_api.py +0 -0
  91. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/models/__init__.py +0 -0
  92. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/models/silero_vad.onnx +0 -0
  93. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/numbers.py +0 -0
  94. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/pidlock.py +0 -0
  95. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/prompt.py +0 -0
  96. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/recorder.py +0 -0
  97. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/retention.py +0 -0
  98. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/service.py +0 -0
  99. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/snapshots.py +0 -0
  100. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/__init__.py +0 -0
  101. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/commit.wav +0 -0
  102. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/start.wav +0 -0
  103. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/stop.wav +0 -0
  104. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/streaming.py +0 -0
  105. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tokens.py +0 -0
  106. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/transcriber.py +0 -0
  107. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/_icons.py +0 -0
  108. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/_indicator.py +0 -0
  109. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/__init__.py +0 -0
  110. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/base.py +0 -0
  111. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/chain.py +0 -0
  112. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/edge_engine.py +0 -0
  113. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/espeak.py +0 -0
  114. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/piper_engine.py +0 -0
  115. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/player.py +0 -0
  116. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/__init__.py +0 -0
  117. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/base.py +0 -0
  118. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/chain.py +0 -0
  119. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/clipboard.py +0 -0
  120. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/ibus.py +0 -0
  121. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/pynput_type.py +0 -0
  122. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/wtype.py +0 -0
  123. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/xdotool.py +0 -0
  124. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/ydotool.py +0 -0
  125. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/vad.py +0 -0
  126. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/vocab_stats.py +0 -0
  127. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/vocabulary.py +0 -0
  128. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/wizard.py +0 -0
  129. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/wordfreq.py +0 -0
  130. {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/worker.py +0 -0
@@ -1,6 +1,6 @@
1
1
  Metadata-Version: 2.4
2
2
  Name: python-voiceio
3
- Version: 0.9.6
3
+ Version: 0.9.8
4
4
  Summary: Voice dictation for Linux. Speak → text, locally, instantly.
5
5
  Author: Hugo Montenegro
6
6
  License-Expression: MIT
@@ -4,7 +4,7 @@ build-backend = "setuptools.build_meta"
4
4
 
5
5
  [project]
6
6
  name = "python-voiceio"
7
- version = "0.9.6"
7
+ version = "0.9.8"
8
8
  description = "Voice dictation for Linux. Speak → text, locally, instantly."
9
9
  readme = "README.md"
10
10
  license = "MIT"
@@ -1,6 +1,6 @@
1
1
  Metadata-Version: 2.4
2
2
  Name: python-voiceio
3
- Version: 0.9.6
3
+ Version: 0.9.8
4
4
  Summary: Voice dictation for Linux. Speak → text, locally, instantly.
5
5
  Author: Hugo Montenegro
6
6
  License-Expression: MIT
@@ -72,3 +72,34 @@ def test_platform_frozen():
72
72
  import pytest
73
73
  with pytest.raises(AttributeError):
74
74
  p.os = "darwin"
75
+
76
+
77
+ class TestOpenInTerminal:
78
+ """The tray menu launches print-and-exit commands (history/doctor/logs);
79
+ hold=True must keep the window open instead of flashing it shut."""
80
+
81
+ def _capture(self, monkeypatch, *, hold):
82
+ import voiceio.platform as plat
83
+ monkeypatch.setattr(plat, "_detect_os", lambda: "linux")
84
+ monkeypatch.setenv("TERMINAL", "myterm")
85
+ monkeypatch.setattr(
86
+ plat.shutil, "which",
87
+ lambda b: "/usr/bin/myterm" if b == "myterm" else None,
88
+ )
89
+ calls = []
90
+ with patch("subprocess.Popen", side_effect=lambda args, *a, **k: calls.append(args)):
91
+ ok = plat.open_in_terminal(["voiceio", "history"], hold=hold)
92
+ assert ok is True
93
+ return calls[0]
94
+
95
+ def test_hold_wraps_in_shell_that_pauses(self, monkeypatch):
96
+ args = self._capture(monkeypatch, hold=True)
97
+ assert args[:2] == ["myterm", "-e"]
98
+ assert "sh" in args and "-c" in args
99
+ script = args[-1]
100
+ assert "voiceio history" in script
101
+ assert "read" in script # pauses so the window survives
102
+
103
+ def test_no_hold_runs_command_directly(self, monkeypatch):
104
+ args = self._capture(monkeypatch, hold=False)
105
+ assert args == ["myterm", "-e", "voiceio", "history"]
@@ -121,6 +121,25 @@ def test_disfluency_mode_rejects_rewording():
121
121
  assert pc.correct(original) == original
122
122
 
123
123
 
124
+ def test_disfluency_mode_rejects_dropped_negation():
125
+ """Dropping 'not' is one deletion under every fraction cap but inverts
126
+ meaning — the negation guard must reject it."""
127
+ pc = PostCorrector(_cfg(remove_disfluencies=True))
128
+ original = "I do not want the feature to ship today at all"
129
+ inverted = "I do want the feature to ship today at all"
130
+ with patch("voiceio.llm_api.chat", return_value=inverted):
131
+ assert pc.correct(original) == original
132
+
133
+
134
+ def test_disfluency_mode_rejects_replaced_negation():
135
+ """Replacing a negation with something else also inverts meaning."""
136
+ pc = PostCorrector(_cfg(remove_disfluencies=True))
137
+ original = "we should never deploy on a friday afternoon"
138
+ changed = "we should always deploy on a friday afternoon"
139
+ with patch("voiceio.llm_api.chat", return_value=changed):
140
+ assert pc.correct(original) == original
141
+
142
+
124
143
  def test_fix_mode_still_rejects_big_deletion():
125
144
  """With disfluency mode OFF, the original conservative guards stand: a big
126
145
  deletion is a length change and must be rejected."""
@@ -24,15 +24,48 @@ class TestStripDisfluencies:
24
24
  assert strip_disfluencies("I had had enough") == "I had had enough"
25
25
 
26
26
  def test_dedups_duplicate_sentence(self):
27
- # The Whisper re-decode artifact: an exact sentence repeated.
28
- out = strip_disfluencies("Do the research. Do the research.")
29
- assert out == "Do the research."
27
+ # The Whisper re-decode artifact: a whole sentence repeated verbatim.
28
+ out = strip_disfluencies("Do the deep research now. Do the deep research now.")
29
+ assert out == "Do the deep research now."
30
30
 
31
31
  def test_keeps_meaningful_words(self):
32
32
  # 'like' as a real verb/preposition and content must survive.
33
33
  text = "I like the design and it works like a charm"
34
34
  assert strip_disfluencies(text) == text
35
35
 
36
+ def test_never_eats_real_words_or_units(self):
37
+ # Default-on runs on everyone's speech: filler patterns must not collide
38
+ # with real words, units, or abbreviations. Case-sensitive matching is
39
+ # what protects the all-caps abbreviations (ER, UM, HM).
40
+ for text in [
41
+ "to err is human",
42
+ "we should err on caution",
43
+ "the bolt is 5 mm wide",
44
+ "set it to 10 mm please",
45
+ "ah yes I remember now",
46
+ "I like the ohm rating",
47
+ "Take him to the ER right now", # ER = emergency room
48
+ "The UM campus in Michigan", # UM = University of Michigan
49
+ "The Er atom is a lanthanide", # Er = erbium
50
+ "we measured 3 hm across", # hm = hectometre (bare, 1 m)
51
+ ]:
52
+ assert strip_disfluencies(text) == text, text
53
+
54
+ def test_preserves_newlines_and_structure(self):
55
+ # A filler on its own line must not swallow the paragraph break.
56
+ assert strip_disfluencies("First para.\n\nSecond para.") == \
57
+ "First para.\n\nSecond para."
58
+ assert "\n\n" in strip_disfluencies("First para.\n\num\n\nSecond para.")
59
+
60
+ def test_uh_huh_removed_whole(self):
61
+ # Regression: ordering bug once stranded "-huh".
62
+ assert strip_disfluencies("uh-huh right") == "right"
63
+ assert strip_disfluencies("uh huh yes") == "yes"
64
+
65
+ def test_preserves_emphatic_short_repeat(self):
66
+ # Short repeats are emphasis, not a re-decode artifact — keep them.
67
+ assert strip_disfluencies("No. No.") == "No. No."
68
+
36
69
  def test_empty(self):
37
70
  assert strip_disfluencies("") == ""
38
71
 
@@ -0,0 +1 @@
1
+ __version__ = "0.9.8"
@@ -93,7 +93,8 @@ class OutputConfig:
93
93
  # and duplicate re-decode sentences on every pass; when [postcorrect] is
94
94
  # also enabled, its LLM pass additionally removes false starts and filler
95
95
  # "like", guarded so it can only delete — never add, rephrase, or reorder.
96
- remove_disfluencies: bool = False
96
+ # On by default: nobody wants "um" in their notes. Set false for verbatim.
97
+ remove_disfluencies: bool = True
97
98
  voice_input_prefix: str = "" # e.g. "[voice]" — empty disables
98
99
  # Incremental finalization: once the un-finalized audio tail grows past
99
100
  # this many seconds, it is cut at the nearest interior speech pause,
@@ -311,8 +311,14 @@ def system_deps_install_cmd(keys: list[str] | None = None) -> str:
311
311
  return f"{prefix} {' '.join(pkgs)}"
312
312
 
313
313
 
314
- def open_in_terminal(cmd: list[str]) -> bool:
315
- """Launch a command in the user's terminal emulator. Returns success."""
314
+ def open_in_terminal(cmd: list[str], *, hold: bool = False) -> bool:
315
+ """Launch a command in the user's terminal emulator. Returns success.
316
+
317
+ ``hold=True`` keeps the window open after the command exits — required for
318
+ print-and-exit commands (history, doctor, logs) that would otherwise flash
319
+ a terminal that closes the instant they finish.
320
+ """
321
+ import shlex
316
322
  import subprocess
317
323
 
318
324
  plat_os = _detect_os()
@@ -334,11 +340,20 @@ def open_in_terminal(cmd: list[str]) -> bool:
334
340
  log.warning("Failed to open cmd.exe")
335
341
  return False
336
342
 
337
- # Linux: try common terminal emulators
343
+ # Linux: try common terminal emulators. When holding, wrap the command in a
344
+ # shell that pauses at the end so the window survives a print-and-exit.
345
+ launch = cmd
346
+ if hold:
347
+ inner = " ".join(shlex.quote(c) for c in cmd)
348
+ launch = [
349
+ "sh", "-c",
350
+ f'{inner}; printf "\\n[voiceio] Press Enter to close…"; read _',
351
+ ]
352
+
338
353
  term = os.environ.get("TERMINAL")
339
354
  if term and shutil.which(term):
340
355
  try:
341
- subprocess.Popen([term, "-e", *cmd])
356
+ subprocess.Popen([term, "-e", *launch])
342
357
  return True
343
358
  except OSError:
344
359
  pass
@@ -346,7 +361,7 @@ def open_in_terminal(cmd: list[str]) -> bool:
346
361
  # x-terminal-emulator (Debian/Ubuntu alternative)
347
362
  if shutil.which("x-terminal-emulator"):
348
363
  try:
349
- subprocess.Popen(["x-terminal-emulator", "-e", *cmd])
364
+ subprocess.Popen(["x-terminal-emulator", "-e", *launch])
350
365
  return True
351
366
  except OSError:
352
367
  pass
@@ -365,7 +380,7 @@ def open_in_terminal(cmd: list[str]) -> bool:
365
380
  for prefix, binary in _TERMINALS:
366
381
  if shutil.which(binary):
367
382
  try:
368
- subprocess.Popen([*prefix, *cmd])
383
+ subprocess.Popen([*prefix, *launch])
369
384
  return True
370
385
  except OSError:
371
386
  continue
@@ -34,6 +34,19 @@ _MAX_INSERTED_WORDS = 0 # adding any word = altering meaning → reject
34
34
  _MAX_REPLACE_FRAC = 0.15 # ASR word-fixes only, never wholesale rewording
35
35
  _MAX_DELETE_FRAC = 0.4 # backstop against deleting real content
36
36
 
37
+ # Words whose deletion/replacement flips meaning — the fraction caps can't catch
38
+ # a single dropped "not". If the edit touches any of these on the original side,
39
+ # reject outright. (Contractions ending in "n't" are handled separately.)
40
+ _MEANING_CRITICAL = frozenset({
41
+ "not", "no", "never", "none", "nor", "neither", "without", "cannot",
42
+ "nothing", "nobody", "nowhere", "n't",
43
+ })
44
+
45
+
46
+ def _is_meaning_critical(word: str) -> bool:
47
+ w = word.lower().strip(".,;:!?\"'()")
48
+ return w in _MEANING_CRITICAL or w.endswith("n't")
49
+
37
50
  _SYSTEM_PROMPT = (
38
51
  "You fix automatic speech recognition errors in dictated text. "
39
52
  "The user dictates about software engineering and their projects. "
@@ -46,18 +59,25 @@ _SYSTEM_PROMPT = (
46
59
  # Disfluency mode: also strip spoken filler, delete-only. The strict rules
47
60
  # mirror the guards — the model is told exactly what the diff check enforces.
48
61
  _SYSTEM_PROMPT_CLEAN = (
49
- "You convert dictated speech into clean written text. The user dictates "
50
- "about software engineering and their projects. Do exactly two things:\n"
62
+ "You convert dictated speech into clean, concise written text. The user "
63
+ "dictates about software engineering and their projects, and wants it to "
64
+ "read like something they wrote, not spoke. Do exactly two things:\n"
51
65
  "1. Fix words the recognizer misheard (wrong proper nouns, homophones, "
52
66
  "garbled technical terms).\n"
53
- "2. Remove speech disfluencies: filler sounds (um, uh, er); filler uses of "
54
- "'like', 'you know', 'I mean'; false starts and self-corrections (keep the "
55
- "corrected version); and stray word repetitions.\n"
56
- "STRICT RULES: Only DELETE disfluencies and FIX misheard words. NEVER add "
57
- "words. NEVER rephrase, reword, reorder, or summarize. NEVER drop real "
58
- "content, meaningful hedges, or negations. If unsure whether something is a "
59
- "disfluency, KEEP it. Preserve the speaker's own wording and punctuation. "
60
- "Return only the cleaned text, nothing else."
67
+ "2. Cut spoken disfluency so it reads tight and clear:\n"
68
+ " - filler sounds (um, uh, er);\n"
69
+ " - filler discourse markers with no content: leading 'so', 'yeah', "
70
+ "'okay', 'well', 'now', 'right', and 'you know', 'I mean', 'like', 'sort "
71
+ "of'/'kind of' when used as filler, 'or something'/'or whatever', a "
72
+ "trailing 'right?' tag;\n"
73
+ " - false starts and self-corrections (keep only the corrected version);\n"
74
+ " - stray word and phrase repetitions.\n"
75
+ "Lean toward the briefer version — remove hesitation and padding freely.\n"
76
+ "HARD LIMITS (these protect meaning): only DELETE and FIX — NEVER add "
77
+ "words, rephrase, reword, reorder, or summarize. NEVER drop or alter actual "
78
+ "content, a negation ('not', 'never', \"n't\"), or a hedge that changes "
79
+ "certainty ('I think', 'maybe', 'probably', 'might'). Keep the speaker's own "
80
+ "wording. Return only the cleaned text, nothing else."
61
81
  )
62
82
 
63
83
  _MAX_RECENT = 3
@@ -315,8 +335,12 @@ class PostCorrector:
315
335
  inserted += j2 - j1
316
336
  elif tag == "replace":
317
337
  replaced += max(i2 - i1, j2 - j1)
338
+ if any(_is_meaning_critical(w) for w in aw[i1:i2]):
339
+ return False, "negation" # e.g. "not" → something else
318
340
  elif tag == "delete":
319
341
  deleted += i2 - i1
342
+ if any(_is_meaning_critical(w) for w in aw[i1:i2]):
343
+ return False, "negation" # dropping "not" inverts meaning
320
344
  n = len(aw)
321
345
  if inserted > _MAX_INSERTED_WORDS:
322
346
  return False, "inserted" # added content — meaning changed
@@ -14,13 +14,25 @@ if TYPE_CHECKING:
14
14
  _NO_CASE_LANGUAGES = frozenset({"zh", "ja", "ko", "ar", "he", "th", "hi", "bn", "ka", "my"})
15
15
 
16
16
  # Filler SOUNDS only — tokens with no lexical meaning, so deleting them can
17
- # never change meaning. Word repetitions ("had had" is valid English) and
18
- # filler uses of "like"/"you know" need judgment and are left to the LLM layer.
19
- # Consumes an adjacent comma on either side so "be, uh, found" → "be found".
17
+ # never change meaning. Deliberately conservative because this runs on every
18
+ # user's speech:
19
+ # * CASE-SENSITIVE (no IGNORECASE): all-caps abbreviations that look like
20
+ # fillers are real words and must survive — "ER" (emergency room), "UM"
21
+ # (University of Michigan), "HM". We match lowercase and Title-case forms
22
+ # ("um", "Um") — Whisper's filler spellings — but never all-caps.
23
+ # * Excluded entirely: "er"/"erm" ("ER"/"Er"=erbium), "mm" (millimetres),
24
+ # "ah" (interjection), bare "hm" (hectometre — require "hmm", 2+ m's).
25
+ # * Surrounding whitespace is horizontal-only ([^\S\n]) so a filler on its
26
+ # own line doesn't swallow the paragraph/list break around it.
27
+ # Word repetitions ("had had" is valid English) and filler "like"/"you know"
28
+ # need judgment and are left to the LLM layer. Order matters: multi-token
29
+ # "uh-huh" before "[Uu]h+" so it isn't clipped to a stray "-huh".
20
30
  _FILLER_RE = re.compile(
21
- r"\s*,?\s*\b(?:u+m+|u+h+m*|e+r+m*|erm+|a+h+|h+m+|mm+|mhm|uh[-\s]?huh)\b\s*,?\s*",
22
- re.IGNORECASE,
31
+ r"[^\S\n]*,?[^\S\n]*\b(?:[Uu]h[-\s]?huh|[Mm]hm|[Uu]h+m*|[Uu]m+|[Hh]m{2,})\b[^\S\n]*,?[^\S\n]*",
23
32
  )
33
+ # A re-decode artifact is a whole duplicated sentence; require this many words
34
+ # so emphatic short repeats ("No. No.", "Stop. Stop.") are preserved.
35
+ _MIN_DEDUP_WORDS = 4
24
36
 
25
37
 
26
38
  def strip_disfluencies(text: str) -> str:
@@ -36,19 +48,33 @@ def strip_disfluencies(text: str) -> str:
36
48
  return text
37
49
  text = _FILLER_RE.sub(" ", text)
38
50
  text = _dedup_adjacent_sentences(text)
39
- text = re.sub(r"\s{2,}", " ", text)
51
+ # Repair the debris the deletions leave, without crossing newlines (so
52
+ # paragraph/list structure survives even when punctuation_cleanup is off).
53
+ text = re.sub(r"[^\S\n]+([,.;:?!])", r"\1", text) # space before punctuation
54
+ text = re.sub(r"[^\S\n]{2,}", " ", text) # collapse runs of spaces
40
55
  return text.strip()
41
56
 
42
57
 
43
58
  def _dedup_adjacent_sentences(text: str) -> str:
44
- """Drop a sentence identical to the one immediately before it."""
45
- parts = re.split(r"(?<=[.?!])\s+", text)
59
+ """Drop a sentence identical to the one immediately before it.
60
+
61
+ Separators are captured and preserved on rejoin so paragraph/list breaks
62
+ survive; only a full (>= _MIN_DEDUP_WORDS) verbatim repeat is removed.
63
+ """
64
+ tokens = re.split(r"((?<=[.?!])\s+)", text) # [sent, sep, sent, sep, …]
46
65
  out: list[str] = []
47
- for p in parts:
48
- if out and p.strip().lower() == out[-1].strip().lower():
49
- continue
50
- out.append(p)
51
- return " ".join(out)
66
+ last_kept: str | None = None
67
+ for i in range(0, len(tokens), 2):
68
+ sentence = tokens[i]
69
+ sep = tokens[i + 1] if i + 1 < len(tokens) else ""
70
+ norm = sentence.strip().lower()
71
+ if (last_kept is not None and norm == last_kept
72
+ and len(sentence.split()) >= _MIN_DEDUP_WORDS):
73
+ continue # drop the duplicate sentence and its separator
74
+ out.append(sentence)
75
+ out.append(sep)
76
+ last_kept = norm
77
+ return "".join(out)
52
78
 
53
79
 
54
80
  def cleanup(text: str, language: str = "en") -> str:
@@ -146,7 +146,7 @@ def _read_stdout(proc: subprocess.Popen, toggle_cb: Callable[[], None]) -> None:
146
146
  action = cmd[5:]
147
147
  cli_cmd = _MENU_COMMANDS.get(action)
148
148
  if cli_cmd:
149
- if not open_in_terminal(cli_cmd):
149
+ if not open_in_terminal(cli_cmd, hold=True):
150
150
  log.warning("Failed to open terminal for: %s", action)
151
151
  except (OSError, ValueError):
152
152
  pass
@@ -45,7 +45,7 @@ def start(
45
45
  from voiceio.platform import open_in_terminal
46
46
  cli_cmd = _MENU_COMMANDS.get(action)
47
47
  if cli_cmd:
48
- open_in_terminal(cli_cmd)
48
+ open_in_terminal(cli_cmd, hold=True)
49
49
  return _handler
50
50
 
51
51
  _icon = pystray.Icon(
@@ -1 +0,0 @@
1
- __version__ = "0.9.6"
File without changes
File without changes
File without changes