python-voiceio 0.9.6__tar.gz → 0.9.8__tar.gz
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- {python_voiceio-0.9.6/python_voiceio.egg-info → python_voiceio-0.9.8}/PKG-INFO +1 -1
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/pyproject.toml +1 -1
- {python_voiceio-0.9.6 → python_voiceio-0.9.8/python_voiceio.egg-info}/PKG-INFO +1 -1
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_platform.py +31 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_postcorrect.py +19 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_postprocess.py +36 -3
- python_voiceio-0.9.8/voiceio/__init__.py +1 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/config.py +2 -1
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/platform.py +21 -6
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/postcorrect.py +34 -10
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/postprocess.py +39 -13
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/__init__.py +1 -1
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/_pystray.py +1 -1
- python_voiceio-0.9.6/voiceio/__init__.py +0 -1
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/LICENSE +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/README.md +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/SOURCES.txt +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/dependency_links.txt +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/entry_points.txt +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/requires.txt +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/python_voiceio.egg-info/top_level.txt +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/setup.cfg +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_adjudicate.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_app_wiring.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_audio_quality.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_audit.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_autocorrect.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_autocorrect_state.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_backend_probes.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_cli.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_clipboard_read.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_commands.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_concurrency_lockdown.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_config.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_correct_batch.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_corrections.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_evaluate.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_fallback.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_health.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_hints.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_history.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_ibus_pending.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_ibus_ping.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_ibus_typer.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_llm.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_llm_api.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_numbers.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_prebuffer.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_prompt.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_recorder_integration.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_retention.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_robustness.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_security_hardening.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_snapshots.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_streaming.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_tokens.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_transcriber.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_tts.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_vad.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_vocabulary.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_wizard.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/tests/test_wordfreq.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/__main__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/app.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/audit.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/autocorrect.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/autocorrect_state.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/backends.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/cli.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/clipboard_read.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/commands.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/consent.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/corrections.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/demo.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/evaluate.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/feedback.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/health.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hints.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/history.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/__init__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/base.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/chain.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/evdev.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/pynput_backend.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/hotkeys/socket_backend.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/ibus/__init__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/ibus/engine.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/ibus/pending.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/llm.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/llm_api.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/models/__init__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/models/silero_vad.onnx +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/numbers.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/pidlock.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/prompt.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/recorder.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/retention.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/service.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/snapshots.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/__init__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/commit.wav +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/start.wav +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/sounds/stop.wav +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/streaming.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tokens.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/transcriber.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/_icons.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tray/_indicator.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/__init__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/base.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/chain.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/edge_engine.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/espeak.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/piper_engine.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/tts/player.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/__init__.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/base.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/chain.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/clipboard.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/ibus.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/pynput_type.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/wtype.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/xdotool.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/typers/ydotool.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/vad.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/vocab_stats.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/vocabulary.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/wizard.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/wordfreq.py +0 -0
- {python_voiceio-0.9.6 → python_voiceio-0.9.8}/voiceio/worker.py +0 -0
|
@@ -72,3 +72,34 @@ def test_platform_frozen():
|
|
|
72
72
|
import pytest
|
|
73
73
|
with pytest.raises(AttributeError):
|
|
74
74
|
p.os = "darwin"
|
|
75
|
+
|
|
76
|
+
|
|
77
|
+
class TestOpenInTerminal:
|
|
78
|
+
"""The tray menu launches print-and-exit commands (history/doctor/logs);
|
|
79
|
+
hold=True must keep the window open instead of flashing it shut."""
|
|
80
|
+
|
|
81
|
+
def _capture(self, monkeypatch, *, hold):
|
|
82
|
+
import voiceio.platform as plat
|
|
83
|
+
monkeypatch.setattr(plat, "_detect_os", lambda: "linux")
|
|
84
|
+
monkeypatch.setenv("TERMINAL", "myterm")
|
|
85
|
+
monkeypatch.setattr(
|
|
86
|
+
plat.shutil, "which",
|
|
87
|
+
lambda b: "/usr/bin/myterm" if b == "myterm" else None,
|
|
88
|
+
)
|
|
89
|
+
calls = []
|
|
90
|
+
with patch("subprocess.Popen", side_effect=lambda args, *a, **k: calls.append(args)):
|
|
91
|
+
ok = plat.open_in_terminal(["voiceio", "history"], hold=hold)
|
|
92
|
+
assert ok is True
|
|
93
|
+
return calls[0]
|
|
94
|
+
|
|
95
|
+
def test_hold_wraps_in_shell_that_pauses(self, monkeypatch):
|
|
96
|
+
args = self._capture(monkeypatch, hold=True)
|
|
97
|
+
assert args[:2] == ["myterm", "-e"]
|
|
98
|
+
assert "sh" in args and "-c" in args
|
|
99
|
+
script = args[-1]
|
|
100
|
+
assert "voiceio history" in script
|
|
101
|
+
assert "read" in script # pauses so the window survives
|
|
102
|
+
|
|
103
|
+
def test_no_hold_runs_command_directly(self, monkeypatch):
|
|
104
|
+
args = self._capture(monkeypatch, hold=False)
|
|
105
|
+
assert args == ["myterm", "-e", "voiceio", "history"]
|
|
@@ -121,6 +121,25 @@ def test_disfluency_mode_rejects_rewording():
|
|
|
121
121
|
assert pc.correct(original) == original
|
|
122
122
|
|
|
123
123
|
|
|
124
|
+
def test_disfluency_mode_rejects_dropped_negation():
|
|
125
|
+
"""Dropping 'not' is one deletion under every fraction cap but inverts
|
|
126
|
+
meaning — the negation guard must reject it."""
|
|
127
|
+
pc = PostCorrector(_cfg(remove_disfluencies=True))
|
|
128
|
+
original = "I do not want the feature to ship today at all"
|
|
129
|
+
inverted = "I do want the feature to ship today at all"
|
|
130
|
+
with patch("voiceio.llm_api.chat", return_value=inverted):
|
|
131
|
+
assert pc.correct(original) == original
|
|
132
|
+
|
|
133
|
+
|
|
134
|
+
def test_disfluency_mode_rejects_replaced_negation():
|
|
135
|
+
"""Replacing a negation with something else also inverts meaning."""
|
|
136
|
+
pc = PostCorrector(_cfg(remove_disfluencies=True))
|
|
137
|
+
original = "we should never deploy on a friday afternoon"
|
|
138
|
+
changed = "we should always deploy on a friday afternoon"
|
|
139
|
+
with patch("voiceio.llm_api.chat", return_value=changed):
|
|
140
|
+
assert pc.correct(original) == original
|
|
141
|
+
|
|
142
|
+
|
|
124
143
|
def test_fix_mode_still_rejects_big_deletion():
|
|
125
144
|
"""With disfluency mode OFF, the original conservative guards stand: a big
|
|
126
145
|
deletion is a length change and must be rejected."""
|
|
@@ -24,15 +24,48 @@ class TestStripDisfluencies:
|
|
|
24
24
|
assert strip_disfluencies("I had had enough") == "I had had enough"
|
|
25
25
|
|
|
26
26
|
def test_dedups_duplicate_sentence(self):
|
|
27
|
-
# The Whisper re-decode artifact:
|
|
28
|
-
out = strip_disfluencies("Do the research. Do the research.")
|
|
29
|
-
assert out == "Do the research."
|
|
27
|
+
# The Whisper re-decode artifact: a whole sentence repeated verbatim.
|
|
28
|
+
out = strip_disfluencies("Do the deep research now. Do the deep research now.")
|
|
29
|
+
assert out == "Do the deep research now."
|
|
30
30
|
|
|
31
31
|
def test_keeps_meaningful_words(self):
|
|
32
32
|
# 'like' as a real verb/preposition and content must survive.
|
|
33
33
|
text = "I like the design and it works like a charm"
|
|
34
34
|
assert strip_disfluencies(text) == text
|
|
35
35
|
|
|
36
|
+
def test_never_eats_real_words_or_units(self):
|
|
37
|
+
# Default-on runs on everyone's speech: filler patterns must not collide
|
|
38
|
+
# with real words, units, or abbreviations. Case-sensitive matching is
|
|
39
|
+
# what protects the all-caps abbreviations (ER, UM, HM).
|
|
40
|
+
for text in [
|
|
41
|
+
"to err is human",
|
|
42
|
+
"we should err on caution",
|
|
43
|
+
"the bolt is 5 mm wide",
|
|
44
|
+
"set it to 10 mm please",
|
|
45
|
+
"ah yes I remember now",
|
|
46
|
+
"I like the ohm rating",
|
|
47
|
+
"Take him to the ER right now", # ER = emergency room
|
|
48
|
+
"The UM campus in Michigan", # UM = University of Michigan
|
|
49
|
+
"The Er atom is a lanthanide", # Er = erbium
|
|
50
|
+
"we measured 3 hm across", # hm = hectometre (bare, 1 m)
|
|
51
|
+
]:
|
|
52
|
+
assert strip_disfluencies(text) == text, text
|
|
53
|
+
|
|
54
|
+
def test_preserves_newlines_and_structure(self):
|
|
55
|
+
# A filler on its own line must not swallow the paragraph break.
|
|
56
|
+
assert strip_disfluencies("First para.\n\nSecond para.") == \
|
|
57
|
+
"First para.\n\nSecond para."
|
|
58
|
+
assert "\n\n" in strip_disfluencies("First para.\n\num\n\nSecond para.")
|
|
59
|
+
|
|
60
|
+
def test_uh_huh_removed_whole(self):
|
|
61
|
+
# Regression: ordering bug once stranded "-huh".
|
|
62
|
+
assert strip_disfluencies("uh-huh right") == "right"
|
|
63
|
+
assert strip_disfluencies("uh huh yes") == "yes"
|
|
64
|
+
|
|
65
|
+
def test_preserves_emphatic_short_repeat(self):
|
|
66
|
+
# Short repeats are emphasis, not a re-decode artifact — keep them.
|
|
67
|
+
assert strip_disfluencies("No. No.") == "No. No."
|
|
68
|
+
|
|
36
69
|
def test_empty(self):
|
|
37
70
|
assert strip_disfluencies("") == ""
|
|
38
71
|
|
|
@@ -0,0 +1 @@
|
|
|
1
|
+
__version__ = "0.9.8"
|
|
@@ -93,7 +93,8 @@ class OutputConfig:
|
|
|
93
93
|
# and duplicate re-decode sentences on every pass; when [postcorrect] is
|
|
94
94
|
# also enabled, its LLM pass additionally removes false starts and filler
|
|
95
95
|
# "like", guarded so it can only delete — never add, rephrase, or reorder.
|
|
96
|
-
|
|
96
|
+
# On by default: nobody wants "um" in their notes. Set false for verbatim.
|
|
97
|
+
remove_disfluencies: bool = True
|
|
97
98
|
voice_input_prefix: str = "" # e.g. "[voice]" — empty disables
|
|
98
99
|
# Incremental finalization: once the un-finalized audio tail grows past
|
|
99
100
|
# this many seconds, it is cut at the nearest interior speech pause,
|
|
@@ -311,8 +311,14 @@ def system_deps_install_cmd(keys: list[str] | None = None) -> str:
|
|
|
311
311
|
return f"{prefix} {' '.join(pkgs)}"
|
|
312
312
|
|
|
313
313
|
|
|
314
|
-
def open_in_terminal(cmd: list[str]) -> bool:
|
|
315
|
-
"""Launch a command in the user's terminal emulator. Returns success.
|
|
314
|
+
def open_in_terminal(cmd: list[str], *, hold: bool = False) -> bool:
|
|
315
|
+
"""Launch a command in the user's terminal emulator. Returns success.
|
|
316
|
+
|
|
317
|
+
``hold=True`` keeps the window open after the command exits — required for
|
|
318
|
+
print-and-exit commands (history, doctor, logs) that would otherwise flash
|
|
319
|
+
a terminal that closes the instant they finish.
|
|
320
|
+
"""
|
|
321
|
+
import shlex
|
|
316
322
|
import subprocess
|
|
317
323
|
|
|
318
324
|
plat_os = _detect_os()
|
|
@@ -334,11 +340,20 @@ def open_in_terminal(cmd: list[str]) -> bool:
|
|
|
334
340
|
log.warning("Failed to open cmd.exe")
|
|
335
341
|
return False
|
|
336
342
|
|
|
337
|
-
# Linux: try common terminal emulators
|
|
343
|
+
# Linux: try common terminal emulators. When holding, wrap the command in a
|
|
344
|
+
# shell that pauses at the end so the window survives a print-and-exit.
|
|
345
|
+
launch = cmd
|
|
346
|
+
if hold:
|
|
347
|
+
inner = " ".join(shlex.quote(c) for c in cmd)
|
|
348
|
+
launch = [
|
|
349
|
+
"sh", "-c",
|
|
350
|
+
f'{inner}; printf "\\n[voiceio] Press Enter to close…"; read _',
|
|
351
|
+
]
|
|
352
|
+
|
|
338
353
|
term = os.environ.get("TERMINAL")
|
|
339
354
|
if term and shutil.which(term):
|
|
340
355
|
try:
|
|
341
|
-
subprocess.Popen([term, "-e", *
|
|
356
|
+
subprocess.Popen([term, "-e", *launch])
|
|
342
357
|
return True
|
|
343
358
|
except OSError:
|
|
344
359
|
pass
|
|
@@ -346,7 +361,7 @@ def open_in_terminal(cmd: list[str]) -> bool:
|
|
|
346
361
|
# x-terminal-emulator (Debian/Ubuntu alternative)
|
|
347
362
|
if shutil.which("x-terminal-emulator"):
|
|
348
363
|
try:
|
|
349
|
-
subprocess.Popen(["x-terminal-emulator", "-e", *
|
|
364
|
+
subprocess.Popen(["x-terminal-emulator", "-e", *launch])
|
|
350
365
|
return True
|
|
351
366
|
except OSError:
|
|
352
367
|
pass
|
|
@@ -365,7 +380,7 @@ def open_in_terminal(cmd: list[str]) -> bool:
|
|
|
365
380
|
for prefix, binary in _TERMINALS:
|
|
366
381
|
if shutil.which(binary):
|
|
367
382
|
try:
|
|
368
|
-
subprocess.Popen([*prefix, *
|
|
383
|
+
subprocess.Popen([*prefix, *launch])
|
|
369
384
|
return True
|
|
370
385
|
except OSError:
|
|
371
386
|
continue
|
|
@@ -34,6 +34,19 @@ _MAX_INSERTED_WORDS = 0 # adding any word = altering meaning → reject
|
|
|
34
34
|
_MAX_REPLACE_FRAC = 0.15 # ASR word-fixes only, never wholesale rewording
|
|
35
35
|
_MAX_DELETE_FRAC = 0.4 # backstop against deleting real content
|
|
36
36
|
|
|
37
|
+
# Words whose deletion/replacement flips meaning — the fraction caps can't catch
|
|
38
|
+
# a single dropped "not". If the edit touches any of these on the original side,
|
|
39
|
+
# reject outright. (Contractions ending in "n't" are handled separately.)
|
|
40
|
+
_MEANING_CRITICAL = frozenset({
|
|
41
|
+
"not", "no", "never", "none", "nor", "neither", "without", "cannot",
|
|
42
|
+
"nothing", "nobody", "nowhere", "n't",
|
|
43
|
+
})
|
|
44
|
+
|
|
45
|
+
|
|
46
|
+
def _is_meaning_critical(word: str) -> bool:
|
|
47
|
+
w = word.lower().strip(".,;:!?\"'()")
|
|
48
|
+
return w in _MEANING_CRITICAL or w.endswith("n't")
|
|
49
|
+
|
|
37
50
|
_SYSTEM_PROMPT = (
|
|
38
51
|
"You fix automatic speech recognition errors in dictated text. "
|
|
39
52
|
"The user dictates about software engineering and their projects. "
|
|
@@ -46,18 +59,25 @@ _SYSTEM_PROMPT = (
|
|
|
46
59
|
# Disfluency mode: also strip spoken filler, delete-only. The strict rules
|
|
47
60
|
# mirror the guards — the model is told exactly what the diff check enforces.
|
|
48
61
|
_SYSTEM_PROMPT_CLEAN = (
|
|
49
|
-
"You convert dictated speech into clean written text. The user
|
|
50
|
-
"about software engineering and their projects
|
|
62
|
+
"You convert dictated speech into clean, concise written text. The user "
|
|
63
|
+
"dictates about software engineering and their projects, and wants it to "
|
|
64
|
+
"read like something they wrote, not spoke. Do exactly two things:\n"
|
|
51
65
|
"1. Fix words the recognizer misheard (wrong proper nouns, homophones, "
|
|
52
66
|
"garbled technical terms).\n"
|
|
53
|
-
"2.
|
|
54
|
-
"
|
|
55
|
-
"
|
|
56
|
-
"
|
|
57
|
-
"
|
|
58
|
-
"
|
|
59
|
-
"
|
|
60
|
-
"
|
|
67
|
+
"2. Cut spoken disfluency so it reads tight and clear:\n"
|
|
68
|
+
" - filler sounds (um, uh, er);\n"
|
|
69
|
+
" - filler discourse markers with no content: leading 'so', 'yeah', "
|
|
70
|
+
"'okay', 'well', 'now', 'right', and 'you know', 'I mean', 'like', 'sort "
|
|
71
|
+
"of'/'kind of' when used as filler, 'or something'/'or whatever', a "
|
|
72
|
+
"trailing 'right?' tag;\n"
|
|
73
|
+
" - false starts and self-corrections (keep only the corrected version);\n"
|
|
74
|
+
" - stray word and phrase repetitions.\n"
|
|
75
|
+
"Lean toward the briefer version — remove hesitation and padding freely.\n"
|
|
76
|
+
"HARD LIMITS (these protect meaning): only DELETE and FIX — NEVER add "
|
|
77
|
+
"words, rephrase, reword, reorder, or summarize. NEVER drop or alter actual "
|
|
78
|
+
"content, a negation ('not', 'never', \"n't\"), or a hedge that changes "
|
|
79
|
+
"certainty ('I think', 'maybe', 'probably', 'might'). Keep the speaker's own "
|
|
80
|
+
"wording. Return only the cleaned text, nothing else."
|
|
61
81
|
)
|
|
62
82
|
|
|
63
83
|
_MAX_RECENT = 3
|
|
@@ -315,8 +335,12 @@ class PostCorrector:
|
|
|
315
335
|
inserted += j2 - j1
|
|
316
336
|
elif tag == "replace":
|
|
317
337
|
replaced += max(i2 - i1, j2 - j1)
|
|
338
|
+
if any(_is_meaning_critical(w) for w in aw[i1:i2]):
|
|
339
|
+
return False, "negation" # e.g. "not" → something else
|
|
318
340
|
elif tag == "delete":
|
|
319
341
|
deleted += i2 - i1
|
|
342
|
+
if any(_is_meaning_critical(w) for w in aw[i1:i2]):
|
|
343
|
+
return False, "negation" # dropping "not" inverts meaning
|
|
320
344
|
n = len(aw)
|
|
321
345
|
if inserted > _MAX_INSERTED_WORDS:
|
|
322
346
|
return False, "inserted" # added content — meaning changed
|
|
@@ -14,13 +14,25 @@ if TYPE_CHECKING:
|
|
|
14
14
|
_NO_CASE_LANGUAGES = frozenset({"zh", "ja", "ko", "ar", "he", "th", "hi", "bn", "ka", "my"})
|
|
15
15
|
|
|
16
16
|
# Filler SOUNDS only — tokens with no lexical meaning, so deleting them can
|
|
17
|
-
# never change meaning.
|
|
18
|
-
#
|
|
19
|
-
#
|
|
17
|
+
# never change meaning. Deliberately conservative because this runs on every
|
|
18
|
+
# user's speech:
|
|
19
|
+
# * CASE-SENSITIVE (no IGNORECASE): all-caps abbreviations that look like
|
|
20
|
+
# fillers are real words and must survive — "ER" (emergency room), "UM"
|
|
21
|
+
# (University of Michigan), "HM". We match lowercase and Title-case forms
|
|
22
|
+
# ("um", "Um") — Whisper's filler spellings — but never all-caps.
|
|
23
|
+
# * Excluded entirely: "er"/"erm" ("ER"/"Er"=erbium), "mm" (millimetres),
|
|
24
|
+
# "ah" (interjection), bare "hm" (hectometre — require "hmm", 2+ m's).
|
|
25
|
+
# * Surrounding whitespace is horizontal-only ([^\S\n]) so a filler on its
|
|
26
|
+
# own line doesn't swallow the paragraph/list break around it.
|
|
27
|
+
# Word repetitions ("had had" is valid English) and filler "like"/"you know"
|
|
28
|
+
# need judgment and are left to the LLM layer. Order matters: multi-token
|
|
29
|
+
# "uh-huh" before "[Uu]h+" so it isn't clipped to a stray "-huh".
|
|
20
30
|
_FILLER_RE = re.compile(
|
|
21
|
-
r"\
|
|
22
|
-
re.IGNORECASE,
|
|
31
|
+
r"[^\S\n]*,?[^\S\n]*\b(?:[Uu]h[-\s]?huh|[Mm]hm|[Uu]h+m*|[Uu]m+|[Hh]m{2,})\b[^\S\n]*,?[^\S\n]*",
|
|
23
32
|
)
|
|
33
|
+
# A re-decode artifact is a whole duplicated sentence; require this many words
|
|
34
|
+
# so emphatic short repeats ("No. No.", "Stop. Stop.") are preserved.
|
|
35
|
+
_MIN_DEDUP_WORDS = 4
|
|
24
36
|
|
|
25
37
|
|
|
26
38
|
def strip_disfluencies(text: str) -> str:
|
|
@@ -36,19 +48,33 @@ def strip_disfluencies(text: str) -> str:
|
|
|
36
48
|
return text
|
|
37
49
|
text = _FILLER_RE.sub(" ", text)
|
|
38
50
|
text = _dedup_adjacent_sentences(text)
|
|
39
|
-
|
|
51
|
+
# Repair the debris the deletions leave, without crossing newlines (so
|
|
52
|
+
# paragraph/list structure survives even when punctuation_cleanup is off).
|
|
53
|
+
text = re.sub(r"[^\S\n]+([,.;:?!])", r"\1", text) # space before punctuation
|
|
54
|
+
text = re.sub(r"[^\S\n]{2,}", " ", text) # collapse runs of spaces
|
|
40
55
|
return text.strip()
|
|
41
56
|
|
|
42
57
|
|
|
43
58
|
def _dedup_adjacent_sentences(text: str) -> str:
|
|
44
|
-
"""Drop a sentence identical to the one immediately before it.
|
|
45
|
-
|
|
59
|
+
"""Drop a sentence identical to the one immediately before it.
|
|
60
|
+
|
|
61
|
+
Separators are captured and preserved on rejoin so paragraph/list breaks
|
|
62
|
+
survive; only a full (>= _MIN_DEDUP_WORDS) verbatim repeat is removed.
|
|
63
|
+
"""
|
|
64
|
+
tokens = re.split(r"((?<=[.?!])\s+)", text) # [sent, sep, sent, sep, …]
|
|
46
65
|
out: list[str] = []
|
|
47
|
-
|
|
48
|
-
|
|
49
|
-
|
|
50
|
-
|
|
51
|
-
|
|
66
|
+
last_kept: str | None = None
|
|
67
|
+
for i in range(0, len(tokens), 2):
|
|
68
|
+
sentence = tokens[i]
|
|
69
|
+
sep = tokens[i + 1] if i + 1 < len(tokens) else ""
|
|
70
|
+
norm = sentence.strip().lower()
|
|
71
|
+
if (last_kept is not None and norm == last_kept
|
|
72
|
+
and len(sentence.split()) >= _MIN_DEDUP_WORDS):
|
|
73
|
+
continue # drop the duplicate sentence and its separator
|
|
74
|
+
out.append(sentence)
|
|
75
|
+
out.append(sep)
|
|
76
|
+
last_kept = norm
|
|
77
|
+
return "".join(out)
|
|
52
78
|
|
|
53
79
|
|
|
54
80
|
def cleanup(text: str, language: str = "en") -> str:
|
|
@@ -146,7 +146,7 @@ def _read_stdout(proc: subprocess.Popen, toggle_cb: Callable[[], None]) -> None:
|
|
|
146
146
|
action = cmd[5:]
|
|
147
147
|
cli_cmd = _MENU_COMMANDS.get(action)
|
|
148
148
|
if cli_cmd:
|
|
149
|
-
if not open_in_terminal(cli_cmd):
|
|
149
|
+
if not open_in_terminal(cli_cmd, hold=True):
|
|
150
150
|
log.warning("Failed to open terminal for: %s", action)
|
|
151
151
|
except (OSError, ValueError):
|
|
152
152
|
pass
|
|
@@ -1 +0,0 @@
|
|
|
1
|
-
__version__ = "0.9.6"
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|
|
File without changes
|