Spaces:
Runtime error
Runtime error
Commit ·
f9732d0
1
Parent(s): d30ef31
Upload 2 files
Browse files- segments_test.py +48 -0
- vad_test.py +66 -0
segments_test.py
ADDED
|
@@ -0,0 +1,48 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
import sys
|
| 2 |
+
import unittest
|
| 3 |
+
|
| 4 |
+
sys.path.append('../whisper-webui')
|
| 5 |
+
|
| 6 |
+
from src.segments import merge_timestamps
|
| 7 |
+
|
| 8 |
+
class TestSegments(unittest.TestCase):
|
| 9 |
+
def __init__(self, *args, **kwargs):
|
| 10 |
+
super(TestSegments, self).__init__(*args, **kwargs)
|
| 11 |
+
|
| 12 |
+
def test_merge_segments(self):
|
| 13 |
+
segments = [
|
| 14 |
+
{'start': 10.0, 'end': 20.0},
|
| 15 |
+
{'start': 22.0, 'end': 27.0},
|
| 16 |
+
{'start': 31.0, 'end': 35.0},
|
| 17 |
+
{'start': 45.0, 'end': 60.0},
|
| 18 |
+
{'start': 61.0, 'end': 65.0},
|
| 19 |
+
{'start': 68.0, 'end': 98.0},
|
| 20 |
+
{'start': 100.0, 'end': 102.0},
|
| 21 |
+
{'start': 110.0, 'end': 112.0}
|
| 22 |
+
]
|
| 23 |
+
|
| 24 |
+
result = merge_timestamps(segments, merge_window=5, max_merge_size=30, padding_left=1, padding_right=1)
|
| 25 |
+
|
| 26 |
+
self.assertListEqual(result, [
|
| 27 |
+
{'start': 9.0, 'end': 36.0},
|
| 28 |
+
{'start': 44.0, 'end': 66.0},
|
| 29 |
+
{'start': 67.0, 'end': 99.0},
|
| 30 |
+
{'start': 99.0, 'end': 103.0},
|
| 31 |
+
{'start': 109.0, 'end': 113.0}
|
| 32 |
+
])
|
| 33 |
+
|
| 34 |
+
def test_overlap_next(self):
|
| 35 |
+
segments = [
|
| 36 |
+
{'start': 5.0, 'end': 39.182},
|
| 37 |
+
{'start': 39.986, 'end': 40.814}
|
| 38 |
+
]
|
| 39 |
+
|
| 40 |
+
result = merge_timestamps(segments, merge_window=5, max_merge_size=30, padding_left=1, padding_right=1)
|
| 41 |
+
|
| 42 |
+
self.assertListEqual(result, [
|
| 43 |
+
{'start': 4.0, 'end': 39.584},
|
| 44 |
+
{'start': 39.584, 'end': 41.814}
|
| 45 |
+
])
|
| 46 |
+
|
| 47 |
+
if __name__ == '__main__':
|
| 48 |
+
unittest.main()
|
vad_test.py
ADDED
|
@@ -0,0 +1,66 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
import pprint
|
| 2 |
+
import unittest
|
| 3 |
+
import numpy as np
|
| 4 |
+
import sys
|
| 5 |
+
|
| 6 |
+
sys.path.append('../whisper-webui')
|
| 7 |
+
|
| 8 |
+
from src.vad import AbstractTranscription, VadSileroTranscription
|
| 9 |
+
|
| 10 |
+
class TestVad(unittest.TestCase):
|
| 11 |
+
def __init__(self, *args, **kwargs):
|
| 12 |
+
super(TestVad, self).__init__(*args, **kwargs)
|
| 13 |
+
self.transcribe_calls = []
|
| 14 |
+
|
| 15 |
+
def test_transcript(self):
|
| 16 |
+
mock = MockVadTranscription()
|
| 17 |
+
|
| 18 |
+
self.transcribe_calls.clear()
|
| 19 |
+
result = mock.transcribe("mock", lambda segment : self.transcribe_segments(segment))
|
| 20 |
+
|
| 21 |
+
self.assertListEqual(self.transcribe_calls, [
|
| 22 |
+
[30, 30],
|
| 23 |
+
[100, 100]
|
| 24 |
+
])
|
| 25 |
+
|
| 26 |
+
self.assertListEqual(result['segments'],
|
| 27 |
+
[{'end': 50.0, 'start': 40.0, 'text': 'Hello world '},
|
| 28 |
+
{'end': 120.0, 'start': 110.0, 'text': 'Hello world '}]
|
| 29 |
+
)
|
| 30 |
+
|
| 31 |
+
def transcribe_segments(self, segment):
|
| 32 |
+
self.transcribe_calls.append(segment.tolist())
|
| 33 |
+
|
| 34 |
+
# Dummy text
|
| 35 |
+
return {
|
| 36 |
+
'text': "Hello world ",
|
| 37 |
+
'segments': [
|
| 38 |
+
{
|
| 39 |
+
"start": 10.0,
|
| 40 |
+
"end": 20.0,
|
| 41 |
+
"text": "Hello world "
|
| 42 |
+
}
|
| 43 |
+
],
|
| 44 |
+
'language': ""
|
| 45 |
+
}
|
| 46 |
+
|
| 47 |
+
class MockVadTranscription(AbstractTranscription):
|
| 48 |
+
def __init__(self):
|
| 49 |
+
super().__init__()
|
| 50 |
+
|
| 51 |
+
def get_audio_segment(self, str, start_time: str = None, duration: str = None):
|
| 52 |
+
start_time_seconds = float(start_time.removesuffix("s"))
|
| 53 |
+
duration_seconds = float(duration.removesuffix("s"))
|
| 54 |
+
|
| 55 |
+
# For mocking, this just returns a simple numppy array
|
| 56 |
+
return np.array([start_time_seconds, duration_seconds], dtype=np.float64)
|
| 57 |
+
|
| 58 |
+
def get_transcribe_timestamps(self, audio: str):
|
| 59 |
+
result = []
|
| 60 |
+
|
| 61 |
+
result.append( { 'start': 30, 'end': 60 } )
|
| 62 |
+
result.append( { 'start': 100, 'end': 200 } )
|
| 63 |
+
return result
|
| 64 |
+
|
| 65 |
+
if __name__ == '__main__':
|
| 66 |
+
unittest.main()
|