Narrow AutoTokenizer type for transformers 5
This commit is contained in:
parent
c9466af986
commit
d9cc9a4eea
1 changed files with 2 additions and 1 deletions
|
|
@ -33,6 +33,7 @@ async def test_local_chunker(qa_corpus: list[dict[str, str]]):
|
||||||
|
|
||||||
# Load tokenizer for verification
|
# Load tokenizer for verification
|
||||||
tokenizer = AutoTokenizer.from_pretrained(chunker.tokenizer_name)
|
tokenizer = AutoTokenizer.from_pretrained(chunker.tokenizer_name)
|
||||||
|
assert tokenizer is not None
|
||||||
|
|
||||||
# Ensure that chunks are reasonably sized (allowing more flexibility for structure-aware chunking)
|
# Ensure that chunks are reasonably sized (allowing more flexibility for structure-aware chunking)
|
||||||
total_tokens = 0
|
total_tokens = 0
|
||||||
|
|
@ -74,7 +75,7 @@ async def test_local_chunker_runs_off_event_loop_thread():
|
||||||
|
|
||||||
original = chunker._chunk_sync
|
original = chunker._chunk_sync
|
||||||
|
|
||||||
def recording_chunk_sync(self, document):
|
def recording_chunk_sync(_self, document):
|
||||||
called_from.append(threading.current_thread())
|
called_from.append(threading.current_thread())
|
||||||
return original(document)
|
return original(document)
|
||||||
|
|
||||||
|
|
|
||||||
Loading…
Reference in a new issue