Merge pull request #1087 from huggingface/fix-warnings

Decode now calls private property instead of public method
This commit is contained in:
Thomas Wolf 2019-08-28 22:22:11 +02:00 committed by GitHub
commit 5f297c7be3
No known key found for this signature in database
GPG Key ID: 4AEE18F83AFDEB23

View File

@ -641,9 +641,9 @@ class PreTrainedTokenizer(object):
filtered_tokens = self.convert_ids_to_tokens(token_ids, skip_special_tokens=skip_special_tokens)
text = self.convert_tokens_to_string(filtered_tokens)
if self.sep_token is not None and self.sep_token in text:
text = text.replace(self.cls_token, self.sep_token)
split_text = list(filter(lambda sentence: len(sentence) > 0, text.split(self.sep_token)))
if self._sep_token is not None and self._sep_token in text:
text = text.replace(self._cls_token, self._sep_token)
split_text = list(filter(lambda sentence: len(sentence) > 0, text.split(self._sep_token)))
if clean_up_tokenization_spaces:
clean_text = [self.clean_up_tokenization(text) for text in split_text]
return clean_text