shivash commited on
Commit
65be92d
·
verified ·
1 Parent(s): af12e95

Upload tokenizer_config.json with huggingface_hub

Browse files
Files changed (1) hide show
  1. tokenizer_config.json +24 -0
tokenizer_config.json ADDED
@@ -0,0 +1,24 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "add_bos_token": false,
3
+ "add_eos_token": false,
4
+ "add_prefix_space": false,
5
+ "bos_token": "<|endoftext|>",
6
+ "clean_up_tokenization_spaces": true,
7
+ "eos_token": "<|endoftext|>",
8
+ "errors": "replace",
9
+ "model_max_length": 4096,
10
+ "pad_token": "<|endoftext|>",
11
+ "tokenizer_class": "AutoTokenizer",
12
+ "unk_token": "<|endoftext|>",
13
+ "auto_map": {
14
+ "AutoTokenizer": [
15
+ "transformers",
16
+ "AutoTokenizer"
17
+ ]
18
+ },
19
+ "_tokenizer_fallbacks": [
20
+ "GPT2Tokenizer",
21
+ "LlamaTokenizer",
22
+ "PreTrainedTokenizerFast"
23
+ ]
24
+ }