kirp commited on
Commit
251efdc
1 Parent(s): 0cc0520

Upload 3 files

Browse files
Files changed (3) hide show
  1. tokenizer.model +3 -0
  2. tokenizer.vocab +0 -0
  3. tokenizer_config.json +12 -0
tokenizer.model ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e2676d4ca29ca1750f6ff203328d73b189321dc5776ceede037cbd36541d70c0
3
+ size 757958
tokenizer.vocab ADDED
The diff for this file is too large to render. See raw diff
 
tokenizer_config.json ADDED
@@ -0,0 +1,12 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token": "<s>",
3
+ "eos_token": "</s>",
4
+ "model_max_length": 1000000000000000019884624838656,
5
+ "tokenizer_class": "LlamaTokenizer",
6
+ "unk_token": "<unk>",
7
+ "sep_token": "<sep>",
8
+ "pad_token": "<pad>",
9
+ "cls_token": "<cls>",
10
+ "mask_token": "<mask>",
11
+ "additional_special_tokens": ["[QUE]", "[DESC]", "[KWD]", "[ANS]"]
12
+ }