first commit

2025-01-21 13:55:11 +08:00 · 2025-01-21 13:55:11 +08:00 · 8daf9bcd00
parent 559b22fdc7
commit 8daf9bcd00
6 changed files with 8330 additions and 0 deletions
--- a/config.json
+++ b/config.json
--- a/preprocessor_config.json
+++ b/preprocessor_config.json
@ -0,0 +1,17 @@
 {
  "do_normalize": true,
  "do_resize": true,
  "feature_extractor_type": "ViTFeatureExtractor",
  "image_mean": [
    0.5,
    0.5,
    0.5
  ],
  "image_std": [
    0.5,
    0.5,
    0.5
  ],
  "resample": 2,
  "size": 224
 }
--- a/pytorch_model.bin
+++ b/pytorch_model.bin
--- a/special_tokens_map.json
+++ b/special_tokens_map.json
@ -0,0 +1 @@
 {"unk_token": "[UNK]", "sep_token": "[SEP]", "pad_token": "[PAD]", "cls_token": "[CLS]", "mask_token": "[MASK]"}
--- a/tokenizer_config.json
+++ b/tokenizer_config.json
@ -0,0 +1 @@
 {"unk_token": "[UNK]", "sep_token": "[SEP]", "pad_token": "[PAD]", "cls_token": "[CLS]", "mask_token": "[MASK]", "do_lower_case": false, "do_word_tokenize": true, "do_subword_tokenize": true, "word_tokenizer_type": "mecab", "subword_tokenizer_type": "character", "never_split": null, "mecab_kwargs": {"mecab_dic": "unidic_lite"}, "special_tokens_map_file": null, "tokenizer_file": null, "name_or_path": "cl-tohoku/bert-base-japanese-char-v2", "tokenizer_class": "BertJapaneseTokenizer"}
--- a/vocab.txt
+++ b/vocab.txt
		`@ -0,0 +1 @@`
							`{"unk_token": "[UNK]", "sep_token": "[SEP]", "pad_token": "[PAD]", "cls_token": "[CLS]", "mask_token": "[MASK]"}`