File size: 924 Bytes
039f477
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
{
    "text_encoder": {
        "tokenizer_class": "bert",
        "model_type": "bert",
        "dim": 768,
        "context_dim": 768,
        "vocab_size": 30522,
        "padding_idx": 0,
        "num_layers": 4,
        "num_heads": 12,
        "embedding_dim": 256,
        "multimodal_layers_ids": [
            2,
            3
        ],
        "head_one_neuron": false,
        "pooling": "cls",
        "max_position_embeddings": 77,
        "dropout_prob": 0.1
    },
    "image_encoder": {
        "normalization_means": [
            0.48145466,
            0.4578275,
            0.40821073
        ],
        "normalization_deviations": [
            0.26862954,
            0.26130258,
            0.27577711
        ],
        "dim": 768,
        "patch_size": 16,
        "image_size": 224,
        "num_layers": 12,
        "num_heads": 12,
        "embedding_dim": 256,
        "pooling": "cls"
    }
}