{ | |
"model_cfg": { | |
"embed_dim": 1024, | |
"init_logit_bias": -10, | |
"custom_text": true, | |
"vision_cfg": { | |
"image_size": 256, | |
"timm_model_name": "vit_large_patch16_siglip_256", | |
"timm_model_pretrained": false, | |
"timm_pool": "map", | |
"timm_proj": "none" | |
}, | |
"text_cfg": { | |
"context_length": 64, | |
"vocab_size": 32000, | |
"tokenizer_kwargs": { | |
"clean": "canonicalize" | |
}, | |
"width": 1024, | |
"heads": 16, | |
"layers": 24, | |
"no_causal_mask": true, | |
"proj_bias": true, | |
"pool_type": "last", | |
"norm_kwargs": { | |
"eps": 1e-06 | |
}, | |
"hf_tokenizer_name": "timm/ViT-L-16-SigLIP-256" | |
} | |
}, | |
"preprocess_cfg": { | |
"mean": [ | |
0.5, | |
0.5, | |
0.5 | |
], | |
"std": [ | |
0.5, | |
0.5, | |
0.5 | |
], | |
"interpolation": "bicubic", | |
"resize_mode": "squash" | |
} | |
} |