{
    "embed_dim": 1024,
    "vision_cfg": {
        "image_size": 336,
        "layers": 40,
        "width": 1408,
        "head_width": 88,
        "patch_size": 14,
        "mlp_ratio": 4.3637,
        "drop_path_rate": 0.4
    },
    "text_cfg": {
        "context_length": 77,
        "vocab_size": 49408,
        "width": 768,
        "heads": 12,
        "layers": 12
    }
}
