@misc{indiciaefe782434bd11, title = {ViTamin: Designing Scalable Vision Models in the Vision-Language Era}, author = {Jieneng Chen and Qihang Yu and Xiaohui Shen and Alan Yuille and Liang-Chieh Chen}, year = {2024}, url = {https://arxiv.org/abs/2404.02132}, note = {Source identifier: 2404.02132} }