@misc{indiciae8f33e9124b7d, title = {VILA: Learning Image Aesthetics from User Comments with Vision-Language Pretraining}, author = {Junjie Ke and Keren Ye and Jiahui Yu and Yonghui Wu and Peyman Milanfar and Feng Yang}, year = {2023}, url = {https://arxiv.org/abs/2303.14302}, note = {Source identifier: 2303.14302} }