@misc{indiciaecc2c13e4667e, title = {Multimodal Large Language Model is a Human-Aligned Annotator for Text-to-Image Generation}, author = {Xun Wu and Shaohan Huang and Furu Wei}, year = {2024}, url = {https://arxiv.org/abs/2404.15100}, note = {Source identifier: 2404.15100} }