@misc{indiciae143355e27f80, title = {ROVI: A VLM-LLM Re-Captioned Dataset for Open-Vocabulary Instance-Grounded Text-to-Image Generation}, author = {Cihang Peng and Qiming Hou and Zhong Ren and Kun Zhou}, year = {2025}, url = {https://arxiv.org/abs/2508.01008}, note = {Source identifier: 2508.01008} }