@misc{indiciae21734e50938a, title = {Feedback-Driven Vision-Language Alignment with Minimal Human Supervision}, author = {Giorgio Giannone and Ruoteng Li and Qianli Feng and Evgeny Perevodchikov and Rui Chen and Aleix Martinez}, year = {2025}, url = {https://arxiv.org/abs/2501.04568}, note = {Source identifier: 2501.04568} }