@misc{indiciaeb28a96f30aec, title = {RelationVLM: Making Large Vision-Language Models Understand Visual Relations}, author = {Zhipeng Huang and Zhizheng Zhang and Zheng-Jun Zha and Yan Lu and Baining Guo}, year = {2024}, url = {https://arxiv.org/abs/2403.12801}, note = {Source identifier: 2403.12801} }