@misc{indiciaeaa876e19de7b, title = {Benchmarking Large Vision-Language Models via Directed Scene Graph for Comprehensive Image Captioning}, author = {Fan Lu and Wei Wu and Kecheng Zheng and Shuailei Ma and Biao Gong and Jiawei Liu and Wei Zhai and Yang Cao and Yujun Shen and Zheng-Jun Zha}, year = {2025}, url = {https://arxiv.org/abs/2412.08614}, note = {Source identifier: 2412.08614} }