@misc{indiciae52e5bd165d1c, title = {LLaVA-SpaceSGG: Visual Instruct Tuning for Open-vocabulary Scene Graph Generation with Enhanced Spatial Relations}, author = {Mingjie Xu and Mengyang Wu and Yuzhi Zhao and Jason Chun Lok Li and Weifeng Ou}, year = {2024}, url = {https://arxiv.org/abs/2412.06322}, note = {Source identifier: 2412.06322} }