@misc{indiciae5cfc7760a84b, title = {Part-Whole Relational Fusion Towards Multi-Modal Scene Understanding}, author = {Yi Liu and Chengxin Li and Shoukun Xu and Jungong Han}, year = {2024}, url = {https://arxiv.org/abs/2410.14944}, note = {Source identifier: 2410.14944} }