@misc{indiciae8e59a6f2f707, title = {Semantic2Graph: Graph-based Multi-modal Feature Fusion for Action Segmentation in Videos}, author = {Junbin Zhang and Pei-Hsuan Tsai and Meng-Hsun Tsai}, year = {2024}, doi = {10.1007/s10489-023-05259-z}, url = {https://arxiv.org/abs/2209.05653}, note = {Source identifier: 2209.05653} }