@misc{indiciaeef45f731feb7, title = {Acoustically Grounded Cost Learning for Open-Vocabulary Audio-Visual Semantic Segmentation}, author = {Tianrui Hui and Shaofei Huang and Qisong Han and Yaxiong Wang and Lechao Cheng and Zhedong Zheng and Zhun Zhong and Richang Hong and Meng Wang}, year = {2026}, url = {https://arxiv.org/abs/2608.29121}, note = {Source identifier: 2608.29121} }