@misc{indiciaee9bcccfbdf41, title = {Can LLM-as-a-Judge Reliably Verify Rubrics in Agentic Scenarios?}, author = {Yangda Peng and Yunjia Qi and Haotian Xia and Guanzhong He and Xintong Shi and Richeng Xuan and Songyuanyi Lu and Yixian Liu and Zhichao Hu and Yuhong Liu and Hao Peng}, year = {2026}, url = {https://arxiv.org/abs/2606.29920}, note = {Source identifier: 2606.29920} }