<?xml version="1.0" encoding="UTF-8"?>
<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9"
        xmlns:news="http://www.google.com/schemas/sitemap-news/0.9">
  <url>
    <loc>https://ainews.surl.tw/article/puda-以受限-cli-與事件日誌連接實驗室代理-將生成式規劃隔離於硬體驅動層-81a0182d</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T16:02:55.113Z</news:publication_date>
      <news:title>PUDA 以受限 CLI 與事件日誌連接實驗室代理，將生成式規劃隔離於硬體驅動層</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/capa-測試程式助理的跨工作階段記憶-同一使用者歷史讓首次成功率平均提高-15-6-點-b07c4a7e</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T16:02:01.232Z</news:publication_date>
      <news:title>CAPA 測試程式助理的跨工作階段記憶：同一使用者歷史讓首次成功率平均提高 15.6 點</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/trek-以可執行規則驗證旅遊代理-最強模型完整可行率僅-46-2-529999d6</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T12:03:24.798Z</news:publication_date>
      <news:title>TREK 以可執行規則驗證旅遊代理，最強模型完整可行率僅 46.2%</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/secrespond-重建入侵後鑑識流程-23-款模型無一完整偵測並修復單一靶場-705dbe43</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T14:02:51.802Z</news:publication_date>
      <news:title>SecRespond 重建入侵後鑑識流程，23 款模型無一完整偵測並修復單一靶場</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/cam-df-把代理工具排序改成成本停止決策-工具暴露量減少-37-14901d01</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T12:02:08.247Z</news:publication_date>
      <news:title>CAM-DF 把代理工具排序改成成本停止決策，工具暴露量減少 37%</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/擴散蒸餾只對齊最終-cfg-速度可能掩蓋分支錯誤-pdm-改為分別約束正向預測與條件方向-ef5d2544</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T10:02:59.242Z</news:publication_date>
      <news:title>擴散蒸餾只對齊最終 CFG 速度可能掩蓋分支錯誤，PDM 改為分別約束正向預測與條件方向</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/memtx-把代理共享記憶改成交易式提交-窮舉-550-萬個協定狀態未見安全規則失效-71276980</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T10:01:44.850Z</news:publication_date>
      <news:title>MemTX 把代理共享記憶改成交易式提交，窮舉 550 萬個協定狀態未見安全規則失效</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/omegause-officeval-為辦公代理加入逐任務成本標籤-最佳模型品質仍落後人類逾三分之一-bbfb5dfe</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T08:02:17.057Z</news:publication_date>
      <news:title>OmegaUse-OfficeVal 為辦公代理加入逐任務成本標籤，最佳模型品質仍落後人類逾三分之一</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/arex-讓研究代理反覆驗證未解條件-4b-模型以壓縮狀態維持長程搜尋-4a389b89</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T04:03:20.727Z</news:publication_date>
      <news:title>AREX 讓研究代理反覆驗證未解條件，4B 模型以壓縮狀態維持長程搜尋</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/ace-rtl-讓-nemotron-反覆編譯與修正-verilog-九類任務平均通過率達-97-1-d46c9474</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T06:03:27.909Z</news:publication_date>
      <news:title>ACE-RTL 讓 Nemotron 反覆編譯與修正 Verilog，九類任務平均通過率達 97.1%</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/poolside-開放-laguna-s-2-1-118b-moe-每個-token-僅啟用-8-5b-參數-0ae9b8cf</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T04:02:18.648Z</news:publication_date>
      <news:title>Poolside 開放 Laguna S 2.1：118B MoE 每個 token 僅啟用 8.5B 參數</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/同一模型兼任工具呼叫預測器-qwen-4b-下一步命中率提高逾-17-點-97c1f6e4</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T02:01:57.800Z</news:publication_date>
      <news:title>同一模型兼任工具呼叫預測器，Qwen 4B 下一步命中率提高逾 17 點</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/skillgate-以規則預篩代理技能-llm-輸入量減少-77-但仍漏掉近四分之一攻擊-bc553366</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T22:03:17.758Z</news:publication_date>
      <news:title>SkillGate 以規則預篩代理技能，LLM 輸入量減少 77% 但仍漏掉近四分之一攻擊</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/shieldstral-將審核政策寫進查詢-3b-多模態模型在文字安全評測平均-f1-達-84-9-9b14b836</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-30T00:02:31.505Z</news:publication_date>
      <news:title>Shieldstral 將審核政策寫進查詢，3B 多模態模型在文字安全評測平均 F1 達 84.9%</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/實測-1-723-個-mcp-應用-近三分之二未在工具執行前要求批准-d2649269</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T22:02:08.236Z</news:publication_date>
      <news:title>實測 1,723 個 MCP 應用：近三分之二未在工具執行前要求批准</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/handbook-md-以百頁企業規章約束代理-30-組模型配置最高嚴格通過率僅-36-2-35a53c90</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T20:02:16.224Z</news:publication_date>
      <news:title>HANDBOOK.md 以百頁企業規章約束代理，30 組模型配置最高嚴格通過率僅 36.2%</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/patientagentbench-以-1-200-場工具對話測試醫療代理-最強模型分流通過率仍僅-88-bff1c3dc</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T18:02:05.727Z</news:publication_date>
      <news:title>PatientAgentBench 以 1,200 場工具對話測試醫療代理，最強模型分流通過率仍僅 88%</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/分層-lora-實驗顯示-模型要學詞彙-事實與行為政策-最佳更新位置並不相同-2389a66a</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T18:02:55.313Z</news:publication_date>
      <news:title>分層 LoRA 實驗顯示：模型要學詞彙、事實與行為政策，最佳更新位置並不相同</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/要求-json-就讓-44-款模型答案趨同-結構化輸出可能改變模型判斷-a2ff1ef1</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T16:03:20.587Z</news:publication_date>
      <news:title>要求 JSON 就讓 44 款模型答案趨同，結構化輸出可能改變模型判斷</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/github-models-7-月-30-日全面關閉-推論-api-與-byok-端點同步失效-f7427198</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T16:02:15.977Z</news:publication_date>
      <news:title>GitHub Models 7 月 30 日全面關閉，推論 API 與 BYOK 端點同步失效</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/ninfer-為兩款-qwen3-6-寫死推論路徑-單張-rtx-5090-長解碼達每秒-543-token-7cd66a68</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T14:03:15.870Z</news:publication_date>
      <news:title>NInfer 為兩款 Qwen3.6 寫死推論路徑，單張 RTX 5090 長解碼達每秒 543 token</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/mlperf-endpoints-0-7-將雲端推論改成持續投稿-首版仍缺成本正規化-039d8f34</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T14:02:16.824Z</news:publication_date>
      <news:title>MLPerf Endpoints 0.7 將雲端推論改成持續投稿，首版仍缺成本正規化</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/specbox-在模型生成期間預熱代理沙箱-p99-延遲最高縮短至三分之一-e1f8adca</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T10:02:56.120Z</news:publication_date>
      <news:title>SpecBox 在模型生成期間預熱代理沙箱，P99 延遲最高縮短至三分之一</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/cotinyvla-以-9-億參數超越-7b-機器人基線-但實驗仍停留在模擬環境-4e398cd0</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T12:02:59.062Z</news:publication_date>
      <news:title>CoTinyVLA 以 9 億參數超越 7B 機器人基線，但實驗仍停留在模擬環境</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/jarvishub-將畫布變成多模態代理的共享狀態-但首版仍缺量化評測-123546d8</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T08:04:07.723Z</news:publication_date>
      <news:title>JarvisHub 將畫布變成多模態代理的共享狀態，但首版仍缺量化評測</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/sol-attn-在-online-softmax-內動態稀疏化注意力-影片生成端到端加速逾-2-倍-43008da4</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T08:02:51.138Z</news:publication_date>
      <news:title>Sol-Attn 在 online softmax 內動態稀疏化注意力，影片生成端到端加速逾 2 倍</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/claude-opus-5-以可調推理強度換取成本-雲端版本支援百萬-token-輸入-af588bf7</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T04:02:03.539Z</news:publication_date>
      <news:title>Claude Opus 5 以可調推理強度換取成本，雲端版本支援百萬 token 輸入</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/nvidia-dynamo-為-kimi-k3-提供多節點部署配方-百萬-token-推論至少需-16-顆-gb200-gb300-061a4d11</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T06:02:28.208Z</news:publication_date>
      <news:title>NVIDIA Dynamo 為 Kimi K3 提供多節點部署配方，百萬 token 推論至少需 16 顆 GB200／GB300</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/微軟以小型資安模型承接九成掃描-mdash-僅將難題升級至-gpt-5-4-9d4c27c8</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T02:01:44.041Z</news:publication_date>
      <news:title>微軟以小型資安模型承接九成掃描，MDASH 僅將難題升級至 GPT-5.4</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/9b-模型以-500-美元-grpo-學會目錄審查-單項工作流超越前沿-api-75e3bca4</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T02:02:30.258Z</news:publication_date>
      <news:title>9B 模型以 500 美元 GRPO 學會目錄審查，單項工作流超越前沿 API</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/grok-4-5-接入-github-copilot-500k-上下文與推理強度可跨-ide-使用-5035a040</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T00:03:26.699Z</news:publication_date>
      <news:title>Grok 4.5 接入 GitHub Copilot，500K 上下文與推理強度可跨 IDE 使用</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/28-款模型的-17-萬次投票顯示-增加代理數量無法消除共同錯誤-83765f51</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-29T00:02:14.689Z</news:publication_date>
      <news:title>28 款模型的 17 萬次投票顯示：增加代理數量無法消除共同錯誤</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/一個語尾詞讓模型同意率擺盪-64-點-反諂媚訓練可能只學到表面句型-df346c5d</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-28T22:01:45.264Z</news:publication_date>
      <news:title>一個語尾詞讓模型同意率擺盪 64 點，反諂媚訓練可能只學到表面句型</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/erunderstand-將-er-圖轉成結構化測試-vlm-在多元關係與弱實體上明顯失準-3946322b</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-28T20:03:14.879Z</news:publication_date>
      <news:title>ERUnderstand 將 ER 圖轉成結構化測試，VLM 在多元關係與弱實體上明顯失準</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://ainews.surl.tw/article/planphys-拆解長程代理訓練-稀疏獎勵下-on-policy-蒸餾比-grpo-更穩定-21d5f267</loc>
    <news:news>
      <news:publication>
        <news:name>SURL AI News</news:name>
        <news:language>zh-tw</news:language>
      </news:publication>
      <news:publication_date>2026-07-28T20:01:57.762Z</news:publication_date>
      <news:title>PlanPhys 拆解長程代理訓練：稀疏獎勵下，on-policy 蒸餾比 GRPO 更穩定</news:title>
    </news:news>
  </url>
</urlset>
