<?xml version="1.0" encoding="utf-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
    <channel>
        <title>KV cache streaming - the llama.cpp fork that enables Qwen 3.8 27B at large contexts for 16GB VRAM GPU owners</title>
        <link>https://video.troed.se/videos/watch/fb44555a-b42a-4170-a79f-3e058130c54a</link>
        <description>Raymond's blog entry: https://medium.com/@raymond860909/running-qwen-27b-on-16g-vram-with-full-context-length-building-adaptive-kv-cache-streaming-for-bf1e819116e9 Their source repo: https://github.com/RaymondHuang210129/llama.cpp-adaptive-kv-streaming Opencode plugin for TG/PP display: https://git.sync.wtf/troed/oc-ls-stats Yes I'm lying down: https://levus.co/</description>
        <lastBuildDate>Sat, 29 Aug 2026 14:39:40 GMT</lastBuildDate>
        <docs>https://validator.w3.org/feed/docs/rss2.html</docs>
        <generator>PeerTube - https://video.troed.se</generator>
        <image>
            <title>KV cache streaming - the llama.cpp fork that enables Qwen 3.8 27B at large contexts for 16GB VRAM GPU owners</title>
            <url>https://video.troed.se/lazy-static/avatars/f128c4cb-0bf7-4ea9-b718-8e51171ac9ce.png</url>
            <link>https://video.troed.se/videos/watch/fb44555a-b42a-4170-a79f-3e058130c54a</link>
        </image>
        <copyright>All rights reserved, unless otherwise specified in the terms specified at https://video.troed.se/about and potential licenses granted by each content's rightholder.</copyright>
        <atom:link href="https://video.troed.se/feeds/video-comments.xml?videoId=fb44555a-b42a-4170-a79f-3e058130c54a" rel="self" type="application/rss+xml"/>
    </channel>
</rss>