<?xml version="1.0" encoding="UTF-8"?><urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9" xmlns:xhtml="http://www.w3.org/1999/xhtml" xmlns:image="http://www.google.com/schemas/sitemap-image/1.1"><url><loc>https://kjallari02.com/blog</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url><url><loc>https://kjallari02.com</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>1.0</priority></url><url><loc>https://kjallari02.com/about</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url><url><loc>https://kjallari02.com/evaluating-long-context-retrieval-beyond-needle-in-a-haystack</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url><url><loc>https://kjallari02.com/why-mixture-of-experts-models-struggle-with-memory-bandwidth</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url><url><loc>https://kjallari02.com/quantization-trade-offs-in-local-model-serving-pipelines</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url><url><loc>https://kjallari02.com/page</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url><url><loc>https://kjallari02.com/contact</loc><lastmod>2026-07-25T02:16:35.000Z</lastmod><priority>0.5</priority></url></urlset>