<?xml version="1.0" encoding="UTF-8"?>
<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9" xmlns:news="http://www.google.com/schemas/sitemap-news/0.9">
  <url>
    <loc>https://bitbytecore.com/article/quantization-explained-what-q4-q8-and-fp16-actually-do-to-a-local-model-2</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-06T05:58:00.000Z</news:publication_date>
      <news:title>Quantization Explained: What Q4, Q8, and FP16 Actually Do to a Local Model</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/cuda-lock-in-is-real-a-precise-cost-accounting-of-what-switching-gpu-vendors-actually-breaks</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-06T05:57:00.000Z</news:publication_date>
      <news:title>CUDA Lock-In Is Real: A Precise Cost Accounting of What Switching GPU Vendors Actually Breaks</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/how-robots-are-really-trained-the-sim-to-real-gap-is-not-a-bug-you-can-patch</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-06T05:56:00.000Z</news:publication_date>
      <news:title>How Robots Are Really Trained: The Sim-to-Real Gap Is Not a Bug You Can Patch</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/mixture-of-experts-models-how-they-work-and-why-they-cut-inference-costs</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-06T05:55:00.000Z</news:publication_date>
      <news:title>Mixture-of-Experts Models: How They Work and Why They Cut Inference Costs</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/the-real-difference-between-mcp-function-calling-and-agent-loops</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-05T05:58:00.000Z</news:publication_date>
      <news:title>The Real Difference Between MCP, Function Calling, and Agent Loops</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/how-much-vram-you-actually-need-to-run-a-local-llm</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-05T05:57:00.000Z</news:publication_date>
      <news:title>How Much VRAM You Actually Need to Run a Local LLM</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/the-real-privacy-audit-what-data-your-ai-coding-assistant-sends-home</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-05T05:56:00.000Z</news:publication_date>
      <news:title>The Real Privacy Audit: What Data Your AI Coding Assistant Sends Home</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://bitbytecore.com/article/how-to-evaluate-a-local-llm-for-a-real-task-a-repeatable-testing-framework</loc>
    <news:news>
      <news:publication>
        <news:name>BitByteCore</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-08-05T05:55:00.000Z</news:publication_date>
      <news:title>How to Evaluate a Local LLM for a Real Task: A Repeatable Testing Framework</news:title>
    </news:news>
  </url>
</urlset>