<?xml version="1.0" encoding="UTF-8"?>
<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9"
        xmlns:news="http://www.google.com/schemas/sitemap-news/0.9">
  <url>
    <loc>https://threatfrontier.com/articles/mlx-vs-llama-cpp-on-apple-silicon-which-local-llm-runtime-to-actually-use</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:22.506Z</news:publication_date>
      <news:title>MLX vs llama.cpp on Apple Silicon: Which Local LLM Runtime to Actually Use</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/how-to-build-an-llm-evaluation-harness-that-catches-real-regressions</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:21.363Z</news:publication_date>
      <news:title>How to Build an LLM Evaluation Harness That Catches Real Regressions</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/prompt-caching-explained-cutting-agent-and-chat-costs-without-touching-the-model</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:20.192Z</news:publication_date>
      <news:title>Prompt Caching Explained: Cutting Agent Costs Without Touching the Model</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/speculative-decoding-and-multi-token-prediction-how-llms-generate-faster</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:19.052Z</news:publication_date>
      <news:title>Speculative Decoding and Multi-Token Prediction: How LLMs Generate Faster</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/structured-output-and-constrained-decoding-making-llms-return-valid-json</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:17.901Z</news:publication_date>
      <news:title>Structured Output and Constrained Decoding: Making LLMs Return Valid JSON</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/reranking-explained-why-your-rag-pipeline-needs-a-second-scoring-pass</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:16.788Z</news:publication_date>
      <news:title>Reranking Explained: Why Your RAG Pipeline Needs a Second Scoring Pass</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/vector-database-comparison-pgvector-qdrant-weaviate-milvus-and-chroma</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:15.759Z</news:publication_date>
      <news:title>Vector Database Comparison: pgvector, Qdrant, Weaviate, Milvus and Chroma</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/best-open-weight-embedding-models-2026-benchmarks-dimensions-and-vram-cost</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:14.725Z</news:publication_date>
      <news:title>Best Open-Weight Embedding Models 2026: Benchmarks, Dimensions and VRAM Cost</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/chunking-strategies-for-rag-how-document-splitting-decides-retrieval-quality</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:13.420Z</news:publication_date>
      <news:title>Chunking Strategies for RAG: How Document Splitting Decides Retrieval Quality</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/production-rag-architecture-retrieval-pipelines-that-actually-work-in-2026</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:37:12.312Z</news:publication_date>
      <news:title>Production RAG Architecture: Retrieval Pipelines That Actually Work</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/chatgpt-alternatives-2026-self-hosted-open-weight-stacks-that-actually-replace-it</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:04:04.612Z</news:publication_date>
      <news:title>Replacing ChatGPT With a Self-Hosted Stack: What Actually Works in 2026</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/rtx-pro-6000-blackwell-96gb-the-local-llm-workstation-buyers-guide</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:04:03.558Z</news:publication_date>
      <news:title>RTX PRO 6000 Blackwell 96GB: The Local LLM Workstation Buyer&apos;s Guide</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/serverless-gpu-inference-compared-baseten-modal-runpod-together-fireworks</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:04:02.542Z</news:publication_date>
      <news:title>Serverless GPU Inference Compared: Baseten, Modal, RunPod, Together and Fireworks</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/dgx-spark-cluster-guide-tensor-parallel-sizing-for-two-and-four-node-rigs</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:04:01.537Z</news:publication_date>
      <news:title>Multi-Node DGX Spark: Tensor-Parallel Sizing for 2x and 4x Clusters</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/agent-harness-architecture-explained-how-coding-agents-actually-run-tools</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:04:00.471Z</news:publication_date>
      <news:title>Agent Harness Architecture: How Coding Agents Actually Execute Tools</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/sglang-vs-vllm-vs-tensorrt-llm-choosing-a-production-inference-engine-in-2026</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:03:59.480Z</news:publication_date>
      <news:title>SGLang vs vLLM vs TensorRT-LLM: Choosing a Production Inference Engine</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/kv-cache-explained-why-context-length-costs-more-vram-than-your-model-weights</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:03:58.464Z</news:publication_date>
      <news:title>KV Cache Explained: Why Long Context Costs More VRAM Than Your Weights</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/does-it-fit-vram-calculator-for-open-weight-llms-across-quantization-formats</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:03:57.520Z</news:publication_date>
      <news:title>Will It Fit? VRAM Math for Open-Weight Models Across Every Quantization Format</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/inference-engineering-explained-the-practical-primer-for-serving-your-own-llms</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:03:56.475Z</news:publication_date>
      <news:title>Inference Engineering: A Practical Primer for Serving Your Own Models</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/nvfp4-vs-fp8-vs-mxfp4-quantization-formats-explained-vram-math-for-local-llms</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T16:03:55.444Z</news:publication_date>
      <news:title>NVFP4 vs FP8 vs MXFP4 vs EXL3: The Quantization Format Guide for Local LLMs</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/fips-140-2-sunset-september-2026-the-migration-guide-to-fips-140-3</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T15:18:10.171Z</news:publication_date>
      <news:title>FIPS 140-2 Sunset September 2026: The Migration Guide to FIPS 140-3</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/eu-cyber-resilience-act-article-14-navigating-mandatory-vulnerability-reporting-requirements-for-enterprise-psirts</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T15:16:00.264Z</news:publication_date>
      <news:title>EU Cyber Resilience Act Article 14: Navigating Mandatory Vulnerability Reporting Requirements for Enterprise PSIRTs</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/beyond-the-patch-evicting-cosmicsting-cve-2024-34102-backdoors-and-stylesmuggler-skimmers-in-magento-and-adobe-commerce</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T15:15:40.264Z</news:publication_date>
      <news:title>Beyond the Patch: Evicting CosmicSting (CVE-2024-34102) Backdoors and StyleSmuggler Skimmers in Magento and Adobe Commerce</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/gitlab-unauthenticated-arbitrary-file-read-exposes-critical-secrets-full-blast-radius-and-incident-recovery-runbook</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T15:15:30.246Z</news:publication_date>
      <news:title>GitLab Unauthenticated Arbitrary File Read Exposes Critical Secrets: Full Blast Radius and Incident Recovery Runbook</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/interlock-ransomware-weaponizes-critical-cisco-fmc-flaw-why-standard-patching-fails</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T15:14:50.224Z</news:publication_date>
      <news:title>Interlock Ransomware Weaponizes Critical Cisco FMC Flaw: Why Standard Patching Fails</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/run-deepseek-v4-1-flash-on-three-dgx-sparks</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T14:58:00.601Z</news:publication_date>
      <news:title>Run DeepSeek V4.1 Flash on Three DGX Sparks</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/best-local-ai-models-to-run-on-your-dgx-sparks</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T14:55:38.555Z</news:publication_date>
      <news:title>Best Local AI Models to Run on Your DGX Sparks</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/hardening-deepseek-harness-architecture-setup-tutorial-and-runtime-security</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T14:51:40.172Z</news:publication_date>
      <news:title>Hardening DeepSeek Harness: Architecture, Setup Tutorial, and Runtime Security</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://threatfrontier.com/articles/deepseek-v4-1-flash-upends-enterprise-ai-economics-via-asymmetric-routing</loc>
    <news:news>
      <news:publication>
        <news:name>ThreatFrontier.com</news:name>
        <news:language>en</news:language>
      </news:publication>
      <news:publication_date>2026-09-12T14:50:10.428Z</news:publication_date>
      <news:title>DeepSeek-V4.1-Flash Upends Enterprise AI Economics via Asymmetric Routing</news:title>
    </news:news>
  </url>
</urlset>
