Staff Site Reliability EngineerEmbedded

External Listing
{
  "source": {
    "name": "linkedin",
    "id": "4411261497",
    "url": "https://www.linkedin.com/jobs/view/staff-site-reliability-engineer-at-obsidian-security-4411261497?_l=en"
  },
  "postedDate": "2026-09-18T15:43:58.216Z",
  "applicationDeadline": null,
  "isActive": true,
  "isExpired": false,
  "matching": {
    "role": {
      "primaryTitle": "Staff Site Reliability Engineer",
      "titleSynonyms": [
        "Site Reliability Engineer",
        "SRE",
        "Production Engineer"
      ],
      "secondaryTitles": [
        "Senior Site Reliability Engineer",
        "Lead Site Reliability Engineer"
      ],
      "function": "ict_professionals",
      "functionConfidence": "high",
      "roleFamily": "u2512_software_developers",
      "roleFamilyConfidence": "medium",
      "roleSubFamily": "Site Reliability Engineering",
      "roleSubFamilyConfidence": "medium",
      "seniority": "mid_senior",
      "industries": [
        "i96_it_services_and_it_consulting"
      ]
    },
    "primarySignals": {
      "tasks": [
        "Define and lead long-term reliability strategy across services",
        "Establish system visibility frameworks for observability, detection, and resilience",
        "Partner across teams to embed reliability and standardize SLI/SLOs",
        "Build intelligent detection systems including anomaly detection and connector health models",
        "Define and evolve incident communication strategy and lead postmortems",
        "Contribute hands-on to system design, monitoring, and debugging across distributed systems and data pipelines"
      ],
      "skills": [
        "Site Reliability Engineering",
        "Production Engineering",
        "AWS",
        "GCP",
        "Kubernetes",
        "Helm",
        "Observability stacks (Prometheus, Grafana)",
        "CI/CD systems (GitLab CI/CD, ArgoCD)",
        "Distributed systems debugging",
        "Systems thinking",
        "Incident management"
      ],
      "tools": [
        "AWS",
        "GCP",
        "Kubernetes",
        "Helm",
        "Prometheus",
        "Grafana",
        "GitLab CI/CD",
        "ArgoCD"
      ],
      "educationLevel": null,
      "educationKeywords": [],
      "certifications": [],
      "languages": [],
      "yearsRelevant": 5
    },
    "secondarySignals": {
      "tasks": [
        "Experience building anomaly detection or intelligent alerting systems",
        "Experience designing customer-facing status pages and incident communication frameworks"
      ],
      "skills": [
        "B2B SaaS experience",
        "Familiarity with third-party SaaS connector architectures and ingestion patterns"
      ],
      "tools": [],
      "educationLevel": null,
      "educationKeywords": [],
      "certifications": [],
      "languages": [],
      "yearsRelevant": 3
    },
    "practical": {
      "locations": [
        "Cheltenham, England, United Kingdom"
      ],
      "locationProvenance": "stated",
      "countries": [
        "GB"
      ],
      "workModes": [
        "unknown"
      ],
      "workModeProvenance": "defaulted",
      "employmentTypes": [
        "full_time"
      ],
      "compensation": {
        "min": 124000,
        "max": 141000,
        "currency": "GBP",
        "period": "year"
      }
    },
    "dataCompleteness": "high"
  },
  "indexing": {
    "function": "ict_professionals",
    "functionConfidence": "high",
    "roleFamily": "u2512_software_developers",
    "roleFamilyConfidence": "medium",
    "countryCode": "GB",
    "countryCodeConfidence": "high",
    "locationBucket": "ttwa_cheltenham",
    "locationBucketConfidence": "high",
    "workMode": "unknown",
    "workModeConfidence": "unknown",
    "employmentType": "full_time",
    "employmentTypeConfidence": "high",
    "isAgency": "direct",
    "salaryMax": 141000,
    "industryGroup": "i96_it_services_and_it_consulting",
    "industryGroupConfidence": "high",
    "educationRequired": null,
    "educationRequiredConfidence": "unknown"
  },
  "display": {
    "title": "Staff Site Reliability Engineer",
    "company": {
      "name": "Obsidian Security"
    },
    "locationDisplay": "Cheltenham, England, United Kingdom",
    "applicationUrl": null
  },
  "roleFamilyEsco": "u2522_systems_administrators",
  "requirementsEssentiality": {
    "items": [
      {
        "text": "5+ years in SRE, Production Engineering, or related roles",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "3+ years operating at a senior or technical leadership level (Staff or equivalent scope)",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Deep expertise in AWS and/or GCP",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Deep expertise in Kubernetes and Helm",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Deep expertise in Observability stacks (Prometheus, Grafana, or equivalent)",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Deep expertise in CI/CD systems (GitLab CI/CD, ArgoCD, etc.)",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Proven experience designing and scaling reliability systems for multi-tenant SaaS platforms",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Strong debugging and systems thinking across distributed microservices and legacy systems",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Demonstrated ability to lead initiatives that improve incident detection, response, and system resilience",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Hands-on engineering approach with a track record of building—not just configuring—reliability systems",
        "category": "skill",
        "essentiality": "compulsory",
        "triggerPhrase": "Required Qualifications",
        "confidence": "high"
      },
      {
        "text": "Experience in B2B SaaS serving enterprise or financial customers",
        "category": "skill",
        "essentiality": "preferred",
        "triggerPhrase": "Preferred Qualifications",
        "confidence": "high"
      },
      {
        "text": "Familiarity with third-party SaaS connector architectures and ingestion patterns",
        "category": "skill",
        "essentiality": "preferred",
        "triggerPhrase": "Preferred Qualifications",
        "confidence": "high"
      },
      {
        "text": "Experience building anomaly detection or intelligent alerting systems",
        "category": "skill",
        "essentiality": "preferred",
        "triggerPhrase": "Preferred Qualifications",
        "confidence": "high"
      },
      {
        "text": "Experience designing customer-facing status pages and incident communication frameworks",
        "category": "skill",
        "essentiality": "preferred",
        "triggerPhrase": "Preferred Qualifications",
        "confidence": "high"
      }
    ]
  }
}