{"schemaVersion":"jobsearcher.job.v1","id":"c68411e25bc0d90ffd46a434","url":"https://jobsearcher.com/jobs/c68411e25bc0d90ffd46a434","canonicalUrl":"https://jobsearcher.com/jobs/c68411e25bc0d90ffd46a434","title":"Senior LLM Inference: GPU Kernel Optimization","description":"NVIDIA is seeking a Sr. Inference Engineer to push LLM inference performance through GPU kernel optimization. You will help develop silicon-measured benchmarking, model-level performance projection tooling, and agentic optimization systems, collaborating across compiler, hardware, and framework teams to surface bottlenecks and deliver measurable gains. You will work on GPU kernel microbenchmarking, end-to-end model performance analysis, and agentic optimization, shaping production inference\n#J-18808-Ljbffr","company":"NVIDIA","rawCompany":"nvidia","city":"Austin","state":"TX","isRemote":false,"isActive":true,"createdAt":"2026-10-03T03:58:10.831Z","occupations":[{"code":"15-1299.08","title":"Computer Systems Engineers/Architects","slug":"computer-systems-engineers-architects"},{"code":"17-2061.00","title":"Computer Hardware Engineers","slug":"computer-hardware-engineers"},{"code":"15-1221.00","title":"Computer and Information Research Scientists","slug":"computer-and-information-research-scientists"}],"industries":[{"code":"334111","title":"Electronic Computer Manufacturing","slug":"electronic-computer-manufacturing"},{"code":"334118","title":"Computer Terminal and Other Computer Peripheral Equipment Manufacturing","slug":"computer-terminal-and-other-computer-peripheral-equipment-manufacturing"},{"code":"513210","title":"Software Publishers","slug":"software-publishers"}],"jobPosting":{"@context":"https://schema.org","@type":"JobPosting","title":"Senior LLM Inference: GPU Kernel Optimization","description":"NVIDIA is seeking a Sr. Inference Engineer to push LLM inference performance through GPU kernel optimization. You will help develop silicon-measured benchmarking, model-level performance projection tooling, and agentic optimization systems, collaborating across compiler, hardware, and framework teams to surface bottlenecks and deliver measurable gains. You will work on GPU kernel microbenchmarking, end-to-end model performance analysis, and agentic optimization, shaping production inference\n#J-18808-Ljbffr","datePosted":"2026-10-03T03:58:10.831Z","dateModified":"2026-10-03T03:58:10.831Z","hiringOrganization":{"@type":"Organization","name":"NVIDIA","sameAs":"https://jobsearcher.com"},"jobLocation":{"@type":"Place","address":{"@type":"PostalAddress","addressLocality":"Austin","addressRegion":"TX","addressCountry":"US"}},"identifier":{"@type":"PropertyValue","name":"JobSearcher","value":"c68411e25bc0d90ffd46a434"},"url":"https://jobsearcher.com/jobs/c68411e25bc0d90ffd46a434"}}