File: llvm.r600.dot4.ll

package info (click to toggle)
llvm-toolchain-6.0 1%3A6.0.1-10
  • links: PTS, VCS
  • area: main
  • in suites: buster
  • size: 598,080 kB
  • sloc: cpp: 3,046,253; ansic: 595,057; asm: 271,965; python: 128,926; objc: 106,554; sh: 21,906; lisp: 10,191; pascal: 6,094; ml: 5,544; perl: 5,265; makefile: 2,227; cs: 2,027; xml: 686; php: 212; csh: 117
file content (11 lines) | stat: -rw-r--r-- 566 bytes parent folder | download | duplicates (15)
1
2
3
4
5
6
7
8
9
10
11
; RUN: llc -march=r600 -mcpu=redwood -verify-machineinstrs < %s

declare float @llvm.r600.dot4(<4 x float>, <4 x float>) nounwind readnone

define amdgpu_kernel void @test_dp4(float addrspace(1)* %out, <4 x float> addrspace(1)* %a, <4 x float> addrspace(1)* %b) nounwind {
  %src0 = load <4 x float>, <4 x float> addrspace(1)* %a, align 16
  %src1 = load <4 x float>, <4 x float> addrspace(1)* %b, align 16
  %dp4 = call float @llvm.r600.dot4(<4 x float> %src0, <4 x float> %src1) nounwind readnone
  store float %dp4, float addrspace(1)* %out, align 4
  ret void
}