tree: e0c39d2c189a40822cd7bbcf29461bc4a9b49254
  1. access-non-generic.ll
  2. activemask.ll
  3. add-sub-128bit.ll
  4. addr-mode.ll
  5. address-folder.ll
  6. address-folder.mir
  7. addrspacecast-cse.ll
  8. addrspacecast-folding.ll
  9. addrspacecast-gvar.ll
  10. addrspacecast-ptx64.ll
  11. addrspacecast.ll
  12. aggr-param.ll
  13. aggregate-return.ll
  14. alias-errors.ll
  15. alias.ll
  16. and-or-setcc.ll
  17. annotations.ll
  18. anonymous-fn-param.ll
  19. APIntLoadStore.ll
  20. APIntParam.ll
  21. APIntSextParam.ll
  22. APIntZextParam.ll
  23. applypriority.ll
  24. arbitrary-fp-to-float.ll
  25. arg-lowering.ll
  26. arithmetic-fp-sm20.ll
  27. arithmetic-int.ll
  28. asm-printer-ptx-module-directives.ll
  29. async-copy.ll
  30. atomic-alignment.err.ll
  31. atomic-lower-local.ll
  32. atomicrmw-allow-ftz-atomics.ll
  33. atomicrmw-expand.err.ll
  34. atomicrmw-sm60.ll
  35. atomicrmw-sm70.ll
  36. atomicrmw-sm90.ll
  37. atomicrmw.py
  38. atomics-b128.ll
  39. atomics-sm60.ll
  40. atomics.ll
  41. b52037.ll
  42. barrier.ll
  43. bf16-instructions.ll
  44. bf16.ll
  45. bf16x2-instructions-approx.ll
  46. bf16x2-instructions.ll
  47. bfe.ll
  48. blocksareclusters-kernel-attr.ll
  49. bmsk.ll
  50. boolean-patterns.ll
  51. branch-fold.ll
  52. branch-fold.mir
  53. brkpt.ll
  54. bswap.ll
  55. bug17709.ll
  56. bug21465.ll
  57. bug22246.ll
  58. bug22322.ll
  59. bug26185-2.ll
  60. bug26185.ll
  61. bug41651.ll
  62. bug52623.ll
  63. bypass-div.ll
  64. byval-arg-vectorize.ll
  65. byval-const-global.ll
  66. call-with-alloca-buffer.ll
  67. call_bitcast_byval.ll
  68. callchain.ll
  69. calling-conv.ll
  70. calls-with-phi.ll
  71. chain-different-as.ll
  72. cluster-dim.ll
  73. clusterlaunchcontrol-multicast.ll
  74. clusterlaunchcontrol.ll
  75. cmpxchg-sm60.ll
  76. cmpxchg-sm70.ll
  77. cmpxchg-sm90.ll
  78. cmpxchg-unsupported-syncscope.err.ll
  79. cmpxchg.ll
  80. cmpxchg.py
  81. combine-mad.ll
  82. combine-min-max.ll
  83. combine-mul-wide-type.ll
  84. combine-wide.ll
  85. common-linkage.ll
  86. compare-int.ll
  87. compute-ptx-value-vts.ll
  88. constant-vectors.ll
  89. convergent-mir-call.ll
  90. convert-call-to-indirect.ll
  91. convert-fp-i128.ll
  92. convert-fp-i8.ll
  93. convert-fp.ll
  94. convert-int-sm20.ll
  95. convert-sm100.ll
  96. convert-sm100a.ll
  97. convert-sm103a.ll
  98. convert-sm80-sf.ll
  99. convert-sm80.ll
  100. convert-sm89.ll
  101. convert-sm90.ll
  102. convert_fp4x2_sm_100f.ll
  103. convert_fp4x2_to_bf16x2.ll
  104. convert_fp6x2_sm_100f.ll
  105. convert_fp6x2_to_bf16x2.ll
  106. convert_fp8x2_sm_100f.ll
  107. convert_fp8x2_to_bf16x2.ll
  108. convert_s2f6x2_sm_100a.ll
  109. copysign.ll
  110. cp-async-bulk-ptx86.ll
  111. cp-async-bulk-s2g-sm100.ll
  112. cp-async-bulk-tensor-g2s-1cta.ll
  113. cp-async-bulk-tensor-g2s-2cta.ll
  114. cp-async-bulk-tensor-g2s-cta-sm100.ll
  115. cp-async-bulk-tensor-g2s-cta-sm100a.ll
  116. cp-async-bulk-tensor-g2s-cta-sm90.ll
  117. cp-async-bulk-tensor-g2s-gather4.ll
  118. cp-async-bulk-tensor-g2s-im2colw.ll
  119. cp-async-bulk-tensor-g2s-im2colw128.ll
  120. cp-async-bulk-tensor-g2s-invalid.ll
  121. cp-async-bulk-tensor-g2s.ll
  122. cp-async-bulk-tensor-prefetch-sm100a.ll
  123. cp-async-bulk-tensor-prefetch.ll
  124. cp-async-bulk-tensor-reduce.ll
  125. cp-async-bulk-tensor-s2g-scatter4.ll
  126. cp-async-bulk-tensor-s2g.ll
  127. cp-async-bulk.ll
  128. cse-mov-sym.ll
  129. ctlz.ll
  130. ctpop.ll
  131. cttz.ll
  132. dag-cse.ll
  133. dead-shfl.ll
  134. default-sm.ll
  135. demote-vars.ll
  136. disable-opt.ll
  137. discard.ll
  138. disjoint-or-addr.ll
  139. distributed-shared-cluster.ll
  140. div-ri.ll
  141. div.ll
  142. divrem-combine.ll
  143. dot-product.ll
  144. dynamic-stackalloc-regression.ll
  145. dynamic_stackalloc.ll
  146. elect.ll
  147. empty-type.ll
  148. envreg.ll
  149. extern-shared-valid-name.ll
  150. extloadv.ll
  151. extractelement.ll
  152. f16-abs.ll
  153. f16-add-sat.ll
  154. f16-ex2.ll
  155. f16-instructions.ll
  156. f16-mul-sat.ll
  157. f16-sub-sat.ll
  158. f16x2-instructions.ll
  159. f32-ex2.ll
  160. f32-lg2.ll
  161. f32x2-convert-i32x2.ll
  162. f32x2-instructions.ll
  163. fabs-intrinsics.ll
  164. fast-math.ll
  165. fcos-no-fast-math.ll
  166. fence-proxy-sm90-ptx86.ll
  167. fence-proxy-sm90.ll
  168. fence-proxy-tensormap-invalid.ll
  169. fence-proxy-tensormap.ll
  170. fence-proxy.ll
  171. fence.ll
  172. fence.py
  173. fexp2.ll
  174. filetype-null.ll
  175. flo.ll
  176. float-to-arbitrary-fp.ll
  177. flog2.ll
  178. fma-assoc.ll
  179. fma-disable.ll
  180. fma-oob.ll
  181. fma-relu-contract.ll
  182. fma-relu-fma-intrinsic.ll
  183. fma-relu-instruction-flag.ll
  184. fma.ll
  185. fmax3.ll
  186. fminimum-fmaximum.ll
  187. fns.ll
  188. fold-movs.ll
  189. forward-ld-param.ll
  190. fp-arith-sat.ll
  191. fp-contract-f32x2.ll
  192. fp-contract.ll
  193. fp-fold-sub.ll
  194. fp-literals.ll
  195. fp128-storage-type.ll
  196. frameindex-lifetime.ll
  197. frem.ll
  198. fsin-no-fast-math.ll
  199. function-align.ll
  200. funnel-shift-clamp.ll
  201. generic-to-nvvm-ir.ll
  202. generic-to-nvvm.ll
  203. global-addrspace.ll
  204. global-ctor-empty.ll
  205. global-incomplete-init.ll
  206. global-ordering.ll
  207. global-variable-big.ll
  208. global-visibility.ll
  209. globals_init.ll
  210. globals_lowering.ll
  211. griddepcontrol.ll
  212. gvar-init.ll
  213. gvn-scalar-pre-reg-pressure.ll
  214. half.ll
  215. i1-array-global.ll
  216. i1-ext-load.ll
  217. i1-global.ll
  218. i1-icmp.ll
  219. i1-int-to-fp.ll
  220. i1-load-lower.ll
  221. i1-param.ll
  222. i1-select.ll
  223. i128-array.ll
  224. i128-global.ll
  225. i128-ld-st.ll
  226. i128-param.ll
  227. i128-retval.ll
  228. i128-struct.ll
  229. i128.ll
  230. i16x2-instructions.ll
  231. i32x2-instructions.ll
  232. i8-param.ll
  233. i8x2-instructions.ll
  234. i8x4-instructions.ll
  235. idioms.ll
  236. imad.ll
  237. indirect_byval.ll
  238. inline-asm-b128-test1.ll
  239. inline-asm-b128-test2.ll
  240. inline-asm-b128-test3.ll
  241. inline-asm-line-info-inlined-at.ll
  242. inline-asm-line-info-per-instruction.ll
  243. inline-asm-line-number-before.ll
  244. inline-asm.ll
  245. inlineasm-output-template.ll
  246. insert-vector-elt-bitcast-legalize.ll
  247. insert-vector-elt-shuffle-i8.ll
  248. insertelt-dynamic.ll
  249. intr-range.ll
  250. intrinsic-immarg-print-mismatched-signature.ll
  251. intrinsic-old.ll
  252. intrinsics-sm90-ptx81.ll
  253. intrinsics-sm90.ll
  254. intrinsics.ll
  255. isspacep.ll
  256. jump-table.ll
  257. kernel-param-align.ll
  258. ld-addrspace.ll
  259. ld-generic.ll
  260. ld-param-sink.ll
  261. ld-st-addrrspace.py
  262. ldg-invariant-256.ll
  263. ldg-invariant.ll
  264. ldparam-v4.ll
  265. ldu-i8.ll
  266. ldu-ldg.ll
  267. ldu-reg-plus-offset.ll
  268. lit.local.cfg
  269. load-sext-i1.ll
  270. load-store-256-addressing-invariant.ll
  271. load-store-256-addressing.ll
  272. load-store-atomic.err.ll
  273. load-store-scalars.ll
  274. load-store-sm-70.ll
  275. load-store-sm-90.ll
  276. load-store-vectors-256.ll
  277. load-store-vectors.ll
  278. load-with-non-coherent-cache.ll
  279. LoadStoreVectorizer.ll
  280. local-stack-frame.ll
  281. loop-vectorize.ll
  282. lower-aggr-copies-shared.ll
  283. lower-aggr-copies.ll
  284. lower-alloca.ll
  285. lower-args-alignment.ll
  286. lower-args-gridconstant.ll
  287. lower-args.ll
  288. lower-byval-args-dbg.ll
  289. lower-byval-args-idempotent.ll
  290. lower-byval-args-mem-attrs.ll
  291. lower-byval-args.ll
  292. lower-ctor-dtor.ll
  293. lower-kernel-ptr-arg.ll
  294. machine-cse-predicate-inversion-multiple-users.ll
  295. machine-cse-predicate-inversion-rollback.mir
  296. machine-cse-predicate-inversion-vector-float.ll
  297. machine-cse-predicate-inversion.ll
  298. machine-cse-predicate-no-inversion.ll
  299. machine-sink.ll
  300. machinelicm-no-preheader.mir
  301. MachineSink-call.ll
  302. MachineSink-convergent.ll
  303. managed.ll
  304. mark-kernel-ptrs-global.ll
  305. masked-load-3xhalf.ll
  306. masked-load-vectors.ll
  307. masked-store-variable-mask.ll
  308. masked-store-vectors-256.ll
  309. match.ll
  310. math-intrins-sm53-ptx42.ll
  311. math-intrins-sm80-ptx70-autoupgrade.ll
  312. math-intrins-sm80-ptx70-instcombine.ll
  313. math-intrins-sm80-ptx70.ll
  314. math-intrins-sm86-ptx72-autoupgrade.ll
  315. math-intrins-sm86-ptx72.ll
  316. math-intrins.ll
  317. max-align.ll
  318. maxclusterrank.ll
  319. mbarrier.ll
  320. mbarrier_arr.ll
  321. mbarrier_arr_relaxed.ll
  322. mbarrier_tx.ll
  323. mbarrier_wait_sm80_ptx70.ll
  324. mbarrier_wait_sm80_ptx71.ll
  325. mbarrier_wait_sm90_ptx78.ll
  326. mbarrier_wait_sm90_ptx80.ll
  327. mbarrier_wait_sm90_ptx86.ll
  328. minmax-negative.ll
  329. misaligned-vector-ldst.ll
  330. misched_func_call.ll
  331. mixed-precision-fp.ll
  332. mma-no-sink-after-laneid-check.ll
  333. module-inline-asm.ll
  334. movmatrix.ll
  335. mulhi-intrins.ll
  336. mulwide.ll
  337. naked-fn-with-frame-pointer.ll
  338. nanosleep.ll
  339. no-extra-parens.ll
  340. no-f32x2.ll
  341. no-stack-protector-libcall-error.ll
  342. noduplicate-syncthreads.ll
  343. nofunc.ll
  344. noreturn.ll
  345. nounroll.ll
  346. nvcl-param-align.ll
  347. nvptx-aa-inline-asm.ll
  348. nvptx-aa.ll
  349. nvptx-fold-fma.ll
  350. nvptx-prec-divf32-flag.ll
  351. NVPTXAA_before_BasicAA.ll
  352. nvvm-abs.ll
  353. nvvm-annotations-D120129.ll
  354. nvvm-reflect-arch-O0.ll
  355. nvvm-reflect-arch.ll
  356. nvvm-reflect-module-flag.ll
  357. nvvm-reflect-ocl.ll
  358. nvvm-reflect-opaque.ll
  359. nvvm-reflect-options.ll
  360. nvvm-reflect.ll
  361. op-fence.ll
  362. packed-aggr.ll
  363. param-add.ll
  364. param-align.ll
  365. param-load-store.ll
  366. param-overalign.ll
  367. param-space-subqualifiers.ll
  368. param-vectorize-device.ll
  369. param-vectorize-kernel.ll
  370. pass-name.ll
  371. pm-event.ll
  372. pow2_mask_cmp.ll
  373. pr126337.ll
  374. pr13291-i1-store.ll
  375. pr16278.ll
  376. pr17529.ll
  377. prefetch-inferas-test.ll
  378. prefetch.ll
  379. prmt-const-folding.ll
  380. prmt.ll
  381. promote-param-align.ll
  382. proxy-reg-erasure-kill-flag.mir
  383. proxy-reg-erasure-ptx.ll
  384. proxy-reg-erasure.mir
  385. ptx-version-validation.ll
  386. rcp-opt.ll
  387. read-global-variable-constant.ll
  388. reduction-intrinsics.ll
  389. redux-sync-f32.ll
  390. redux-sync.ll
  391. refl1.ll
  392. reg-copy.ll
  393. reg-types.ll
  394. reqnctapercluster-const-fold.ll
  395. reqntid-const-fold.ll
  396. reserved-smem-offset.ll
  397. ret-align-mismatch.ll
  398. rotate-add.ll
  399. rotate.ll
  400. rotate_64.ll
  401. rsqrt-opt.ll
  402. rsqrt.ll
  403. sad-intrins.ll
  404. scalar-to-vector.ll
  405. scalarize-non-coalescable-v2f32.ll
  406. sched1.ll
  407. sched2.ll
  408. setmaxnreg-sm100a.ll
  409. setmaxnreg.ll
  410. sext-in-reg.ll
  411. sext-params.ll
  412. sext-setcc.ll
  413. shfl-p.ll
  414. shfl-sync-p.ll
  415. shfl-sync.ll
  416. shfl.ll
  417. shift-opt.ll
  418. shift-parts.ll
  419. short-ptr.ll
  420. shuffle-vec-undef-init.ll
  421. simple-call.ll
  422. sm-version.ll
  423. speculative-execution-divergent-target.ll
  424. sqrt-approx.ll
  425. srl-bitcast-bv.ll
  426. st-addrspace.ll
  427. st-generic.ll
  428. st-param-imm.ll
  429. st_async_mbarrier.ll
  430. st_async_mbarrier_b128.ll
  431. st_async_release.ll
  432. st_async_release_multimem.ll
  433. st_bulk.ll
  434. stackaddress.ll
  435. stacksaverestore.ll
  436. store-retval.ll
  437. store-undef.ll
  438. sub-byte-constant-vector-convert.ll
  439. sub-byte-constant-vectors-i4-i2.ll
  440. surf-read-cuda.ll
  441. surf-read.ll
  442. surf-tex.py
  443. surf-write-cuda.ll
  444. surf-write.ll
  445. switch-loop-header.mir
  446. switch.ll
  447. symbol-naming.ll
  448. szext.ll
  449. tag-invariant-loads.ll
  450. TailDuplication-convergent.ll
  451. tanhf.ll
  452. tcgen05-alloc.ll
  453. tcgen05-commit.ll
  454. tcgen05-cp.ll
  455. tcgen05-fence.ll
  456. tcgen05-ld-red.ll
  457. tcgen05-ld.ll
  458. tcgen05-mma-block-scale-invalid.ll
  459. tcgen05-mma-block-scale-ptx88-aa.ll
  460. tcgen05-mma-block-scale-ptx88.ll
  461. tcgen05-mma-block-scale.ll
  462. tcgen05-mma-disable-output-lane-i8.ll
  463. tcgen05-mma-disable-output-lane.ll
  464. tcgen05-mma-i8.ll
  465. tcgen05-mma-invalid.ll
  466. tcgen05-mma-scale-d-invalid.ll
  467. tcgen05-mma-scale-d.ll
  468. tcgen05-mma-tensor-formatted.ll
  469. tcgen05-mma-ws-i8.ll
  470. tcgen05-mma-ws.ll
  471. tcgen05-mma.ll
  472. tcgen05-shift.ll
  473. tcgen05-st.ll
  474. tensormap_replace.ll
  475. tensormap_replace_invalid.ll
  476. tensormap_replace_sm_100a.ll
  477. tensormap_replace_sm_103a.ll
  478. tex-read-cuda.ll
  479. tex-read.ll
  480. texsurf-queries.ll
  481. thread-fence.ll
  482. tid-range.ll
  483. trunc-setcc.ll
  484. trunc-tofp.ll
  485. unaligned-param-load-store.ll
  486. unfold-masked-merge-vector-variablemask.ll
  487. unknown-intrinsic.ll
  488. unreachable.ll
  489. unrecognized-sm1x.ll
  490. upgrade-nvvm-annotations.ll
  491. used-bytes-mask.ll
  492. vaargs.ll
  493. variadics-backend.ll
  494. variadics-lowering.ll
  495. vec-param-load.ll
  496. vec8.ll
  497. vector-args.ll
  498. vector-call.ll
  499. vector-compare.ll
  500. vector-global.ll
  501. vector-loads.ll
  502. vector-returns.ll
  503. vector-select.ll
  504. vector-stores.ll
  505. vectorize-misaligned.ll
  506. vote.ll
  507. weak-global.ll
  508. weak-linkage.ll
  509. wgmma-sm90a-fence.ll
  510. wmma-ptx60-sm70.py
  511. wmma-ptx61-sm70.py
  512. wmma-ptx63-sm72.py
  513. wmma-ptx63-sm75.py
  514. wmma-ptx64-sm70.py
  515. wmma-ptx65-sm75.py
  516. wmma-ptx71-sm80.py
  517. wmma-ptx78-sm90.py
  518. wmma-ptx86-sm100a.py
  519. wmma-ptx86-sm101a.py
  520. wmma-ptx87-sm120a.py
  521. wmma-ptx88-sm100f.py
  522. wmma-ptx88-sm120a.py
  523. wmma-ptx88-sm120f.py
  524. wmma-ptx90-sm110f.py
  525. wmma-ptx91-sm120a.py
  526. wmma-ptx91-sm120f.py
  527. wmma.py
  528. zeroext-32bit.ll