tree: 6675ccc0cbe27047021a10db13cf1ac63a50e1fa
  1. access-non-generic.ll
  2. activemask.ll
  3. add-sub-128bit.ll
  4. addr-mode.ll
  5. address-folder.ll
  6. address-folder.mir
  7. addrspacecast-cse.ll
  8. addrspacecast-folding.ll
  9. addrspacecast-gvar.ll
  10. addrspacecast-ptx64.ll
  11. addrspacecast.ll
  12. aggr-param.ll
  13. aggregate-return.ll
  14. alias-errors.ll
  15. alias.ll
  16. and-or-setcc.ll
  17. annotations.ll
  18. anonymous-fn-param.ll
  19. APIntLoadStore.ll
  20. APIntParam.ll
  21. APIntSextParam.ll
  22. APIntZextParam.ll
  23. applypriority.ll
  24. arbitrary-fp-to-float.ll
  25. arg-lowering.ll
  26. arithmetic-fp-sm20.ll
  27. arithmetic-int.ll
  28. asm-printer-ptx-module-directives.ll
  29. async-copy.ll
  30. atomic-alignment.err.ll
  31. atomic-lower-local.ll
  32. atomicrmw-allow-ftz-atomics.ll
  33. atomicrmw-expand.err.ll
  34. atomicrmw-sm60.ll
  35. atomicrmw-sm70.ll
  36. atomicrmw-sm90.ll
  37. atomicrmw.py
  38. atomics-b128.ll
  39. atomics-sm60.ll
  40. atomics.ll
  41. b52037.ll
  42. barrier.ll
  43. bf16-instructions.ll
  44. bf16-neg-noftz.ll
  45. bf16-setp-noftz.ll
  46. bf16.ll
  47. bf16x2-instructions-approx.ll
  48. bf16x2-instructions.ll
  49. bfe.ll
  50. blocksareclusters-kernel-attr.ll
  51. bmsk.ll
  52. boolean-patterns.ll
  53. branch-fold.ll
  54. branch-fold.mir
  55. brkpt.ll
  56. bswap.ll
  57. bug17709.ll
  58. bug21465.ll
  59. bug22246.ll
  60. bug22322.ll
  61. bug26185-2.ll
  62. bug26185.ll
  63. bug41651.ll
  64. bug52623.ll
  65. bypass-div.ll
  66. byval-arg-vectorize.ll
  67. byval-const-global.ll
  68. cache-hint-atomics.ll
  69. cache-hint-cache-policy.ll
  70. cache-hint-intrinsics.ll
  71. cache-hint-invalid.ll
  72. cache-hint-load-store.ll
  73. cache-hint-sm-version.ll
  74. cache-hint-transforms.ll
  75. call-with-alloca-buffer.ll
  76. call_bitcast_byval.ll
  77. callchain.ll
  78. calling-conv.ll
  79. calls-with-phi.ll
  80. chain-different-as.ll
  81. cluster-dim.ll
  82. clusterlaunchcontrol-multicast.ll
  83. clusterlaunchcontrol.ll
  84. cmpxchg-sm60.ll
  85. cmpxchg-sm70.ll
  86. cmpxchg-sm90.ll
  87. cmpxchg-unsupported-syncscope.err.ll
  88. cmpxchg.ll
  89. cmpxchg.py
  90. combine-mad.ll
  91. combine-min-max.ll
  92. combine-mul-wide-type.ll
  93. combine-wide.ll
  94. common-linkage.ll
  95. compare-int.ll
  96. compute-ptx-value-vts.ll
  97. constant-vectors.ll
  98. convergent-mir-call.ll
  99. convert-call-to-indirect.ll
  100. convert-fp-i128.ll
  101. convert-fp-i8.ll
  102. convert-fp.ll
  103. convert-int-sm20.ll
  104. convert-sm100.ll
  105. convert-sm100a.ll
  106. convert-sm103a.ll
  107. convert-sm80-sf.ll
  108. convert-sm80.ll
  109. convert-sm89.ll
  110. convert-sm90.ll
  111. convert-ue5m3x2.ll
  112. convert_fp4x2_sm_100f.ll
  113. convert_fp4x2_to_bf16x2.ll
  114. convert_fp6x2_sm_100f.ll
  115. convert_fp6x2_to_bf16x2.ll
  116. convert_fp8x2_sm_100f.ll
  117. convert_fp8x2_to_bf16x2.ll
  118. convert_s2f6x2_sm_100a.ll
  119. copysign.ll
  120. cp-async-bulk-ptx86.ll
  121. cp-async-bulk-s2g-sm100.ll
  122. cp-async-bulk-tensor-g2s-1cta.ll
  123. cp-async-bulk-tensor-g2s-2cta.ll
  124. cp-async-bulk-tensor-g2s-cta-sm100.ll
  125. cp-async-bulk-tensor-g2s-cta-sm100a.ll
  126. cp-async-bulk-tensor-g2s-cta-sm90.ll
  127. cp-async-bulk-tensor-g2s-gather4.ll
  128. cp-async-bulk-tensor-g2s-im2colw.ll
  129. cp-async-bulk-tensor-g2s-im2colw128.ll
  130. cp-async-bulk-tensor-g2s-invalid.ll
  131. cp-async-bulk-tensor-g2s.ll
  132. cp-async-bulk-tensor-prefetch-sm100a.ll
  133. cp-async-bulk-tensor-prefetch.ll
  134. cp-async-bulk-tensor-reduce-im2colw.ll
  135. cp-async-bulk-tensor-reduce-override.ll
  136. cp-async-bulk-tensor-reduce.ll
  137. cp-async-bulk-tensor-s2g-im2colw.ll
  138. cp-async-bulk-tensor-s2g-override.ll
  139. cp-async-bulk-tensor-s2g-scatter4.ll
  140. cp-async-bulk-tensor-s2g.ll
  141. cp-async-bulk.ll
  142. cse-mov-sym.ll
  143. ctlz.ll
  144. ctpop.ll
  145. cttz.ll
  146. dag-cse.ll
  147. dead-shfl.ll
  148. default-sm.ll
  149. demote-vars.ll
  150. disable-opt.ll
  151. discard.ll
  152. disjoint-or-addr.ll
  153. distributed-shared-cluster.ll
  154. div-ri.ll
  155. div.ll
  156. divrem-combine.ll
  157. dot-product.ll
  158. dynamic-stackalloc-regression.ll
  159. dynamic_stackalloc.ll
  160. elect.ll
  161. empty-type.ll
  162. envreg.ll
  163. extern-shared-valid-name.ll
  164. extloadv.ll
  165. extractelement.ll
  166. f128-no-libcall-error.ll
  167. f16-abs.ll
  168. f16-add-sat.ll
  169. f16-ex2.ll
  170. f16-instructions.ll
  171. f16-mul-sat.ll
  172. f16-sub-sat.ll
  173. f16x2-instructions.ll
  174. f32-ex2.ll
  175. f32-lg2.ll
  176. f32x2-convert-i32x2.ll
  177. f32x2-instructions.ll
  178. fabs-intrinsics.ll
  179. fast-math.ll
  180. fcos-no-fast-math.ll
  181. fence-proxy-sm90-ptx86.ll
  182. fence-proxy-sm90.ll
  183. fence-proxy-tensormap-invalid.ll
  184. fence-proxy-tensormap.ll
  185. fence-proxy.ll
  186. fence.ll
  187. fence.py
  188. fexp2.ll
  189. filetype-null.ll
  190. flo.ll
  191. float-to-arbitrary-fp.ll
  192. flog2.ll
  193. fma-assoc.ll
  194. fma-disable.ll
  195. fma-oob.ll
  196. fma-relu-contract.ll
  197. fma-relu-fma-intrinsic.ll
  198. fma-relu-instruction-flag.ll
  199. fma.ll
  200. fmax3.ll
  201. fminimum-fmaximum.ll
  202. fns.ll
  203. fold-movs.ll
  204. forward-ld-param.ll
  205. fp-arith-sat.ll
  206. fp-contract-f32x2.ll
  207. fp-contract.ll
  208. fp-fold-sub.ll
  209. fp-literals.ll
  210. fp128-conv-no-libcall-error.ll
  211. fp128-global.ll
  212. fp128-storage-type.ll
  213. frameindex-lifetime.ll
  214. frem.ll
  215. fsin-no-fast-math.ll
  216. function-align.ll
  217. funnel-shift-clamp.ll
  218. generic-to-nvvm-ir.ll
  219. generic-to-nvvm.ll
  220. global-addrspace.ll
  221. global-ctor-empty.ll
  222. global-cycle-alias.ll
  223. global-cycle-internal-subcycle.ll
  224. global-cycle-internal.ll
  225. global-cycle.ll
  226. global-incomplete-init.ll
  227. global-ordering.ll
  228. global-variable-big.ll
  229. global-visibility.ll
  230. globals_init.ll
  231. globals_lowering.ll
  232. griddepcontrol.ll
  233. gvar-init.ll
  234. gvn-scalar-pre-reg-pressure.ll
  235. half.ll
  236. i1-array-global.ll
  237. i1-ext-load.ll
  238. i1-global.ll
  239. i1-icmp.ll
  240. i1-int-to-fp.ll
  241. i1-load-lower.ll
  242. i1-param.ll
  243. i1-select.ll
  244. i128-array.ll
  245. i128-global.ll
  246. i128-ld-st.ll
  247. i128-param.ll
  248. i128-retval.ll
  249. i128-struct.ll
  250. i128.ll
  251. i16x2-instructions.ll
  252. i32x2-instructions.ll
  253. i8-param.ll
  254. i8x2-instructions.ll
  255. i8x4-instructions.ll
  256. idioms.ll
  257. imad.ll
  258. indirect_byval.ll
  259. inline-asm-b128-test1.ll
  260. inline-asm-b128-test2.ll
  261. inline-asm-b128-test3.ll
  262. inline-asm-line-info-inlined-at.ll
  263. inline-asm-line-info-per-instruction.ll
  264. inline-asm-line-number-before.ll
  265. inline-asm.ll
  266. inlineasm-output-template.ll
  267. insert-vector-elt-bitcast-legalize.ll
  268. insert-vector-elt-shuffle-i8.ll
  269. insertelt-dynamic.ll
  270. intr-range.ll
  271. intrinsic-immarg-print-mismatched-signature.ll
  272. intrinsic-old.ll
  273. intrinsics-sm90-ptx81.ll
  274. intrinsics-sm90.ll
  275. intrinsics.ll
  276. isspacep.ll
  277. jump-table.ll
  278. kernel-param-align.ll
  279. ld-addrspace.ll
  280. ld-generic.ll
  281. ld-param-sink.ll
  282. ld-st-addrrspace.py
  283. ldg-invariant-256.ll
  284. ldg-invariant.ll
  285. ldparam-v4.ll
  286. ldu-i8.ll
  287. ldu-ldg.ll
  288. ldu-reg-plus-offset.ll
  289. lit.local.cfg
  290. llc-pipeline-npm.ll
  291. load-sext-i1.ll
  292. load-store-256-addressing-invariant.ll
  293. load-store-256-addressing.ll
  294. load-store-atomic.err.ll
  295. load-store-scalars.ll
  296. load-store-sm-70.ll
  297. load-store-sm-90.ll
  298. load-store-vectors-256.ll
  299. load-store-vectors.ll
  300. load-with-non-coherent-cache.ll
  301. LoadStoreVectorizer.ll
  302. local-stack-frame.ll
  303. loop-vectorize.ll
  304. lower-aggr-copies-shared.ll
  305. lower-aggr-copies.ll
  306. lower-alloca.ll
  307. lower-args-alignment.ll
  308. lower-args-gridconstant.ll
  309. lower-args.ll
  310. lower-byval-args-dbg.ll
  311. lower-byval-args-idempotent.ll
  312. lower-byval-args-mem-attrs.ll
  313. lower-byval-args.ll
  314. lower-ctor-dtor.ll
  315. lower-kernel-ptr-arg.ll
  316. machine-cse-predicate-inversion-multiple-users.ll
  317. machine-cse-predicate-inversion-rollback.mir
  318. machine-cse-predicate-inversion-vector-float.ll
  319. machine-cse-predicate-inversion.ll
  320. machine-cse-predicate-no-inversion.ll
  321. machine-sink.ll
  322. machinelicm-no-preheader.mir
  323. MachineSink-call.ll
  324. MachineSink-convergent.ll
  325. managed.ll
  326. mark-kernel-ptrs-global.ll
  327. masked-divrem.ll
  328. masked-load-3xhalf.ll
  329. masked-load-vectors.ll
  330. masked-store-variable-mask.ll
  331. masked-store-vectors-256.ll
  332. match.ll
  333. math-intrins-sm53-ptx42.ll
  334. math-intrins-sm80-ptx70-autoupgrade.ll
  335. math-intrins-sm80-ptx70-instcombine.ll
  336. math-intrins-sm80-ptx70.ll
  337. math-intrins-sm86-ptx72-autoupgrade.ll
  338. math-intrins-sm86-ptx72.ll
  339. math-intrins.ll
  340. max-align.ll
  341. maxclusterrank.ll
  342. mbarrier.ll
  343. mbarrier_arr.ll
  344. mbarrier_arr_relaxed.ll
  345. mbarrier_tx.ll
  346. mbarrier_wait_sm80_ptx70.ll
  347. mbarrier_wait_sm80_ptx71.ll
  348. mbarrier_wait_sm90_ptx78.ll
  349. mbarrier_wait_sm90_ptx80.ll
  350. mbarrier_wait_sm90_ptx86.ll
  351. memcpy-alloca-align.ll
  352. minmax-negative.ll
  353. misaligned-vector-ldst.ll
  354. misched_func_call.ll
  355. mixed-precision-fp.ll
  356. mma-no-sink-after-laneid-check.ll
  357. module-inline-asm.ll
  358. movmatrix.ll
  359. mulhi-intrins.ll
  360. mulwide.ll
  361. naked-fn-with-frame-pointer.ll
  362. nanosleep.ll
  363. no-extra-parens.ll
  364. no-f32x2.ll
  365. no-stack-protector-libcall-error.ll
  366. noduplicate-syncthreads.ll
  367. nofunc.ll
  368. noreturn.ll
  369. nounroll.ll
  370. nvcl-param-align.ll
  371. nvptx-aa-inline-asm.ll
  372. nvptx-aa.ll
  373. nvptx-fold-fma.ll
  374. nvptx-prec-divf32-flag.ll
  375. NVPTXAA_before_BasicAA.ll
  376. nvvm-abs.ll
  377. nvvm-annotations-D120129.ll
  378. nvvm-reflect-arch-O0.ll
  379. nvvm-reflect-arch.ll
  380. nvvm-reflect-module-flag.ll
  381. nvvm-reflect-ocl.ll
  382. nvvm-reflect-opaque.ll
  383. nvvm-reflect-options.ll
  384. nvvm-reflect.ll
  385. op-fence.ll
  386. packed-aggr-self-ptx70.ll
  387. packed-aggr.ll
  388. param-add.ll
  389. param-align.ll
  390. param-load-store.ll
  391. param-overalign.ll
  392. param-space-subqualifiers.ll
  393. param-vectorize-device.ll
  394. param-vectorize-kernel.ll
  395. pass-name.ll
  396. peephole-cvta-local-short-ptr.mir
  397. pm-event.ll
  398. pow2_mask_cmp.ll
  399. powi.ll
  400. pr126337.ll
  401. pr13291-i1-store.ll
  402. pr16278.ll
  403. pr17529.ll
  404. prefetch-inferas-test.ll
  405. prefetch.ll
  406. prmt-const-folding.ll
  407. prmt.ll
  408. promote-param-align.ll
  409. proxy-reg-erasure-kill-flag.mir
  410. proxy-reg-erasure-ptx.ll
  411. proxy-reg-erasure.mir
  412. ptx-version-validation.ll
  413. rcp-opt.ll
  414. read-global-variable-constant.ll
  415. reduction-intrinsics.ll
  416. redux-sync-f32.ll
  417. redux-sync.ll
  418. refl1.ll
  419. reg-copy.ll
  420. reg-types.ll
  421. reqnctapercluster-const-fold.ll
  422. reqntid-const-fold.ll
  423. reserved-smem-offset.ll
  424. ret-align-mismatch.ll
  425. rotate-add.ll
  426. rotate.ll
  427. rotate_64.ll
  428. rsqrt-opt.ll
  429. rsqrt.ll
  430. sad-intrins.ll
  431. scalar-to-vector.ll
  432. scalarize-non-coalescable-v2f32.ll
  433. sched1.ll
  434. sched2.ll
  435. setmaxnreg-sm100a.ll
  436. setmaxnreg.ll
  437. sext-in-reg.ll
  438. sext-params.ll
  439. sext-setcc.ll
  440. shfl-p.ll
  441. shfl-sync-p.ll
  442. shfl-sync.ll
  443. shfl.ll
  444. shift-opt.ll
  445. shift-parts.ll
  446. short-ptr.ll
  447. shuffle-vec-undef-init.ll
  448. simple-call.ll
  449. sm-version.ll
  450. speculative-execution-divergent-target.ll
  451. sqrt-approx.ll
  452. srl-bitcast-bv.ll
  453. st-addrspace.ll
  454. st-generic.ll
  455. st-param-imm.ll
  456. st_async_mbarrier.ll
  457. st_async_mbarrier_b128.ll
  458. st_async_release.ll
  459. st_async_release_multimem.ll
  460. st_bulk.ll
  461. stackaddress.ll
  462. stacksaverestore.ll
  463. store-retval.ll
  464. store-undef.ll
  465. sub-byte-constant-vector-convert.ll
  466. sub-byte-constant-vectors-i4-i2.ll
  467. surf-read-cuda.ll
  468. surf-read.ll
  469. surf-tex.py
  470. surf-write-cuda.ll
  471. surf-write.ll
  472. switch-loop-header.mir
  473. switch.ll
  474. symbol-naming.ll
  475. szext.ll
  476. tag-invariant-loads.ll
  477. TailDuplication-convergent.ll
  478. tanhf.ll
  479. tcgen05-alloc-dealloc-exclusive.ll
  480. tcgen05-alloc.ll
  481. tcgen05-commit-sm107.ll
  482. tcgen05-commit.ll
  483. tcgen05-cp.ll
  484. tcgen05-fence.ll
  485. tcgen05-ld-red.ll
  486. tcgen05-ld.ll
  487. tcgen05-mma-block-scale-invalid.ll
  488. tcgen05-mma-block-scale-ptx88-aa.ll
  489. tcgen05-mma-block-scale-ptx88.ll
  490. tcgen05-mma-block-scale.ll
  491. tcgen05-mma-collector-b-i8-invalid.ll
  492. tcgen05-mma-collector-b.ll
  493. tcgen05-mma-disable-output-lane-collector-b.ll
  494. tcgen05-mma-disable-output-lane-i8.ll
  495. tcgen05-mma-disable-output-lane.ll
  496. tcgen05-mma-i8.ll
  497. tcgen05-mma-invalid.ll
  498. tcgen05-mma-scale-d-invalid.ll
  499. tcgen05-mma-scale-d.ll
  500. tcgen05-mma-sp-collector-b-i8-invalid.ll
  501. tcgen05-mma-sp-collector-b.ll
  502. tcgen05-mma-sp-disable-output-lane-collector-b.ll
  503. tcgen05-mma-sp-mxf4-mxf4nvf4-kind.ll
  504. tcgen05-mma-ti16-kind.ll
  505. tcgen05-mma-ws-i8.ll
  506. tcgen05-mma-ws.ll
  507. tcgen05-mma.ll
  508. tcgen05-shift.ll
  509. tcgen05-st.ll
  510. tensormap_replace.ll
  511. tensormap_replace_invalid.ll
  512. tensormap_replace_sm_100a.ll
  513. tensormap_replace_sm_103a.ll
  514. tex-read-cuda.ll
  515. tex-read.ll
  516. texsurf-queries.ll
  517. thread-fence.ll
  518. tid-range.ll
  519. trunc-setcc.ll
  520. trunc-tofp.ll
  521. unaligned-param-load-store.ll
  522. unfold-masked-merge-vector-variablemask.ll
  523. unknown-intrinsic.ll
  524. unreachable.ll
  525. unrecognized-sm1x.ll
  526. upgrade-nvvm-annotations.ll
  527. used-bytes-mask.ll
  528. vaargs.ll
  529. variadics-backend.ll
  530. variadics-lowering.ll
  531. vec-param-load.ll
  532. vec8.ll
  533. vector-args.ll
  534. vector-call.ll
  535. vector-compare.ll
  536. vector-global.ll
  537. vector-loads.ll
  538. vector-returns.ll
  539. vector-select.ll
  540. vector-stores.ll
  541. vectorize-misaligned.ll
  542. vote.ll
  543. weak-global.ll
  544. weak-linkage.ll
  545. wgmma-sm90a-fence.ll
  546. wmma-ptx60-sm70.py
  547. wmma-ptx61-sm70.py
  548. wmma-ptx63-sm72.py
  549. wmma-ptx63-sm75.py
  550. wmma-ptx64-sm70.py
  551. wmma-ptx65-sm75.py
  552. wmma-ptx71-sm80.py
  553. wmma-ptx78-sm90.py
  554. wmma-ptx86-sm100a.py
  555. wmma-ptx86-sm101a.py
  556. wmma-ptx87-sm120a.py
  557. wmma-ptx88-sm100f.py
  558. wmma-ptx88-sm120a.py
  559. wmma-ptx88-sm120f.py
  560. wmma-ptx90-sm110f.py
  561. wmma-ptx91-sm120a.py
  562. wmma-ptx91-sm120f.py
  563. wmma.py
  564. zeroext-32bit.ll