[INFRA] Import NVIDIA/CCCL upstream as optimization reference library
CCCL (CUDA C++ Core Libraries) provides: - CUB: device/block/warp-level GPU primitives (reduce, scan, sort, topk) - Thrust: high-level parallel algorithms (transform_reduce, sort, scan) - libcudacxx: CUDA C++ standard library (atomics, barriers, memory) - cudax: experimental features (memory resources, allocators) - Tuning policies: per-SM hardware-specific algorithm parameters Competition optimization vectors mapped to CCCL: - Output TPS (83% weight): warp_reduce, block_reduce, device_topk - Input TPS (14% weight): device_scan, block_load, prefetch - Cache TPS (3% weight): prefix caching strategy patterns - Memory (0.9 util): pooled/cached/buddy allocators Source: https://github.com/NVIDIA/cccl (shallow clone, HEAD only) License: Apache-2.0
This commit is contained in:
94
cccl_upstream/cudax/CMakeLists.txt
Normal file
94
cccl_upstream/cudax/CMakeLists.txt
Normal file
@@ -0,0 +1,94 @@
|
||||
if (NOT CCCL_ENABLE_CUDAX)
|
||||
include(cmake/cudaxAddSubdir.cmake)
|
||||
return()
|
||||
endif()
|
||||
|
||||
cmake_minimum_required(VERSION 3.21)
|
||||
project(cudax LANGUAGES CXX CUDA)
|
||||
|
||||
option(
|
||||
cudax_ENABLE_HEADER_TESTING
|
||||
"Test that CUDA Experimental's public headers compile."
|
||||
ON
|
||||
)
|
||||
option(cudax_ENABLE_TESTING "Build CUDA Experimental's tests." ON)
|
||||
option(cudax_ENABLE_EXAMPLES "Build CUDA Experimental's examples." ON)
|
||||
option(cudax_ENABLE_PLACES "Enable standalone Places subproject" ON)
|
||||
option(cudax_ENABLE_CUDASTF "Enable CUDASTF subproject" ON)
|
||||
option(
|
||||
cudax_ENABLE_CUDASTF_CODE_GENERATION
|
||||
"Enable code generation using STF's parallel_for or launch with CUDA compiler."
|
||||
ON
|
||||
)
|
||||
option(
|
||||
cudax_ENABLE_CUDASTF_BOUNDSCHECK
|
||||
"Enable bounds checks for STF targets. Requires debug build."
|
||||
OFF
|
||||
)
|
||||
option(
|
||||
cudax_ENABLE_CUDASTF_MATHLIBS
|
||||
"Enable STF tests/examples that use cublas/cusolver."
|
||||
OFF
|
||||
)
|
||||
option(cudax_ENABLE_CUFILE "Enable cuFile in CUDA Experimental" ON)
|
||||
|
||||
if (cudax_ENABLE_CUFILE)
|
||||
if (WIN32)
|
||||
message(FATAL_ERROR "cuFile is not available on Windows.")
|
||||
endif()
|
||||
|
||||
if (CMAKE_VERSION VERSION_LESS "3.25.0")
|
||||
message(
|
||||
FATAL_ERROR
|
||||
"cuFile is not available before cmake 3.25.0, please, use newer cmake."
|
||||
)
|
||||
endif()
|
||||
|
||||
cccl_get_cudatoolkit()
|
||||
|
||||
if (CUDAToolkit_VERSION VERSION_LESS "12.9.0")
|
||||
message(FATAL_ERROR "cuFile support requires at least CUDA 12.9.")
|
||||
endif()
|
||||
|
||||
if (NOT TARGET CUDA::cuFile)
|
||||
message(
|
||||
FATAL_ERROR
|
||||
"CUDA::cuFile target required for requested cuFile support was not found."
|
||||
)
|
||||
endif()
|
||||
endif()
|
||||
|
||||
if (
|
||||
cudax_ENABLE_CUDASTF_BOUNDSCHECK
|
||||
AND NOT CMAKE_BUILD_TYPE MATCHES "Debug"
|
||||
AND NOT CMAKE_BUILD_TYPE MATCHES "RelWithDebInfo"
|
||||
)
|
||||
message(
|
||||
FATAL_ERROR
|
||||
"cudax_ENABLE_CUDASTF_BOUNDSCHECK requires a Debug build."
|
||||
)
|
||||
endif()
|
||||
|
||||
include(cmake/cudaxBuildCompilerTargets.cmake)
|
||||
if (cudax_ENABLE_PLACES)
|
||||
include(cmake/cudaxPlacesConfigureTarget.cmake)
|
||||
endif()
|
||||
if (cudax_ENABLE_CUDASTF)
|
||||
include(cmake/cudaxSTFConfigureTarget.cmake)
|
||||
endif()
|
||||
|
||||
if (cudax_ENABLE_HEADER_TESTING)
|
||||
include(cmake/cudaxHeaderTesting.cmake)
|
||||
endif()
|
||||
|
||||
if (cudax_ENABLE_TESTING)
|
||||
add_subdirectory(test)
|
||||
endif()
|
||||
|
||||
if (cudax_ENABLE_EXAMPLES)
|
||||
add_subdirectory(examples)
|
||||
endif()
|
||||
|
||||
if (CCCL_ENABLE_BENCHMARKS)
|
||||
add_subdirectory(benchmarks)
|
||||
endif()
|
||||
Reference in New Issue
Block a user