revert: restore to 8c8c0286 (last confirmed build success)

Revert LD_PRELOAD addition and patch_xformers profiling skip.
Need to identify which change caused build failure before re-adding.
This commit is contained in:
project6-dev
2026-08-13 15:08:40 +00:00
parent bed1fc4d54
commit aebc660a10
2 changed files with 0 additions and 6 deletions

View File

@@ -49,6 +49,4 @@ env:
value: '1'
- name: PYTORCH_CUDA_ALLOC_CONF
value: max_split_size_mb:512
- name: LD_PRELOAD
value: /workspace/qwen3_6_scripts/cccl_preload/libcccl_allocator.so

View File

@@ -200,10 +200,6 @@ FALLBACK_METHOD = '''
max_seqlen = max(seq_lens_list)
try:
# Skip flash_attn during profiling — OOMs on large dummy batch
import os
if os.environ.get("BI100_IN_STARTUP_PROFILE") == "1":
raise RuntimeError("skip flash_attn during profiling")
out = _ixf.flash_attn_varlen_func(
q_flat.to(torch.float16),
k_flat.to(torch.float16),