1st

Browse files

Files changed (12) hide show

.gitattributes +10 -0
losses.py +103 -0
step_144783 +3 -0
step_217174 +3 -0
step_289565 +3 -0
step_361957 +3 -0
step_434348 +3 -0
step_506739 +3 -0
step_579130 +3 -0
step_651522 +3 -0
step_72391 +3 -0
step_723914 +3 -0

.gitattributes CHANGED Viewed

@@ -33,3 +33,13 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
 *.zip filter=lfs diff=lfs merge=lfs -text
 *.zst filter=lfs diff=lfs merge=lfs -text
 *tfevents* filter=lfs diff=lfs merge=lfs -text

 *.zip filter=lfs diff=lfs merge=lfs -text
 *.zst filter=lfs diff=lfs merge=lfs -text
 *tfevents* filter=lfs diff=lfs merge=lfs -text
+step_144783 filter=lfs diff=lfs merge=lfs -text
+step_217174 filter=lfs diff=lfs merge=lfs -text
+step_289565 filter=lfs diff=lfs merge=lfs -text
+step_361957 filter=lfs diff=lfs merge=lfs -text
+step_434348 filter=lfs diff=lfs merge=lfs -text
+step_506739 filter=lfs diff=lfs merge=lfs -text
+step_579130 filter=lfs diff=lfs merge=lfs -text
+step_651522 filter=lfs diff=lfs merge=lfs -text
+step_72391 filter=lfs diff=lfs merge=lfs -text
+step_723914 filter=lfs diff=lfs merge=lfs -text

losses.py ADDED Viewed

	@@ -0,0 +1,103 @@

+from typing import Any, Tuple, Dict, Sequence, Optional
+import torch
+import torch.nn.functional as F
+from torch import nn
+import math
+IGNORE_LABEL_ID = -100
+def s(x, epsilon=1e-30):
+    return torch.where(
+        x<0,
+        1/(1-x+ epsilon),
+        x + 1
+    )
+def log_stablemax(x, dim=-1):
+    s_x = s(x)
+    return torch.log(s_x/torch.sum(s_x, dim=dim, keepdim=True))
+def stablemax_cross_entropy(logits, labels, ignore_index: int = -100, valid_mask=None):
+    logprobs = log_stablemax(logits.to(torch.float64), dim=-1)
+    if valid_mask is None:
+        valid_mask = (labels != ignore_index)
+    transformed_labels = torch.where(valid_mask, labels, 0)
+    prediction_logprobs = torch.gather(logprobs, index=transformed_labels.to(torch.long).unsqueeze(-1), dim=-1).squeeze(-1)
+    return -torch.where(valid_mask, prediction_logprobs, 0)
+def softmax_cross_entropy(logits, labels, ignore_index: int = -100):
+    # Cast logits to f32
+    # Flatten logits
+    return F.cross_entropy(logits.to(torch.float32).view(-1, logits.shape[-1]), labels.to(torch.long).view(-1), ignore_index=ignore_index, reduction="none").view(labels.shape)
+class ACTLossHead(nn.Module):
+    def __init__(self, model: nn.Module, loss_type: str):
+        super().__init__()
+        self.model = model
+        self.loss_fn = globals()[loss_type]
+    def initial_carry(self, *args, **kwargs):
+        return self.model.initial_carry(*args, **kwargs)  # type: ignore
+    def forward(
+        self,
+        return_keys: Sequence[str],
+        # Model args
+        **model_kwargs,
+    ) -> Tuple[Any, torch.Tensor, Dict[str, torch.Tensor], Optional[Dict[str, torch.Tensor]], torch.Tensor]:
+        # Model logits
+        # B x SeqLen x D
+        new_carry, outputs = self.model(**model_kwargs)
+        labels = new_carry.current_data["labels"]
+        with torch.no_grad():
+            # Preds
+            outputs["preds"] = torch.argmax(outputs["logits"], dim=-1)
+            # Correctness
+            mask = (labels != IGNORE_LABEL_ID)
+            loss_counts = mask.sum(-1)
+            loss_divisor = loss_counts.clamp_min(1).unsqueeze(-1)  # Avoid NaNs in division
+            is_correct = mask & (torch.argmax(outputs["logits"], dim=-1) == labels)
+            seq_is_correct = is_correct.sum(-1) == loss_counts
+            # Metrics (halted)
+            valid_metrics = new_carry.halted & (loss_counts > 0)
+            metrics = {
+                "count": valid_metrics.sum(),
+                "accuracy":       torch.where(valid_metrics, (is_correct.to(torch.float32) / loss_divisor).sum(-1), 0).sum(),
+                "exact_accuracy": (valid_metrics & seq_is_correct).sum(),
+                "q_halt_accuracy": (valid_metrics & ((outputs["q_halt_logits"] >= 0) == seq_is_correct)).sum(),
+                "steps":          torch.where(valid_metrics, new_carry.steps, 0).sum(),
+            }
+        # Losses
+        lm_loss = (self.loss_fn(outputs["logits"], labels, ignore_index=IGNORE_LABEL_ID, valid_mask=mask) / loss_divisor).sum()
+        q_halt_loss = F.binary_cross_entropy_with_logits(outputs["q_halt_logits"], seq_is_correct.to(outputs["q_halt_logits"].dtype), reduction="sum")
+        metrics.update({
+            "lm_loss": lm_loss.detach(),
+            "q_halt_loss": q_halt_loss.detach(),
+        })
+        # Q continue (bootstrapping target loss); Alexia: This fits Q-learning, but seems totally unecessary
+        q_continue_loss = 0
+        if "target_q_continue" in outputs:
+            q_continue_loss = F.binary_cross_entropy_with_logits(outputs["q_continue_logits"], outputs["target_q_continue"], reduction="sum")
+            metrics["q_continue_loss"] = q_continue_loss.detach()
+        # Filter outputs for return
+        detached_outputs = {k: outputs[k].detach() for k in return_keys if k in outputs}
+        return new_carry, lm_loss + 0.5 * (q_halt_loss + q_continue_loss), metrics, detached_outputs, new_carry.halted.all()

step_144783 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:9dcc0b0bf3fb7d64c567b59df6ac196706d224b6498f379df945ac3ed6905a9b
+size 2467988405

step_217174 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:af2831ec79ccef44854b775ec65adc92eaa05c58a41fa2aabf9f6c5117c05a41
+size 2467988405

step_289565 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:26dc32b7d05c8da50de0eb489d4e5c72e6348f78bfbedbfad347c80f5f95507f
+size 2467988405

step_361957 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:04083d75970bbdaf4f6649ba7fa3d09136eeebe955149a4e371e24f4e453be7f
+size 2467988405

step_434348 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d371fc16816ea73db48eaade70290375feec060e98f211280da69d3b29792ad9
+size 2467988405

step_506739 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:efa317c0d741402096f58d0833ab42f619630929f7aebb7cc3e8bbda3899e586
+size 2467988405

step_579130 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:0c9b2be79abe8123de922fc34266e52324220dda4ab8addfc88652fa2188a04e
+size 2467988405

step_651522 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:0cf3dfcbaa96d69867db31ef68c7d5015fdf40cad956c973ab34725b10a3fb4b
+size 2467988405

step_72391 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:27025e2f399c310ba22bc3711396c0d67e664e245ebd515737aca7e110fbdc40
+size 2467988386

step_723914 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:15a35887067ad27145bf8e8551b0c0deb30c12b19c7087082556e5358d5d0f78
+size 2467988405